//==--- sycl_esimd_mixed_unnamed.cpp - DPC++ ESIMD on-device test ---------==//
//
// Part of the LLVM Project, under the Apache License v2.0 with LLVM Exceptions.
// See https://llvm.org/LICENSE.txt for license information.
// SPDX-License-Identifier: Apache-2.0 WITH LLVM-exception
//
//===----------------------------------------------------------------------===//
// This is basic test for mixing unnamed SYCL and ESIMD kernels in the same
// source and in the same program .

// RUN: %{build} -o %t.out
// RUN: %{run} %t.out

#include <iostream>
#include <sycl/ext/intel/esimd.hpp>
#include <sycl/detail/core.hpp>

using namespace ::sycl;

bool checkResult(const std::vector<float> &A, int Inc) {
  int err_cnt = 0;
  unsigned Size = A.size();

  for (unsigned i = 0; i < Size; ++i) {
    if (A[i] != i + Inc)
      if (++err_cnt < 10)
        std::cerr << "failed at A[" << i << "]: " << A[i] << " != " << i + Inc
                  << "\n";
  }

  if (err_cnt > 0) {
    std::cout << "  pass rate: "
              << ((float)(Size - err_cnt) / (float)Size) * 100.0f << "% ("
              << (Size - err_cnt) << "/" << Size << ")\n";
    return false;
  }
  return true;
}

int main(void) {
  constexpr unsigned Size = 32;
  constexpr unsigned VL = 16;

  std::vector<float> A(Size);

  for (unsigned i = 0; i < Size; ++i) {
    A[i] = i;
  }

  try {
    buffer<float, 1> bufa(A.data(), range<1>(Size));

    // We need that many workgroups
    ::sycl::range<1> GlobalRange{Size};
    // We need that many threads in each group
    ::sycl::range<1> LocalRange{1};

    queue q;

    auto dev = q.get_device();
    std::cout << "Running on " << dev.get_info<info::device::name>() << "\n";

    auto e = q.submit([&](handler &cgh) {
      auto PA = bufa.get_access<access::mode::read_write>(cgh);
      cgh.parallel_for(GlobalRange * LocalRange,
                       [=](id<1> i) { PA[i] = PA[i] + 1; });
    });
    e.wait();
  } catch (::sycl::exception const &e) {
    std::cout << "SYCL exception caught: " << e.what() << '\n';
    return 2;
  }

  if (checkResult(A, 1)) {
    std::cout << "SYCL kernel passed\n";
  } else {
    std::cout << "SYCL kernel failed\n";
    return 1;
  }

  try {
    buffer<float, 1> bufa(A.data(), range<1>(Size));

    // We need that many workgroups
    ::sycl::range<1> GlobalRange{Size / VL};
    // We need that many threads in each group
    ::sycl::range<1> LocalRange{1};

    queue q;

    auto dev = q.get_device();
    std::cout << "Running on " << dev.get_info<info::device::name>() << "\n";

    auto e = q.submit([&](handler &cgh) {
      auto PA = bufa.get_access<access::mode::read_write>(cgh);
      cgh.parallel_for(GlobalRange * LocalRange,
                       [=](id<1> i) SYCL_ESIMD_KERNEL {
                         using namespace sycl::ext::intel::esimd;
                         unsigned int offset = i * VL * sizeof(float);
                         simd<float, VL> va;
                         va.copy_from(PA, offset);
                         simd<float, VL> vc = va + 1;
                         vc.copy_to(PA, offset);
                       });
    });
    e.wait();
  } catch (::sycl::exception const &e) {
    std::cout << "SYCL exception caught: " << e.what() << '\n';
    return 2;
  }

  if (checkResult(A, 2)) {
    std::cout << "ESIMD kernel passed\n";
  } else {
    std::cout << "ESIMD kernel failed\n";
    return 1;
  }
  return 0;
}
