Add block evaluationto CwiseUnaryOp and add PreferBlockAccess enum to all evaluators

author: Eugene Zhulenev <ezhulenev@google.com> 2018-08-10 16:53:36 -0700
committer: Eugene Zhulenev <ezhulenev@google.com> 2018-08-10 16:53:36 -0700
commit: f2209d06e428e0691de71f30fc2db4cb29191cd2 (patch)
tree: 37d7294a61f80c87389e8e930700a549554afe51 /unsupported/test/cxx11_tensor_executor.cpp
parent: cfaedb38cd662def3b5684a20965b3bc1b0d6a3f (diff)
1 files changed, 31 insertions, 0 deletions
diff --git a/unsupported/test/cxx11_tensor_executor.cpp b/unsupported/test/cxx11_tensor_executor.cpp
index d16ae4d11..1dc18220c 100644
--- a/unsupported/test/cxx11_tensor_executor.cpp
+++ b/unsupported/test/cxx11_tensor_executor.cpp
@@ -31,6 +31,33 @@ static array<Index, NumDims> RandomDims(int min_dim = 1, int max_dim = 20) {
 
 template <typename T, int NumDims, typename Device, bool Vectorizable,
           bool Tileable, int Layout>
+static void test_execute_unary_expr(Device d) {
+  static constexpr int Options = 0 | Layout;
+
+  // Pick a large enough tensor size to bypass small tensor block evaluation
+  // optimization.
+  auto dims = RandomDims<NumDims>(50 / NumDims, 100 / NumDims);
+
+  Tensor<T, NumDims, Options, Index> src(dims);
+  Tensor<T, NumDims, Options, Index> dst(dims);
+
+  src.setRandom();
+  const auto expr = src.square();
+
+  using Assign = TensorAssignOp<decltype(dst), const decltype(expr)>;
+  using Executor =
+      internal::TensorExecutor<const Assign, Device, Vectorizable, Tileable>;
+
+  Executor::run(Assign(dst, expr), d);
+
+  for (Index i = 0; i < dst.dimensions().TotalSize(); ++i) {
+    T square = src.coeff(i) * src.coeff(i);
+    VERIFY_IS_EQUAL(square, dst.coeff(i));
+  }
+}
+
+template <typename T, int NumDims, typename Device, bool Vectorizable,
+          bool Tileable, int Layout>
 static void test_execute_binary_expr(Device d)
 {
   static constexpr int Options = 0 | Layout;
@@ -445,6 +472,10 @@ EIGEN_DECLARE_TEST(cxx11_tensor_executor) {
   Eigen::ThreadPool tp(num_threads);
   Eigen::ThreadPoolDevice tp_device(&tp, num_threads);
 
+  CALL_SUBTEST_COMBINATIONS(test_execute_unary_expr, float, 3);
+  CALL_SUBTEST_COMBINATIONS(test_execute_unary_expr, float, 4);
+  CALL_SUBTEST_COMBINATIONS(test_execute_unary_expr, float, 5);
+
   CALL_SUBTEST_COMBINATIONS(test_execute_binary_expr, float, 3);
   CALL_SUBTEST_COMBINATIONS(test_execute_binary_expr, float, 4);
   CALL_SUBTEST_COMBINATIONS(test_execute_binary_expr, float, 5);
author	Eugene Zhulenev <ezhulenev@google.com>	2018-08-10 16:53:36 -0700
committer	Eugene Zhulenev <ezhulenev@google.com>	2018-08-10 16:53:36 -0700
commit	f2209d06e428e0691de71f30fc2db4cb29191cd2 (patch)
tree	37d7294a61f80c87389e8e930700a549554afe51 /unsupported/test/cxx11_tensor_executor.cpp
parent	cfaedb38cd662def3b5684a20965b3bc1b0d6a3f (diff)