Middle-end: Adjust decrement IV style partial vectorization COST model

author Juzhe-Zhong <juzhe.zhong@rivai.ai>

Wed, 13 Dec 2023 09:21:07 +0000 (17:21 +0800)

committer Pan Li <pan2.li@intel.com>

Wed, 13 Dec 2023 11:51:59 +0000 (19:51 +0800)
author Juzhe-Zhong <juzhe.zhong@rivai.ai>
Wed, 13 Dec 2023 09:21:07 +0000 (17:21 +0800)
committer Pan Li <pan2.li@intel.com>
Wed, 13 Dec 2023 11:51:59 +0000 (19:51 +0800)
diff --git a/gcc/testsuite/gcc.dg/vect/costmodel/riscv/rvv/pr111317.c b/gcc/testsuite/gcc.dg/vect/costmodel/riscv/rvv/pr111317.c

new file mode 100644 (file)

index 0000000..d4bea24
--- /dev/null
+++ b/gcc/testsuite/gcc.dg/vect/costmodel/riscv/rvv/pr111317.c
@@ -0,0 +1,12 @@
+/* { dg-do compile } */
+/* { dg-options "-march=rv64gcv -mabi=lp64d -O3 -ftree-vectorize --param=riscv-autovec-lmul=m1" } */
+
+void
+foo (char *__restrict a, short *__restrict b, int n)
+{
+  for (int i = 0; i < n; i++)
+    b[i] = (short) a[i];
+}
+
+/* { dg-final { scan-assembler-times {vsetvli\s+[a-x0-9]+,\s*[a-x0-9]+,\s*e16,\s*m1,\s*t[au],\s*m[au]} 1 } } */
+/* { dg-final { scan-assembler-times {vsetvli} 1 } } */
diff --git a/gcc/tree-vect-loop.cc b/gcc/tree-vect-loop.cc

index 6261cd1be1ddbf6ebbf95796313011e4934d7a87..19e38b8637b0535a45b50f10fd924de4435fbb72 100644 (file)
--- a/gcc/tree-vect-loop.cc
+++ b/gcc/tree-vect-loop.cc
@@ -4870,10 +4870,21 @@ vect_estimate_min_profitable_iters (loop_vec_info loop_vinfo,
             if (partial_load_store_bias != 0)
               body_stmts += 1;
  
-           /* Each may need two MINs and one MINUS to update lengths in body
-              for next iteration.  */
+           unsigned int length_update_cost = 0;
+           if (LOOP_VINFO_USING_DECREMENTING_IV_P (loop_vinfo))
+             /* For decrement IV style, we use a single SELECT_VL since
+                beginning to calculate the number of elements need to be
+                processed in current iteration, and a SHIFT operation to
+                compute the next memory address instead of adding vectorization
+                factor.  */
+             length_update_cost = 2;
+           else
+             /* For increment IV stype, Each may need two MINs and one MINUS to
+                update lengths in body for next iteration.  */
+             length_update_cost = 3;
+
             if (need_iterate_p)
-             body_stmts += 3 * num_vectors;
+             body_stmts += length_update_cost * num_vectors;
           }
  
        (void) add_stmt_cost (target_cost_data, prologue_stmts,
author	Juzhe-Zhong <juzhe.zhong@rivai.ai>
	Wed, 13 Dec 2023 09:21:07 +0000 (17:21 +0800)
committer	Pan Li <pan2.li@intel.com>
	Wed, 13 Dec 2023 11:51:59 +0000 (19:51 +0800)
gcc/testsuite/gcc.dg/vect/costmodel/riscv/rvv/pr111317.c	[new file with mode: 0644]	patch \| blob
gcc/tree-vect-loop.cc		patch \| blob \| blame \| history