aboutsummaryrefslogtreecommitdiff
path: root/gcc
diff options
context:
space:
mode:
authorRichard Biener <rguenther@suse.de>2023-08-21 14:09:48 +0200
committerRichard Biener <rguenther@suse.de>2023-08-21 15:10:06 +0200
commit2eaebcf3df4dba0bfa379b7adde455710e9b3e41 (patch)
tree70d1cee3060c03f25f4adc695f69ec4fb099d412 /gcc
parente4e6a92407432cc19db73f8ce108243007d614f0 (diff)
downloadgcc-2eaebcf3df4dba0bfa379b7adde455710e9b3e41.zip
gcc-2eaebcf3df4dba0bfa379b7adde455710e9b3e41.tar.gz
gcc-2eaebcf3df4dba0bfa379b7adde455710e9b3e41.tar.bz2
Fix FAIL: gcc.target/i386/pr87007-5.c
The following fixes the gcc.target/i386/pr87007-5.c testcase which changed code generation again after the recent sinking improvements. We now have vxorps %xmm0, %xmm0, %xmm0 vsqrtsd d2(%rip), %xmm0, %xmm0 and a necessary xor again in one case, the other vsqrtsd has a register source and a properly zeroing load: vmovsd d3(%rip), %xmm0 testl %esi, %esi jg .L11 .L3: vsqrtsd %xmm0, %xmm0, %xmm0 the following patch adjusts the scan. * gcc.target/i386/pr87007-5.c: Update comment, adjust subtest.
Diffstat (limited to 'gcc')
-rw-r--r--gcc/testsuite/gcc.target/i386/pr87007-5.c6
1 files changed, 4 insertions, 2 deletions
diff --git a/gcc/testsuite/gcc.target/i386/pr87007-5.c b/gcc/testsuite/gcc.target/i386/pr87007-5.c
index a6cdf11..8f2dc94 100644
--- a/gcc/testsuite/gcc.target/i386/pr87007-5.c
+++ b/gcc/testsuite/gcc.target/i386/pr87007-5.c
@@ -1,6 +1,8 @@
/* { dg-do compile } */
/* { dg-options "-Ofast -march=skylake-avx512 -mfpmath=sse -fno-tree-vectorize -fdump-tree-cddce3-details -fdump-tree-lsplit-optimized" } */
-/* Load of d2/d3 is hoisted out, vrndscalesd will reuse loades register to avoid partial dependence. */
+/* Load of d2/d3 is hoisted out, the loop is split, store of d1 and sqrt
+ are sunk out of the loop and the loop is elided. One vsqrtsd with
+ memory operand needs a xor to avoid partial dependence. */
#include<math.h>
@@ -17,4 +19,4 @@ foo (int n, int k)
/* { dg-final { scan-tree-dump "optimized: loop split" "lsplit" } } */
/* { dg-final { scan-tree-dump-times "removing loop" 2 "cddce3" } } */
-/* { dg-final { scan-assembler-times "vxorps\[^\n\r\]*xmm\[0-9\]" 0 } } */
+/* { dg-final { scan-assembler-times "vxorps\[^\n\r\]*xmm\[0-9\]" 1 } } */