|
MultiSource/Benchmarks/Trimaran/enc-md5/enc-md5
Profile
|
-75.99% |
12.643 |
3.035 |
0.000 |
-75.99% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_int64_t_
|
-27.67% |
1044685.999 |
755617.828 |
1229.059 |
-27.68% |
1229.059 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_int64_t_
|
-27.62% |
1044201.728 |
755829.272 |
1042.917 |
-27.63% |
1042.917 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_int32_t_
|
-26.92% |
515789.113 |
376939.778 |
349.669 |
-26.91% |
349.669 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_int32_t_
|
-26.87% |
515563.776 |
377041.890 |
459.943 |
-26.91% |
459.943 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint64_t>/10
|
-26.65% |
38.566 |
28.287 |
0.003 |
-26.79% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint64_t>/16
|
-25.78% |
53.203 |
39.486 |
0.064 |
-25.80% |
0.064 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_uint8_t_
|
-25.40% |
125023.006 |
93273.332 |
162.975 |
-25.33% |
162.975 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_uint8_t_
|
-25.39% |
125003.557 |
93262.458 |
165.758 |
-25.34% |
165.758 |
|
SingleSource/Benchmarks/Shootout/Shootout-ary3
Profile
|
-23.48% |
0.984 |
0.753 |
0.001 |
-23.52% |
0.001 |
|
MultiSource/Benchmarks/Ptrdist/yacr2/yacr2
Profile
|
-23.08% |
2.837 |
2.182 |
0.001 |
-23.05% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint64_t>/16
|
-20.76% |
32.542 |
25.785 |
0.011 |
-20.77% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint16_t>/10
|
-20.00% |
9.386 |
7.508 |
0.017 |
-20.00% |
0.017 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint16_t>/16
|
-19.80% |
9.385 |
7.527 |
0.010 |
-19.80% |
0.010 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint8_t>/16
|
-19.03% |
60.092 |
48.659 |
0.097 |
-19.03% |
0.097 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint8_t>/10
|
-18.87% |
41.302 |
33.510 |
0.149 |
-18.87% |
0.149 |
|
MultiSource/Benchmarks/Olden/em3d/em3d
Profile
|
-17.93% |
18.299 |
15.018 |
0.009 |
-17.95% |
0.009 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint64_t>/10
|
-17.62% |
10.017 |
8.252 |
0.009 |
-17.58% |
0.009 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint64_t>/16
|
-17.48% |
10.012 |
8.262 |
0.006 |
-17.55% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint64_t>/10
|
-17.45% |
24.393 |
20.136 |
0.027 |
-17.38% |
0.027 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/28
|
-17.06% |
104.489 |
86.662 |
0.000 |
-17.07% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/16
|
-15.41% |
66.944 |
56.626 |
0.001 |
-15.43% |
0.001 |
|
MultiSource/Benchmarks/SciMark2-C/scimark2
Profile
|
-15.04% |
166.584 |
141.536 |
0.045 |
-15.09% |
0.045 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint16_t>/10
|
-13.81% |
35.617 |
30.700 |
0.004 |
-13.89% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/10
|
-13.69% |
48.181 |
41.587 |
0.011 |
-13.69% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint8_t>/16
|
-10.23% |
34.012 |
30.533 |
0.392 |
-10.18% |
0.392 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_novec_int64_t_
|
-9.86% |
1044836.010 |
941793.968 |
1044.869 |
-9.93% |
1044.869 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_autovec_int64_t_
|
-9.77% |
1044678.309 |
942568.112 |
939.055 |
-9.91% |
939.055 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_autovec_int32_t_
|
-8.74% |
515667.865 |
470599.875 |
463.183 |
-8.68% |
463.183 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_novec_int32_t_
|
-8.65% |
515501.258 |
470929.460 |
254.732 |
-8.58% |
254.732 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint16_t>/16
|
-8.54% |
49.964 |
45.697 |
0.005 |
-8.51% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/28
|
-8.50% |
157.949 |
144.529 |
0.004 |
-8.51% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/256
|
-8.39% |
1207.584 |
1106.243 |
0.017 |
-8.40% |
0.017 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/51
|
-8.38% |
272.480 |
249.642 |
0.004 |
-8.38% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/999
|
-8.32% |
2364.895 |
2168.127 |
0.049 |
-8.33% |
0.049 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/999
|
-8.30% |
4642.751 |
4257.631 |
0.432 |
-8.38% |
0.432 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/256
|
-8.24% |
1245.773 |
1143.107 |
0.009 |
-8.24% |
0.009 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/999
|
-8.23% |
4804.538 |
4409.293 |
0.616 |
-8.24% |
0.616 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/256
|
-8.23% |
642.920 |
590.038 |
0.005 |
-8.22% |
0.005 |
|
MultiSource/Benchmarks/ASC_Sequoia/AMGmk/AMGmk
Profile
|
-8.20% |
43.392 |
39.833 |
0.011 |
-8.34% |
0.011 |
|
SingleSource/Benchmarks/Adobe-C++/simple_types_constant_folding
Profile
|
-8.12% |
2.365 |
2.173 |
0.001 |
-8.17% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/51
|
-8.11% |
266.240 |
244.651 |
0.004 |
-8.11% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/999
|
-8.09% |
2443.245 |
2245.577 |
0.053 |
-8.09% |
0.053 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/256
|
-8.07% |
662.260 |
608.813 |
0.026 |
-8.08% |
0.026 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C++/HPCCG/HPCCG
Profile
|
-8.01% |
4.652 |
4.280 |
0.053 |
-8.00% |
0.053 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/28
|
-7.85% |
155.497 |
143.284 |
0.004 |
-7.85% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/16
|
-7.50% |
100.112 |
92.604 |
0.001 |
-7.50% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/51
|
-7.42% |
173.008 |
160.172 |
0.004 |
-7.37% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/10
|
-7.28% |
81.651 |
75.711 |
0.001 |
-7.28% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/51
|
-7.21% |
169.253 |
157.053 |
0.004 |
-7.21% |
0.004 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C++/miniFE/miniFE
Profile
|
-6.90% |
17.838 |
16.607 |
0.055 |
-6.96% |
0.055 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/256
|
-6.69% |
51.634 |
48.180 |
0.007 |
-6.66% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/28
|
-6.68% |
51.629 |
48.179 |
0.001 |
-6.63% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/16
|
-6.67% |
89.163 |
83.218 |
0.001 |
-6.67% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/28
|
-6.63% |
122.634 |
114.504 |
0.000 |
-6.63% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/51
|
-6.61% |
51.588 |
48.177 |
0.004 |
-6.59% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/999
|
-6.55% |
51.556 |
48.180 |
0.003 |
-6.68% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/28
|
-6.25% |
120.135 |
112.630 |
0.011 |
-6.25% |
0.011 |
|
External/SPEC/CINT2017rate/525.x264_r/525.x264_r
Profile
|
-5.99% |
144.568 |
135.910 |
0.075 |
-4.41% |
0.075 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/999
|
-5.78% |
1453.679 |
1369.715 |
0.194 |
-5.78% |
0.194 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/10
|
-5.69% |
76.961 |
72.582 |
0.000 |
-5.69% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint64_t>/999
|
-5.67% |
5517.533 |
5204.603 |
0.712 |
-5.67% |
0.712 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/256
|
-5.64% |
1087.497 |
1026.207 |
0.031 |
-5.64% |
0.031 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/999
|
-5.57% |
2847.697 |
2689.168 |
0.193 |
-5.57% |
0.193 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/51
|
-5.56% |
241.826 |
228.377 |
0.025 |
-5.57% |
0.025 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/256
|
-5.55% |
744.590 |
703.242 |
0.043 |
-5.55% |
0.043 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/256
|
-5.55% |
745.302 |
703.927 |
0.015 |
-5.55% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/256
|
-5.55% |
580.970 |
548.744 |
0.005 |
-5.55% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/999
|
-5.54% |
2848.364 |
2690.520 |
0.026 |
-5.54% |
0.026 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/256
|
-5.54% |
745.184 |
703.919 |
0.012 |
-5.54% |
0.012 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/999
|
-5.53% |
1450.989 |
1370.821 |
0.626 |
-5.53% |
0.626 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/51
|
-5.52% |
170.193 |
160.801 |
0.003 |
-5.52% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/28
|
-5.50% |
142.350 |
134.525 |
0.001 |
-5.50% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/51
|
-5.50% |
170.818 |
161.430 |
0.004 |
-5.50% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/51
|
-5.50% |
170.812 |
161.426 |
0.002 |
-5.50% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/256
|
-5.48% |
399.842 |
377.911 |
0.012 |
-5.48% |
0.012 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/256
|
-5.47% |
400.462 |
378.544 |
0.005 |
-5.48% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/256
|
-5.47% |
400.448 |
378.545 |
0.039 |
-5.47% |
0.039 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/999
|
-5.46% |
1450.729 |
1371.572 |
0.131 |
-5.45% |
0.131 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/999
|
-5.45% |
2846.974 |
2691.862 |
0.053 |
-5.45% |
0.053 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint64_t>/51
|
-5.45% |
310.350 |
293.444 |
0.005 |
-5.45% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/16
|
-5.41% |
92.604 |
87.595 |
0.003 |
-5.41% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/999
|
-5.39% |
52.247 |
49.430 |
0.000 |
-5.39% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/16
|
-5.39% |
52.245 |
49.428 |
0.002 |
-5.39% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/10
|
-5.39% |
52.247 |
49.429 |
0.000 |
-5.39% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/256
|
-5.39% |
52.247 |
49.430 |
0.001 |
-5.39% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/28
|
-5.39% |
52.245 |
49.429 |
0.002 |
-5.39% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/51
|
-5.39% |
52.246 |
49.431 |
0.001 |
-5.39% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint64_t>/999
|
-5.38% |
2804.270 |
2653.482 |
0.061 |
-5.39% |
0.061 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/999
|
-5.35% |
2131.638 |
2017.499 |
0.022 |
-5.37% |
0.022 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/999
|
-5.32% |
3006.023 |
2846.199 |
0.529 |
-5.32% |
0.529 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/51
|
-5.12% |
110.130 |
104.490 |
0.002 |
-5.13% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/999
|
-5.11% |
4156.189 |
3943.813 |
1.678 |
-5.98% |
1.678 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/51
|
-5.08% |
110.746 |
105.118 |
0.002 |
-5.08% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/10
|
-5.05% |
16.893 |
16.041 |
0.093 |
-5.05% |
0.093 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/51
|
-5.01% |
156.107 |
148.292 |
0.005 |
-5.01% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/999
|
-4.99% |
1528.786 |
1452.559 |
0.040 |
-5.01% |
0.040 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/51
|
-4.92% |
178.323 |
169.546 |
0.008 |
-4.92% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/16
|
-4.88% |
83.534 |
79.461 |
0.002 |
-4.87% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint64_t>/51
|
-4.78% |
196.465 |
187.082 |
0.007 |
-4.78% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/999
|
-4.76% |
39.419 |
37.541 |
0.002 |
-4.76% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/51
|
-4.76% |
39.420 |
37.542 |
0.001 |
-4.76% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/256
|
-4.76% |
39.418 |
37.541 |
0.001 |
-4.77% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/999
|
-4.69% |
40.046 |
38.167 |
0.001 |
-4.69% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/51
|
-4.69% |
40.045 |
38.167 |
0.001 |
-4.69% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/999
|
-4.69% |
40.046 |
38.168 |
0.001 |
-4.69% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/256
|
-4.69% |
40.045 |
38.167 |
0.001 |
-4.69% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/51
|
-4.69% |
40.044 |
38.167 |
0.001 |
-4.69% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/256
|
-4.69% |
40.043 |
38.167 |
0.001 |
-4.69% |
0.001 |
|
MultiSource/Benchmarks/TSVC/IndirectAddressing-flt/IndirectAddressing-flt
Profile
|
-4.48% |
15.694 |
14.991 |
0.016 |
-4.80% |
0.016 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/16
|
-4.47% |
16.893 |
16.138 |
0.075 |
-4.47% |
0.075 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/28
|
-4.45% |
112.625 |
107.618 |
0.003 |
-4.45% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/51
|
-4.41% |
113.878 |
108.858 |
0.004 |
-4.42% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/256
|
-4.35% |
50.368 |
48.178 |
0.001 |
-4.36% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/999
|
-4.35% |
50.368 |
48.178 |
0.001 |
-4.35% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/51
|
-4.35% |
50.369 |
48.179 |
0.000 |
-4.35% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/16
|
-4.35% |
50.368 |
48.178 |
0.001 |
-4.35% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/28
|
-4.35% |
50.368 |
48.180 |
0.000 |
-4.35% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/28
|
-4.03% |
16.893 |
16.212 |
0.067 |
-4.04% |
0.067 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/51
|
-4.00% |
109.501 |
105.116 |
0.002 |
-4.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint8_t>/10
|
-3.93% |
24.539 |
23.574 |
0.346 |
-3.93% |
0.346 |
|
MultiSource/Applications/JM/lencod/lencod
Profile
|
-3.86% |
35.701 |
34.324 |
0.032 |
-3.98% |
0.032 |
|
SingleSource/Benchmarks/Misc-C++/oopack_v1p8
Profile
|
-3.83% |
0.478 |
0.460 |
0.000 |
-3.89% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/51
|
-3.40% |
40.045 |
38.682 |
0.044 |
-3.40% |
0.044 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint64_t>/999
|
-3.30% |
56.940 |
55.060 |
0.002 |
-3.30% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint64_t>/51
|
-3.30% |
56.938 |
55.061 |
0.001 |
-3.30% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/256
|
-3.27% |
20.021 |
19.367 |
0.134 |
-3.27% |
0.134 |
|
MultiSource/Benchmarks/MiBench/consumer-lame/consumer-lame
Profile
|
-3.24% |
0.682 |
0.660 |
0.000 |
-3.27% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/999
|
-3.17% |
40.044 |
38.776 |
0.002 |
-3.18% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC15
|
-2.94% |
21.274 |
20.648 |
0.001 |
-2.94% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/10
|
-2.89% |
40.674 |
39.500 |
0.024 |
-2.89% |
0.024 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC16
|
-2.88% |
21.260 |
20.649 |
0.001 |
-0.12% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/10
|
-2.74% |
41.295 |
40.164 |
0.008 |
-2.74% |
0.008 |
|
MultiSource/Benchmarks/VersaBench/8b10b/8b10b
Profile
|
-2.66% |
9.299 |
9.052 |
0.001 |
-2.66% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint16_t>/16
|
-2.53% |
29.557 |
28.808 |
0.167 |
-2.75% |
0.167 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC16
|
-2.52% |
21.183 |
20.648 |
0.000 |
-2.67% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4BigLoopTC16
|
-2.42% |
273.908 |
267.291 |
2.238 |
0.64% |
2.238 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC16
|
-2.41% |
273.894 |
267.294 |
2.342 |
0.56% |
2.342 |
|
External/SPEC/CINT2017rate/523.xalancbmk_r/523.xalancbmk_r
Profile
|
-2.37% |
211.504 |
206.487 |
0.747 |
-1.46% |
0.747 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC16
|
-2.37% |
273.838 |
267.362 |
2.302 |
0.67% |
2.302 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC32
|
-2.18% |
28.784 |
28.157 |
0.002 |
-2.17% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC31
|
-2.18% |
28.784 |
28.157 |
0.000 |
-2.17% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC32
|
-2.17% |
28.782 |
28.157 |
0.000 |
-2.18% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC31
|
-2.17% |
28.783 |
28.158 |
0.000 |
-2.17% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC31
|
-2.17% |
28.782 |
28.157 |
0.000 |
-2.17% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC31
|
-2.17% |
28.781 |
28.158 |
0.001 |
-2.16% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC32
|
-2.16% |
28.779 |
28.157 |
0.001 |
-2.12% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC32
|
-2.16% |
28.780 |
28.157 |
0.000 |
-2.15% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC31
|
-2.16% |
28.778 |
28.157 |
0.001 |
-2.16% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC31
|
-2.16% |
28.779 |
28.158 |
0.000 |
-2.16% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC31
|
-2.16% |
28.778 |
28.156 |
0.001 |
-2.15% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC32
|
-2.12% |
28.767 |
28.157 |
0.002 |
-2.14% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC16
|
-2.10% |
21.093 |
20.649 |
0.001 |
-1.09% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/16
|
-2.10% |
55.687 |
54.517 |
0.031 |
-2.11% |
0.031 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/16
|
-2.07% |
56.312 |
55.149 |
0.006 |
-2.07% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC32
|
-2.06% |
28.748 |
28.157 |
0.001 |
-1.98% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/10
|
-2.05% |
30.560 |
29.934 |
0.014 |
-2.37% |
0.014 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC15
|
-2.03% |
21.076 |
20.648 |
0.002 |
-2.05% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/256
|
-2.02% |
353.548 |
346.408 |
0.027 |
-2.50% |
0.027 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC32
|
-1.88% |
28.696 |
28.157 |
0.001 |
-1.99% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC15
|
-1.86% |
21.040 |
20.648 |
0.000 |
-1.36% |
0.000 |
|
MultiSource/Applications/sgefa/sgefa
Profile
|
-1.85% |
0.641 |
0.629 |
0.002 |
-2.07% |
0.002 |
|
SingleSource/Benchmarks/CoyoteBench/huffbench
Profile
|
-1.78% |
50.634 |
49.733 |
0.003 |
-1.77% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC15
|
-1.75% |
272.992 |
268.208 |
1.586 |
-1.37% |
1.586 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/16
|
-1.66% |
38.085 |
37.452 |
0.001 |
-1.54% |
0.001 |
|
SingleSource/Benchmarks/Shootout/Shootout-matrix
Profile
|
-1.60% |
4.692 |
4.617 |
0.000 |
-1.61% |
0.000 |
|
External/SPEC/CFP2017rate/511.povray_r/511.povray_r
Profile
|
-1.59% |
28.677 |
28.220 |
0.084 |
-1.45% |
0.084 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/10
|
-1.57% |
48.940 |
48.170 |
0.004 |
-1.50% |
0.004 |
|
MultiSource/Applications/hexxagon/hexxagon
Profile
|
-1.57% |
8.512 |
8.379 |
0.000 |
-1.58% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC63
|
-1.43% |
43.801 |
43.173 |
0.001 |
-1.43% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC64
|
-1.43% |
43.801 |
43.173 |
0.000 |
-1.43% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC63
|
-1.43% |
43.800 |
43.173 |
0.002 |
-1.43% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC64
|
-1.43% |
43.800 |
43.173 |
0.001 |
-1.43% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC64
|
-1.43% |
43.800 |
43.173 |
0.001 |
-1.43% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC64
|
-1.43% |
43.799 |
43.173 |
0.000 |
-1.43% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC63
|
-1.43% |
43.800 |
43.174 |
0.001 |
-1.43% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC63
|
-1.43% |
43.800 |
43.174 |
0.002 |
-1.43% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.000 |
-1.43% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC63
|
-1.43% |
43.800 |
43.174 |
0.001 |
-1.43% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC63
|
-1.43% |
43.799 |
43.173 |
0.002 |
-1.43% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.002 |
-1.43% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.001 |
-1.43% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC63
|
-1.43% |
43.799 |
43.174 |
0.001 |
-1.43% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/28
|
-1.40% |
85.724 |
84.524 |
0.024 |
-1.40% |
0.024 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/28
|
-1.38% |
86.351 |
85.156 |
0.006 |
-1.38% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/28
|
-1.22% |
53.128 |
52.477 |
0.010 |
-1.09% |
0.010 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint8_t>/16
|
-1.05% |
9.829 |
9.725 |
0.036 |
0.96% |
0.036 |