|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_int64_t_
|
-27.65% |
1044878.546 |
755945.148 |
1742.983 |
0.12% |
1742.983 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_int64_t_
|
-27.52% |
1044463.436 |
757073.129 |
1338.797 |
0.10% |
1338.797 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_int32_t_
|
-26.98% |
515829.116 |
376664.800 |
1101.959 |
0.07% |
1101.959 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_int32_t_
|
-26.97% |
515728.518 |
376627.207 |
724.429 |
-0.10% |
724.429 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint64_t>/10
|
-26.83% |
38.640 |
28.275 |
0.007 |
-0.01% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint64_t>/16
|
-25.80% |
53.215 |
39.487 |
0.011 |
-0.02% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_uint8_t_
|
-25.13% |
124920.394 |
93533.753 |
172.107 |
0.22% |
172.107 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_uint8_t_
|
-25.12% |
124916.399 |
93536.409 |
171.494 |
0.22% |
171.494 |
|
SingleSource/Benchmarks/Shootout/Shootout-ary3
Profile
|
-23.54% |
0.985 |
0.753 |
0.001 |
-0.06% |
0.001 |
|
MultiSource/Benchmarks/Ptrdist/yacr2/yacr2
Profile
|
-23.02% |
2.836 |
2.183 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint64_t>/16
|
-20.87% |
32.544 |
25.752 |
0.015 |
0.03% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint16_t>/10
|
-19.88% |
9.385 |
7.520 |
0.005 |
-0.07% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint16_t>/16
|
-19.84% |
9.386 |
7.524 |
0.018 |
0.09% |
0.018 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint8_t>/16
|
-18.69% |
60.094 |
48.862 |
0.029 |
-0.05% |
0.029 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint8_t>/10
|
-18.03% |
41.303 |
33.858 |
0.016 |
-0.11% |
0.016 |
|
MultiSource/Benchmarks/Olden/em3d/em3d
Profile
|
-17.93% |
18.304 |
15.022 |
0.009 |
-0.05% |
0.009 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint64_t>/16
|
-17.80% |
10.020 |
8.237 |
0.003 |
-0.33% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint64_t>/10
|
-17.48% |
10.012 |
8.262 |
0.015 |
-0.00% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint64_t>/10
|
-17.31% |
24.373 |
20.154 |
0.015 |
0.02% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/28
|
-17.07% |
104.498 |
86.658 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/16
|
-15.44% |
66.954 |
56.619 |
0.006 |
-0.01% |
0.006 |
|
MultiSource/Benchmarks/SciMark2-C/scimark2
Profile
|
-15.06% |
166.680 |
141.579 |
0.026 |
0.00% |
0.026 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint16_t>/10
|
-13.88% |
35.652 |
30.702 |
0.056 |
0.03% |
0.056 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/10
|
-13.64% |
48.181 |
41.612 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_autovec_int64_t_
|
-9.90% |
1046245.511 |
942673.250 |
467.032 |
-0.00% |
467.032 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_novec_int64_t_
|
-9.82% |
1045645.924 |
942952.546 |
1036.285 |
-0.03% |
1036.285 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint8_t>/16
|
-9.73% |
33.992 |
30.685 |
0.201 |
0.12% |
0.201 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint16_t>/16
|
-8.53% |
49.949 |
45.686 |
0.028 |
-0.02% |
0.028 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/28
|
-8.50% |
157.975 |
144.541 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_autovec_int32_t_
|
-8.47% |
515323.384 |
471684.330 |
266.010 |
-0.03% |
266.010 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_novec_int32_t_
|
-8.46% |
515135.584 |
471551.970 |
162.037 |
-0.00% |
162.037 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/256
|
-8.39% |
1207.632 |
1106.282 |
0.091 |
0.00% |
0.091 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/999
|
-8.38% |
4647.048 |
4257.548 |
0.224 |
-0.00% |
0.224 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/51
|
-8.38% |
272.488 |
249.656 |
0.006 |
0.00% |
0.006 |
|
MultiSource/Benchmarks/ASC_Sequoia/AMGmk/AMGmk
Profile
|
-8.33% |
43.459 |
39.837 |
0.004 |
-0.01% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/999
|
-8.33% |
2365.086 |
2168.182 |
0.094 |
-0.00% |
0.094 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/256
|
-8.24% |
1245.804 |
1143.121 |
0.253 |
0.00% |
0.253 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/999
|
-8.23% |
4805.013 |
4409.684 |
0.291 |
-0.00% |
0.291 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/256
|
-8.22% |
642.905 |
590.066 |
0.010 |
0.00% |
0.010 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/51
|
-8.10% |
266.237 |
244.664 |
0.001 |
-0.00% |
0.001 |
|
SingleSource/Benchmarks/Adobe-C++/simple_types_constant_folding
Profile
|
-8.09% |
2.366 |
2.175 |
0.002 |
-0.03% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/256
|
-8.08% |
662.339 |
608.805 |
0.430 |
0.00% |
0.430 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/999
|
-8.08% |
2443.259 |
2245.803 |
0.213 |
-0.00% |
0.213 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C++/HPCCG/HPCCG
Profile
|
-7.88% |
4.652 |
4.286 |
0.056 |
-0.09% |
0.056 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/28
|
-7.85% |
155.494 |
143.289 |
0.005 |
0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/16
|
-7.50% |
100.115 |
92.608 |
0.006 |
-0.00% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/51
|
-7.37% |
172.925 |
160.181 |
0.047 |
0.00% |
0.047 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/10
|
-7.28% |
81.654 |
75.712 |
0.009 |
0.00% |
0.009 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/51
|
-7.21% |
169.256 |
157.059 |
0.002 |
0.00% |
0.002 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C++/miniFE/miniFE
Profile
|
-6.94% |
17.849 |
16.610 |
0.034 |
-0.10% |
0.034 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/999
|
-6.68% |
51.627 |
48.177 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/16
|
-6.66% |
89.163 |
83.221 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/256
|
-6.66% |
51.617 |
48.180 |
0.007 |
0.00% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/28
|
-6.63% |
122.638 |
114.507 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/28
|
-6.63% |
51.601 |
48.180 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/51
|
-6.59% |
51.577 |
48.180 |
0.000 |
0.01% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/28
|
-6.24% |
120.135 |
112.635 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/999
|
-5.89% |
4194.518 |
3947.375 |
0.122 |
0.00% |
0.122 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/999
|
-5.77% |
1453.725 |
1369.774 |
0.178 |
-0.01% |
0.178 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/10
|
-5.69% |
76.960 |
72.583 |
0.003 |
0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint64_t>/999
|
-5.64% |
5517.621 |
5206.220 |
0.283 |
0.00% |
0.283 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/256
|
-5.64% |
1087.542 |
1026.190 |
0.054 |
0.00% |
0.054 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/51
|
-5.56% |
241.837 |
228.391 |
0.008 |
-0.00% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/999
|
-5.56% |
2847.792 |
2689.590 |
0.041 |
-0.00% |
0.041 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/256
|
-5.55% |
581.001 |
548.737 |
0.023 |
-0.00% |
0.023 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/256
|
-5.54% |
744.568 |
703.290 |
0.049 |
0.00% |
0.049 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/256
|
-5.54% |
745.255 |
703.943 |
0.036 |
0.00% |
0.036 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/256
|
-5.54% |
745.223 |
703.918 |
0.037 |
-0.00% |
0.037 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/999
|
-5.54% |
2848.455 |
2690.604 |
0.062 |
-0.01% |
0.062 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/999
|
-5.54% |
1451.132 |
1370.777 |
0.630 |
0.00% |
0.630 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/51
|
-5.52% |
170.196 |
160.802 |
0.005 |
-0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/51
|
-5.49% |
170.824 |
161.439 |
0.004 |
-0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/28
|
-5.49% |
142.352 |
134.533 |
0.003 |
0.01% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/51
|
-5.49% |
170.815 |
161.436 |
0.015 |
-0.00% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/28
|
-5.49% |
16.894 |
15.967 |
0.140 |
-1.30% |
0.140 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/256
|
-5.48% |
400.495 |
378.553 |
0.033 |
-0.00% |
0.033 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/256
|
-5.48% |
399.822 |
377.925 |
0.014 |
0.00% |
0.014 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/256
|
-5.47% |
400.440 |
378.554 |
0.011 |
-0.00% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/999
|
-5.45% |
2846.931 |
2691.633 |
0.234 |
-0.00% |
0.234 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/999
|
-5.45% |
1450.661 |
1371.659 |
0.052 |
-0.00% |
0.052 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint64_t>/51
|
-5.44% |
310.351 |
293.466 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/16
|
-5.40% |
92.607 |
87.602 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/999
|
-5.39% |
52.248 |
49.430 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/28
|
-5.39% |
52.247 |
49.431 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/10
|
-5.39% |
52.248 |
49.431 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/16
|
-5.39% |
52.246 |
49.431 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/51
|
-5.39% |
52.248 |
49.433 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/256
|
-5.39% |
52.246 |
49.432 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint64_t>/999
|
-5.38% |
2804.568 |
2653.744 |
0.178 |
0.00% |
0.178 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/999
|
-5.37% |
2131.884 |
2017.337 |
0.149 |
-0.01% |
0.149 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/999
|
-5.32% |
3006.073 |
2846.115 |
0.192 |
0.00% |
0.192 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/10
|
-5.22% |
16.894 |
16.012 |
0.013 |
-0.04% |
0.013 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/16
|
-5.21% |
16.894 |
16.014 |
0.114 |
-0.94% |
0.114 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/51
|
-5.13% |
110.143 |
104.495 |
0.003 |
0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/51
|
-5.08% |
110.746 |
105.118 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/999
|
-5.03% |
1529.178 |
1452.267 |
0.227 |
-0.03% |
0.227 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/51
|
-5.01% |
156.117 |
148.293 |
0.008 |
-0.00% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/51
|
-4.92% |
178.325 |
169.553 |
0.003 |
0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/16
|
-4.87% |
83.532 |
79.466 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint64_t>/51
|
-4.77% |
196.478 |
187.098 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/999
|
-4.76% |
39.419 |
37.542 |
0.003 |
0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/51
|
-4.76% |
39.420 |
37.544 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/256
|
-4.76% |
39.419 |
37.543 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/51
|
-4.69% |
40.046 |
38.167 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/51
|
-4.69% |
40.046 |
38.168 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/256
|
-4.69% |
40.045 |
38.169 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/256
|
-4.68% |
40.046 |
38.170 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/999
|
-4.68% |
40.046 |
38.170 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/999
|
-4.68% |
40.046 |
38.171 |
0.001 |
0.00% |
0.001 |
|
MultiSource/Benchmarks/TSVC/IndirectAddressing-flt/IndirectAddressing-flt
Profile
|
-4.58% |
15.747 |
15.026 |
0.017 |
-0.08% |
0.017 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC15
|
-4.49% |
273.341 |
261.067 |
0.577 |
-1.15% |
0.577 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/28
|
-4.45% |
112.632 |
107.625 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/51
|
-4.41% |
113.886 |
108.867 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/256
|
-4.35% |
50.371 |
48.179 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/16
|
-4.35% |
50.371 |
48.179 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/51
|
-4.35% |
50.371 |
48.180 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/28
|
-4.35% |
50.370 |
48.180 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/999
|
-4.35% |
50.370 |
48.180 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4BigLoopTC15
|
-4.27% |
272.396 |
260.751 |
0.694 |
-1.15% |
0.694 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/51
|
-4.00% |
109.500 |
105.120 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC15
|
-3.93% |
271.926 |
261.235 |
0.440 |
-1.15% |
0.440 |
|
SingleSource/Benchmarks/Misc-C++/oopack_v1p8
Profile
|
-3.88% |
0.479 |
0.460 |
0.000 |
0.01% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/999
|
-3.50% |
40.048 |
38.647 |
0.055 |
0.14% |
0.055 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint8_t>/10
|
-3.38% |
24.538 |
23.710 |
0.206 |
0.21% |
0.206 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint64_t>/51
|
-3.30% |
56.940 |
55.062 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint64_t>/999
|
-3.29% |
56.939 |
55.063 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/51
|
-3.20% |
40.045 |
38.764 |
0.007 |
-0.04% |
0.007 |
|
MultiSource/Benchmarks/TSVC/IndirectAddressing-dbl/IndirectAddressing-dbl
Profile
|
-3.01% |
18.523 |
17.966 |
0.117 |
-1.18% |
0.117 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/10
|
-2.94% |
40.674 |
39.477 |
0.046 |
0.08% |
0.046 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC15
|
-2.94% |
21.274 |
20.649 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/10
|
-2.79% |
41.296 |
40.145 |
0.020 |
-0.01% |
0.020 |
|
MultiSource/Benchmarks/TSVC/Symbolics-flt/Symbolics-flt
Profile
|
-2.75% |
4.797 |
4.665 |
0.030 |
-2.50% |
0.030 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint16_t>/16
|
-2.75% |
29.623 |
28.810 |
0.006 |
0.00% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC16
|
-2.66% |
21.214 |
20.649 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC7
|
-2.58% |
161.315 |
157.158 |
1.076 |
-0.24% |
1.076 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/256
|
-2.49% |
355.305 |
346.451 |
0.033 |
0.00% |
0.033 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/10
|
-2.42% |
30.662 |
29.919 |
0.013 |
-0.05% |
0.013 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC7
|
-2.36% |
160.984 |
157.180 |
1.120 |
-0.22% |
1.120 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC32
|
-2.17% |
28.783 |
28.157 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC31
|
-2.17% |
28.782 |
28.156 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC31
|
-2.17% |
28.783 |
28.157 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC32
|
-2.17% |
28.782 |
28.158 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC31
|
-2.17% |
28.782 |
28.157 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC31
|
-2.17% |
28.779 |
28.156 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4BigLoopTC7
|
-2.16% |
160.668 |
157.192 |
1.024 |
-0.20% |
1.024 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC31
|
-2.16% |
28.779 |
28.157 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC31
|
-2.16% |
28.778 |
28.157 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/16
|
-2.16% |
55.694 |
54.493 |
0.048 |
-0.05% |
0.048 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC31
|
-2.15% |
28.776 |
28.157 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC32
|
-2.15% |
28.777 |
28.158 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC32
|
-2.13% |
28.771 |
28.160 |
0.000 |
0.01% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC32
|
-2.12% |
28.768 |
28.157 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/16
|
-2.08% |
56.315 |
55.143 |
0.010 |
0.02% |
0.010 |
|
MultiSource/Applications/sgefa/sgefa
Profile
|
-2.06% |
0.642 |
0.629 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC15
|
-2.05% |
21.081 |
20.649 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC32
|
-1.98% |
28.728 |
28.158 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC32
|
-1.98% |
28.726 |
28.157 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/256
|
-1.95% |
20.022 |
19.632 |
0.113 |
1.55% |
0.113 |
|
SingleSource/Benchmarks/CoyoteBench/huffbench
Profile
|
-1.76% |
50.629 |
49.739 |
0.004 |
0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/16
|
-1.67% |
38.037 |
37.402 |
0.037 |
-0.13% |
0.037 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_asin_novec_double_
|
-1.64% |
1184.843 |
1165.436 |
5.632 |
-1.29% |
5.632 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/10
|
-1.53% |
48.905 |
48.159 |
0.011 |
-0.03% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC63
|
-1.43% |
43.799 |
43.174 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC63
|
-1.43% |
43.800 |
43.176 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC64
|
-1.43% |
43.800 |
43.175 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC63
|
-1.43% |
43.799 |
43.175 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC64
|
-1.43% |
43.798 |
43.174 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC64
|
-1.43% |
43.799 |
43.175 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC64
|
-1.42% |
43.799 |
43.176 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC63
|
-1.42% |
43.799 |
43.175 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC63
|
-1.42% |
43.800 |
43.177 |
0.011 |
0.01% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC63
|
-1.42% |
43.799 |
43.176 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC64
|
-1.42% |
43.798 |
43.176 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC63
|
-1.42% |
43.799 |
43.177 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC4
|
-1.41% |
107.017 |
105.512 |
0.334 |
-1.64% |
0.334 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/28
|
-1.40% |
86.349 |
85.140 |
0.015 |
0.00% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC1
|
-1.40% |
108.216 |
106.703 |
0.284 |
-0.89% |
0.284 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/28
|
-1.38% |
85.726 |
84.539 |
0.008 |
0.00% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC15
|
-1.35% |
20.932 |
20.649 |
0.001 |
-0.00% |
0.001 |
|
MultiSource/Applications/lua/lua
Profile
|
-1.35% |
70.915 |
69.957 |
0.035 |
-0.13% |
0.035 |
|
External/SPEC/CFP2017rate/511.povray_r/511.povray_r
Profile
|
-1.24% |
28.635 |
28.281 |
0.023 |
0.16% |
0.023 |
|
MultiSource/Applications/lambda-0.1.3/lambda
Profile
|
-1.14% |
11.808 |
11.674 |
0.011 |
-0.11% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/28
|
-1.11% |
53.057 |
52.470 |
0.025 |
0.03% |
0.025 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC16
|
-1.09% |
20.877 |
20.648 |
0.001 |
-0.00% |
0.001 |
|
MultiSource/Benchmarks/TSVC/ControlFlow-flt/ControlFlow-flt
Profile
|
-1.08% |
24.408 |
24.145 |
0.091 |
2.23% |
0.091 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4BigLoopTC1
|
-1.07% |
108.324 |
107.166 |
0.117 |
-0.52% |
0.117 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC16
|
-1.07% |
265.793 |
262.960 |
0.410 |
0.65% |
0.410 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC4
|
-1.05% |
107.133 |
106.010 |
0.261 |
-0.86% |
0.261 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C/RSBench/rsbench
Profile
|
-1.03% |
1.222 |
1.210 |
0.003 |
-0.11% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC16
|
-1.02% |
265.574 |
262.855 |
0.423 |
0.46% |
0.423 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4BigLoopTC2
|
-1.02% |
108.197 |
107.094 |
0.288 |
-0.60% |
0.288 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC3
|
-1.01% |
107.884 |
106.799 |
0.292 |
-0.41% |
0.292 |