|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_int64_t_
|
-27.54% |
1044878.546 |
757134.792 |
409.904 |
-0.10% |
409.904 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_int64_t_
|
-27.44% |
1044463.436 |
757895.086 |
267.298 |
0.21% |
267.298 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_int32_t_
|
-26.91% |
515829.116 |
377026.536 |
401.695 |
-0.18% |
401.695 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_int32_t_
|
-26.78% |
515728.518 |
377636.305 |
664.916 |
0.02% |
664.916 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint64_t>/10
|
-26.76% |
38.640 |
28.298 |
0.007 |
-0.03% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint64_t>/16
|
-25.77% |
53.215 |
39.502 |
0.104 |
-0.01% |
0.104 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_uint8_t_
|
-25.33% |
124920.394 |
93276.605 |
354.054 |
-0.18% |
354.054 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_uint8_t_
|
-25.32% |
124916.399 |
93285.607 |
351.860 |
-0.14% |
351.860 |
|
SingleSource/Benchmarks/Shootout/Shootout-ary3
Profile
|
-23.61% |
0.985 |
0.752 |
0.001 |
-0.14% |
0.001 |
|
MultiSource/Benchmarks/Ptrdist/yacr2/yacr2
Profile
|
-22.99% |
2.836 |
2.184 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint64_t>/16
|
-20.87% |
32.544 |
25.753 |
0.018 |
-0.04% |
0.018 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint16_t>/16
|
-19.91% |
9.386 |
7.517 |
0.014 |
-0.02% |
0.014 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint16_t>/10
|
-19.86% |
9.385 |
7.521 |
0.003 |
-0.08% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint8_t>/16
|
-18.68% |
60.094 |
48.868 |
0.014 |
-0.02% |
0.014 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint64_t>/16
|
-17.98% |
10.020 |
8.218 |
0.035 |
-0.05% |
0.035 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint8_t>/10
|
-17.96% |
41.303 |
33.886 |
0.008 |
-0.00% |
0.008 |
|
MultiSource/Benchmarks/Olden/em3d/em3d
Profile
|
-17.94% |
18.304 |
15.021 |
0.009 |
-0.06% |
0.009 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint64_t>/10
|
-17.53% |
10.012 |
8.257 |
0.029 |
-0.21% |
0.029 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint64_t>/10
|
-17.26% |
24.373 |
20.167 |
0.025 |
0.04% |
0.025 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/28
|
-17.07% |
104.498 |
86.658 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/16
|
-15.46% |
66.954 |
56.602 |
0.012 |
-0.05% |
0.012 |
|
MultiSource/Benchmarks/SciMark2-C/scimark2
Profile
|
-15.13% |
166.680 |
141.463 |
0.023 |
-0.05% |
0.023 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint16_t>/10
|
-13.97% |
35.652 |
30.673 |
0.009 |
-0.04% |
0.009 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/10
|
-13.65% |
48.181 |
41.605 |
0.002 |
-0.01% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint8_t>/16
|
-10.21% |
33.992 |
30.523 |
0.098 |
-1.80% |
0.098 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_autovec_int64_t_
|
-9.95% |
1046245.511 |
942102.924 |
1011.915 |
-0.06% |
1011.915 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_novec_int64_t_
|
-9.90% |
1045645.924 |
942079.012 |
1104.741 |
-0.17% |
1104.741 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_autovec_int32_t_
|
-8.60% |
515323.384 |
471024.233 |
237.908 |
-0.14% |
237.908 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_novec_int32_t_
|
-8.58% |
515135.584 |
470955.105 |
453.341 |
0.03% |
453.341 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint16_t>/16
|
-8.52% |
49.949 |
45.695 |
0.004 |
0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/28
|
-8.51% |
157.975 |
144.538 |
0.563 |
-0.00% |
0.563 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/256
|
-8.40% |
1207.632 |
1106.224 |
0.021 |
0.00% |
0.021 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/999
|
-8.39% |
4647.048 |
4257.316 |
0.110 |
-0.00% |
0.110 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/51
|
-8.38% |
272.488 |
249.647 |
0.588 |
-0.00% |
0.588 |
|
MultiSource/Benchmarks/ASC_Sequoia/AMGmk/AMGmk
Profile
|
-8.34% |
43.459 |
39.833 |
0.013 |
-0.01% |
0.013 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/999
|
-8.33% |
2365.086 |
2167.998 |
0.058 |
-0.01% |
0.058 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/256
|
-8.24% |
1245.804 |
1143.127 |
0.012 |
0.00% |
0.012 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/999
|
-8.24% |
4805.013 |
4409.316 |
0.594 |
-0.00% |
0.594 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/256
|
-8.22% |
642.905 |
590.032 |
0.009 |
-0.00% |
0.009 |
|
SingleSource/Benchmarks/Adobe-C++/simple_types_constant_folding
Profile
|
-8.15% |
2.366 |
2.174 |
0.001 |
-0.13% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/51
|
-8.11% |
266.237 |
244.651 |
0.004 |
0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/999
|
-8.08% |
2443.259 |
2245.724 |
0.651 |
-0.00% |
0.651 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/256
|
-8.08% |
662.339 |
608.820 |
0.590 |
-0.00% |
0.590 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C++/HPCCG/HPCCG
Profile
|
-7.91% |
4.652 |
4.284 |
0.040 |
-0.05% |
0.040 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/28
|
-7.85% |
155.494 |
143.282 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/16
|
-7.51% |
100.115 |
92.601 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/51
|
-7.37% |
172.925 |
160.178 |
0.004 |
-0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/10
|
-7.28% |
81.654 |
75.709 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/51
|
-7.21% |
169.256 |
157.053 |
0.009 |
-0.00% |
0.009 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C++/miniFE/miniFE
Profile
|
-6.90% |
17.849 |
16.616 |
0.031 |
0.03% |
0.031 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/999
|
-6.68% |
51.627 |
48.178 |
0.020 |
-0.00% |
0.020 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/16
|
-6.67% |
89.163 |
83.217 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/256
|
-6.66% |
51.617 |
48.178 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/28
|
-6.64% |
51.601 |
48.176 |
0.118 |
-0.00% |
0.118 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/28
|
-6.64% |
122.638 |
114.500 |
0.209 |
0.00% |
0.209 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/51
|
-6.59% |
51.577 |
48.180 |
0.007 |
0.00% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/28
|
-6.25% |
120.135 |
112.627 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/256
|
-6.06% |
20.022 |
18.808 |
0.464 |
-4.59% |
0.464 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/999
|
-5.90% |
4194.518 |
3947.228 |
0.116 |
-0.01% |
0.116 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/999
|
-5.78% |
1453.725 |
1369.739 |
0.776 |
-0.02% |
0.776 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/10
|
-5.69% |
76.960 |
72.579 |
0.005 |
-0.01% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/256
|
-5.65% |
1087.542 |
1026.145 |
0.036 |
-0.00% |
0.036 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint64_t>/999
|
-5.64% |
5517.621 |
5206.159 |
0.373 |
-0.00% |
0.373 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/51
|
-5.57% |
241.837 |
228.376 |
0.008 |
-0.00% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/999
|
-5.55% |
2847.792 |
2689.643 |
0.057 |
0.00% |
0.057 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/256
|
-5.55% |
581.001 |
548.739 |
0.023 |
0.00% |
0.023 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/256
|
-5.55% |
745.255 |
703.889 |
0.018 |
-0.01% |
0.018 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/256
|
-5.55% |
744.568 |
703.271 |
0.027 |
-0.00% |
0.027 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/256
|
-5.55% |
745.223 |
703.900 |
0.135 |
-0.00% |
0.135 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/999
|
-5.54% |
2848.455 |
2690.631 |
0.125 |
0.00% |
0.125 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/51
|
-5.52% |
170.196 |
160.803 |
0.004 |
-0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/28
|
-5.50% |
142.352 |
134.522 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/999
|
-5.50% |
1451.132 |
1371.351 |
0.524 |
-0.00% |
0.524 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/51
|
-5.50% |
170.824 |
161.432 |
0.004 |
0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/51
|
-5.50% |
170.815 |
161.426 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/256
|
-5.48% |
399.822 |
377.912 |
0.017 |
-0.00% |
0.017 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/256
|
-5.48% |
400.495 |
378.564 |
0.003 |
0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/256
|
-5.47% |
400.440 |
378.533 |
0.015 |
-0.00% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/999
|
-5.46% |
2846.931 |
2691.409 |
0.287 |
-0.01% |
0.287 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/999
|
-5.45% |
1450.661 |
1371.577 |
0.042 |
-0.00% |
0.042 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint64_t>/51
|
-5.44% |
310.351 |
293.463 |
0.019 |
0.00% |
0.019 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/16
|
-5.41% |
92.607 |
87.596 |
0.002 |
-0.01% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/28
|
-5.39% |
52.247 |
49.429 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/10
|
-5.39% |
52.248 |
49.430 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/999
|
-5.39% |
52.248 |
49.430 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/51
|
-5.39% |
52.248 |
49.430 |
0.002 |
-0.01% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/256
|
-5.39% |
52.246 |
49.431 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/16
|
-5.39% |
52.246 |
49.431 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint64_t>/999
|
-5.38% |
2804.568 |
2653.672 |
0.181 |
0.00% |
0.181 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/999
|
-5.37% |
2131.884 |
2017.451 |
0.107 |
-0.00% |
0.107 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/16
|
-5.36% |
16.894 |
15.989 |
0.126 |
-0.67% |
0.126 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/999
|
-5.33% |
3006.073 |
2845.956 |
0.487 |
-0.01% |
0.487 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/51
|
-5.12% |
110.143 |
104.499 |
0.008 |
-0.00% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/51
|
-5.08% |
110.746 |
105.117 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/51
|
-5.01% |
156.117 |
148.290 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/999
|
-5.01% |
1529.178 |
1452.568 |
0.141 |
-0.00% |
0.141 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/51
|
-4.92% |
178.325 |
169.555 |
0.003 |
0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/16
|
-4.86% |
83.532 |
79.469 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint64_t>/51
|
-4.78% |
196.478 |
187.092 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/51
|
-4.77% |
39.420 |
37.540 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/256
|
-4.77% |
39.419 |
37.540 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/999
|
-4.76% |
39.419 |
37.542 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/256
|
-4.69% |
40.046 |
38.166 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/256
|
-4.69% |
40.045 |
38.166 |
0.004 |
-0.01% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/999
|
-4.69% |
40.046 |
38.167 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/999
|
-4.69% |
40.046 |
38.167 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/51
|
-4.69% |
40.046 |
38.167 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/51
|
-4.69% |
40.046 |
38.167 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC15
|
-4.60% |
273.341 |
260.768 |
3.752 |
-5.21% |
3.752 |
|
MultiSource/Benchmarks/TSVC/IndirectAddressing-flt/IndirectAddressing-flt
Profile
|
-4.58% |
15.747 |
15.027 |
0.042 |
0.06% |
0.042 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/28
|
-4.51% |
16.894 |
16.133 |
0.076 |
0.42% |
0.076 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/28
|
-4.45% |
112.632 |
107.624 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/51
|
-4.41% |
113.886 |
108.858 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/999
|
-4.35% |
50.370 |
48.177 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/51
|
-4.35% |
50.371 |
48.178 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/256
|
-4.35% |
50.371 |
48.179 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/16
|
-4.35% |
50.371 |
48.179 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/28
|
-4.35% |
50.370 |
48.180 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/10
|
-4.24% |
16.894 |
16.179 |
0.085 |
-0.23% |
0.085 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4BigLoopTC15
|
-4.11% |
272.396 |
261.196 |
3.205 |
-4.96% |
3.205 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint16_t>/10
|
-4.11% |
22.245 |
21.332 |
0.013 |
-0.47% |
0.013 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/51
|
-4.00% |
109.500 |
105.117 |
0.001 |
-0.00% |
0.001 |
|
MultiSource/Applications/JM/lencod/lencod
Profile
|
-3.92% |
35.747 |
34.345 |
0.031 |
-4.08% |
0.031 |
|
SingleSource/Benchmarks/Misc-C++/oopack_v1p8
Profile
|
-3.91% |
0.479 |
0.460 |
0.000 |
-0.01% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC15
|
-3.88% |
271.926 |
261.376 |
3.456 |
-4.90% |
3.456 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/999
|
-3.77% |
40.048 |
38.538 |
0.107 |
-0.60% |
0.107 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/51
|
-3.45% |
40.045 |
38.665 |
0.048 |
0.02% |
0.048 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint64_t>/51
|
-3.30% |
56.940 |
55.061 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint64_t>/999
|
-3.30% |
56.939 |
55.061 |
0.006 |
-0.00% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC7
|
-3.13% |
161.315 |
156.265 |
1.560 |
-4.59% |
1.560 |
|
MultiSource/Benchmarks/MiBench/consumer-lame/consumer-lame
Profile
|
-3.13% |
0.682 |
0.661 |
0.001 |
-2.23% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC15
|
-2.94% |
21.274 |
20.648 |
0.001 |
-0.00% |
0.001 |
|
External/SPEC/CINT2017rate/525.x264_r/525.x264_r
Profile
|
-2.90% |
142.176 |
138.047 |
0.135 |
-3.12% |
0.135 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC7
|
-2.90% |
160.984 |
156.317 |
1.594 |
-4.59% |
1.594 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/10
|
-2.85% |
40.674 |
39.514 |
0.017 |
0.10% |
0.017 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint16_t>/16
|
-2.76% |
29.623 |
28.807 |
0.010 |
-0.06% |
0.010 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/10
|
-2.76% |
41.296 |
40.158 |
0.008 |
-0.02% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4BigLoopTC7
|
-2.74% |
160.668 |
156.262 |
1.669 |
-4.54% |
1.669 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC16
|
-2.67% |
21.214 |
20.648 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/256
|
-2.51% |
355.305 |
346.394 |
0.041 |
-0.01% |
0.041 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/10
|
-2.38% |
30.662 |
29.932 |
0.023 |
-0.03% |
0.023 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C/RSBench/rsbench
Profile
|
-2.33% |
1.222 |
1.194 |
0.006 |
-1.24% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC31
|
-2.18% |
28.783 |
28.156 |
0.001 |
-0.01% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC31
|
-2.17% |
28.782 |
28.157 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC32
|
-2.17% |
28.782 |
28.157 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC32
|
-2.17% |
28.783 |
28.158 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC31
|
-2.17% |
28.782 |
28.157 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC31
|
-2.16% |
28.779 |
28.156 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC31
|
-2.16% |
28.779 |
28.157 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC31
|
-2.16% |
28.778 |
28.156 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC31
|
-2.15% |
28.776 |
28.157 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC32
|
-2.15% |
28.777 |
28.158 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/16
|
-2.14% |
55.694 |
54.505 |
0.016 |
0.02% |
0.016 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC32
|
-2.13% |
28.771 |
28.157 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC32
|
-2.12% |
28.768 |
28.157 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/16
|
-2.08% |
56.315 |
55.143 |
0.030 |
-0.02% |
0.030 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC15
|
-2.05% |
21.081 |
20.649 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC32
|
-1.99% |
28.728 |
28.157 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC32
|
-1.98% |
28.726 |
28.158 |
0.000 |
-0.00% |
0.000 |
|
MultiSource/Applications/sgefa/sgefa
Profile
|
-1.96% |
0.642 |
0.630 |
0.003 |
-0.36% |
0.003 |
|
SingleSource/Benchmarks/CoyoteBench/huffbench
Profile
|
-1.76% |
50.629 |
49.740 |
0.007 |
0.00% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint8_t>/10
|
-1.66% |
24.538 |
24.131 |
0.078 |
0.34% |
0.078 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/16
|
-1.62% |
38.037 |
37.422 |
0.020 |
-0.03% |
0.020 |
|
SingleSource/Benchmarks/Shootout/Shootout-matrix
Profile
|
-1.61% |
4.693 |
4.617 |
0.000 |
-1.60% |
0.000 |
|
MultiSource/Benchmarks/Rodinia/hotspot/hotspot
Profile
|
-1.51% |
2.312 |
2.277 |
0.011 |
0.07% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/10
|
-1.49% |
48.905 |
48.177 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC63
|
-1.43% |
43.800 |
43.173 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC63
|
-1.43% |
43.800 |
43.174 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC63
|
-1.43% |
43.799 |
43.173 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC64
|
-1.43% |
43.799 |
43.173 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC64
|
-1.43% |
43.799 |
43.173 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC64
|
-1.43% |
43.800 |
43.174 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC64
|
-1.43% |
43.798 |
43.173 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC63
|
-1.43% |
43.799 |
43.173 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC63
|
-1.43% |
43.799 |
43.174 |
0.004 |
-0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC63
|
-1.43% |
43.799 |
43.174 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC64
|
-1.43% |
43.798 |
43.173 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC63
|
-1.42% |
43.799 |
43.175 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/28
|
-1.39% |
86.349 |
85.146 |
0.009 |
-0.01% |
0.009 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/28
|
-1.39% |
85.726 |
84.538 |
0.000 |
0.05% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC15
|
-1.36% |
20.932 |
20.648 |
0.001 |
-0.00% |
0.001 |
|
MultiSource/Applications/lua/lua
Profile
|
-1.19% |
70.915 |
70.072 |
0.003 |
0.01% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/28
|
-1.15% |
53.057 |
52.446 |
0.020 |
-0.01% |
0.020 |
|
MultiSource/Applications/lambda-0.1.3/lambda
Profile
|
-1.12% |
11.808 |
11.676 |
0.007 |
-0.05% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC16
|
-1.09% |
20.877 |
20.649 |
0.001 |
-0.00% |
0.001 |
|
External/SPEC/CFP2017rate/511.povray_r/511.povray_r
Profile
|
-1.07% |
28.635 |
28.329 |
0.021 |
0.11% |
0.021 |