|
MultiSource/Benchmarks/Trimaran/enc-md5/enc-md5
Profile
|
-75.99% |
12.643 |
3.035 |
0.000 |
-0.01% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_int64_t_
|
-27.65% |
1044463.436 |
755716.357 |
1001.087 |
-0.12% |
1001.087 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_int64_t_
|
-27.64% |
1044878.546 |
756054.860 |
683.175 |
-0.16% |
683.175 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_int32_t_
|
-27.03% |
515728.518 |
376342.556 |
792.942 |
-0.15% |
792.942 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint64_t>/10
|
-26.89% |
38.640 |
28.250 |
0.053 |
-0.01% |
0.053 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_int32_t_
|
-26.86% |
515829.116 |
377278.442 |
247.668 |
0.02% |
247.668 |
|
MultiSource/Benchmarks/Ptrdist/yacr2/yacr2
Profile
|
-26.69% |
2.836 |
2.079 |
0.001 |
-4.85% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint64_t>/16
|
-25.76% |
53.215 |
39.509 |
0.015 |
-0.05% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_novec_uint8_t_
|
-25.25% |
124916.399 |
93370.368 |
185.706 |
0.09% |
185.706 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_with_cond_arith_autovec_uint8_t_
|
-25.25% |
124920.394 |
93374.533 |
169.710 |
0.08% |
169.710 |
|
SingleSource/Benchmarks/Shootout/Shootout-ary3
Profile
|
-23.43% |
0.985 |
0.754 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint64_t>/16
|
-20.88% |
32.544 |
25.749 |
0.024 |
-0.01% |
0.024 |
|
MultiSource/Benchmarks/Prolangs-C/gnugo/gnugo
Profile
|
-20.57% |
0.137 |
0.109 |
0.000 |
-20.77% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint16_t>/16
|
-19.88% |
9.386 |
7.520 |
0.013 |
-0.01% |
0.013 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint16_t>/10
|
-19.83% |
9.385 |
7.524 |
0.023 |
-0.04% |
0.023 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint8_t>/16
|
-18.79% |
60.094 |
48.805 |
0.036 |
-0.12% |
0.036 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint8_t>/10
|
-18.04% |
41.303 |
33.852 |
0.014 |
-0.18% |
0.014 |
|
MultiSource/Benchmarks/Olden/em3d/em3d
Profile
|
-17.91% |
18.304 |
15.026 |
0.006 |
0.03% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint64_t>/10
|
-17.80% |
10.012 |
8.230 |
0.015 |
-0.30% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_single_load<uint64_t>/16
|
-17.64% |
10.020 |
8.252 |
0.027 |
-0.06% |
0.027 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint64_t>/10
|
-17.35% |
24.373 |
20.143 |
0.019 |
-0.15% |
0.019 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/28
|
-17.35% |
104.498 |
86.367 |
0.138 |
-0.25% |
0.138 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/16
|
-15.43% |
66.954 |
56.624 |
0.011 |
-0.01% |
0.011 |
|
MultiSource/Benchmarks/SciMark2-C/scimark2
Profile
|
-15.11% |
166.680 |
141.492 |
0.013 |
-0.03% |
0.013 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint16_t>/10
|
-13.93% |
35.652 |
30.685 |
0.007 |
0.01% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/10
|
-13.64% |
48.181 |
41.610 |
0.002 |
0.00% |
0.002 |
|
MultiSource/Benchmarks/TSVC/LoopRestructuring-dbl/LoopRestructuring-dbl
Profile
|
-11.40% |
55.687 |
49.339 |
2.428 |
-16.22% |
2.428 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint8_t>/16
|
-10.11% |
33.992 |
30.557 |
0.358 |
-0.84% |
0.358 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_novec_int64_t_
|
-9.98% |
1045645.924 |
941315.977 |
505.256 |
-0.06% |
505.256 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_autovec_int64_t_
|
-9.95% |
1046245.511 |
942101.379 |
557.608 |
-0.01% |
557.608 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_autovec_int32_t_
|
-8.66% |
515323.384 |
470706.171 |
186.005 |
-0.07% |
186.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:BENCHMARK_multi_csa_only_novec_int32_t_
|
-8.65% |
515135.584 |
470582.603 |
560.793 |
-0.05% |
560.793 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_single_load<uint16_t>/16
|
-8.52% |
49.949 |
45.695 |
0.049 |
-0.02% |
0.049 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/28
|
-8.51% |
157.975 |
144.532 |
0.005 |
0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/256
|
-8.40% |
1207.632 |
1106.213 |
0.070 |
-0.00% |
0.070 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/999
|
-8.39% |
4647.048 |
4257.363 |
0.109 |
-0.01% |
0.109 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/51
|
-8.38% |
272.488 |
249.649 |
0.011 |
-0.00% |
0.011 |
|
MultiSource/Benchmarks/ASC_Sequoia/AMGmk/AMGmk
Profile
|
-8.36% |
43.459 |
39.826 |
0.013 |
-0.02% |
0.013 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/999
|
-8.33% |
2365.086 |
2168.129 |
0.043 |
-0.00% |
0.043 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/256
|
-8.24% |
1245.804 |
1143.151 |
0.122 |
-0.00% |
0.122 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint8_t>/999
|
-8.23% |
4805.013 |
4409.522 |
0.214 |
-0.01% |
0.214 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/256
|
-8.22% |
642.905 |
590.047 |
0.009 |
0.00% |
0.009 |
|
SingleSource/Benchmarks/Adobe-C++/simple_types_constant_folding
Profile
|
-8.19% |
2.366 |
2.173 |
0.001 |
-0.02% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/51
|
-8.11% |
266.237 |
244.642 |
0.010 |
-0.00% |
0.010 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/999
|
-8.09% |
2443.259 |
2245.598 |
0.126 |
-0.00% |
0.126 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/256
|
-8.08% |
662.339 |
608.789 |
0.075 |
-0.01% |
0.075 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C++/HPCCG/HPCCG
Profile
|
-7.85% |
4.652 |
4.287 |
0.034 |
0.04% |
0.034 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/28
|
-7.85% |
155.494 |
143.284 |
0.005 |
-0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/16
|
-7.50% |
100.115 |
92.603 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/51
|
-7.37% |
172.925 |
160.177 |
0.004 |
-0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint16_t>/10
|
-7.28% |
81.654 |
75.711 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/51
|
-7.21% |
169.256 |
157.060 |
0.003 |
0.00% |
0.003 |
|
MultiSource/Benchmarks/DOE-ProxyApps-C++/miniFE/miniFE
Profile
|
-6.87% |
17.849 |
16.622 |
0.062 |
-0.04% |
0.062 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/999
|
-6.68% |
51.627 |
48.179 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/16
|
-6.66% |
89.163 |
83.221 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/256
|
-6.66% |
51.617 |
48.178 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/28
|
-6.63% |
51.601 |
48.179 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint8_t>/28
|
-6.63% |
122.638 |
114.508 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint8_t>/51
|
-6.59% |
51.577 |
48.178 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/28
|
-6.25% |
120.135 |
112.622 |
0.010 |
-0.01% |
0.010 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/999
|
-5.92% |
4194.518 |
3946.157 |
0.686 |
-0.04% |
0.686 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/999
|
-5.78% |
1453.725 |
1369.737 |
0.346 |
0.00% |
0.346 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint16_t>/10
|
-5.69% |
76.960 |
72.582 |
0.007 |
-0.00% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint64_t>/999
|
-5.67% |
5517.621 |
5204.657 |
0.859 |
-0.02% |
0.859 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/256
|
-5.65% |
1087.542 |
1026.136 |
0.094 |
-0.01% |
0.094 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/51
|
-5.57% |
241.837 |
228.371 |
0.015 |
-0.00% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/999
|
-5.55% |
2847.792 |
2689.649 |
0.093 |
-0.01% |
0.093 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/256
|
-5.55% |
745.255 |
703.902 |
0.021 |
-0.01% |
0.021 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/256
|
-5.55% |
20.022 |
18.912 |
0.403 |
-3.61% |
0.403 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/999
|
-5.54% |
2848.455 |
2690.565 |
0.361 |
-0.00% |
0.361 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/256
|
-5.54% |
745.223 |
703.923 |
0.026 |
-0.00% |
0.026 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/256
|
-5.54% |
744.568 |
703.305 |
0.008 |
0.00% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/256
|
-5.54% |
581.001 |
548.802 |
0.001 |
0.01% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/51
|
-5.52% |
170.196 |
160.805 |
0.005 |
0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint16_t>/51
|
-5.50% |
170.824 |
161.430 |
0.007 |
-0.00% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/28
|
-5.50% |
142.352 |
134.528 |
0.008 |
0.00% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/51
|
-5.50% |
170.815 |
161.428 |
0.006 |
0.00% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/999
|
-5.49% |
2846.931 |
2690.592 |
0.461 |
-0.02% |
0.461 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/256
|
-5.48% |
400.495 |
378.545 |
0.009 |
0.00% |
0.009 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/256
|
-5.48% |
399.822 |
377.915 |
0.006 |
-0.00% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/256
|
-5.47% |
400.440 |
378.546 |
0.011 |
0.00% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/999
|
-5.47% |
1451.132 |
1371.808 |
0.457 |
-0.02% |
0.457 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/999
|
-5.45% |
1450.661 |
1371.594 |
0.079 |
-0.00% |
0.079 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint64_t>/51
|
-5.45% |
310.351 |
293.448 |
0.008 |
-0.00% |
0.008 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_three_loads<uint32_t>/16
|
-5.41% |
92.607 |
87.600 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/10
|
-5.39% |
52.248 |
49.430 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/51
|
-5.39% |
52.248 |
49.430 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/28
|
-5.39% |
52.247 |
49.430 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/999
|
-5.39% |
52.248 |
49.431 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/256
|
-5.39% |
52.246 |
49.429 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint16_t>/16
|
-5.39% |
52.246 |
49.429 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint64_t>/999
|
-5.37% |
2804.568 |
2653.850 |
0.035 |
-0.01% |
0.035 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/999
|
-5.37% |
2131.884 |
2017.401 |
0.068 |
-0.00% |
0.068 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/999
|
-5.33% |
3006.073 |
2845.999 |
0.239 |
-0.00% |
0.239 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/28
|
-5.30% |
16.894 |
15.998 |
0.123 |
-0.37% |
0.123 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint32_t>/51
|
-5.13% |
110.143 |
104.496 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint8_t>/51
|
-5.09% |
110.746 |
105.114 |
0.005 |
-0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/999
|
-5.03% |
1529.178 |
1452.322 |
0.280 |
-0.02% |
0.280 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/51
|
-5.01% |
156.117 |
148.288 |
0.010 |
-0.00% |
0.010 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/16
|
-5.01% |
16.894 |
16.048 |
0.062 |
-0.18% |
0.062 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/51
|
-4.93% |
178.325 |
169.529 |
0.011 |
-0.01% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/16
|
-4.87% |
83.532 |
79.465 |
0.006 |
0.00% |
0.006 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/10
|
-4.84% |
16.894 |
16.076 |
0.016 |
0.08% |
0.016 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint64_t>/51
|
-4.78% |
196.478 |
187.085 |
0.005 |
-0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/51
|
-4.77% |
39.420 |
37.540 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/256
|
-4.77% |
39.419 |
37.541 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint8_t>/999
|
-4.76% |
39.419 |
37.542 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/51
|
-4.69% |
40.046 |
38.166 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/256
|
-4.69% |
40.046 |
38.166 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/51
|
-4.69% |
40.046 |
38.168 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint32_t>/999
|
-4.69% |
40.046 |
38.168 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/256
|
-4.69% |
40.045 |
38.168 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint16_t>/999
|
-4.69% |
40.046 |
38.169 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_three_loads<uint32_t>/28
|
-4.45% |
112.632 |
107.620 |
0.005 |
-0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/51
|
-4.42% |
113.886 |
108.855 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/256
|
-4.35% |
50.371 |
48.178 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/16
|
-4.35% |
50.371 |
48.178 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/999
|
-4.35% |
50.370 |
48.177 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/51
|
-4.35% |
50.371 |
48.179 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint32_t>/28
|
-4.34% |
50.370 |
48.182 |
0.001 |
-0.00% |
0.001 |
|
External/SPEC/CINT2017rate/525.x264_r/525.x264_r
Profile
|
-4.29% |
142.176 |
136.075 |
0.219 |
0.27% |
0.219 |
|
MultiSource/Benchmarks/TSVC/IndirectAddressing-flt/IndirectAddressing-flt
Profile
|
-4.28% |
15.747 |
15.073 |
0.016 |
-0.01% |
0.016 |
|
MultiSource/Applications/JM/lencod/lencod
Profile
|
-4.23% |
35.747 |
34.236 |
0.053 |
0.02% |
0.053 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint16_t>/51
|
-4.00% |
109.500 |
105.115 |
0.003 |
-0.00% |
0.003 |
|
SingleSource/Benchmarks/Misc-C++/oopack_v1p8
Profile
|
-3.94% |
0.479 |
0.460 |
0.000 |
0.01% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC15
|
-3.49% |
273.341 |
263.805 |
2.743 |
0.48% |
2.743 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint64_t>/51
|
-3.30% |
56.940 |
55.060 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_three_loads<uint64_t>/999
|
-3.30% |
56.939 |
55.060 |
0.005 |
-0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/51
|
-3.19% |
40.045 |
38.767 |
0.012 |
-0.02% |
0.012 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_first_two_loads<uint64_t>/999
|
-3.19% |
40.048 |
38.770 |
0.003 |
-0.01% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4BigLoopTC15
|
-3.19% |
272.396 |
263.710 |
2.786 |
0.47% |
2.786 |
|
MultiSource/Benchmarks/MiBench/consumer-lame/consumer-lame
Profile
|
-3.05% |
0.682 |
0.661 |
0.001 |
-0.12% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC15
|
-3.00% |
271.926 |
263.776 |
2.689 |
0.38% |
2.689 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/10
|
-2.96% |
40.674 |
39.471 |
0.035 |
-0.12% |
0.035 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC15
|
-2.94% |
21.274 |
20.648 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/10
|
-2.77% |
41.296 |
40.150 |
0.015 |
-0.04% |
0.015 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint16_t>/16
|
-2.74% |
29.623 |
28.812 |
0.003 |
-0.03% |
0.003 |
|
MultiSource/Benchmarks/VersaBench/8b10b/8b10b
Profile
|
-2.68% |
9.299 |
9.050 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC16
|
-2.67% |
21.214 |
20.648 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4BigLoopTC8
|
-2.55% |
165.902 |
161.665 |
0.483 |
-0.62% |
0.483 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/256
|
-2.49% |
355.305 |
346.453 |
0.016 |
0.00% |
0.016 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/10
|
-2.46% |
30.662 |
29.907 |
0.036 |
0.15% |
0.036 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4BigLoopTC8
|
-2.46% |
166.167 |
162.081 |
0.323 |
-0.61% |
0.323 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC32
|
-2.17% |
28.783 |
28.157 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC31
|
-2.17% |
28.783 |
28.158 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC32
|
-2.17% |
28.782 |
28.157 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC31
|
-2.17% |
28.782 |
28.158 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC31
|
-2.17% |
28.782 |
28.158 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC31
|
-2.16% |
28.779 |
28.157 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC31
|
-2.16% |
28.779 |
28.158 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC31
|
-2.16% |
28.778 |
28.158 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC32
|
-2.15% |
28.777 |
28.157 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC31
|
-2.15% |
28.776 |
28.158 |
0.003 |
0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC32
|
-2.14% |
28.771 |
28.156 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC32
|
-2.12% |
28.768 |
28.158 |
0.000 |
0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/16
|
-2.11% |
55.694 |
54.518 |
0.007 |
0.03% |
0.007 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/16
|
-2.07% |
56.315 |
55.150 |
0.004 |
-0.01% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC15
|
-2.05% |
21.081 |
20.649 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:benchForTruncOrZextVecInLoopWithVW16From_uint32_t_To_uint64_t_
|
-2.05% |
17785.479 |
17421.022 |
120.536 |
0.20% |
120.536 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC32
|
-1.99% |
28.728 |
28.157 |
0.000 |
-0.00% |
0.000 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC32
|
-1.97% |
28.726 |
28.158 |
0.000 |
-0.00% |
0.000 |
|
MultiSource/Benchmarks/TSVC/IndirectAddressing-dbl/IndirectAddressing-dbl
Profile
|
-1.94% |
18.523 |
18.163 |
0.112 |
-0.30% |
0.112 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4BigLoopTC8
|
-1.91% |
165.014 |
161.862 |
0.539 |
-0.34% |
0.539 |
|
MultiSource/Applications/hexxagon/hexxagon
Profile
|
-1.80% |
8.514 |
8.361 |
0.002 |
-0.24% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/16
|
-1.67% |
38.037 |
37.402 |
0.025 |
0.02% |
0.025 |
|
External/SPEC/CFP2017rate/511.povray_r/511.povray_r
Profile
|
-1.63% |
28.635 |
28.167 |
0.117 |
-0.22% |
0.117 |
|
SingleSource/Benchmarks/Shootout/Shootout-matrix
Profile
|
-1.61% |
4.693 |
4.617 |
0.000 |
-0.01% |
0.000 |
|
MultiSource/Applications/sgefa/sgefa
Profile
|
-1.59% |
0.642 |
0.632 |
0.001 |
-0.19% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint8_t>/10
|
-1.49% |
48.905 |
48.175 |
0.003 |
-0.01% |
0.003 |
|
MultiSource/Applications/sqlite3/sqlite3
Profile
|
-1.47% |
12.283 |
12.102 |
0.010 |
-1.17% |
0.010 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC63
|
-1.43% |
43.799 |
43.173 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC64
|
-1.43% |
43.800 |
43.173 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC64
|
-1.43% |
43.799 |
43.173 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW4LoopWithReductionTC63
|
-1.43% |
43.799 |
43.173 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC63
|
-1.43% |
43.799 |
43.174 |
0.003 |
-0.00% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC63
|
-1.43% |
43.800 |
43.175 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW1LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC64
|
-1.43% |
43.798 |
43.173 |
0.002 |
-0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC64
|
-1.43% |
43.798 |
43.173 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC4VW1LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.005 |
-0.00% |
0.005 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC64
|
-1.43% |
43.799 |
43.174 |
0.002 |
0.00% |
0.002 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForLoopWithReductionAutoVecTC63
|
-1.43% |
43.800 |
43.175 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW4LoopWithReductionTC63
|
-1.42% |
43.799 |
43.175 |
0.001 |
-0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC1VW4LoopWithReductionTC63
|
-1.42% |
43.799 |
43.176 |
0.001 |
0.00% |
0.001 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint64_t>/28
|
-1.42% |
85.726 |
84.511 |
0.040 |
-0.03% |
0.040 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_single_load<uint8_t>/10
|
-1.41% |
24.538 |
24.193 |
0.011 |
-0.50% |
0.011 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_no_early_exit_two_loads<uint32_t>/28
|
-1.38% |
86.349 |
85.155 |
0.004 |
-0.00% |
0.004 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC15
|
-1.35% |
20.932 |
20.649 |
0.001 |
0.00% |
0.001 |
|
SingleSource/Benchmarks/CoyoteBench/huffbench
Profile
|
-1.17% |
50.629 |
50.037 |
0.006 |
0.58% |
0.006 |
|
MultiSource/Applications/lambda-0.1.3/lambda
Profile
|
-1.13% |
11.808 |
11.675 |
0.003 |
-0.04% |
0.003 |
|
MicroBenchmarks/LoopVectorization/LoopVectorizationBenchmarks.test:autovec_early_exit_taken_mid_two_loads<uint64_t>/28
|
-1.12% |
53.057 |
52.461 |
0.010 |
0.00% |
0.010 |
|
External/SPEC/CINT2017rate/523.xalancbmk_r/523.xalancbmk_r
Profile
|
-1.12% |
209.538 |
207.189 |
0.154 |
-2.07% |
0.154 |
|
MicroBenchmarks/LoopVectorization/LoopInterleavingBenchmarks.test:benchForIC2VW1LoopWithReductionTC16
|
-1.09% |
20.877 |
20.649 |
0.000 |
0.00% |
0.000 |