test_alignment.py |
12 |
12 |
0 |
0 |
0 |
0 |
0 |
0 |
test_analysis.py |
28 |
28 |
0 |
0 |
0 |
0 |
0 |
0 |
test_aot_inductor.py |
950 |
455 |
480 |
3 |
7 |
5 |
0 |
0 |
test_aot_inductor_arrayref.py |
314 |
128 |
182 |
3 |
1 |
0 |
0 |
0 |
test_aot_inductor_custom_ops.py |
35 |
35 |
0 |
0 |
0 |
0 |
0 |
0 |
test_aot_inductor_package.py |
88 |
59 |
25 |
0 |
4 |
0 |
0 |
0 |
test_async_compile.py |
8 |
8 |
0 |
0 |
0 |
0 |
0 |
0 |
test_augmented_graph_helper.py |
20 |
20 |
0 |
0 |
0 |
0 |
0 |
0 |
test_auto_chunker.py |
9 |
9 |
0 |
0 |
0 |
0 |
0 |
0 |
test_auto_functionalize.py |
39 |
38 |
1 |
0 |
0 |
0 |
0 |
0 |
test_benchmark_fusion.py |
16 |
10 |
4 |
0 |
0 |
2 |
0 |
0 |
test_benchmarking.py |
14 |
8 |
0 |
6 |
0 |
0 |
0 |
0 |
test_best_config.py |
1 |
1 |
0 |
0 |
0 |
0 |
0 |
0 |
test_binary_folding.py |
6 |
5 |
1 |
0 |
0 |
0 |
0 |
0 |
test_block_analysis.py |
10 |
10 |
0 |
0 |
0 |
0 |
0 |
0 |
test_cache.py |
725 |
725 |
0 |
0 |
0 |
0 |
0 |
0 |
test_caching.py |
212 |
212 |
0 |
0 |
0 |
0 |
0 |
0 |
test_ck_backend.py |
34 |
0 |
34 |
0 |
0 |
0 |
0 |
0 |
test_codecache.py |
257 |
193 |
59 |
1 |
3 |
1 |
0 |
0 |
test_codegen_triton.py |
1 |
1 |
0 |
0 |
0 |
0 |
0 |
0 |
test_collective_autotuning.py |
2 |
0 |
2 |
0 |
0 |
0 |
0 |
0 |
test_combo_kernels.py |
77 |
70 |
5 |
0 |
0 |
2 |
0 |
0 |
test_compile.py |
10 |
10 |
0 |
0 |
0 |
0 |
0 |
0 |
test_compile_subprocess.py |
936 |
868 |
40 |
0 |
8 |
20 |
0 |
0 |
test_compile_worker.py |
16 |
16 |
0 |
0 |
0 |
0 |
0 |
0 |
test_compiled_autograd.py |
875 |
865 |
4 |
5 |
1 |
0 |
0 |
0 |
test_compiled_optimizers.py |
682 |
677 |
3 |
0 |
0 |
2 |
0 |
0 |
test_config.py |
14 |
14 |
0 |
0 |
0 |
0 |
0 |
0 |
test_control_deps.py |
4 |
4 |
0 |
0 |
0 |
0 |
0 |
0 |
test_control_flow.py |
741 |
739 |
2 |
0 |
0 |
0 |
0 |
0 |
test_cooperative_reductions.py |
163 |
160 |
0 |
0 |
0 |
3 |
0 |
0 |
test_coordinate_descent_tuner.py |
5 |
5 |
0 |
0 |
0 |
0 |
0 |
0 |
test_cpp_wrapper_hipify.py |
3 |
3 |
0 |
0 |
0 |
0 |
0 |
0 |
test_cpu_repro.py |
752 |
747 |
5 |
0 |
0 |
0 |
0 |
0 |
test_cuda_repro.py |
98 |
78 |
3 |
1 |
3 |
13 |
0 |
0 |
test_cudacodecache.py |
3 |
0 |
0 |
0 |
3 |
0 |
0 |
0 |
test_cudagraph_trees.py |
189 |
180 |
8 |
0 |
0 |
1 |
0 |
0 |
test_cudagraph_trees_expandable_segments.py |
158 |
151 |
6 |
0 |
0 |
1 |
0 |
0 |
test_custom_lowering.py |
6 |
4 |
1 |
0 |
1 |
0 |
0 |
0 |
test_custom_op_autotune.py |
4 |
3 |
0 |
0 |
1 |
0 |
0 |
0 |
test_custom_partitioner_fn.py |
1 |
1 |
0 |
0 |
0 |
0 |
0 |
0 |
test_custom_post_grad_passes.py |
6 |
6 |
0 |
0 |
0 |
0 |
0 |
0 |
test_cutedsl_grouped_mm.py |
24 |
0 |
24 |
0 |
0 |
0 |
0 |
0 |
test_cutedsl_template.py |
13 |
0 |
13 |
0 |
0 |
0 |
0 |
0 |
test_cutlass_backend.py |
181 |
0 |
181 |
0 |
0 |
0 |
0 |
0 |
test_cutlass_evt.py |
8 |
0 |
8 |
0 |
0 |
0 |
0 |
0 |
test_debug_trace.py |
3 |
3 |
0 |
0 |
0 |
0 |
0 |
0 |
test_decompose_mem_bound_mm.py |
37 |
35 |
2 |
0 |
0 |
0 |
0 |
0 |
test_dependencies.py |
5 |
5 |
0 |
0 |
0 |
0 |
0 |
0 |
test_deterministic.py |
32 |
7 |
0 |
0 |
24 |
1 |
0 |
0 |
test_device_assert.py |
8 |
8 |
0 |
0 |
0 |
0 |
0 |
0 |
test_distributed_patterns.py |
20 |
20 |
0 |
0 |
0 |
0 |
0 |
0 |
test_efficient_conv_bn_eval.py |
2 |
1 |
0 |
0 |
0 |
1 |
0 |
0 |
test_exc_lowering_stack_trace.py |
2 |
2 |
0 |
0 |
0 |
0 |
0 |
0 |
test_extension_backend.py |
1 |
0 |
0 |
0 |
1 |
0 |
0 |
0 |
test_external_callables.py |
3 |
2 |
0 |
0 |
0 |
1 |
0 |
0 |
test_flex_attention.py |
786 |
750 |
27 |
1 |
2 |
6 |
0 |
0 |
test_flex_decoding.py |
557 |
554 |
1 |
2 |
0 |
0 |
0 |
0 |
test_flex_flash.py |
210 |
6 |
204 |
0 |
0 |
0 |
0 |
0 |
test_foreach.py |
607 |
581 |
26 |
0 |
0 |
0 |
0 |
0 |
test_fp8.py |
202 |
98 |
55 |
0 |
44 |
5 |
0 |
0 |
test_fused_attention.py |
116 |
114 |
2 |
0 |
0 |
0 |
0 |
0 |
test_fusion_regions.py |
6 |
6 |
0 |
0 |
0 |
0 |
0 |
0 |
test_fuzzer.py |
11 |
10 |
1 |
0 |
0 |
0 |
0 |
0 |
test_fx_fusion.py |
4 |
4 |
0 |
0 |
0 |
0 |
0 |
0 |
test_fxir_backend.py |
75 |
74 |
0 |
0 |
0 |
1 |
0 |
0 |
test_gpu_cpp_wrapper.py |
297 |
289 |
1 |
0 |
3 |
4 |
0 |
0 |
test_gpu_select_algorithm.py |
58 |
58 |
0 |
0 |
0 |
0 |
0 |
0 |
test_graph_transform_observer.py |
1 |
1 |
0 |
0 |
0 |
0 |
0 |
0 |
test_group_batch_fusion.py |
13 |
11 |
2 |
0 |
0 |
0 |
0 |
0 |
test_halide.py |
4 |
0 |
4 |
0 |
0 |
0 |
0 |
0 |
test_helion_kernels.py |
2 |
0 |
2 |
0 |
0 |
0 |
0 |
0 |
test_indexing.py |
22 |
22 |
0 |
0 |
0 |
0 |
0 |
0 |
test_inductor_annotations.py |
2 |
2 |
0 |
0 |
0 |
0 |
0 |
0 |
test_inductor_freezing.py |
48 |
33 |
1 |
0 |
1 |
2 |
0 |
11 |
test_inductor_scheduler.py |
10 |
10 |
0 |
0 |
0 |
0 |
0 |
0 |
test_inductor_utils.py |
2 |
2 |
0 |
0 |
0 |
0 |
0 |
0 |
test_inplace_padding.py |
9 |
9 |
0 |
0 |
0 |
0 |
0 |
0 |
test_inplacing_pass.py |
23 |
23 |
0 |
0 |
0 |
0 |
0 |
0 |
test_kernel_benchmark.py |
19 |
16 |
0 |
0 |
0 |
3 |
0 |
0 |
test_kernel_optimization.py |
1 |
1 |
0 |
0 |
0 |
0 |
0 |
0 |
test_lookup_table.py |
37 |
29 |
4 |
0 |
0 |
0 |
0 |
4 |
test_loop_ordering.py |
70 |
68 |
0 |
0 |
0 |
2 |
0 |
0 |
test_max_autotune.py |
472 |
182 |
173 |
0 |
17 |
95 |
0 |
5 |
test_max_autotune_blackwell.py |
106 |
2 |
104 |
0 |
0 |
0 |
0 |
0 |
test_mem_estimation.py |
4 |
4 |
0 |
0 |
0 |
0 |
0 |
0 |
test_memory.py |
8 |
8 |
0 |
0 |
0 |
0 |
0 |
0 |
test_memory_planning.py |
4 |
4 |
0 |
0 |
0 |
0 |
0 |
0 |
test_metrics.py |
6 |
6 |
0 |
0 |
0 |
0 |
0 |
0 |
test_minifier.py |
14 |
14 |
0 |
0 |
0 |
0 |
0 |
0 |
test_minifier_isolate.py |
2 |
1 |
1 |
0 |
0 |
0 |
0 |
0 |
test_minifier_utils.py |
3 |
3 |
0 |
0 |
0 |
0 |
0 |
0 |
test_mix_order_reduction.py |
485 |
248 |
219 |
0 |
16 |
1 |
0 |
1 |
test_mkldnn_pattern_matcher.py |
20 |
16 |
4 |
0 |
0 |
0 |
0 |
0 |
test_mmdecomp.py |
28 |
28 |
0 |
0 |
0 |
0 |
0 |
0 |
test_move_constructors_to_gpu.py |
7 |
6 |
1 |
0 |
0 |
0 |
0 |
0 |
test_mps_basic.py |
46 |
24 |
0 |
0 |
20 |
2 |
0 |
0 |
test_multi_kernel.py |
19 |
16 |
2 |
0 |
0 |
1 |
0 |
0 |
test_native_matmul.py |
14 |
2 |
0 |
0 |
4 |
8 |
0 |
0 |
test_needs_exact_strides.py |
2 |
2 |
0 |
0 |
0 |
0 |
0 |
0 |
test_nv_universal_gemm.py |
23 |
3 |
20 |
0 |
0 |
0 |
0 |
0 |
test_online_softmax.py |
31 |
31 |
0 |
0 |
0 |
0 |
0 |
0 |
test_op_completeness.py |
5 |
4 |
1 |
0 |
0 |
0 |
0 |
0 |
test_op_dtype_prop.py |
581 |
581 |
0 |
0 |
0 |
0 |
0 |
0 |
test_ordered_set.py |
401 |
386 |
15 |
0 |
0 |
0 |
0 |
0 |
test_pad_mm.py |
19 |
8 |
1 |
0 |
0 |
10 |
0 |
0 |
test_padding.py |
55 |
46 |
9 |
0 |
0 |
0 |
0 |
0 |
test_pattern_matcher.py |
63 |
56 |
0 |
0 |
0 |
7 |
0 |
0 |
test_perf.py |
66 |
66 |
0 |
0 |
0 |
0 |
0 |
0 |
test_profiler.py |
8 |
8 |
0 |
0 |
0 |
0 |
0 |
0 |
test_provenance_tracing.py |
16 |
16 |
0 |
0 |
0 |
0 |
0 |
0 |
test_quantization.py |
2 |
2 |
0 |
0 |
0 |
0 |
0 |
0 |
test_remote_cache.py |
3 |
3 |
0 |
0 |
0 |
0 |
0 |
0 |
test_scatter_optimization.py |
8 |
8 |
0 |
0 |
0 |
0 |
0 |
0 |
test_segmented_tree.py |
12 |
12 |
0 |
0 |
0 |
0 |
0 |
0 |
test_select_algorithm.py |
27 |
14 |
1 |
0 |
2 |
9 |
0 |
1 |
test_selective_lowering.py |
2 |
2 |
0 |
0 |
0 |
0 |
0 |
0 |
test_smoke.py |
3 |
3 |
0 |
0 |
0 |
0 |
0 |
0 |
test_snode_runtime.py |
22 |
22 |
0 |
0 |
0 |
0 |
0 |
0 |
test_split_cat_fx_aten_passes.py |
5 |
5 |
0 |
0 |
0 |
0 |
0 |
0 |
test_split_cat_fx_passes.py |
11 |
11 |
0 |
0 |
0 |
0 |
0 |
0 |
test_static_triton_launcher.py |
17 |
17 |
0 |
0 |
0 |
0 |
0 |
0 |
test_subgraph_choice.py |
2 |
2 |
0 |
0 |
0 |
0 |
0 |
0 |
test_template_heuristics_registry.py |
7 |
7 |
0 |
0 |
0 |
0 |
0 |
0 |
test_torchbind.py |
16 |
16 |
0 |
0 |
0 |
0 |
0 |
0 |
test_torchinductor.py |
1041 |
962 |
42 |
0 |
3 |
22 |
0 |
12 |
test_torchinductor_codegen_config_overrides.py |
4 |
4 |
0 |
0 |
0 |
0 |
0 |
0 |
test_torchinductor_codegen_dynamic_shapes.py |
1864 |
1466 |
200 |
195 |
3 |
0 |
0 |
0 |
test_torchinductor_dynamic_shapes.py |
1933 |
1808 |
112 |
3 |
8 |
2 |
0 |
0 |
test_torchinductor_opinfo.py |
3674 |
1555 |
691 |
38 |
0 |
1390 |
0 |
0 |
test_torchinductor_strided_blocks.py |
304 |
86 |
194 |
0 |
1 |
1 |
0 |
22 |
test_triton_extension_backend.py |
3 |
3 |
0 |
0 |
0 |
0 |
0 |
0 |
test_triton_helpers.py |
2 |
2 |
0 |
0 |
0 |
0 |
0 |
0 |
test_triton_heuristics.py |
13 |
11 |
1 |
0 |
1 |
0 |
0 |
0 |
test_triton_kernels.py |
372 |
326 |
44 |
0 |
0 |
2 |
0 |
0 |
test_triton_syntax.py |
1 |
1 |
0 |
0 |
0 |
0 |
0 |
0 |
test_triton_wrapper.py |
2 |
2 |
0 |
0 |
0 |
0 |
0 |
0 |
test_unbacked_symints.py |
34 |
32 |
0 |
0 |
0 |
2 |
0 |
0 |
test_utils.py |
11 |
11 |
0 |
0 |
0 |
0 |
0 |
0 |
test_xpu_basic.py |
4 |
4 |
0 |
0 |
0 |
0 |
0 |
0 |
| Total |
22962 |
17580 |
3258 |
258 |
182 |
1628 |
0 |
56 |
gfx1250 PyTorch Inductor Merged Outcome
Updated:
2026-08-05Important
This is the latest exact-node merge, not a fully completed targeted rerun.
All seven high-MISS full suites completed, but the targeted phase completed
27 / 1,848nodes and left1,821reruns pending.Overall Result
22,962PASSEDoutcomes:17,58021,096 / 22,962 = 91.87%1,866 / 22,962 = 8.13%Suite Summary
test_alignment.pytest_analysis.pytest_aot_inductor.pytest_aot_inductor_arrayref.pytest_aot_inductor_custom_ops.pytest_aot_inductor_package.pytest_async_compile.pytest_augmented_graph_helper.pytest_auto_chunker.pytest_auto_functionalize.pytest_benchmark_fusion.pytest_benchmarking.pytest_best_config.pytest_binary_folding.pytest_block_analysis.pytest_cache.pytest_caching.pytest_ck_backend.pytest_codecache.pytest_codegen_triton.pytest_collective_autotuning.pytest_combo_kernels.pytest_compile.pytest_compile_subprocess.pytest_compile_worker.pytest_compiled_autograd.pytest_compiled_optimizers.pytest_config.pytest_control_deps.pytest_control_flow.pytest_cooperative_reductions.pytest_coordinate_descent_tuner.pytest_cpp_wrapper_hipify.pytest_cpu_repro.pytest_cuda_repro.pytest_cudacodecache.pytest_cudagraph_trees.pytest_cudagraph_trees_expandable_segments.pytest_custom_lowering.pytest_custom_op_autotune.pytest_custom_partitioner_fn.pytest_custom_post_grad_passes.pytest_cutedsl_grouped_mm.pytest_cutedsl_template.pytest_cutlass_backend.pytest_cutlass_evt.pytest_debug_trace.pytest_decompose_mem_bound_mm.pytest_dependencies.pytest_deterministic.pytest_device_assert.pytest_distributed_patterns.pytest_efficient_conv_bn_eval.pytest_exc_lowering_stack_trace.pytest_extension_backend.pytest_external_callables.pytest_flex_attention.pytest_flex_decoding.pytest_flex_flash.pytest_foreach.pytest_fp8.pytest_fused_attention.pytest_fusion_regions.pytest_fuzzer.pytest_fx_fusion.pytest_fxir_backend.pytest_gpu_cpp_wrapper.pytest_gpu_select_algorithm.pytest_graph_transform_observer.pytest_group_batch_fusion.pytest_halide.pytest_helion_kernels.pytest_indexing.pytest_inductor_annotations.pytest_inductor_freezing.pytest_inductor_scheduler.pytest_inductor_utils.pytest_inplace_padding.pytest_inplacing_pass.pytest_kernel_benchmark.pytest_kernel_optimization.pytest_lookup_table.pytest_loop_ordering.pytest_max_autotune.pytest_max_autotune_blackwell.pytest_mem_estimation.pytest_memory.pytest_memory_planning.pytest_metrics.pytest_minifier.pytest_minifier_isolate.pytest_minifier_utils.pytest_mix_order_reduction.pytest_mkldnn_pattern_matcher.pytest_mmdecomp.pytest_move_constructors_to_gpu.pytest_mps_basic.pytest_multi_kernel.pytest_native_matmul.pytest_needs_exact_strides.pytest_nv_universal_gemm.pytest_online_softmax.pytest_op_completeness.pytest_op_dtype_prop.pytest_ordered_set.pytest_pad_mm.pytest_padding.pytest_pattern_matcher.pytest_perf.pytest_profiler.pytest_provenance_tracing.pytest_quantization.pytest_remote_cache.pytest_scatter_optimization.pytest_segmented_tree.pytest_select_algorithm.pytest_selective_lowering.pytest_smoke.pytest_snode_runtime.pytest_split_cat_fx_aten_passes.pytest_split_cat_fx_passes.pytest_static_triton_launcher.pytest_subgraph_choice.pytest_template_heuristics_registry.pytest_torchbind.pytest_torchinductor.pytest_torchinductor_codegen_config_overrides.pytest_torchinductor_codegen_dynamic_shapes.pytest_torchinductor_dynamic_shapes.pytest_torchinductor_opinfo.pytest_torchinductor_strided_blocks.pytest_triton_extension_backend.pytest_triton_helpers.pytest_triton_heuristics.pytest_triton_kernels.pytest_triton_syntax.pytest_triton_wrapper.pytest_unbacked_symints.pytest_utils.pytest_xpu_basic.pyExecution and Merge Notes
Continuation coverage
6,254 / 6,254completed.27 / 1,848completed;1,821pending.6,281 / 8,102.16,681latest outcomes from ROCm/HIP7.15.0and6,281from ROCm/HIP7.15.26305.Stop-condition correction
2.5immediately adjacent executed tests trigger SIGABRT; any other completed test resets the streak.Notable outcomes
test_builtin_score_mods_different_block_size_score_mod6_BLOCK_SIZE3_cuda_float16passed in the continuation.test_layer_norm_bwd_with_dynamic_shape_dynamic_dims2remainsMISSEDand is still pending in the targeted phase.Largest unresolved clusters
test_torchinductor_opinfo.py(1,390) andtest_max_autotune.py(95).test_fp8.py(44),test_deterministic.py(24), andtest_mps_basic.py(20).test_torchinductor_strided_blocks.py(22),test_torchinductor.py(12), andtest_inductor_freezing.py(11).Environment
gfx12502.11.0+rocm7.15.0a20260721, ROCm/HIP7.15.0, Triton3.8.02.11.0+rocm10.1.0a20260803, ROCm/HIP7.15.26305, Triton3.8.0Detailed reproduction steps and high-risk test notes are in this issue comment.
The comments below are intermediate checkpoints and can be ignored.