Tune MI300X fused MoE Triton kernel JSON config. (#3492)

This commit is contained in:
Wen-Heng (Jack) Chung
2025-02-11 12:27:25 -06:00
committed by GitHub
parent bb418ced80
commit cadd5dbe6a

View File

@@ -72,10 +72,10 @@
"waves_per_eu": 0
},
"64": {
"BLOCK_SIZE_M": 256,
"BLOCK_SIZE_M": 32,
"BLOCK_SIZE_N": 128,
"BLOCK_SIZE_K": 128,
"GROUP_SIZE_M": 1,
"GROUP_SIZE_M": 4,
"num_warps": 4,
"num_stages": 2,
"waves_per_eu": 0