Tune MI300X fused MoE Triton kernel JSON config. (#3492)
This commit is contained in:
committed by
GitHub
parent
bb418ced80
commit
cadd5dbe6a
@@ -72,10 +72,10 @@
|
|||||||
"waves_per_eu": 0
|
"waves_per_eu": 0
|
||||||
},
|
},
|
||||||
"64": {
|
"64": {
|
||||||
"BLOCK_SIZE_M": 256,
|
"BLOCK_SIZE_M": 32,
|
||||||
"BLOCK_SIZE_N": 128,
|
"BLOCK_SIZE_N": 128,
|
||||||
"BLOCK_SIZE_K": 128,
|
"BLOCK_SIZE_K": 128,
|
||||||
"GROUP_SIZE_M": 1,
|
"GROUP_SIZE_M": 4,
|
||||||
"num_warps": 4,
|
"num_warps": 4,
|
||||||
"num_stages": 2,
|
"num_stages": 2,
|
||||||
"waves_per_eu": 0
|
"waves_per_eu": 0
|
||||||
|
|||||||
Reference in New Issue
Block a user