Fused moe triton cfg opt for rocm (#2612)
Co-authored-by: wunhuang <wunhuang@amd.com>
This commit is contained in:
@@ -55,35 +55,35 @@
|
|||||||
"kpack": 2
|
"kpack": 2
|
||||||
},
|
},
|
||||||
"128": {
|
"128": {
|
||||||
"BLOCK_SIZE_M": 16,
|
"BLOCK_SIZE_M": 64,
|
||||||
"BLOCK_SIZE_N": 64,
|
"BLOCK_SIZE_N": 128,
|
||||||
"BLOCK_SIZE_K": 256,
|
"BLOCK_SIZE_K": 128,
|
||||||
"GROUP_SIZE_M": 1,
|
"GROUP_SIZE_M": 1,
|
||||||
"num_warps": 4,
|
"num_warps": 8,
|
||||||
"num_stages": 0,
|
"num_stages": 0,
|
||||||
"waves_per_eu": 1,
|
"waves_per_eu": 0,
|
||||||
"matrix_instr_nonkdim": 16,
|
"matrix_instr_nonkdim": 16,
|
||||||
"kpack": 1
|
"kpack": 1
|
||||||
},
|
},
|
||||||
"256": {
|
"256": {
|
||||||
"BLOCK_SIZE_M": 16,
|
"BLOCK_SIZE_M": 64,
|
||||||
"BLOCK_SIZE_N": 64,
|
"BLOCK_SIZE_N": 128,
|
||||||
"BLOCK_SIZE_K": 256,
|
"BLOCK_SIZE_K": 128,
|
||||||
"GROUP_SIZE_M": 1,
|
"GROUP_SIZE_M": 1,
|
||||||
"num_warps": 4,
|
"num_warps": 8,
|
||||||
"num_stages": 0,
|
"num_stages": 0,
|
||||||
"waves_per_eu": 1,
|
"waves_per_eu": 0,
|
||||||
"matrix_instr_nonkdim": 16,
|
"matrix_instr_nonkdim": 16,
|
||||||
"kpack": 1
|
"kpack": 1
|
||||||
},
|
},
|
||||||
"512": {
|
"512": {
|
||||||
"BLOCK_SIZE_M": 64,
|
"BLOCK_SIZE_M": 64,
|
||||||
"BLOCK_SIZE_N": 64,
|
"BLOCK_SIZE_N": 128,
|
||||||
"BLOCK_SIZE_K": 256,
|
"BLOCK_SIZE_K": 128,
|
||||||
"GROUP_SIZE_M": 1,
|
"GROUP_SIZE_M": 1,
|
||||||
"num_warps": 4,
|
"num_warps": 8,
|
||||||
"num_stages": 0,
|
"num_stages": 0,
|
||||||
"waves_per_eu": 2,
|
"waves_per_eu": 0,
|
||||||
"matrix_instr_nonkdim": 16,
|
"matrix_instr_nonkdim": 16,
|
||||||
"kpack": 2
|
"kpack": 2
|
||||||
},
|
},
|
||||||
|
|||||||
Reference in New Issue
Block a user