state: generation 10907 (intent)

This commit is contained in:
2026-09-20 20:50:51 +00:00
parent e7d526346a
commit 1b4aaa4104
9 changed files with 3090 additions and 2295 deletions

View File

@@ -303,6 +303,44 @@
"targetGpu": "Ascend_910-b3",
"taskType": "text-generation"
},
"ascend_910-b4|vllm_tokenizer_patch|text-generation|model_type:gemma4": {
"architectureSignature": "model_type:gemma4",
"architectures": [],
"evidenceCount": 1,
"expiresAt": "2026-10-20T02:52:43.762431+00:00",
"framework": "vllm_tokenizer_patch",
"latestFailureAt": "2026-09-20T02:52:43.762431+00:00",
"latestSuccessfulAt": null,
"matchType": "model_type",
"modelType": "gemma4",
"sourceModelIds": [
"mlx-community/gemma-4-e4b-it-OptiQ-4bit"
],
"sourceTaskIds": [
"4968999"
],
"targetGpu": "Ascend_910-b4",
"taskType": "text-generation"
},
"ascend_910-b4|vllm_tokenizer_patch|text-generation|model_type:qwen3_5": {
"architectureSignature": "model_type:qwen3_5",
"architectures": [],
"evidenceCount": 1,
"expiresAt": "2026-10-20T02:52:43.746995+00:00",
"framework": "vllm_tokenizer_patch",
"latestFailureAt": "2026-09-20T02:52:43.746995+00:00",
"latestSuccessfulAt": null,
"matchType": "model_type",
"modelType": "qwen3_5",
"sourceModelIds": [
"Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed"
],
"sourceTaskIds": [
"4968998"
],
"targetGpu": "Ascend_910-b4",
"taskType": "text-generation"
},
"ascend_910-b4|vllm|text-generation|architectures:nanbeigeforcausallm": {
"architectureSignature": "architectures:nanbeigeforcausallm",
"architectures": [
@@ -430,7 +468,7 @@
"ascend_910-b4|vllm|text-generation|model_type:qwen3_5": {
"architectureSignature": "model_type:qwen3_5",
"architectures": [],
"evidenceCount": 6,
"evidenceCount": 4,
"expiresAt": "2026-10-17T03:21:21+00:00",
"framework": "vllm",
"latestFailureAt": "2026-09-17T03:21:21+00:00",
@@ -441,15 +479,13 @@
"r0b0tlab/Qwen3.8-27B-NVFP4-MTP-sm121",
"ornith-ai/Ornith-1.5-9B-MLX-8bit",
"cyankiwi/Ornith-1.5-9B-AWQ-INT4",
"Jackrong/DeepSeek-V4-Pro-Qwen3.5-9B",
"cyankiwi/Ornith-1.5-9B-AWQ-FP8"
"Jackrong/DeepSeek-V4-Pro-Qwen3.5-9B"
],
"sourceTaskIds": [
"4610340",
"4610335",
"4610341",
"4610339",
"4610336"
"4610339"
],
"targetGpu": "Ascend_910-b4",
"taskType": "text-generation"
@@ -948,7 +984,7 @@
"hygon_k100-ai|vllm|text-generation|model_type:qwen3_5": {
"architectureSignature": "model_type:qwen3_5",
"architectures": [],
"evidenceCount": 5,
"evidenceCount": 4,
"expiresAt": "2026-10-20T18:51:21+00:00",
"framework": "vllm",
"latestFailureAt": "2026-09-20T18:51:21+00:00",
@@ -959,15 +995,13 @@
"logic65/Qwen3.8-Whittle-w50-18.3B-unrepaired",
"ewinregirgojr/Qwen3.8-14B-Instruct-Turbo",
"VerStella/Qwen3.8-27B-heretic-ara-W8A8-DCU",
"EschaLabs/Qwen3.8-27B-Escha-W2",
"douyamv/Qwen3.8-27B-FP8"
"EschaLabs/Qwen3.8-27B-Escha-W2"
],
"sourceTaskIds": [
"4523697",
"4490374",
"4460362",
"4595432",
"4609095"
"4595432"
],
"targetGpu": "hygon_k100-ai",
"taskType": "text-generation"
@@ -1378,19 +1412,23 @@
"kunlunxin_p-800|vllm_tokenizer_patch|text-generation|model_type:qwen3_5_moe": {
"architectureSignature": "model_type:qwen3_5_moe",
"architectures": [],
"evidenceCount": 3,
"expiresAt": "2026-10-20T02:52:51.436311+00:00",
"evidenceCount": 5,
"expiresAt": "2026-10-20T03:09:17.835766+00:00",
"framework": "vllm_tokenizer_patch",
"latestFailureAt": "2026-09-20T02:52:51.436311+00:00",
"latestFailureAt": "2026-09-20T03:09:17.835766+00:00",
"latestSuccessfulAt": null,
"matchType": "model_type",
"modelType": "qwen3_5_moe",
"sourceModelIds": [
"apodex/Apodex-1.1-mini-GPTQ-Int4",
"cyankiwi/Apodex-1.1-mini-AWQ-INT4",
"mlx-community/Qwen3.5-122B-A10B-OptiQ-2bit-REAP-63B",
"primitive-ai/Nex-N2.5-mini-mixed-NVFP4-FP8",
"primitive-ai/Nex-N2.5-mini-NVFP4"
],
"sourceTaskIds": [
"4969320",
"4969047",
"4969042",
"4969040",
"4969041"
@@ -1398,6 +1436,25 @@
"targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation"
},
"kunlunxin_p-800|vllm|text-generation|model_type:deepseek_v4": {
"architectureSignature": "model_type:deepseek_v4",
"architectures": [],
"evidenceCount": 1,
"expiresAt": "2026-10-20T02:52:51.889837+00:00",
"framework": "vllm",
"latestFailureAt": "2026-09-20T02:52:51.889837+00:00",
"latestSuccessfulAt": null,
"matchType": "model_type",
"modelType": "deepseek_v4",
"sourceModelIds": [
"JANGQ-AI/DeepSeek-V4-Flash-JANGTQ-K"
],
"sourceTaskIds": [
"4969069"
],
"targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation"
},
"metax_c-500|vllm|text-generation|model_type:mellum": {
"architectureSignature": "model_type:mellum",
"architectures": [],
@@ -1630,7 +1687,7 @@
"mthreads_s4000|vllm|text-generation|model_type:qwen3_5_moe": {
"architectureSignature": "model_type:qwen3_5_moe",
"architectures": [],
"evidenceCount": 2,
"evidenceCount": 1,
"expiresAt": "2026-10-15T07:13:14.093936+00:00",
"framework": "vllm",
"latestFailureAt": "2026-09-15T07:13:14.093936+00:00",
@@ -1638,12 +1695,10 @@
"matchType": "model_type",
"modelType": "qwen3_5_moe",
"sourceModelIds": [
"mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit",
"mlx-community/XYZ-Aquila-mini-OptiQ-4bit"
"mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit"
],
"sourceTaskIds": [
"4868226",
"4868227"
"4868226"
],
"targetGpu": "Mthreads_s4000",
"taskType": "text-generation"
@@ -2209,19 +2264,21 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-09-20T20:31:52.456913+00:00",
"generatedAt": "2026-09-20T20:49:20.618975+00:00",
"summary": {
"activeBlockCount": 108,
"activeBlockCount": 111,
"byGpuFramework": {
"Ascend_910-b3|vllm": 11,
"Ascend_910-b3|vllm_tokenizer_patch": 4,
"Ascend_910-b4|vllm": 12,
"Ascend_910-b4|vllm_tokenizer_patch": 2,
"Biren_166m|vllm": 3,
"Cambricon_mlu-370-x8|vllm": 6,
"Cambricon_mlu-370-x8|vllm-mlu": 5,
"Iluvatar_bi-150|transformers": 4,
"Iluvatar_bi-150|vllm": 7,
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0": 3,
"Kunlunxin_p-800|vllm": 1,
"Kunlunxin_p-800|vllm_tokenizer_patch": 2,
"MetaX_c-500|vllm": 7,
"Mthreads_s4000|vllm": 9,

View File

@@ -425,7 +425,7 @@
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-09-20T20:48:18.654035+00:00",
"generatedAt": "2026-09-20T20:49:36.415249+00:00",
"gpuStats": {
"Ascend_910-b3": {
"available": true,

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -14,13 +14,13 @@
100
],
"accounts": 12,
"activeScanned": 1160,
"activeScanned": 1154,
"ageCleanupMode": "admission_only",
"agePolicySkipped": {
"cleanupDisabled": true,
"reason": "admission_only"
},
"architectureBlockCount": 108,
"architectureBlockCount": 111,
"architectureFrameworkCatalog": {
"ascend_910-b3|text-generation": [
"llamacpp",
@@ -67,90 +67,153 @@
]
},
"architectureFrameworkCatalogErrors": {},
"architectureIncompatibleCount": 4,
"architectureIncompatibleCount": 7,
"architectureIncompatibleTasks": [
{
"accountIndex": 3,
"architectureBlockEvidenceCount": 3,
"architectureBlockExpiresAt": "2026-10-20T02:52:51.436311+00:00",
"accountIndex": 2,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5_moe",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5_moe"
"model_type:qwen3_5"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Kunlunxin_p-800",
"modelId": "mlx-community/Qwen3.6-35B-A3B-OptiQ-4bit-REAP-19B",
"gpuType": "Ascend_910-b4",
"modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed-FP16",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969638,
"taskId": 4969016,
"taskType": "text-generation"
},
{
"accountIndex": 3,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Ascend_910-b4",
"modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4782245,
"taskType": "text-generation"
},
{
"accountIndex": 3,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Ascend_910-b4",
"modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969005,
"taskType": "text-generation"
},
{
"accountIndex": 3,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.762431+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:gemma4",
"architectureSignatures": [
"model_type:gemma4"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Ascend_910-b4",
"modelId": "mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969017,
"taskType": "text-generation"
},
{
"accountIndex": 5,
"architectureBlockEvidenceCount": 3,
"architectureBlockExpiresAt": "2026-10-20T02:52:51.436311+00:00",
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5_moe",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5_moe"
"model_type:qwen3_5"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Kunlunxin_p-800",
"modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit",
"gpuType": "Ascend_910-b4",
"modelId": "rapid-mlx/Qwen3.8-27B-4bit-MTP-MLX",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969802,
"taskId": 4899358,
"taskType": "text-generation"
},
{
"accountIndex": 7,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Ascend_910-b4",
"modelId": "Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969013,
"taskType": "text-generation"
},
{
"accountIndex": 8,
"architectureBlockEvidenceCount": 3,
"architectureBlockExpiresAt": "2026-10-20T02:52:51.436311+00:00",
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5_moe",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5_moe"
"model_type:qwen3_5"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Kunlunxin_p-800",
"modelId": "ornith-ai/Ornith-1.5-35B-A3B-NVFP4",
"gpuType": "Ascend_910-b4",
"modelId": "ornith-ai/Ornith-1.5-9B-MLX-6bit",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969619,
"taskType": "text-generation"
},
{
"accountIndex": 9,
"architectureBlockEvidenceCount": 3,
"architectureBlockExpiresAt": "2026-10-20T02:52:51.436311+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5_moe",
"architectureSignatures": [
"model_type:qwen3_5_moe"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Kunlunxin_p-800",
"modelId": "mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969335,
"taskId": 4895729,
"taskType": "text-generation"
}
],
@@ -176,29 +239,29 @@
"unsloth/LFM2-2.6B-Exp-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/unsloth/LFM2-2.6B-Exp-GGUF/resolve/master/config.json (status=404)",
"z-lab/Qwen3.8-27B-DFlash2-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/z-lab/Qwen3.8-27B-DFlash2-GGUF/resolve/master/config.json (status=404)"
},
"architectureModelConfigsComplete": 64,
"architectureModelConfigsComplete": 62,
"architectureOnly": true,
"architecturePolicySkipped": {
"frameworkCatalogUnknown": 0,
"frameworkContextUnknown": 76,
"frameworkContextUnknown": 74,
"modelArchitectureUnknown": 134,
"noMatchingBlock": 997,
"noMatchingBlock": 988,
"partiallyBlockedFrameworkSet": 25,
"runningMatchedProtected": 0,
"submissionContextMismatch": 0,
"submissionContextUnknown": 0
},
"cancelledCount": 4,
"cancelledCount": 7,
"cancelledTasks": [
{
"accountIndex": 3,
"architectureBlockEvidenceCount": 3,
"architectureBlockExpiresAt": "2026-10-20T02:52:51.436311+00:00",
"accountIndex": 2,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5_moe",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5_moe"
"model_type:qwen3_5"
],
"cleanupReasons": [
"known_framework_architecture_incompatible"
@@ -207,22 +270,94 @@
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Kunlunxin_p-800",
"modelId": "mlx-community/Qwen3.6-35B-A3B-OptiQ-4bit-REAP-19B",
"gpuType": "Ascend_910-b4",
"modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed-FP16",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969638,
"taskId": 4969016,
"taskType": "text-generation"
},
{
"accountIndex": 3,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5"
],
"cleanupReasons": [
"known_framework_architecture_incompatible"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Ascend_910-b4",
"modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4782245,
"taskType": "text-generation"
},
{
"accountIndex": 3,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5"
],
"cleanupReasons": [
"known_framework_architecture_incompatible"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Ascend_910-b4",
"modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969005,
"taskType": "text-generation"
},
{
"accountIndex": 3,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.762431+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:gemma4",
"architectureSignatures": [
"model_type:gemma4"
],
"cleanupReasons": [
"known_framework_architecture_incompatible"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Ascend_910-b4",
"modelId": "mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969017,
"taskType": "text-generation"
},
{
"accountIndex": 5,
"architectureBlockEvidenceCount": 3,
"architectureBlockExpiresAt": "2026-10-20T02:52:51.436311+00:00",
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5_moe",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5_moe"
"model_type:qwen3_5"
],
"cleanupReasons": [
"known_framework_architecture_incompatible"
@@ -231,22 +366,46 @@
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Kunlunxin_p-800",
"modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit",
"gpuType": "Ascend_910-b4",
"modelId": "rapid-mlx/Qwen3.8-27B-4bit-MTP-MLX",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969802,
"taskId": 4899358,
"taskType": "text-generation"
},
{
"accountIndex": 7,
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5"
],
"cleanupReasons": [
"known_framework_architecture_incompatible"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Ascend_910-b4",
"modelId": "Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969013,
"taskType": "text-generation"
},
{
"accountIndex": 8,
"architectureBlockEvidenceCount": 3,
"architectureBlockExpiresAt": "2026-10-20T02:52:51.436311+00:00",
"architectureBlockEvidenceCount": 1,
"architectureBlockExpiresAt": "2026-10-20T02:52:43.746995+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5_moe",
"architectureSignature": "model_type:qwen3_5",
"architectureSignatures": [
"model_type:qwen3_5_moe"
"model_type:qwen3_5"
],
"cleanupReasons": [
"known_framework_architecture_incompatible"
@@ -255,41 +414,17 @@
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Kunlunxin_p-800",
"modelId": "ornith-ai/Ornith-1.5-35B-A3B-NVFP4",
"gpuType": "Ascend_910-b4",
"modelId": "ornith-ai/Ornith-1.5-9B-MLX-6bit",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969619,
"taskType": "text-generation"
},
{
"accountIndex": 9,
"architectureBlockEvidenceCount": 3,
"architectureBlockExpiresAt": "2026-10-20T02:52:51.436311+00:00",
"architectureMatchScope": "exact_framework",
"architectureMatchType": "model_type",
"architectureSignature": "model_type:qwen3_5_moe",
"architectureSignatures": [
"model_type:qwen3_5_moe"
],
"cleanupReasons": [
"known_framework_architecture_incompatible"
],
"evaluatedFrameworks": [
"vllm_tokenizer_patch"
],
"framework": "vllm_tokenizer_patch",
"gpuType": "Kunlunxin_p-800",
"modelId": "mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B",
"reason": "known_framework_architecture_incompatible",
"status": "waiting",
"taskId": 4969335,
"taskId": 4895729,
"taskType": "text-generation"
}
],
"certainOomCount": 0,
"certainOomTasks": [],
"cleanupCandidateCount": 4,
"cleanupCandidateCount": 7,
"dryRun": false,
"listingErrors": {},
"modelAgeErrors": {},
@@ -314,7 +449,7 @@
],
"oldOverflowCount": 0,
"oldOverflowTasks": [],
"policyCancelledRecorded": 4,
"policyCancelledRecorded": 7,
"policyNoLongerAppliesCount": 0,
"policyNoLongerAppliesTasks": [],
"recentModelDays": 7,
@@ -327,5 +462,5 @@
"repositorySizeUnknown": 0
},
"stopErrors": [],
"uniqueModels": 341
"uniqueModels": 338
}

View File

@@ -55,6 +55,7 @@
{"failReason": "memory_capacity", "failureAction": "reject_if_estimated_load_exceeds_memory", "failureCategory": "memory_capacity", "failureClassificationReason": "structured_oom", "failureCode": "PREFLIGHT_OOM", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureObservedGpuMemoryGiB": 32.0, "failureScope": "model_gpu", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T13:30:42.156001+00:00", "modelId": "mlx-community/Qwen3.5-122B-A10B-OptiQ-2bit-REAP-63B", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 25593158704, "estimatedRequiredGiB": 28.626, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 25613691403, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7626590448, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:expert-pruning", "custom_tag:reap", "custom_tag:moe", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 25613691403}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T05:16:13.139470+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4971297", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T13:30:42.155948+00:00", "modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21930054875, "estimatedRequiredGiB": 24.532, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 21950415614, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6024588416, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21950415614}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T05:16:13.138764+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4971298", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T13:30:42.155984+00:00", "modelId": "mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23524076260, "estimatedRequiredGiB": 26.327, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 23556606333, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6649116524, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 23556606333}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T05:16:13.136785+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4971300", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "参数/模板问题", "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T20:49:19.683254+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-20T04:46:07.939821+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970952", "taskType": "text-generation", "verifyResult": null}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:58:28.353121+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:42:19.137834+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970925", "taskType": "text-generation", "verifyResult": -1}
{"failReason": null, "framework": "", "lastSyncTime": "2026-09-20T04:34:35.568144+00:00", "modelId": "AI-ModelScope/granite-20b-code-instruct", "modelProfile": {}, "outcome": "success", "status": "success", "submitTime": "2026-09-20T04:31:21+00:00", "targetGpu": "MetaX_c-500", "taskId": "4079070", "taskType": "text-generation", "verifyResult": 1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:46:14.462925+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:31:05.666143+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970790", "taskType": "text-generation", "verifyResult": -1}
@@ -66,6 +67,7 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:34:29.856389+00:00", "modelId": "BAAI/AquilaMed-RL", "modelProfile": {"architectures": ["AquilaForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6968, "estimatedRequiredGiB": 18.386, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "aquila3", "modelscopeFileSize": 16451477782, "modelscopeLicense": "other", "modelscopeParams": 8223748096, "modelscopeTags": ["license:other", "model_type:aquila3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:arxiv:2406.12182"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16451477782}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:14:41.855366+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970587", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-20T14:34:57.856723+00:00", "modelId": "RedHatAI/Meta-Llama-3.1-8B-Instruct-FP8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9081287016, "estimatedRequiredGiB": 10.159, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9090504850, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9090504850}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:14:41.848764+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970580", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-20T12:34:29.856361+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955963132}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:14:41.846486+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970590", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T20:49:19.683371+00:00", "modelId": "neuralmagic/Qwen2-7B-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8708787584, "estimatedRequiredGiB": 9.746, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 8720442075, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7615616512, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 8720442075}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:14:41.836936+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970582", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:58:28.353096+00:00", "modelId": "RedHatAI/starcoder2-15b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785269848, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788663591, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788663591}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:09:56.550329+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970503", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-20T15:32:19.355892+00:00", "modelId": "neuralmagic/Mistral-Nemo-Instruct-2407-quantized.w4a16", "modelProfile": {"architectures": ["MistralForCausalLM"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8308282608, "estimatedRequiredGiB": 9.319, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral", "modelscopeFileSize": 8338221588, "modelscopeLicense": "llama2", "modelscopeParams": 12247782400, "modelscopeTags": ["license:llama2", "model_type:mistral", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 8338221588}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:09:56.548188+00:00", "targetGpu": "Biren_166m", "taskId": "4970505", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-20T19:47:52.978254+00:00", "modelId": "aisingapore/SEA-LION-v1-7B-IT-GPTQ", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "f7d2aee0bcf1396cd8ab3b1064b6a032eccc04f2d655a336023e13528574903d", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5469228368, "estimatedRequiredGiB": 6.118, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 5473985594, "modelscopeLicense": "mit", "modelscopeParams": 1922048000, "modelscopeTags": ["license:mit", "model_type:mpt", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5473985594}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:09:56.541899+00:00", "targetGpu": "Biren_166m", "taskId": "4970509", "taskType": "text-generation", "verifyResult": -1}
@@ -87,6 +89,7 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:08:06.464776+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-bf16", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2340697867, "estimatedRequiredGiB": 2.621, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 2345554722, "modelscopeLicense": "other", "modelscopeParams": 1170340608, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2345554722}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:37.837081+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970332", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:08:06.464814+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248516007, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248516007}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:37.834872+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970331", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:08:06.464803+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Instruct-MLX-bf16", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2340697867, "estimatedRequiredGiB": 2.622, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 2345693224, "modelscopeLicense": "other", "modelscopeParams": 1170340608, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2345693224}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:37.819101+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970330", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-20T20:49:19.683358+00:00", "modelId": "RedHatAI/Meta-Llama-3-8B-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084013360, "estimatedRequiredGiB": 10.163, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093325064, "modelscopeLicense": "llama3", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9093325064}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:37.806607+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970326", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T20:31:51.861095+00:00", "modelId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-4BIT", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 19935028960, "estimatedRequiredGiB": 22.299, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 19953127685, "modelscopeLicense": "mit", "modelscopeParams": 32762123264, "modelscopeTags": ["license:mit", "model_type:qwen3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 19953127685}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:37.777020+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970324", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-20T14:14:39.678925+00:00", "modelId": "aisingapore/SEA-LION-v1-7B", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "0335c7e4668aaa963303ebe91f80d6eb1e7e8d0299011bd85f7e05a4df031dc7", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361968, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008115847, "modelscopeLicense": "mit", "modelscopeParams": 7501651968, "modelscopeTags": ["license:mit", "model_type:mpt", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008115847}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:46:35.801267+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970119", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-20T17:09:45.266093+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29274355224, "estimatedRequiredGiB": 32.761, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 29314389092, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 29314389092}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:46:35.789758+00:00", "targetGpu": "Biren_166m", "taskId": "4970118", "taskType": "text-generation", "verifyResult": -1}
@@ -127,6 +130,7 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T20:31:51.861054+00:00", "modelId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34327989512, "estimatedRequiredGiB": 38.385, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 34346039951, "modelscopeLicense": "mit", "modelscopeParams": 32762123264, "modelscopeTags": ["license:mit", "model_type:qwen3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 34346039951}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.543315+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969362", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["zaya"], "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T11:30:54.600365+00:00", "modelId": "JANGQ-AI/Osaurus-AppleScript-8B-JANG_4M", "modelProfile": {"architectures": ["ZayaForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7915599888, "estimatedRequiredGiB": 8.885, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "zaya", "modelscopeFileSize": 7950463599, "modelscopeLicense": "other", "modelscopeParams": 2274126328, "modelscopeTags": ["license:other", "model_type:zaya", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:applescript", "custom_tag:macos", "custom_tag:computer-use", "custom_tag:agent", "custom_tag:tool-calling", "custom_tag:function-calling", "custom_tag:mlx", "custom_tag:moe", "custom_tag:osaurus"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7950463599}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.452113+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4969361", "taskType": "text-generation", "verifyResult": -1}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T14:32:15.260751+00:00", "modelId": "dickeyli/llmrec-onereason-8b-sh2-grpo4-step160", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16779959304, "estimatedRequiredGiB": 18.777, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 16801630203, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 8389956608, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16801630203}, "outcome": "success", "status": "success", "submitTime": "2026-09-20T03:09:18.440547+00:00", "targetGpu": "Biren_166m", "taskId": "4969359", "taskType": "text-generation", "verifyResult": 1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:49:19.683330+00:00", "modelId": "apodex/Apodex-1.1-mini-GPTQ-Int4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 24651300904, "estimatedRequiredGiB": 27.587, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 24684544516, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "gptq", "repositoryOnDiskBytes": 24684544516}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:17.835766+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969320", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T18:18:38.666111+00:00", "modelId": "neuralmagic/SmolLM-135M-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 219850592, "estimatedRequiredGiB": 0.249, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 223236742, "modelscopeLicense": "apache-2.0", "modelscopeParams": 162826560, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 223236742}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:11.539809+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969300", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T18:18:38.666090+00:00", "modelId": "LiquidAI/LFM2.5-2.6B", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5394427456, "estimatedRequiredGiB": 6.049, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 5412393192, "modelscopeLicense": "other", "modelscopeParams": 2697198592, "modelscopeTags": ["license:other", "model_type:lfm2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5412393192}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:11.537727+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969303", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-20T19:47:52.978239+00:00", "modelId": "dickeyli/llmrec-onereason-8b-sh2-grpo4-step160", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "8b3c0a82b858cf26b1a0f6075074234a66487d5734741dbbfd4743963b967faa", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16779959304, "estimatedRequiredGiB": 18.777, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 16801630203, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 8389956608, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16801630203}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.295134+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969090", "taskType": "text-generation", "verifyResult": -1}
@@ -134,6 +138,7 @@
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T18:53:01.473075+00:00", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 37184031394, "estimatedRequiredGiB": 41.579, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 37204398600, "modelscopeLicense": "apache-2.0", "modelscopeParams": 10061358960, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:mixed-precision", "custom_tag:qwen3.6", "custom_tag:agentic-search"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 37204398600}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.043508+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969075", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "runtime_memory", "failureAction": "lower_memory_risk_or_reject_combination", "failureCategory": "runtime_memory", "failureClassificationReason": "structured_device_oom", "failureCode": "DEVICE_OOM", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T19:56:29.962443+00:00", "modelId": "prithivMLmods/CEERS-2112-14B-Instruct", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29540133904, "estimatedRequiredGiB": 33.022, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 29547221374, "modelscopeLicense": "apache-2.0", "modelscopeParams": 14770033664, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:Code", "custom_tag:Math", "custom_tag:Reasoning", "custom_tag:text-generation-inference", "custom_tag:Reinforcement-learning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 29547221374}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.939142+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969070", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedArchitectures": ["Cohere2ForCausalLM"], "framework": "vllm", "lastSyncTime": "2026-09-20T17:09:45.266177+00:00", "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730845002, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730845002}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.937617+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4969071", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["deepseek_v4"], "framework": "vllm", "lastSyncTime": "2026-09-20T20:49:19.683316+00:00", "modelId": "JANGQ-AI/DeepSeek-V4-Flash-JANGTQ-K", "modelProfile": {"architectures": ["DeepseekV4ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 85870266317, "estimatedRequiredGiB": 95.98, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "deepseek_v4", "modelscopeFileSize": 85881510108, "modelscopeLicense": "mit", "modelscopeParams": 21758832856, "modelscopeTags": ["license:mit", "model_type:deepseek_v4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:deepseek", "custom_tag:deepseek-v4", "custom_tag:dsv4", "custom_tag:mixture-of-experts", "custom_tag:mla", "custom_tag:mhc", "custom_tag:sparse-indexer", "custom_tag:million-token-context", "custom_tag:reasoning", "custom_tag:tool-use", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:quantized", "custom_tag:jangtq", "custom_tag:jang"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 85881510108}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.889837+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969069", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T16:44:23.353846+00:00", "modelId": "aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 42709318472, "estimatedRequiredGiB": 47.753, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 42728540075, "modelscopeLicense": "llama3.1", "modelscopeParams": 70553707056, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 42728540075}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.842110+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969064", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:16:30.961259+00:00", "modelId": "mlx-community/Qwen3.6-35B-A3B-OptiQ-4bit-REAP-19B", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 14864536516, "estimatedRequiredGiB": 16.635, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 14884986381, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3818458992, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:expert-pruning", "custom_tag:reap", "custom_tag:moe", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 14884986381}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.748956+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969058", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:16:30.961280+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed-FP16", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683239637, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703971805, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:m1", "custom_tag:m2", "custom_tag:fp16", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703971805}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.741993+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969059", "taskType": "text-generation", "verifyResult": -1}
@@ -146,6 +151,7 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-20T15:18:36.848582+00:00", "modelId": "solidrust/Llama-3-11B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7541230432, "estimatedRequiredGiB": 8.438, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 7550622334, "modelscopeLicense": "other", "modelscopeParams": 11520053248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 7550622334}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.540705+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969050", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-20T15:18:36.848613+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-v0.4-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727938576, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737300809, "modelscopeLicense": "other", "modelscopeParams": 8030261248, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737300809}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.537909+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969049", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm", "lastSyncTime": "2026-09-20T17:09:45.266125+00:00", "modelId": "OpenBMB/MiniCPM5-2B-MLX", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1416035216, "estimatedRequiredGiB": 1.594, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 1426233716, "modelscopeLicense": "apache-2.0", "modelscopeParams": 393390080, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1426233716}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.449778+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969045", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:49:19.683231+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26119812080, "estimatedRequiredGiB": 29.233, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26156887523, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26156887523}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.448250+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969047", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:31:51.860975+00:00", "modelId": "mlx-community/Qwen3.5-122B-A10B-OptiQ-2bit-REAP-63B", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 25593158704, "estimatedRequiredGiB": 28.626, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 25613691403, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7626590448, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:expert-pruning", "custom_tag:reap", "custom_tag:moe", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 25613691403}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.436311+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969042", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:31:51.860998+00:00", "modelId": "primitive-ai/Nex-N2.5-mini-mixed-NVFP4-FP8", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26057989872, "estimatedRequiredGiB": 29.158, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26090095180, "modelscopeLicense": "apache-2.0", "modelscopeParams": 21416999792, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:nvfp4", "custom_tag:fp8", "custom_tag:mixed-precision", "custom_tag:quantized", "custom_tag:moe", "custom_tag:nex-n2.5", "custom_tag:agentic", "custom_tag:tool-calling", "custom_tag:single-gpu", "custom_tag:blackwell"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26090095180}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.351734+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969040", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:31:51.861046+00:00", "modelId": "primitive-ai/Nex-N2.5-mini-NVFP4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23926484304, "estimatedRequiredGiB": 26.777, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 23959625409, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35107212656, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:nvfp4", "custom_tag:quantized", "custom_tag:moe", "custom_tag:nex-n2.5", "custom_tag:agentic", "custom_tag:tool-calling", "custom_tag:single-gpu", "custom_tag:blackwell"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 23959625409}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.349521+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969041", "taskType": "text-generation", "verifyResult": -1}
@@ -153,8 +159,12 @@
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T18:18:38.666148+00:00", "modelId": "primitive-ai/Qwen3.8-27B-mixed-NVFP4-FP8", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 22210552448, "estimatedRequiredGiB": 24.849, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 22234091413, "modelscopeLicense": "apache-2.0", "modelscopeParams": 17463440388, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:nvfp4", "custom_tag:fp8", "custom_tag:mixed-precision", "custom_tag:quantized", "custom_tag:speculative-decoding", "custom_tag:blackwell", "custom_tag:a100"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 22234091413}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.336010+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969039", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T18:18:38.666126+00:00", "modelId": "naver-hyperclovax/HyperCLOVAX-SEED-Think-32B", "modelProfile": {"architectures": ["HyperCLOVAXVisionV2ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 66626957488, "estimatedRequiredGiB": 74.478, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "hyperclovax_vision_v2", "modelscopeFileSize": 66642228907, "modelscopeLicense": "other", "modelscopeParams": 33313410304, "modelscopeTags": ["license:other", "model_type:hyperclovax_vision_v2", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 66642228907}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.307162+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969037", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T17:54:13.170257+00:00", "modelId": "logic65/Qwen3.8-Whittle-w50-18.3B-unrepaired", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 36679335352, "estimatedRequiredGiB": 41.028, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 36711610638, "modelscopeLicense": "apache-2.0", "modelscopeParams": 18339618304, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:pruning", "custom_tag:width-pruning", "custom_tag:zero-training", "custom_tag:qwen3_5", "custom_tag:gated-deltanet"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 36711610638}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.297854+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969036", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T20:49:19.683345+00:00", "modelId": "JANGQ-AI/Hy3-preview-JANGTQ", "modelProfile": {"architectures": ["HYV3ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 85146629560, "estimatedRequiredGiB": 95.18, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "hy_v3", "modelscopeFileSize": 85166002853, "modelscopeLicense": "other", "modelscopeParams": 21562957504, "modelscopeTags": ["license:other", "model_type:hy_v3", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:jang", "custom_tag:jangtq", "custom_tag:hy3", "custom_tag:hunyuan", "custom_tag:hy_v3", "custom_tag:moe", "custom_tag:apple-silicon", "custom_tag:2bit", "custom_tag:295b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 85166002853}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.248597+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969031", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T20:31:51.861068+00:00", "modelId": "BSC-LT/ALIA-40b", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 80867821352, "estimatedRequiredGiB": 90.437, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 80921299838, "modelscopeLicense": "apache-2.0", "modelscopeParams": 40433885184, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 80921299838}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.237797+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969032", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T18:53:01.473054+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v3-9B", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 18483466448, "estimatedRequiredGiB": 20.702, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 18523972484, "modelscopeLicense": "gemma", "modelscopeParams": 9241705984, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 18523972484}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:50.898445+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969015", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:49:19.683277+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-bf16", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2340697867, "estimatedRequiredGiB": 2.621, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 2345554722, "modelscopeLicense": "other", "modelscopeParams": 1170340608, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2345554722}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:43.837507+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969003", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["gemma4"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:49:19.683300+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:43.762431+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4968999", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:49:19.683286+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683241105, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703487237, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703487237}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:43.746995+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4968998", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-19T23:13:47.695492+00:00", "modelId": "RedHatAI/SmolLM-135M-Instruct-quantized.w8a16", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-19T23:07:21+00:00", "targetGpu": "Vastai_va16", "taskId": "4477839", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-19T23:13:47.695526+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-NVFP4", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-19T23:07:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4458016", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm", "lastSyncTime": "2026-09-19T23:04:02.598845+00:00", "modelId": "Jackrong/DeepSeek-V4-Pro-Qwen3.5-9B", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-19T23:03:21+00:00", "targetGpu": "Vastai_va16", "taskId": "4609086", "taskType": "text-generation", "verifyResult": -1}
@@ -288,13 +298,3 @@
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm", "lastSyncTime": "2026-09-17T01:19:27.449250+00:00", "modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21930054875, "estimatedRequiredGiB": 24.532, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 21950415614, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6024588416, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21950415614}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T07:13:14.093936+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868226", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-17T01:10:41.230625+00:00", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelProfile": {"architectures": ["Mistral3ForConditionalGeneration"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17367755740, "estimatedRequiredGiB": 19.429, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral3", "modelscopeFileSize": 17385072379, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4929899520, "modelscopeTags": ["license:apache-2.0", "model_type:mistral3", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:mistral", "custom_tag:mistral3", "custom_tag:devstral"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17385072379}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T07:13:14.092802+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868221", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["gemma4"], "framework": "vllm", "lastSyncTime": "2026-09-17T01:10:41.230600+00:00", "modelId": "mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23524076260, "estimatedRequiredGiB": 26.327, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 23556606333, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6649116524, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 23556606333}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T07:13:14.091353+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868229", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "参数/模板问题", "framework": "vllm", "lastSyncTime": "2026-09-17T04:25:23.963394+00:00", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-15T07:13:14.090218+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868222", "taskType": "text-generation", "verifyResult": null}
{"failReason": "参数/模板问题", "framework": "vllm", "lastSyncTime": "2026-09-17T04:25:23.963532+00:00", "modelId": "mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-15T07:13:14.088705+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868224", "taskType": "text-generation", "verifyResult": null}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm", "lastSyncTime": "2026-09-16T20:22:19.008738+00:00", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 37184031394, "estimatedRequiredGiB": 41.579, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 37204398600, "modelscopeLicense": "apache-2.0", "modelscopeParams": 10061358960, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:mixed-precision", "custom_tag:qwen3.6", "custom_tag:agentic-search"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 37204398600}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T07:13:14.086293+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868227", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["nemotron_h"], "framework": "vllm", "lastSyncTime": "2026-09-15T06:27:13.910681+00:00", "modelId": "primitive-ai/Nemotron-3.5-Lightning-30B-A3B-mixed-INT4-INT8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T06:23:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4544266", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm", "lastSyncTime": "2026-09-15T06:27:13.910666+00:00", "modelId": "douyamv/Qwen3.8-27B-FP8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T06:19:21+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4609095", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "platform_infrastructure", "failureAction": "open_short_gpu_framework_circuit_and_retry_other_models", "failureCategory": "platform_infrastructure", "failureClassificationReason": "no_idle_device", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm", "lastSyncTime": "2026-09-15T06:16:21.605532+00:00", "modelId": "mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T06:15:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4490279", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm", "lastSyncTime": "2026-09-15T06:05:16.606079+00:00", "modelId": "cyankiwi/Ornith-1.5-9B-AWQ-FP8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T05:59:21+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4610336", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "", "lastSyncTime": "2026-09-15T05:54:37.536152+00:00", "modelId": "douyamv/Qwen3.8-27B-FP8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T05:49:21+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4332445", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm", "lastSyncTime": "2026-09-15T05:34:43.803047+00:00", "modelId": "EschaLabs/Qwen3.8-27B-Escha-W2", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T05:29:21+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4610338", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-15T05:15:31.609401+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-v0.4-AWQ", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-15T05:09:21+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4594447", "taskType": "text-generation", "verifyResult": -1}

View File

@@ -1818,6 +1818,59 @@
{"batchId": "3c04fe78082844058f61faa6168bf58e", "completedAt": "2026-09-20T20:45:26.450856+00:00", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:39:11.627762+00:00", "framework": "vllm-customized", "intentId": "5942bdc4b3af4592ba269ba17fbbdd5c", "lastModified": "2026-08-24T20:26:05+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Qwen2-7B-Instruct-quantized.w8a8", "reason": null, "reconciledAt": "2026-09-20T20:48:15.520062+00:00", "repoId": "neuralmagic/Qwen2-7B-Instruct-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4987372", "taskType": "text-generation"}
{"batchId": "aa07bc4f2da94ea6a5d431c9d766a356", "completedAt": "2026-09-20T20:48:14.078056+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:48:05.263302+00:00", "framework": "vllm", "intentId": "6c0149290d594fcf90255005b71655e0", "lastModified": "2026-09-15T15:08:08+00:00", "modelAddress": "https://modelscope.cn/models/mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "reason": null, "reconciledAt": "2026-09-20T20:48:15.521971+00:00", "repoId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Kunlunxin_p-800", "taskId": "4987441", "taskType": "text-generation"}
{"batchId": "aa07bc4f2da94ea6a5d431c9d766a356", "completedAt": "2026-09-20T20:48:14.078068+00:00", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:48:05.263468+00:00", "framework": "vllm-customized", "intentId": "77773a2afbb540be8d09c7020aaf3ba8", "lastModified": "2026-09-16T17:02:01+00:00", "modelAddress": "https://modelscope.cn/models/mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "reason": null, "reconciledAt": "2026-09-20T20:48:15.519786+00:00", "repoId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4987442", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.332978+00:00", "framework": "vllm_tokenizer_patch", "intentId": "b51e94235ac64da9afd83b99d549c1e2", "lastModified": "2026-08-24T20:05:18+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "repoId": "RedHatAI/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333079+00:00", "framework": "vllm_tokenizer_patch", "intentId": "eef1ddf74299485f9ab0a1b1fc171a66", "lastModified": "2026-08-24T19:59:13+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-FP8", "repoId": "RedHatAI/starcoder2-7b-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333122+00:00", "framework": "vllm_tokenizer_patch", "intentId": "c95783446abd44358dd9515ed4fbcd22", "lastModified": "2026-08-24T20:02:01+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-2b-quantized.w8a16", "repoId": "RedHatAI/gemma-2-2b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333162+00:00", "framework": "vllm", "intentId": "c2ed754cb92f43008e122ab83d0fdf7d", "lastModified": "2026-09-10T07:52:22+00:00", "modelAddress": "https://modelscope.cn/models/TokenRhythm/NeoHorse-1-4B", "repoId": "TokenRhythm/NeoHorse-1-4B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333226+00:00", "framework": "vllm", "intentId": "bdc3b98ec6834623a95bcc6b2d4b18d3", "lastModified": "2026-08-26T12:24:21+00:00", "modelAddress": "https://modelscope.cn/models/ornith-ai/Ornith-1.5-35B-A3B-NVFP4", "repoId": "ornith-ai/Ornith-1.5-35B-A3B-NVFP4", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333286+00:00", "framework": "vllm", "intentId": "f08b2c1f402143ed9df6062cc9978031", "lastModified": "2026-08-31T14:36:00+00:00", "modelAddress": "https://modelscope.cn/models/ewinregirgojr/Qwen3.8-9B-Instruct-Turbo", "repoId": "ewinregirgojr/Qwen3.8-9B-Instruct-Turbo", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333345+00:00", "framework": "vllm", "intentId": "56dfc7e73234487a95bca82c0defe177", "lastModified": "2026-08-24T20:06:34+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-3b-quantized.w8a8", "repoId": "RedHatAI/starcoder2-3b-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333403+00:00", "framework": "vllm", "intentId": "2c6773370a0d4d4aa585479be930a2ea", "lastModified": "2026-08-28T09:38:42+00:00", "modelAddress": "https://modelscope.cn/models/mlx-community/Qwen3.6-35B-A3B-OptiQ-4bit-REAP-19B", "repoId": "mlx-community/Qwen3.6-35B-A3B-OptiQ-4bit-REAP-19B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333461+00:00", "framework": "vllm", "intentId": "3f9751e903fc46be93e5040a48c0df18", "lastModified": "2026-09-08T15:45:26+00:00", "modelAddress": "https://modelscope.cn/models/JANGQ-AI/AppleScript-8B-JANG_4M", "repoId": "JANGQ-AI/AppleScript-8B-JANG_4M", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333519+00:00", "framework": "vllm", "intentId": "3730047a5ba04ebab23c0884b2800c24", "lastModified": "2026-08-26T20:08:51+00:00", "modelAddress": "https://modelscope.cn/models/VerStella/Qwen3.8-27B-heretic-ara-W8A8-DCU", "repoId": "VerStella/Qwen3.8-27B-heretic-ara-W8A8-DCU", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333577+00:00", "framework": "vllm", "intentId": "153fcdd5fa3d455280f960ddcb19eec9", "lastModified": "2026-08-29T01:07:23+00:00", "modelAddress": "https://modelscope.cn/models/mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B", "repoId": "mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333634+00:00", "framework": "vllm", "intentId": "7ccd49aad11e4e1094fe65d92f7b2af1", "lastModified": "2026-09-03T07:00:50+00:00", "modelAddress": "https://modelscope.cn/models/EschaLabs/Qwen3.8-27B-Escha-W2", "repoId": "EschaLabs/Qwen3.8-27B-Escha-W2", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333702+00:00", "framework": "vllm_fix_tokenizer", "intentId": "caafd10176a44648bec76516ec9377f4", "lastModified": "2026-09-12T12:30:14+00:00", "modelAddress": "https://modelscope.cn/models/OpenBMB/MiniCPM5-2B-MLX", "repoId": "OpenBMB/MiniCPM5-2B-MLX", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333747+00:00", "framework": "vllm_fix_tokenizer", "intentId": "50262f57e582434eb470824bd30cfd8b", "lastModified": "2026-09-04T12:06:20+00:00", "modelAddress": "https://modelscope.cn/models/inclusionAI/Ling-3.0-tiny", "repoId": "inclusionAI/Ling-3.0-tiny", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333791+00:00", "framework": "vllm_fix_tokenizer", "intentId": "8e1592caf2fb4b2ea5f6cb83764385ff", "lastModified": "2026-08-24T20:31:44+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Meta-Llama-3-8B-Instruct-quantized.w8a8", "repoId": "neuralmagic/Meta-Llama-3-8B-Instruct-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333834+00:00", "framework": "vllm_fix_tokenizer", "intentId": "e75bb6d915894f9c95594e0ee077a8f4", "lastModified": "2026-08-24T20:20:16+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-15b-FP8", "repoId": "neuralmagic/starcoder2-15b-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333878+00:00", "framework": "vllm_fix_tokenizer", "intentId": "059f1f134f784985b0795423c3e61e00", "lastModified": "2026-08-24T20:22:00+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-2b-it-quantized.w8a16", "repoId": "neuralmagic/gemma-2-2b-it-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333921+00:00", "framework": "vllm_fix_tokenizer", "intentId": "429949036c894fc79c4f0cba6c3aa980", "lastModified": "2026-08-26T17:18:42+00:00", "modelAddress": "https://modelscope.cn/models/nv-community/NVIDIA-Nemotron-3-Nano-30B-A3B-NVFP4", "repoId": "nv-community/NVIDIA-Nemotron-3-Nano-30B-A3B-NVFP4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.333964+00:00", "framework": "vllm_fix_tokenizer", "intentId": "374ee31681434d5aa5462eeaa53da558", "lastModified": "2026-09-08T03:51:51+00:00", "modelAddress": "https://modelscope.cn/models/XHToken/Spark-X2.5-4B", "repoId": "XHToken/Spark-X2.5-4B", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334012+00:00", "framework": "vllm_fix_tokenizer", "intentId": "1313d0f8a183447e9f5921b0a6dcb97c", "lastModified": "2026-09-10T05:31:36+00:00", "modelAddress": "https://modelscope.cn/models/primitive-ai/Nex-N2.5-mini-NVFP4", "repoId": "primitive-ai/Nex-N2.5-mini-NVFP4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334056+00:00", "framework": "vllm_fix_tokenizer", "intentId": "dd45307edec94a8d951ec6d56e8c4406", "lastModified": "2026-09-12T12:32:06+00:00", "modelAddress": "https://modelscope.cn/models/OpenBMB/MiniCPM5-2B-GPTQ", "repoId": "OpenBMB/MiniCPM5-2B-GPTQ", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334099+00:00", "framework": "vllm_fix_tokenizer", "intentId": "82d1fd36bd2a4d27b6949b0d566fd9e5", "lastModified": "2026-09-10T05:36:14+00:00", "modelAddress": "https://modelscope.cn/models/primitive-ai/Nex-N2.5-mini-mixed-NVFP4-FP8", "repoId": "primitive-ai/Nex-N2.5-mini-mixed-NVFP4-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334142+00:00", "framework": "vllm_fix_tokenizer", "intentId": "8eda320191e5405da2bd9b49ff891ff8", "lastModified": "2026-08-31T14:27:40+00:00", "modelAddress": "https://modelscope.cn/models/sbintuitions/sarashina2.2-3b-instruct-v0.1", "repoId": "sbintuitions/sarashina2.2-3b-instruct-v0.1", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334185+00:00", "framework": "vllm_fix_tokenizer", "intentId": "e79bc9934bfe41a492947f155395122a", "lastModified": "2026-08-25T12:06:29+00:00", "modelAddress": "https://modelscope.cn/models/LiquidAI/LFM2.5-8B-A1B", "repoId": "LiquidAI/LFM2.5-8B-A1B", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334228+00:00", "framework": "vllm_fix_tokenizer", "intentId": "63bdda7200a642d4b40d478ae4666a76", "lastModified": "2026-09-03T13:37:32+00:00", "modelAddress": "https://modelscope.cn/models/XHToken/Spark-X2.5-1.7B", "repoId": "XHToken/Spark-X2.5-1.7B", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334270+00:00", "framework": "vllm_fix_tokenizer", "intentId": "e2fed4b7988e47d19bf14176d5c06c5e", "lastModified": "2026-08-26T16:41:11+00:00", "modelAddress": "https://modelscope.cn/models/apodex/Apodex-1.1-mini-GPTQ-Int4", "repoId": "apodex/Apodex-1.1-mini-GPTQ-Int4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334313+00:00", "framework": "vllm_fix_tokenizer", "intentId": "33ea05829df046478757a21736991b13", "lastModified": "2026-08-26T12:23:32+00:00", "modelAddress": "https://modelscope.cn/models/LiquidAI/LFM2.5-1.2B-Thinking-MLX-4bit", "repoId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-4bit", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334355+00:00", "framework": "vllm_fix_tokenizer", "intentId": "dd36e13321fc4d549f6d5067bf102ef4", "lastModified": "2026-08-26T17:23:58+00:00", "modelAddress": "https://modelscope.cn/models/LiquidAI/LFM2.5-1.2B-Thinking-MLX-8bit", "repoId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-8bit", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334399+00:00", "framework": "vllm_fix_tokenizer", "intentId": "f71f5427dc134649aa2b71dcbc81c3bc", "lastModified": "2026-08-26T20:28:29+00:00", "modelAddress": "https://modelscope.cn/models/LiquidAI/LFM2.5-1.2B-Thinking-MLX-5bit", "repoId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-5bit", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334442+00:00", "framework": "vllm_fix_tokenizer", "intentId": "afa8878ec2264856adb5b19d1ff166f4", "lastModified": "2026-08-24T20:33:29+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-3b-FP8", "repoId": "RedHatAI/starcoder2-3b-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334485+00:00", "framework": "vllm_fix_tokenizer", "intentId": "e508c171d09243698e51a0dac9977260", "lastModified": "2026-08-24T20:38:11+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-15b-quantized.w8a16", "repoId": "RedHatAI/starcoder2-15b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334528+00:00", "framework": "vllm_fix_tokenizer", "intentId": "f1d7d017e65d440c86acb17276f118bb", "lastModified": "2026-08-24T20:26:53+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-9b-it-quantized.w8a16", "repoId": "RedHatAI/gemma-2-9b-it-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334571+00:00", "framework": "vllm_fix_tokenizer", "intentId": "16de8f05e24846e8a728cfa0d5c20ccb", "lastModified": "2026-09-01T03:45:04+00:00", "modelAddress": "https://modelscope.cn/models/primitive-ai/Qwen3.8-27B-mixed-NVFP4-FP8", "repoId": "primitive-ai/Qwen3.8-27B-mixed-NVFP4-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334613+00:00", "framework": "vllm_fix_tokenizer", "intentId": "4c7fa18679224b1d93ab9c3e2eba3322", "lastModified": "2026-09-04T12:06:11+00:00", "modelAddress": "https://modelscope.cn/models/sapientinc/HRM-Text-1B", "repoId": "sapientinc/HRM-Text-1B", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334655+00:00", "framework": "vllm_fix_tokenizer", "intentId": "b3ffb8db71344b44a1ce31d67af57ebc", "lastModified": "2026-09-03T18:32:29+00:00", "modelAddress": "https://modelscope.cn/models/cyankiwi/Apodex-1.1-mini-AWQ-INT4", "repoId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334702+00:00", "framework": "vllm_fix_tokenizer", "intentId": "82a0aa54bf464381970b69300c246bf4", "lastModified": "2026-09-10T07:49:50+00:00", "modelAddress": "https://modelscope.cn/models/TokenRhythm/NeoHorse-1-9B", "repoId": "TokenRhythm/NeoHorse-1-9B", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334745+00:00", "framework": "vllm_fix_tokenizer", "intentId": "692a9882f37d4d139b329937848af53c", "lastModified": "2026-08-31T23:02:41+00:00", "modelAddress": "https://modelscope.cn/models/ewinregirgojr/Qwen3.8-14B-Instruct-Turbo", "repoId": "ewinregirgojr/Qwen3.8-14B-Instruct-Turbo", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334787+00:00", "framework": "vllm_fix_tokenizer", "intentId": "95577269ac114610838e2885c85fa522", "lastModified": "2026-09-04T07:05:43+00:00", "modelAddress": "https://modelscope.cn/models/XHToken/Spark-X2.5-4B-FP8", "repoId": "XHToken/Spark-X2.5-4B-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334830+00:00", "framework": "vllm_fix_tokenizer", "intentId": "003de52c46134c0089dce26c46f0ade1", "lastModified": "2026-09-09T06:31:08+00:00", "modelAddress": "https://modelscope.cn/models/ddalcu/Qwen3.8-27B-MLX-Serve-6bit", "repoId": "ddalcu/Qwen3.8-27B-MLX-Serve-6bit", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334873+00:00", "framework": "vllm_fix_tokenizer", "intentId": "ab6b35f5201a4b25befaac2165f8da16", "lastModified": "2026-09-02T05:56:49+00:00", "modelAddress": "https://modelscope.cn/models/Mia-AiLab/Qwen3.8-27B-EXL3-3.5bpw", "repoId": "Mia-AiLab/Qwen3.8-27B-EXL3-3.5bpw", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334916+00:00", "framework": "vllm_fix_tokenizer", "intentId": "e324c6188c114fc890d4f10cbf4e516d", "lastModified": "2026-08-26T13:13:18+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/Gemma-SEA-LION-v4-27B-IT-FP8-Dynamic", "repoId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-FP8-Dynamic", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.334959+00:00", "framework": "vllm_fix_tokenizer", "intentId": "8a6b42dad9f3462b8dfff8fd98965477", "lastModified": "2026-09-03T06:52:45+00:00", "modelAddress": "https://modelscope.cn/models/bharatgenai/Param2-17B-A2.4B-Thinking", "repoId": "bharatgenai/Param2-17B-A2.4B-Thinking", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335002+00:00", "framework": "vllm_fix_tokenizer", "intentId": "df1dc8a8c105418999078015a93c13b2", "lastModified": "2026-08-26T16:41:36+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "repoId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335045+00:00", "framework": "vllm_fix_tokenizer", "intentId": "2f6c35bf034b49d4b52dbc2bd0b09325", "lastModified": "2026-09-08T15:40:32+00:00", "modelAddress": "https://modelscope.cn/models/JANGQ-AI/Ling-2.6-flash-JANGTQ", "repoId": "JANGQ-AI/Ling-2.6-flash-JANGTQ", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335089+00:00", "framework": "vllm_fix_tokenizer", "intentId": "ea834e19ea5d42718bd727d6b3994bd3", "lastModified": "2026-08-26T15:57:06+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "repoId": "aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335132+00:00", "framework": "vllm_fix_tokenizer", "intentId": "1bf64249c193429fa6c618203ef3ea6a", "lastModified": "2026-09-09T06:27:24+00:00", "modelAddress": "https://modelscope.cn/models/ddalcu/Muse-Glimmer-30B-MLX-Serve-4bit", "repoId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-4bit", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "d8df464812ff2fd3c46dc89b0fc6a96370e694fab8e52179f1ada2cdd0ea92b9", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335174+00:00", "framework": "vllm_fix_tokenizer", "intentId": "f3e88b274b8e4b27a0fb882d9d0fd397", "lastModified": "2026-08-26T16:48:25+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/SEA-LION-v1-7B-IT", "repoId": "aisingapore/SEA-LION-v1-7B-IT", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335217+00:00", "framework": "vllm_fix_tokenizer", "intentId": "e590c1ee99d54e9cbecff58230239fc9", "lastModified": "2026-08-28T05:30:04+00:00", "modelAddress": "https://modelscope.cn/models/whcl412/mlx-LycheeAI-coder-1.7b", "repoId": "whcl412/mlx-LycheeAI-coder-1.7b", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "d8df464812ff2fd3c46dc89b0fc6a96370e694fab8e52179f1ada2cdd0ea92b9", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335260+00:00", "framework": "vllm_fix_tokenizer", "intentId": "efe57bd2c02b44a2b28bb3e2d61884d7", "lastModified": "2026-08-26T17:29:18+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/SEA-LION-v1-7B", "repoId": "aisingapore/SEA-LION-v1-7B", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335312+00:00", "framework": "vllm_fix_tokenizer", "intentId": "8b6733fc38c84da89045d9927926318b", "lastModified": "2026-08-26T19:32:38+00:00", "modelAddress": "https://modelscope.cn/models/pfnet/plamo-3-nict-8b-base", "repoId": "pfnet/plamo-3-nict-8b-base", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335356+00:00", "framework": "vllm_fix_tokenizer", "intentId": "8288db30140647e997cb24187a8b5c9d", "lastModified": "2026-08-26T19:43:40+00:00", "modelAddress": "https://modelscope.cn/models/inceptionai/Jais-2-8B-Chat", "repoId": "inceptionai/Jais-2-8B-Chat", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335400+00:00", "framework": "vllm_fix_tokenizer", "intentId": "2c3e1ed730694f4ab45bb21365e0de1f", "lastModified": "2026-08-26T12:21:34+00:00", "modelAddress": "https://modelscope.cn/models/LiquidAI/LFM2.5-1.2B-Instruct-MLX-bf16", "repoId": "LiquidAI/LFM2.5-1.2B-Instruct-MLX-bf16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "1ee6be24c46b46d0bece26f38a907139", "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:50:51.335442+00:00", "framework": "vllm", "intentId": "a2d6f65f294b4ab8b710602d1d5e2b30", "lastModified": "2026-08-24T20:29:58+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-quantized.w8a16", "repoId": "RedHatAI/starcoder2-7b-quantized.w8a16", "safeConfigVector": {"dtype": "-", "gpuMemoryUtilization": 0.9, "gpuNum": 1, "maxModelLen": 4096, "tensorParallel": 1}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
{"batchId": "4bbac00140124bb88bb1375aa2a426b7", "completedAt": "2026-09-20T20:47:39.262351+00:00", "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:45:27.369105+00:00", "framework": "vllm_fix_tokenizer", "intentId": "ea3c83bd5ecc4b41bcccc5e6281dee46", "repoId": "mlx-community/Spark-X2.5-4B-OptiQ-4bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Sunrise_pt-200-x1", "taskId": null, "taskType": "text-generation"}
{"batchId": "4bbac00140124bb88bb1375aa2a426b7", "completedAt": "2026-09-20T20:47:39.262348+00:00", "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:45:27.369050+00:00", "framework": "vllm_fix_tokenizer", "intentId": "6ea0f4aa546b47048561649ccae0af32", "repoId": "aisingapore/Llama-SEA-LION-v3-8B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Sunrise_pt-200-x1", "taskId": null, "taskType": "text-generation"}
{"batchId": "4bbac00140124bb88bb1375aa2a426b7", "completedAt": "2026-09-20T20:47:39.262345+00:00", "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configSource": "modelhub_live", "createdAt": "2026-09-20T20:45:27.368993+00:00", "framework": "vllm_fix_tokenizer", "intentId": "206ae750dc974d0e9032ac1fb0aac1fb", "repoId": "primitive-ai/Nex-N2.5-mini-FP8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Sunrise_pt-200-x1", "taskId": null, "taskType": "text-generation"}