state: generation 13304 (intent)
This commit is contained in:
@@ -3044,7 +3044,7 @@
|
||||
"taskType": "text-generation"
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-23T02:39:33.342313+00:00",
|
||||
"generatedAt": "2026-09-23T02:43:28.891356+00:00",
|
||||
"summary": {
|
||||
"activeBlockCount": 153,
|
||||
"byGpuFramework": {
|
||||
|
||||
@@ -396,7 +396,7 @@
|
||||
}
|
||||
},
|
||||
"frameworkUpdatedAt": null,
|
||||
"generatedAt": "2026-09-23T02:42:24.271501+00:00",
|
||||
"generatedAt": "2026-09-23T02:43:41.289571+00:00",
|
||||
"gpuStats": {
|
||||
"Ascend_910-b4": {
|
||||
"available": true,
|
||||
|
||||
@@ -1,5 +1,5 @@
|
||||
{
|
||||
"catalogUpdatedAt": "2026-09-23T02:42:24.271501+00:00",
|
||||
"catalogUpdatedAt": "2026-09-23T02:43:41.289571+00:00",
|
||||
"configuredTaskTypes": [
|
||||
"text-generation"
|
||||
],
|
||||
@@ -56,7 +56,7 @@
|
||||
"time-series-forecasting"
|
||||
],
|
||||
"errors": [],
|
||||
"generatedAt": "2026-09-23T02:42:24.271501+00:00",
|
||||
"generatedAt": "2026-09-23T02:43:49.953286+00:00",
|
||||
"gpuCatalog": {
|
||||
"Ascend_910-b3": {
|
||||
"canVerify": true,
|
||||
@@ -1761,6 +1761,27 @@
|
||||
],
|
||||
"updatedAt": "2026-09-22T12:27:29.290285+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-FP8|2026-09-18T09:46:17+00:00|Kunlunxin_p-800": {
|
||||
"taskTypes": [
|
||||
"text-generation"
|
||||
],
|
||||
"updatedAt": "2026-09-23T02:43:49.885260+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-FP8|2026-09-18T09:46:17+00:00|MetaX_c-500": {
|
||||
"taskTypes": [
|
||||
"feature_emb",
|
||||
"text-generation",
|
||||
"visual-multi-modal"
|
||||
],
|
||||
"updatedAt": "2026-09-23T02:43:49.953286+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-FP8|2026-09-18T09:46:17+00:00|Mthreads_s4000": {
|
||||
"taskTypes": [
|
||||
"text-generation",
|
||||
"visual-multi-modal"
|
||||
],
|
||||
"updatedAt": "2026-09-23T02:43:49.848061+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-GGUF|2026-09-21T03:33:10+00:00|Ascend_910-b4": {
|
||||
"taskTypes": [
|
||||
"text-generation",
|
||||
@@ -2342,19 +2363,6 @@
|
||||
],
|
||||
"updatedAt": "2026-09-22T16:43:42.121158+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/chenqi2026/7a-8elite|2026-09-22T00:58:21+00:00|Iluvatar_bi-150": {
|
||||
"taskTypes": [
|
||||
"asr",
|
||||
"question_answering",
|
||||
"reinforcement_learning",
|
||||
"text-generation",
|
||||
"text-to-image-generation",
|
||||
"text_classification",
|
||||
"vision_classification",
|
||||
"visual-multi-modal"
|
||||
],
|
||||
"updatedAt": "2026-09-22T01:11:09.614434+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/chenqi2026/7a-8elite|2026-09-22T00:58:21+00:00|Vastai_va16": {
|
||||
"taskTypes": [
|
||||
"text-generation"
|
||||
@@ -4397,28 +4405,12 @@
|
||||
],
|
||||
"updatedAt": "2026-09-22T03:14:36.209893+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-32-v2_D|2026-09-21T18:34:41+00:00|hygon_k100-ai": {
|
||||
"taskTypes": [
|
||||
"text-generation",
|
||||
"text-to-image-generation",
|
||||
"visual-multi-modal"
|
||||
],
|
||||
"updatedAt": "2026-09-22T01:54:36.850717+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-8-v2_D|2026-09-21T18:33:37+00:00|Vastai_va16": {
|
||||
"taskTypes": [
|
||||
"text-generation"
|
||||
],
|
||||
"updatedAt": "2026-09-22T05:08:38.677526+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-8-v2_D|2026-09-21T18:33:37+00:00|hygon_k100-ai": {
|
||||
"taskTypes": [
|
||||
"text-generation",
|
||||
"text-to-image-generation",
|
||||
"visual-multi-modal"
|
||||
],
|
||||
"updatedAt": "2026-09-22T01:54:36.949758+00:00"
|
||||
},
|
||||
"https://modelscope.cn/models/nkkbr/Mini-K3-1H-mamba2-n64-g4-v1_B|2026-09-21T18:24:35+00:00|hygon_k100-ai": {
|
||||
"taskTypes": [
|
||||
"text-generation",
|
||||
@@ -6693,6 +6685,6 @@
|
||||
"updateTime": "2025-12-22 08:59:53"
|
||||
}
|
||||
],
|
||||
"taskTreeUpdatedAt": "2026-09-23T02:42:24.271501+00:00",
|
||||
"taskTreeUpdatedAt": "2026-09-23T02:43:41.289571+00:00",
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -1,6 +1,6 @@
|
||||
{
|
||||
"generatedAt": "2026-09-23T02:39:33.263628+00:00",
|
||||
"lastSyncTime": "2026-09-23T02:39:29.759281+00:00",
|
||||
"generatedAt": "2026-09-23T02:43:28.816899+00:00",
|
||||
"lastSyncTime": "2026-09-23T02:43:25.562726+00:00",
|
||||
"recentLimit": 300,
|
||||
"report": {
|
||||
"architectureCompatibilityBlocks": {
|
||||
@@ -3700,14 +3700,14 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 12,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 35,
|
||||
"ambiguous_runtime": 36,
|
||||
"framework_architecture_unsupported": 8,
|
||||
"model_load": 3,
|
||||
"platform_infrastructure": 2,
|
||||
"tokenizer_compatibility": 1,
|
||||
"参数/模板问题": 5
|
||||
},
|
||||
"failureCount": 54,
|
||||
"failureCount": 55,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm-customized",
|
||||
"pendingCount": 0,
|
||||
@@ -3717,8 +3717,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Cambricon_mlu-370-x8",
|
||||
"taskType": "text-generation",
|
||||
"total": 54,
|
||||
"unresolvedFailureCount": 40
|
||||
"total": 55,
|
||||
"unresolvedFailureCount": 41
|
||||
},
|
||||
"Cambricon_mlu-370-x8|vllm-mlu|text-generation": {
|
||||
"attributableFailureCount": 14,
|
||||
@@ -5747,22 +5747,22 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 12,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 35,
|
||||
"ambiguous_runtime": 36,
|
||||
"framework_architecture_unsupported": 8,
|
||||
"model_load": 3,
|
||||
"platform_infrastructure": 2,
|
||||
"tokenizer_compatibility": 1,
|
||||
"参数/模板问题": 5
|
||||
},
|
||||
"failureCount": 54,
|
||||
"failureCount": 55,
|
||||
"failureRate": 1.0,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 2,
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"total": 54,
|
||||
"unresolvedFailureCount": 40
|
||||
"total": 55,
|
||||
"unresolvedFailureCount": 41
|
||||
},
|
||||
"vllm-mlu": {
|
||||
"attributableFailureCount": 25,
|
||||
@@ -5888,7 +5888,7 @@
|
||||
"unresolvedFailureCount": 175
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-23T02:39:33.251720+00:00",
|
||||
"generatedAt": "2026-09-23T02:43:28.804947+00:00",
|
||||
"gpuSummaries": {
|
||||
"Ascend_910-b3": {
|
||||
"attributableFailureCount": 102,
|
||||
@@ -6009,7 +6009,7 @@
|
||||
"decisionSuccessRate": 0.1892,
|
||||
"decisionTotal": 74,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 142,
|
||||
"ambiguous_runtime": 143,
|
||||
"framework_architecture_unsupported": 48,
|
||||
"memory_capacity": 4,
|
||||
"model_load": 5,
|
||||
@@ -6018,15 +6018,15 @@
|
||||
"参数/模板问题": 41,
|
||||
"验证失败": 22
|
||||
},
|
||||
"failureCount": 267,
|
||||
"failureRate": 0.9502,
|
||||
"failureCount": 268,
|
||||
"failureRate": 0.9504,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 2,
|
||||
"successCount": 14,
|
||||
"successRate": 0.0498,
|
||||
"total": 281,
|
||||
"unresolvedFailureCount": 205
|
||||
"successRate": 0.0496,
|
||||
"total": 282,
|
||||
"unresolvedFailureCount": 206
|
||||
},
|
||||
"Iluvatar_bi-100": {
|
||||
"attributableFailureCount": 194,
|
||||
@@ -11321,6 +11321,29 @@
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Cambricon_mlu-370-x8|vllm-customized|text-generation|qwen3_vl|none": {
|
||||
"attributableFailureCount": 0,
|
||||
"decisionFailureRate": 0.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm-customized",
|
||||
"modelType": "qwen3_vl",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "none",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Cambricon_mlu-370-x8",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Cambricon_mlu-370-x8|vllm-customized|text-generation|sophia_hybrid|none": {
|
||||
"attributableFailureCount": 0,
|
||||
"decisionFailureRate": 0.0,
|
||||
@@ -20888,9 +20911,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 13
|
||||
"ambiguous_runtime": 12
|
||||
},
|
||||
"failureCount": 13,
|
||||
"failureCount": 12,
|
||||
"failureRate": 1.0,
|
||||
"framework": "transformers",
|
||||
"lastPlatformFailureAt": null,
|
||||
@@ -20902,8 +20925,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Iluvatar_bi-100",
|
||||
"taskType": "text-generation",
|
||||
"total": 13,
|
||||
"unresolvedFailureCount": 13
|
||||
"total": 12,
|
||||
"unresolvedFailureCount": 12
|
||||
},
|
||||
"Iluvatar_bi-150|transformers|text-generation": {
|
||||
"attributableFailureCount": 5,
|
||||
@@ -22798,6 +22821,31 @@
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 2
|
||||
},
|
||||
"Cambricon_mlu-370-x8|vllm-customized|text-generation|qwen3_vl|none": {
|
||||
"attributableFailureCount": 0,
|
||||
"consecutiveFailures": 0,
|
||||
"decisionFailureRate": 0.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm-customized",
|
||||
"lastTerminalAt": "2026-09-23T02:43:25.562726+00:00",
|
||||
"modelType": "qwen3_vl",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "none",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Cambricon_mlu-370-x8",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Cambricon_mlu-370-x8|vllm-customized|text-generation|sophia_hybrid|none": {
|
||||
"attributableFailureCount": 0,
|
||||
"consecutiveFailures": 0,
|
||||
@@ -22980,9 +23028,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 5
|
||||
"ambiguous_runtime": 4
|
||||
},
|
||||
"failureCount": 5,
|
||||
"failureCount": 4,
|
||||
"failureRate": 1.0,
|
||||
"framework": "transformers",
|
||||
"lastTerminalAt": "2026-09-21T05:42:03.443548+00:00",
|
||||
@@ -22995,8 +23043,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Iluvatar_bi-100",
|
||||
"taskType": "text-generation",
|
||||
"total": 5,
|
||||
"unresolvedFailureCount": 5
|
||||
"total": 4,
|
||||
"unresolvedFailureCount": 4
|
||||
},
|
||||
"Iluvatar_bi-100|transformers|text-generation|lfm2|none": {
|
||||
"attributableFailureCount": 0,
|
||||
@@ -32080,6 +32128,30 @@
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Cambricon_mlu-370-x8|vllm-customized|text-generation|qwen3_vl|none|34": {
|
||||
"attributableFailureCount": 0,
|
||||
"decisionFailureRate": 0.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm-customized",
|
||||
"loadSizeLog2Bucket": 34,
|
||||
"modelType": "qwen3_vl",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "none",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Cambricon_mlu-370-x8",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Cambricon_mlu-370-x8|vllm-customized|text-generation|sophia_hybrid|none|31": {
|
||||
"attributableFailureCount": 0,
|
||||
"decisionFailureRate": 0.0,
|
||||
@@ -45825,15 +45897,15 @@
|
||||
"unresolvedFailureCount": 0
|
||||
}
|
||||
},
|
||||
"terminalRecords": 16798,
|
||||
"totalRecords": 17016,
|
||||
"terminalRecords": 16799,
|
||||
"totalRecords": 17017,
|
||||
"totals": {
|
||||
"attributableFailureCount": 6013,
|
||||
"decisionFailureRate": 0.8622,
|
||||
"decisionSuccessRate": 0.1378,
|
||||
"decisionTotal": 6974,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 4285,
|
||||
"ambiguous_runtime": 4286,
|
||||
"architecture_compatibility": 212,
|
||||
"attention_backend": 3,
|
||||
"backend_operator": 115,
|
||||
@@ -45849,15 +45921,15 @@
|
||||
"日志缺失": 719,
|
||||
"验证失败": 676
|
||||
},
|
||||
"failureCount": 15837,
|
||||
"failureCount": 15838,
|
||||
"failureRate": 0.9428,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 931,
|
||||
"successCount": 961,
|
||||
"successRate": 0.0572,
|
||||
"total": 16798,
|
||||
"unresolvedFailureCount": 8893
|
||||
"total": 16799,
|
||||
"unresolvedFailureCount": 8894
|
||||
},
|
||||
"warnings": [
|
||||
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
@@ -45870,8 +45942,8 @@
|
||||
"GPU hygon_k100-ai 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Iluvatar_bi-150 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Iluvatar_bi-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
@@ -45896,7 +45968,6 @@
|
||||
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
@@ -45910,10 +45981,11 @@
|
||||
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
||||
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
||||
]
|
||||
},
|
||||
"storageMode": "decision_state_only",
|
||||
"summarizedRecords": 17016,
|
||||
"summarizedRecords": 17017,
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -88,6 +88,7 @@
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T20:22:53.859210+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 4}, "failureCount": 4, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-20T12:46:14.462925+00:00", "modelType": "lfm2", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Iluvatar_bi-150", "taskType": "text-generation", "total": 4, "unresolvedFailureCount": 4}, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.438280+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000549", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859255+00:00", "modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21930054875, "estimatedRequiredGiB": 24.532, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 21950415614, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6024588416, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21950415614}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.436494+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000547", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859295+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 1, "consecutiveFailures": 1, "decisionFailureRate": 1.0, "decisionSuccessRate": 0.0, "decisionTotal": 1, "failureBreakdown": {"model_load": 1}, "failureCount": 1, "failureRate": 1.0, "framework": "vllm_tokenizer_patch", "lastTerminalAt": "2026-09-21T04:36:00.569307+00:00", "modelType": "lfm2", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Iluvatar_bi-150", "taskType": "text-generation", "total": 1, "unresolvedFailureCount": 0}, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.435532+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000550", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-23T02:43:25.562726+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.405767+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000548", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "transformers", "lastSyncTime": "2026-09-21T20:22:53.859247+00:00", "modelId": "mlx-community/gemma-4-12B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9003966484, "estimatedRequiredGiB": 10.099, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 9036402901, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2463563824, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9036402901}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.400088+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000555", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-22T20:46:09.364327+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:54:31.274098+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000347", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-22T18:31:46.771079+00:00", "modelId": "RedHatAI/starcoder2-7b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7857298896, "estimatedRequiredGiB": 8.785, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 7860673720, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 7400416256, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7860673720}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:48:23.283753+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000251", "taskType": "text-generation", "verifyResult": -1}
|
||||
@@ -297,4 +298,3 @@
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T05:14:26.178403+00:00", "modelId": "Arain119/Sophia", "modelProfile": {"architectures": ["SophiaForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2226652224, "estimatedRequiredGiB": 9.768, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "sophia_hybrid", "modelscopeFileSize": 8740644030, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 1113293776, "modelscopeTags": ["license:Apache License 2.0", "model_type:sophia_hybrid", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8740644030}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:24.171826+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986869", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:42:03.443575+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed-FP16", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293473597, "estimatedRequiredGiB": 18.233, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16314185171, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:m1", "custom_tag:m2", "custom_tag:fp16", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:chat", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b", "custom_tag:coding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16314185171}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:24.169691+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986868", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457611+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed-FP16", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683239637, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703971805, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:m1", "custom_tag:m2", "custom_tag:fp16", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703971805}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.997557+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986862", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457647+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.994993+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986861", "taskType": "text-generation", "verifyResult": -1}
|
||||
|
||||
@@ -2458,6 +2458,7 @@
|
||||
{"batchId": "b9814996278e427981ad061a19fc68ec", "completedAt": "2026-09-22T19:09:59.529473+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-22T19:09:58.386064+00:00", "framework": "vllm", "intentId": "18422aa9a8834f06836946ed45997424", "lastModified": "2026-09-22T18:53:47+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-16-v2", "reason": null, "reconciledAt": "2026-09-23T02:39:48.693480+00:00", "repoId": "nkkbr/Mini-K3-1H-kda-kernel-16-v2", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Ascend_910-b4", "taskId": "5023594", "taskType": "text-generation"}
|
||||
{"batchId": "b9814996278e427981ad061a19fc68ec", "completedAt": "2026-09-22T19:09:59.529490+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-22T19:09:58.386188+00:00", "framework": "vllm", "intentId": "3d436f1c2bb84bf2a9dc7a8b998e2b65", "lastModified": "2026-09-22T18:42:06+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-2-v2", "reason": null, "reconciledAt": "2026-09-23T02:39:48.691925+00:00", "repoId": "nkkbr/Mini-K3-1H-kda-kernel-2-v2", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Ascend_910-b4", "taskId": "5023595", "taskType": "text-generation"}
|
||||
{"batchId": "b9814996278e427981ad061a19fc68ec", "completedAt": "2026-09-22T19:09:59.529495+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-22T19:09:58.386250+00:00", "framework": "vllm", "intentId": "cf52124b45d34e5eb09db6f16c17a8d7", "lastModified": "2026-09-22T18:40:12+00:00", "modelAddress": "https://modelscope.cn/models/ProCreations/BetterWright-K2-Horizon-7B-Uno", "reason": null, "reconciledAt": "2026-09-23T02:39:48.692689+00:00", "repoId": "ProCreations/BetterWright-K2-Horizon-7B-Uno", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Ascend_910-b4", "taskId": "5023596", "taskType": "text-generation"}
|
||||
{"batchId": "95bc5a65d4aa4bac9c19706aa528d821", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-23T02:44:28.518459+00:00", "framework": "vllm_fix_tokenizer", "intentId": "374addc9531642c6841513830e93c1e7", "lastModified": "2026-09-18T09:46:17+00:00", "modelAddress": "https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-FP8", "repoId": "XingChen-AGI/Xing4.0-29B-A4B-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "7d0e472f09fd4e38b39dc8933c7ad231", "completedAt": "2026-09-22T13:00:50.920772+00:00", "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configSource": "modelhub_live", "createdAt": "2026-09-22T13:00:20.969224+00:00", "framework": "vllm-mlu", "intentId": "c81e5ad5eee94c7fbfbb4bd7d31b62a7", "repoId": "empero-ai/Qwen3.8-9B-Distill", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "age_policy_deferred", "targetGpu": "Cambricon_mlu-370-x4", "taskId": null, "taskType": "text-generation"}
|
||||
{"batchId": "afe52d157d494615819b937d74c0906c", "completedAt": "2026-09-22T12:31:41.319873+00:00", "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configSource": "modelhub_live", "createdAt": "2026-09-22T12:31:29.032434+00:00", "framework": "vllm-mlu", "intentId": "fe2c75eab7864dca87d4316ef8729c63", "repoId": "empero-ai/Qwen3.8-9B-Distill", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "age_policy_deferred", "targetGpu": "Cambricon_mlu-370-x4", "taskId": null, "taskType": "text-generation"}
|
||||
{"batchId": "2aba49d879674e828cf2c5c72b6897da", "completedAt": "2026-09-22T12:28:38.227933+00:00", "configFingerprint": "415f9524ad0e8e6c9c3d7532026f319100b3d292f6c5a09cdc37467f1144a744", "configSource": "modelhub_live", "createdAt": "2026-09-22T12:27:44.722968+00:00", "framework": "llamacpp", "intentId": "3cc699b425b64bb98e2bad28c8e2317a", "repoId": "LiquidAI/LFM2.5-2.6B-DSpark-GGUF", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Iluvatar_bi-150", "taskId": null, "taskType": "text-generation"}
|
||||
|
||||
@@ -2,24 +2,24 @@
|
||||
"agentVersion": "2026.09.22.1",
|
||||
"checksums": {
|
||||
".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1",
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "a95daeec6c8a78cc817fce3fe1ea7da45ab2c8d024bd7bdc5e75d12ba3a04452",
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "02d159d8ddbf8eeacdc5a2c48b960bfeb4c3f05a764c27cc4de69c2fd77c8244",
|
||||
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
||||
".modelhub_state/market_intelligence.json": "ad4dcb71583f176a838b4ca409bb1094a51f96520feca99053e9e0cead2c2ad7",
|
||||
".modelhub_state/official_capabilities.json": "b129e5b7fb6555914467c832386ccfa69d55b8cbcb90b87140af5b10d6d9ba0f",
|
||||
".modelhub_state/outcome_checkpoint.json": "806c570058bb1af9560a71b1891f93167876293a1d8981600134344976619753",
|
||||
".modelhub_state/market_intelligence.json": "307a858d23a12106244a6e9898f1d81f23459dd6215ff97a14a397539debbe64",
|
||||
".modelhub_state/official_capabilities.json": "902bff43b70f0f076dd088597f4a6d9aa88be8b342a309377f8a985f6e3f378b",
|
||||
".modelhub_state/outcome_checkpoint.json": "55743fdfd5f46db594c372d4535774209ab9afba22b19b91c6e7d64cf2fe9956",
|
||||
".modelhub_state/queue_cleanup_latest.json": "abeedea8ce654c253250bcfd506d9b5eab392b53faf9867483b77b9a5c85233b",
|
||||
".modelhub_state/recent_outcomes.jsonl": "742107c27fa4bac8b0df79e335eed2028fea1719e1ff953459208ee56c339b97",
|
||||
".modelhub_state/recent_outcomes.jsonl": "ef30e009ef3f131ccb91f0e3765f80d2af5c61777924c86d28a680581208505c",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "bac8f0c51a1f1728a6dd549182e7e6976158931a9b7c0eb0ffcf80f88f2a9dff",
|
||||
".modelhub_state/recovery_intents.jsonl": "5bb2f68b8cf9dfc31476561ae31bff166b7b5554a1da0e893ace9ec58162ecd8",
|
||||
".modelhub_state/recovery_intents.jsonl": "275e072b3cc5ceada34bf82da2042914700ccf5eff82c19fcf441cf082cdd581",
|
||||
".modelhub_state/routing_intelligence.json": "39ec7b5e77e11b96887ed828d68b094ab47d4d4e68b51e1ea92f4df7a81ab3e3",
|
||||
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
|
||||
".modelhub_state/worker_crashes.jsonl": "619eac16f716b2420672872ff4ce92d452f5fced4f96e8dbb37f9ce511b53c5b",
|
||||
"ledger/submissions.jsonl": "b97e3ee6089ccc093850c7e2d26ecd715079fefb580a266f9b30ef6c29a32188",
|
||||
"outcomes/submissions.jsonl": "b06470d385200a00d73a98385d752172d09647d6ef320f89cfc109a32e96c638"
|
||||
"outcomes/submissions.jsonl": "adb5d998b9ae00ebd887713726d7124395f261614332d9a1af3dcaad721b294c"
|
||||
},
|
||||
"generation": 13303,
|
||||
"phase": "cycle",
|
||||
"generation": 13304,
|
||||
"phase": "intent",
|
||||
"schemaVersion": 1,
|
||||
"updatedAt": "2026-09-23T02:42:24.915231+00:00",
|
||||
"updatedAt": "2026-09-23T02:44:28.662310+00:00",
|
||||
"writerId": "328f98096427428498653b568f0d5041"
|
||||
}
|
||||
|
||||
@@ -669,7 +669,6 @@
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T19:58:46.051766+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29274355224, "estimatedRequiredGiB": 32.761, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 29314389092, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 29314389092}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:57:55.143712+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000420", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T19:58:46.051786+00:00", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w4a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7686303632, "estimatedRequiredGiB": 8.592, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 7688299781, "modelscopeLicense": "llama2", "modelscopeParams": 13960238080, "modelscopeTags": ["license:llama2", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7688299781}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:57:55.174530+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000425", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-21T20:00:55.754427+00:00", "modelId": "neuralmagic/Llama-2-7b-chat-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7003604832, "estimatedRequiredGiB": 7.829, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 7005626643, "modelscopeLicense": "llama2", "modelscopeParams": 6738415616, "modelscopeTags": ["license:llama2", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7005626643}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:58:50.375505+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5000455", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": "2026-09-21T20:11:13.961614+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:10:32.405767+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000548", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-21T20:50:04.063126+00:00", "modelId": "RedHatAI/SmolLM-1.7B-Instruct-quantized.w8a16", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "b3ef6bcf57e6bb4fc9ede989e67a765b0e40b19fbd407f1f8452833c12bd5c38", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2014810120, "estimatedRequiredGiB": 2.256, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 2018195467, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1812039680, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 2018195467}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.243460+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "5001060", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063136+00:00", "modelId": "RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w4a16", "modelProfile": {"architectures": ["MistralForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8308282608, "estimatedRequiredGiB": 9.319, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral", "modelscopeFileSize": 8338221450, "modelscopeLicense": "llama2", "modelscopeParams": 12247782400, "modelscopeTags": ["license:llama2", "model_type:mistral", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 8338221450}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.251589+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5001057", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063169+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248516007, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248516007}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.247906+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001064", "taskType": "text-generation", "verifyResult": null}
|
||||
|
||||
Reference in New Issue
Block a user