state: generation 5245 (cycle)
This commit is contained in:
@@ -1311,7 +1311,7 @@
|
||||
"taskType": "text-generation"
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-10T15:55:46.473596+00:00",
|
||||
"generatedAt": "2026-09-10T15:58:18.525503+00:00",
|
||||
"summary": {
|
||||
"activeBlockCount": 66,
|
||||
"byGpuFramework": {
|
||||
|
||||
@@ -35,7 +35,7 @@
|
||||
"2": {
|
||||
"complete": false,
|
||||
"lastError": "ModelHubAPIError: 系统错误",
|
||||
"listingErrors": 373,
|
||||
"listingErrors": 374,
|
||||
"nextPage": 1,
|
||||
"recordsScanned": 0,
|
||||
"uniqueRecords": 0
|
||||
@@ -102,12 +102,12 @@
|
||||
"cutoffAt": "2026-09-04T03:55:51.365685+00:00",
|
||||
"failureLogsInspected": 0,
|
||||
"mode": "incremental_decision_only",
|
||||
"nextAccountIndex": 2,
|
||||
"nextAccountIndex": 3,
|
||||
"recordsScanned": 0,
|
||||
"seenTaskIds": [],
|
||||
"startedAt": "2026-09-04T03:55:51.365685+00:00",
|
||||
"terminalRecords": 0,
|
||||
"uniqueRecords": 0,
|
||||
"updatedAt": "2026-09-10T15:55:46.449926+00:00",
|
||||
"updatedAt": "2026-09-10T15:58:18.502283+00:00",
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -416,7 +416,7 @@
|
||||
}
|
||||
},
|
||||
"frameworkUpdatedAt": null,
|
||||
"generatedAt": "2026-09-10T15:54:06.263567+00:00",
|
||||
"generatedAt": "2026-09-10T15:56:55.969367+00:00",
|
||||
"gpuStats": {
|
||||
"Ascend_910-b3": {
|
||||
"available": true,
|
||||
|
||||
File diff suppressed because it is too large
Load Diff
@@ -1,6 +1,6 @@
|
||||
{
|
||||
"generatedAt": "2026-09-10T15:48:18.040272+00:00",
|
||||
"lastSyncTime": "2026-09-10T15:48:17.745945+00:00",
|
||||
"generatedAt": "2026-09-10T15:56:47.793655+00:00",
|
||||
"lastSyncTime": "2026-09-10T15:56:47.546297+00:00",
|
||||
"recentLimit": 300,
|
||||
"report": {
|
||||
"architectureCompatibilityBlocks": {
|
||||
@@ -1905,29 +1905,29 @@
|
||||
"unresolvedFailureCount": 56
|
||||
},
|
||||
"MetaX_c-500|vllm|text-generation": {
|
||||
"attributableFailureCount": 45,
|
||||
"decisionFailureRate": 0.9574,
|
||||
"decisionSuccessRate": 0.0426,
|
||||
"decisionTotal": 47,
|
||||
"attributableFailureCount": 46,
|
||||
"decisionFailureRate": 0.9583,
|
||||
"decisionSuccessRate": 0.0417,
|
||||
"decisionTotal": 48,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 7,
|
||||
"backend_operator": 26,
|
||||
"backend_operator": 27,
|
||||
"framework_architecture_unsupported": 14,
|
||||
"memory_capacity": 1,
|
||||
"model_load": 4,
|
||||
"参数/模板问题": 9
|
||||
},
|
||||
"failureCount": 61,
|
||||
"failureRate": 0.9683,
|
||||
"failureCount": 62,
|
||||
"failureRate": 0.9688,
|
||||
"framework": "vllm",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"successCount": 2,
|
||||
"successRate": 0.0317,
|
||||
"successRate": 0.0312,
|
||||
"targetGpu": "MetaX_c-500",
|
||||
"taskType": "text-generation",
|
||||
"total": 63,
|
||||
"total": 64,
|
||||
"unresolvedFailureCount": 16
|
||||
},
|
||||
"Mthreads_s4000|llamacpp|text-generation": {
|
||||
@@ -2316,13 +2316,13 @@
|
||||
"unresolvedFailureCount": 846
|
||||
},
|
||||
"vllm": {
|
||||
"attributableFailureCount": 239,
|
||||
"attributableFailureCount": 240,
|
||||
"decisionFailureRate": 0.9917,
|
||||
"decisionSuccessRate": 0.0083,
|
||||
"decisionTotal": 241,
|
||||
"decisionTotal": 242,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 174,
|
||||
"backend_operator": 32,
|
||||
"backend_operator": 33,
|
||||
"framework_architecture_unsupported": 169,
|
||||
"memory_capacity": 6,
|
||||
"model_load": 18,
|
||||
@@ -2332,14 +2332,14 @@
|
||||
"tokenizer_compatibility": 3,
|
||||
"参数/模板问题": 24
|
||||
},
|
||||
"failureCount": 438,
|
||||
"failureCount": 439,
|
||||
"failureRate": 0.9955,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 1,
|
||||
"successCount": 2,
|
||||
"successRate": 0.0045,
|
||||
"total": 440,
|
||||
"total": 441,
|
||||
"unresolvedFailureCount": 198
|
||||
},
|
||||
"vllm-mlu": {
|
||||
@@ -2443,7 +2443,7 @@
|
||||
"unresolvedFailureCount": 1
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-10T15:48:18.035267+00:00",
|
||||
"generatedAt": "2026-09-10T15:56:47.789236+00:00",
|
||||
"gpuSummaries": {
|
||||
"Ascend_910-b3": {
|
||||
"attributableFailureCount": 26,
|
||||
@@ -2650,27 +2650,27 @@
|
||||
"unresolvedFailureCount": 89
|
||||
},
|
||||
"MetaX_c-500": {
|
||||
"attributableFailureCount": 45,
|
||||
"decisionFailureRate": 0.8036,
|
||||
"decisionSuccessRate": 0.1964,
|
||||
"decisionTotal": 56,
|
||||
"attributableFailureCount": 46,
|
||||
"decisionFailureRate": 0.807,
|
||||
"decisionSuccessRate": 0.193,
|
||||
"decisionTotal": 57,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 7,
|
||||
"backend_operator": 26,
|
||||
"backend_operator": 27,
|
||||
"framework_architecture_unsupported": 14,
|
||||
"memory_capacity": 1,
|
||||
"model_load": 4,
|
||||
"参数/模板问题": 17,
|
||||
"验证失败": 48
|
||||
},
|
||||
"failureCount": 117,
|
||||
"failureRate": 0.9141,
|
||||
"failureCount": 118,
|
||||
"failureRate": 0.9147,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"successCount": 11,
|
||||
"successRate": 0.0859,
|
||||
"total": 128,
|
||||
"successRate": 0.0853,
|
||||
"total": 129,
|
||||
"unresolvedFailureCount": 72
|
||||
},
|
||||
"Mthreads_s4000": {
|
||||
@@ -3568,15 +3568,15 @@
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"MetaX_c-500|vllm|text-generation|llama|compressed-tensors": {
|
||||
"attributableFailureCount": 13,
|
||||
"decisionFailureRate": 0.9286,
|
||||
"decisionSuccessRate": 0.0714,
|
||||
"decisionTotal": 14,
|
||||
"attributableFailureCount": 14,
|
||||
"decisionFailureRate": 0.9333,
|
||||
"decisionSuccessRate": 0.0667,
|
||||
"decisionTotal": 15,
|
||||
"failureBreakdown": {
|
||||
"backend_operator": 13
|
||||
"backend_operator": 14
|
||||
},
|
||||
"failureCount": 13,
|
||||
"failureRate": 0.9286,
|
||||
"failureCount": 14,
|
||||
"failureRate": 0.9333,
|
||||
"framework": "vllm",
|
||||
"modelType": "llama",
|
||||
"pendingCount": 0,
|
||||
@@ -3584,10 +3584,10 @@
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "compressed-tensors",
|
||||
"successCount": 1,
|
||||
"successRate": 0.0714,
|
||||
"successRate": 0.0667,
|
||||
"targetGpu": "MetaX_c-500",
|
||||
"taskType": "text-generation",
|
||||
"total": 14,
|
||||
"total": 15,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"MetaX_c-500|vllm|text-generation|phi3small|compressed-tensors": {
|
||||
@@ -6545,14 +6545,14 @@
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"MetaX_c-500|vllm|text-generation|llama|compressed-tensors|33": {
|
||||
"attributableFailureCount": 6,
|
||||
"attributableFailureCount": 7,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 6,
|
||||
"decisionTotal": 7,
|
||||
"failureBreakdown": {
|
||||
"backend_operator": 6
|
||||
"backend_operator": 7
|
||||
},
|
||||
"failureCount": 6,
|
||||
"failureCount": 7,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm",
|
||||
"loadSizeLog2Bucket": 33,
|
||||
@@ -6565,7 +6565,7 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "MetaX_c-500",
|
||||
"taskType": "text-generation",
|
||||
"total": 6,
|
||||
"total": 7,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"MetaX_c-500|vllm|text-generation|phi3small|compressed-tensors|33": {
|
||||
@@ -7121,16 +7121,16 @@
|
||||
"unresolvedFailureCount": 0
|
||||
}
|
||||
},
|
||||
"terminalRecords": 1602,
|
||||
"totalRecords": 1689,
|
||||
"terminalRecords": 1603,
|
||||
"totalRecords": 1690,
|
||||
"totals": {
|
||||
"attributableFailureCount": 419,
|
||||
"decisionFailureRate": 0.8896,
|
||||
"decisionSuccessRate": 0.1104,
|
||||
"decisionTotal": 471,
|
||||
"attributableFailureCount": 420,
|
||||
"decisionFailureRate": 0.8898,
|
||||
"decisionSuccessRate": 0.1102,
|
||||
"decisionTotal": 472,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 357,
|
||||
"backend_operator": 39,
|
||||
"backend_operator": 40,
|
||||
"framework_architecture_unsupported": 272,
|
||||
"memory_capacity": 11,
|
||||
"model_load": 48,
|
||||
@@ -7141,30 +7141,31 @@
|
||||
"参数/模板问题": 98,
|
||||
"验证失败": 673
|
||||
},
|
||||
"failureCount": 1550,
|
||||
"failureRate": 0.9675,
|
||||
"failureCount": 1551,
|
||||
"failureRate": 0.9676,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 3,
|
||||
"successCount": 52,
|
||||
"successRate": 0.0325,
|
||||
"total": 1602,
|
||||
"successRate": 0.0324,
|
||||
"total": 1603,
|
||||
"unresolvedFailureCount": 1128
|
||||
},
|
||||
"warnings": [
|
||||
"GPU Ascend_910-b3 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Sunrise_pt-200-x1 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Iluvatar_bi-150 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU hygon_k100-ai 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Iluvatar_bi-150 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"组合 Iluvatar_mrv-100|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Cambricon_mlu-370-x4|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Cambricon_mlu-370-x8|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
@@ -7179,11 +7180,10 @@
|
||||
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Iluvatar_bi-150|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Vastai_va16|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Cambricon_mlu-370-x4|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 hygon_k100-ai|vllm-patch-tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
||||
]
|
||||
},
|
||||
"storageMode": "decision_state_only",
|
||||
"summarizedRecords": 1689,
|
||||
"summarizedRecords": 1690,
|
||||
"version": 1
|
||||
}
|
||||
|
||||
File diff suppressed because it is too large
Load Diff
File diff suppressed because it is too large
Load Diff
@@ -1,24 +1,24 @@
|
||||
{
|
||||
"agentVersion": "2026.09.04.4",
|
||||
"checksums": {
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "5e1315a73de2faad3a72522f87cddef9b6425cafcf52ae24d6bed322696d4d51",
|
||||
".modelhub_state/architecture_history_backfill.json": "01278664c6738d5ea88a60de2dc03463efe354a2ea9446e41d73549aee78d71e",
|
||||
".modelhub_state/market_intelligence.json": "8de1ea1de66870e0fe8bf604e995493c8c549bff722aa6c0a1ac331e33716c33",
|
||||
".modelhub_state/official_capabilities.json": "b2cf0bd003ede2d90d3134f76e8edae3b4197c98e4371b52f6e0e16064a8d61f",
|
||||
".modelhub_state/outcome_checkpoint.json": "ca8829a87164af09127727f6dfb8984cd62c7f3aa4237736893f7009a6070187",
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "c30a623d586f52fb36ed15771c33ef2dfd5a156917cf6da00ccf0facb2845d86",
|
||||
".modelhub_state/architecture_history_backfill.json": "f1349579110bf8c7be96fee1d886639a1c47d5448280aa9d663681334fadbddb",
|
||||
".modelhub_state/market_intelligence.json": "4039e2de80c7e7f6aff356ffe384f2c4eb3734e672532146e54cfbdc37833e71",
|
||||
".modelhub_state/official_capabilities.json": "b2082a0356253dfd28e249ca852f673771374dd6faf61f6e1bcbcf11cd8b4f39",
|
||||
".modelhub_state/outcome_checkpoint.json": "8d4aa81d52779bcf4ff4803a38022c1ae5b9e4833596ac9333126a5d8ed0c1c4",
|
||||
".modelhub_state/queue_cleanup_latest.json": "c4c0c2f745124fbabf68b2aa082c5aeee0f45dbb41d40369e8064e674a33adbe",
|
||||
".modelhub_state/recent_outcomes.jsonl": "2d9c6f0ff443f4eee5090318f02dcc2cc42792959df656ead9cafc98c5ebaf0d",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "93c1ac093940dcf307c3e07beadf0c6531631bf3a534b33cb591b2ed9ae2ac31",
|
||||
".modelhub_state/recovery_intents.jsonl": "e144dc1d248e807589d01635285b3974678fe0d958233e5f0df726a824030793",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "b866cb42511741d322fafea55e178e1941f4e992d28e9ab293fdf50a0e807acc",
|
||||
".modelhub_state/recovery_intents.jsonl": "11203f7a04914916a888e120f91caf2342d144e55d09e256d3171c1f4164a5b9",
|
||||
".modelhub_state/routing_intelligence.json": "07412703d74f144906be81443e9a62c63c16be09020d45073f303d32bb738925",
|
||||
".modelhub_state/submission_exclusions.jsonl": "b2de1125bf1f3e0597cd50cc30b563a5bfee39459e843767163eb72101cebc03",
|
||||
".modelhub_state/worker_crashes.jsonl": "7fa483580e87374493226e80bcf40bf8044847acba70045c96fa175966a1a8ba",
|
||||
"ledger/submissions.jsonl": "f33b133e8b25d698947530491a83aadf25705b3248f050bf303ed51744c603b5",
|
||||
"outcomes/submissions.jsonl": "77be1a4fc6a7c7b12ea68a3994f9f2730126081209a2fe599940d37f0cbdba1c"
|
||||
"outcomes/submissions.jsonl": "5fe831b4a8c3a7e13035cf6d883ceccd800b4cced52bee45b9a82546adf1ca14"
|
||||
},
|
||||
"generation": 5244,
|
||||
"generation": 5245,
|
||||
"phase": "cycle",
|
||||
"schemaVersion": 1,
|
||||
"updatedAt": "2026-09-10T15:55:46.542125+00:00",
|
||||
"updatedAt": "2026-09-10T15:58:19.111070+00:00",
|
||||
"writerId": "18252d8a8ef94333abd55823b2d81c42"
|
||||
}
|
||||
|
||||
@@ -118,7 +118,6 @@
|
||||
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-05T17:46:46.041763+00:00", "modelId": "Jackrong/DeepSeek-V4-Pro-Qwen3.5-9B", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 19306310880, "estimatedRequiredGiB": 21.599, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 19326449341, "modelscopeLicense": "apache-2.0", "modelscopeParams": 9653104368, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:text-generation-inference", "custom_tag:transformers", "custom_tag:unsloth", "custom_tag:qwen3_5", "custom_tag:reasoning", "custom_tag:distillation", "custom_tag:deepseek", "custom_tag:sft", "custom_tag:rl", "custom_tag:gspo", "custom_tag:math", "custom_tag:stem", "custom_tag:tool-use", "custom_tag:function-calling", "custom_tag:mtp"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 19326449341}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T09:43:33.758884+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4645816", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T18:34:58.235126+00:00", "modelId": "zstack/qwen3-4b-fake", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 25165824, "estimatedRequiredGiB": 0.028, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 25169485, "modelscopeLicense": null, "modelscopeParams": 25165435, "modelscopeTags": ["model_type:qwen3", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 25169485}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T10:33:51.484756+00:00", "targetGpu": "MetaX_c-500", "taskId": "4646505", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-05T20:12:28.913845+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "modelProfile": {"architectures": [], "configFingerprint": "e023b149992af42872de3c0c0c07dafd6a98d965e5bdb5a29ec5869e58d45725", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 356491104, "estimatedRequiredGiB": 1.363, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 1219879026, "modelscopeLicense": "other", "modelscopeParams": 327707521, "modelscopeTags": ["license:other", "library:gguf", "library:pytorch", "task:text-generation", "custom_tag:speculative-decoding", "custom_tag:dspark", "custom_tag:lfm2", "custom_tag:draft-model", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1219879026}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T12:03:37.084154+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4648146", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T20:12:28.913909+00:00", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W8A8-FP8-Channelwise-compressed-tensors", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084014480, "estimatedRequiredGiB": 10.162, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093198689, "modelscopeLicense": null, "modelscopeParams": 8030261248, "modelscopeTags": ["model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "recentProfileFeedback": {"attributableFailureCount": 1, "consecutiveFailures": 1, "decisionFailureRate": 1.0, "decisionSuccessRate": 0.0, "decisionTotal": 1, "failureBreakdown": {"backend_operator": 1}, "failureCount": 1, "failureRate": 1.0, "framework": "vllm", "lastTerminalAt": "2026-09-05T03:28:46.602358+00:00", "modelType": "llama", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "compressed-tensors", "successCount": 0, "successRate": 0.0, "targetGpu": "MetaX_c-500", "taskType": "text-generation", "total": 1, "unresolvedFailureCount": 0}, "repositoryOnDiskBytes": 9093198689}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T12:03:37.097444+00:00", "targetGpu": "MetaX_c-500", "taskId": "4648145", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T20:12:28.913926+00:00", "modelId": "siliconflow/Qwen3-Coder-30B-A3B-Instruct-fp8", "modelProfile": {"architectures": ["Qwen3MoeForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 31263441624, "estimatedRequiredGiB": 34.958, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_moe", "modelscopeFileSize": 31280205657, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 30554505408, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_moe", "library:safetensors", "library:other", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 31280205657}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T12:03:37.105603+00:00", "targetGpu": "MetaX_c-500", "taskId": "4648141", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-05T21:15:28.523439+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "modelProfile": {"architectures": [], "configFingerprint": "0dce7c4be445254429760db90c7969177ecdca460563778d8b3371a37cc1a50e", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 356491104, "estimatedRequiredGiB": 1.363, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 1219879026, "modelscopeLicense": "other", "modelscopeParams": 327707521, "modelscopeTags": ["license:other", "library:gguf", "library:pytorch", "task:text-generation", "custom_tag:speculative-decoding", "custom_tag:dspark", "custom_tag:lfm2", "custom_tag:draft-model", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1219879026}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T13:07:23.099409+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4649238", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-05T21:15:28.523328+00:00", "modelId": "zstack/qwen3-4b-fake", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 25165824, "estimatedRequiredGiB": 0.028, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 25169485, "modelscopeLicense": null, "modelscopeParams": 25165435, "modelscopeTags": ["model_type:qwen3", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 25169485}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T13:07:23.114152+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4649233", "taskType": "text-generation", "verifyResult": null}
|
||||
@@ -570,7 +569,7 @@
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-10T15:19:40.009069+00:00", "modelId": "BAAI/AREX-Turbo", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9078620368, "estimatedRequiredGiB": 10.18, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9109259594, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4539265536, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:deep-research", "custom_tag:reasoning", "custom_tag:tool-use", "custom_tag:long-context", "custom_tag:qwen3.5", "custom_tag:dense"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 1}, "failureCount": 1, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-09T14:21:21.949441+00:00", "modelType": "qwen3_5", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation", "total": 1, "unresolvedFailureCount": 1}, "repositoryOnDiskBytes": 9109259594}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-10T07:12:26.214438+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4755598", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-10T15:27:09.820539+00:00", "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-10T07:23:10.180003+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4755805", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-10T15:40:28.508808+00:00", "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-10T07:33:41.787375+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4756078", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-10T07:50:25.385741+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4756398", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-10T15:56:47.546297+00:00", "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-10T07:50:25.385741+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4756398", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-10T09:52:16.471862+00:00", "targetGpu": "MetaX_c-500", "taskId": "4758203", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "primitive-ai/Nex-N2.5-mini-NVFP4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23926484304, "estimatedRequiredGiB": 26.777, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 23959625409, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35107212656, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:nvfp4", "custom_tag:quantized", "custom_tag:moe", "custom_tag:nex-n2.5", "custom_tag:agentic", "custom_tag:tool-calling", "custom_tag:single-gpu", "custom_tag:blackwell"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 23959625409}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-10T13:27:23.407832+00:00", "targetGpu": "Biren_166m", "taskId": "4760701", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "primitive-ai/Nex-N2.5-mini-mixed-NVFP4-FP8", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26057989872, "estimatedRequiredGiB": 29.158, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26090095180, "modelscopeLicense": "apache-2.0", "modelscopeParams": 21416999792, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:nvfp4", "custom_tag:fp8", "custom_tag:mixed-precision", "custom_tag:quantized", "custom_tag:moe", "custom_tag:nex-n2.5", "custom_tag:agentic", "custom_tag:tool-calling", "custom_tag:single-gpu", "custom_tag:blackwell"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26090095180}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-10T13:27:23.401229+00:00", "targetGpu": "Biren_166m", "taskId": "4760703", "taskType": "text-generation", "verifyResult": null}
|
||||
|
||||
Reference in New Issue
Block a user