state: generation 8300 (cycle)

This commit is contained in:
2026-09-16 17:10:51 +00:00
parent 6e4b403f25
commit 09703b5337
10 changed files with 1948 additions and 1930 deletions

View File

@@ -1695,7 +1695,7 @@
"taskType": "text-generation" "taskType": "text-generation"
} }
}, },
"generatedAt": "2026-09-16T17:08:10.465742+00:00", "generatedAt": "2026-09-16T17:10:50.840293+00:00",
"summary": { "summary": {
"activeBlockCount": 82, "activeBlockCount": 82,
"byGpuFramework": { "byGpuFramework": {

View File

@@ -75,7 +75,7 @@
"7": { "7": {
"complete": false, "complete": false,
"lastError": "ModelHubAPIError: 系统错误", "lastError": "ModelHubAPIError: 系统错误",
"listingErrors": 605, "listingErrors": 606,
"nextPage": 1, "nextPage": 1,
"recordsScanned": 0, "recordsScanned": 0,
"uniqueRecords": 0 "uniqueRecords": 0
@@ -102,12 +102,12 @@
"cutoffAt": "2026-09-04T03:55:51.365685+00:00", "cutoffAt": "2026-09-04T03:55:51.365685+00:00",
"failureLogsInspected": 0, "failureLogsInspected": 0,
"mode": "incremental_decision_only", "mode": "incremental_decision_only",
"nextAccountIndex": 7, "nextAccountIndex": 8,
"recordsScanned": 0, "recordsScanned": 0,
"seenTaskIds": [], "seenTaskIds": [],
"startedAt": "2026-09-04T03:55:51.365685+00:00", "startedAt": "2026-09-04T03:55:51.365685+00:00",
"terminalRecords": 0, "terminalRecords": 0,
"uniqueRecords": 0, "uniqueRecords": 0,
"updatedAt": "2026-09-16T17:08:10.441138+00:00", "updatedAt": "2026-09-16T17:10:50.815806+00:00",
"version": 1 "version": 1
} }

View File

@@ -416,7 +416,7 @@
} }
}, },
"frameworkUpdatedAt": null, "frameworkUpdatedAt": null,
"generatedAt": "2026-09-16T17:06:36.491173+00:00", "generatedAt": "2026-09-16T17:09:17.698289+00:00",
"gpuStats": { "gpuStats": {
"Ascend_910-b3": { "Ascend_910-b3": {
"available": true, "available": true,

File diff suppressed because it is too large Load Diff

View File

@@ -1,6 +1,6 @@
{ {
"generatedAt": "2026-09-16T16:51:06.081512+00:00", "generatedAt": "2026-09-16T17:09:11.469983+00:00",
"lastSyncTime": "2026-09-16T16:51:05.849377+00:00", "lastSyncTime": "2026-09-16T17:09:11.209777+00:00",
"recentLimit": 300, "recentLimit": 300,
"report": { "report": {
"architectureCompatibilityBlocks": { "architectureCompatibilityBlocks": {
@@ -2181,25 +2181,25 @@
"decisionSuccessRate": 0.1154, "decisionSuccessRate": 0.1154,
"decisionTotal": 26, "decisionTotal": 26,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 11, "ambiguous_runtime": 12,
"backend_operator": 4, "backend_operator": 4,
"framework_architecture_unsupported": 6, "framework_architecture_unsupported": 6,
"model_load": 4, "model_load": 4,
"runtime_memory": 9, "runtime_memory": 9,
"参数/模板问题": 1 "参数/模板问题": 1
}, },
"failureCount": 35, "failureCount": 36,
"failureRate": 0.9211, "failureRate": 0.9231,
"framework": "vllm_0_17_0_corex_4_4_0", "framework": "vllm_0_17_0_corex_4_4_0",
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 3, "successCount": 3,
"successRate": 0.0789, "successRate": 0.0769,
"targetGpu": "Iluvatar_bi-150", "targetGpu": "Iluvatar_bi-150",
"taskType": "text-generation", "taskType": "text-generation",
"total": 38, "total": 39,
"unresolvedFailureCount": 12 "unresolvedFailureCount": 13
}, },
"Iluvatar_bi-150|vllm|text-generation": { "Iluvatar_bi-150|vllm|text-generation": {
"attributableFailureCount": 26, "attributableFailureCount": 26,
@@ -2859,22 +2859,22 @@
"decisionSuccessRate": 0.1154, "decisionSuccessRate": 0.1154,
"decisionTotal": 26, "decisionTotal": 26,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 11, "ambiguous_runtime": 12,
"backend_operator": 4, "backend_operator": 4,
"framework_architecture_unsupported": 6, "framework_architecture_unsupported": 6,
"model_load": 4, "model_load": 4,
"runtime_memory": 9, "runtime_memory": 9,
"参数/模板问题": 1 "参数/模板问题": 1
}, },
"failureCount": 35, "failureCount": 36,
"failureRate": 0.9211, "failureRate": 0.9231,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 3, "successCount": 3,
"successRate": 0.0789, "successRate": 0.0769,
"total": 38, "total": 39,
"unresolvedFailureCount": 12 "unresolvedFailureCount": 13
}, },
"vllm_fix_tokenizer": { "vllm_fix_tokenizer": {
"attributableFailureCount": 57, "attributableFailureCount": 57,
@@ -2920,7 +2920,7 @@
"unresolvedFailureCount": 5 "unresolvedFailureCount": 5
} }
}, },
"generatedAt": "2026-09-16T16:51:06.076246+00:00", "generatedAt": "2026-09-16T17:09:11.465074+00:00",
"gpuSummaries": { "gpuSummaries": {
"Ascend_910-b3": { "Ascend_910-b3": {
"attributableFailureCount": 49, "attributableFailureCount": 49,
@@ -3066,7 +3066,7 @@
"decisionSuccessRate": 0.1525, "decisionSuccessRate": 0.1525,
"decisionTotal": 59, "decisionTotal": 59,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 19, "ambiguous_runtime": 20,
"backend_operator": 7, "backend_operator": 7,
"framework_architecture_unsupported": 20, "framework_architecture_unsupported": 20,
"model_load": 11, "model_load": 11,
@@ -3075,15 +3075,15 @@
"参数/模板问题": 3, "参数/模板问题": 3,
"验证失败": 30 "验证失败": 30
}, },
"failureCount": 102, "failureCount": 103,
"failureRate": 0.9189, "failureRate": 0.9196,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 9, "successCount": 9,
"successRate": 0.0811, "successRate": 0.0804,
"total": 111, "total": 112,
"unresolvedFailureCount": 52 "unresolvedFailureCount": 53
}, },
"Iluvatar_mrv-100": { "Iluvatar_mrv-100": {
"attributableFailureCount": 111, "attributableFailureCount": 111,
@@ -4450,10 +4450,10 @@
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 0, "decisionTotal": 0,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 2, "ambiguous_runtime": 3,
"参数/模板问题": 1 "参数/模板问题": 1
}, },
"failureCount": 3, "failureCount": 4,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm_0_17_0_corex_4_4_0", "framework": "vllm_0_17_0_corex_4_4_0",
"modelType": "qwen3_5", "modelType": "qwen3_5",
@@ -4465,8 +4465,8 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Iluvatar_bi-150", "targetGpu": "Iluvatar_bi-150",
"taskType": "text-generation", "taskType": "text-generation",
"total": 3, "total": 4,
"unresolvedFailureCount": 3 "unresolvedFailureCount": 4
}, },
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|qwen3|none": { "Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|qwen3|none": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
@@ -8410,6 +8410,30 @@
"total": 1, "total": 1,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|qwen3_5|none|12": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm_0_17_0_corex_4_4_0",
"loadSizeLog2Bucket": 12,
"modelType": "qwen3_5",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Iluvatar_bi-150",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
},
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|qwen3_5|none|33": { "Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|qwen3_5|none|33": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
"decisionFailureRate": 0.0, "decisionFailureRate": 0.0,
@@ -10969,15 +10993,15 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
} }
}, },
"terminalRecords": 2032, "terminalRecords": 2033,
"totalRecords": 2126, "totalRecords": 2127,
"totals": { "totals": {
"attributableFailureCount": 654, "attributableFailureCount": 654,
"decisionFailureRate": 0.8922, "decisionFailureRate": 0.8922,
"decisionSuccessRate": 0.1078, "decisionSuccessRate": 0.1078,
"decisionTotal": 733, "decisionTotal": 733,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 492, "ambiguous_runtime": 493,
"backend_operator": 45, "backend_operator": 45,
"framework_architecture_unsupported": 429, "framework_architecture_unsupported": 429,
"memory_capacity": 12, "memory_capacity": 12,
@@ -10989,28 +11013,28 @@
"参数/模板问题": 124, "参数/模板问题": 124,
"验证失败": 673 "验证失败": 673
}, },
"failureCount": 1953, "failureCount": 1954,
"failureRate": 0.9611, "failureRate": 0.9611,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 10, "platformFailureCount": 10,
"successCount": 79, "successCount": 79,
"successRate": 0.0389, "successRate": 0.0389,
"total": 2032, "total": 2033,
"unresolvedFailureCount": 1289 "unresolvedFailureCount": 1290
}, },
"warnings": [ "warnings": [
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。", "GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。", "GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Sunrise_pt-200-x1 本地统计失败率偏高≥50%),建议重点关注。", "GPU Sunrise_pt-200-x1 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。", "GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。", "GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。", "GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。", "GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。", "GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。", "GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。", "GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。", "GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。",
"组合 Biren_166m|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Biren_166m|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Biren_166m|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Biren_166m|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -11036,6 +11060,6 @@
] ]
}, },
"storageMode": "decision_state_only", "storageMode": "decision_state_only",
"summarizedRecords": 2126, "summarizedRecords": 2127,
"version": 1 "version": 1
} }

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -573,11 +573,9 @@
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "submitTime": "2026-09-15T06:30:01.351500+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4867728", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-sunrise-pt-200-x1"} {"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "submitTime": "2026-09-15T06:30:01.351500+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4867728", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-sunrise-pt-200-x1"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "modelId": "mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "submitTime": "2026-09-15T06:30:01.357590+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4867730", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-sunrise-pt-200-x1"} {"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "modelId": "mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "submitTime": "2026-09-15T06:30:01.357590+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4867730", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-sunrise-pt-200-x1"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.150584+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868223", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.150584+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868223", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "modelId": "mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.091353+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868229", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.093936+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868226", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.093936+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868226", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "modelId": "mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.095741+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868225", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "modelId": "mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.095741+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868225", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/yujianboisme/qwen2.5-1.5b-darkhorse-code-fine-tuning", "modelId": "yujianboisme/qwen2.5-1.5b-darkhorse-code-fine-tuning", "submitTime": "2026-09-15T07:13:14.040459+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868228", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/yujianboisme/qwen2.5-1.5b-darkhorse-code-fine-tuning", "modelId": "yujianboisme/qwen2.5-1.5b-darkhorse-code-fine-tuning", "submitTime": "2026-09-15T07:13:14.040459+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868228", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.092802+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868221", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/hcnote/SparkMuse-4B", "modelId": "hcnote/SparkMuse-4B", "submitTime": "2026-09-15T07:13:14.087477+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868220", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/hcnote/SparkMuse-4B", "modelId": "hcnote/SparkMuse-4B", "submitTime": "2026-09-15T07:13:14.087477+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868220", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "submitTime": "2026-09-15T07:13:14.090218+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868222", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "submitTime": "2026-09-15T07:13:14.090218+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868222", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "modelId": "mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "submitTime": "2026-09-15T07:13:14.088705+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868224", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "modelId": "mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "submitTime": "2026-09-15T07:13:14.088705+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868224", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
@@ -623,6 +621,8 @@
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/r0b0tlab/Qwen3.8-27B-NVFP4-MTP-sm121", "modelId": "r0b0tlab/Qwen3.8-27B-NVFP4-MTP-sm121", "submitTime": "2026-09-09T21:08:32.653546+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4746450", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"} {"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/r0b0tlab/Qwen3.8-27B-NVFP4-MTP-sm121", "modelId": "r0b0tlab/Qwen3.8-27B-NVFP4-MTP-sm121", "submitTime": "2026-09-09T21:08:32.653546+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4746450", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/cyankiwi/Ornith-1.5-9B-AWQ-INT4", "modelId": "cyankiwi/Ornith-1.5-9B-AWQ-INT4", "submitTime": "2026-09-09T21:08:32.655185+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4746453", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"} {"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/cyankiwi/Ornith-1.5-9B-AWQ-INT4", "modelId": "cyankiwi/Ornith-1.5-9B-AWQ-INT4", "submitTime": "2026-09-09T21:08:32.655185+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4746453", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/XYZ-Aquila-mini-OptiQ-4bit", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.086293+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868227", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/XYZ-Aquila-mini-OptiQ-4bit", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.086293+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868227", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "modelId": "mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.091353+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868229", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "submitTime": "2026-09-15T07:13:14.092802+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4868221", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mthreads-s4000"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/ornith-ai/Ornith-1.5-35B-A3B-MLX-6bit", "modelId": "ornith-ai/Ornith-1.5-35B-A3B-MLX-6bit", "submitTime": "2026-09-09T21:08:32.684861+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4746455", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"} {"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/ornith-ai/Ornith-1.5-35B-A3B-MLX-6bit", "modelId": "ornith-ai/Ornith-1.5-35B-A3B-MLX-6bit", "submitTime": "2026-09-09T21:08:32.684861+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4746455", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "submitTime": "2026-09-14T19:47:58.577906+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4856415", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "submitTime": "2026-09-14T19:47:58.577906+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4856415", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "submitTime": "2026-09-14T20:04:19.387527+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4856586", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "submitTime": "2026-09-14T20:04:19.387527+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4856586", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}

View File

@@ -1,24 +1,24 @@
{ {
"agentVersion": "2026.09.04.4", "agentVersion": "2026.09.04.4",
"checksums": { "checksums": {
".modelhub_state/architecture_compatibility_blacklist.json": "0e7cffddd876f718188b94fd09f83af651c5a39ab2745aeff39e568e028d2e32", ".modelhub_state/architecture_compatibility_blacklist.json": "bf9673663908e3d1959a001a81393823b81aa4a03b4eb2d6a87de667e0c281af",
".modelhub_state/architecture_history_backfill.json": "0101f5f602aff093e8d6f1fb3f4a2ee9b84b671b614b0b3340cb33b4764d4eb2", ".modelhub_state/architecture_history_backfill.json": "71643834cf563dc78785c66106acd3462d982c9517930dd7def8f86980de0840",
".modelhub_state/market_intelligence.json": "eee42ce2ee3d29609e8228ae418a442720d8fe39b7a46ffb3694ff75e5745a0e", ".modelhub_state/market_intelligence.json": "365fdcfb91bc5492d0962cb1c85513046f6fd2a29cdfd315cda576b8db6f1475",
".modelhub_state/official_capabilities.json": "640087831487810ba9e2c0718b0ffff87ed9c95820662b9ae36061d67c31f5b6", ".modelhub_state/official_capabilities.json": "b25cc334dd1c64222a63bf219d78ee4fd048231f1c113ef31d64dc2ee67bb588",
".modelhub_state/outcome_checkpoint.json": "f983354ad02e23d87ba4f866eb5aeaae0429b91eea0a7cc9db3f9e488c01f001", ".modelhub_state/outcome_checkpoint.json": "43c4e41958121aaaa4be1068a257e7708aad644d00b9a9a3852ef2c99c1d6635",
".modelhub_state/queue_cleanup_latest.json": "057079282bc277b3a97c1805915c45464c12312f8ee3ba54d0366d05f34f4466", ".modelhub_state/queue_cleanup_latest.json": "057079282bc277b3a97c1805915c45464c12312f8ee3ba54d0366d05f34f4466",
".modelhub_state/recent_outcomes.jsonl": "4647f63c856ce0f5a014c36cc12aa304266c858bf6c218b953ef85e3c811ce8e", ".modelhub_state/recent_outcomes.jsonl": "4647f63c856ce0f5a014c36cc12aa304266c858bf6c218b953ef85e3c811ce8e",
".modelhub_state/recovery_active_tasks.jsonl": "c3ecb82117b136c3bae95718af855f7cf9248852472c54f3bc4c2839ea0a35f8", ".modelhub_state/recovery_active_tasks.jsonl": "e7169a2d2c049cf92e85c094faf988c7482348f73012aade982f8f34cb29d25d",
".modelhub_state/recovery_intents.jsonl": "44f7eb486f7e485d912e9914e97198391d796031b4011a2b2405fee0814a3c99", ".modelhub_state/recovery_intents.jsonl": "aa55c02cf43e3ef2238fe7c865cac2d78d62031913de20edae086f9e67c3b5ae",
".modelhub_state/routing_intelligence.json": "c2bbbf3f23b3ed081a6174109100aea36dc5727a8ca34ef635323840e25553dd", ".modelhub_state/routing_intelligence.json": "c2bbbf3f23b3ed081a6174109100aea36dc5727a8ca34ef635323840e25553dd",
".modelhub_state/submission_exclusions.jsonl": "632739c6cd6a2dd2237572ba18e6a390c5aa37ce56a18a2c1341d78569421552", ".modelhub_state/submission_exclusions.jsonl": "632739c6cd6a2dd2237572ba18e6a390c5aa37ce56a18a2c1341d78569421552",
".modelhub_state/worker_crashes.jsonl": "b41d1dda7bab11bedd96de40fecd2d416e47396901b3283f48a56de5d6348983", ".modelhub_state/worker_crashes.jsonl": "b41d1dda7bab11bedd96de40fecd2d416e47396901b3283f48a56de5d6348983",
"ledger/submissions.jsonl": "ac842bf32275c8854b018cb57f408f8e2a76e32f9dd3c2c8de07e1a61638d1fc", "ledger/submissions.jsonl": "c938ca15c5640f3b1da917b9a3aa64885a4b79417a81b64dddb1ec17913d1eb4",
"outcomes/submissions.jsonl": "c2ed8d956e387d936d38a2b0c029e78dbc99e2574faac4b8474935c7fe14d848" "outcomes/submissions.jsonl": "447222a889b3b07abb26783bfed2d971707f3a5e9921410ee126634a216e17ac"
}, },
"generation": 8299, "generation": 8300,
"phase": "cycle", "phase": "cycle",
"schemaVersion": 1, "schemaVersion": 1,
"updatedAt": "2026-09-16T17:08:10.553221+00:00", "updatedAt": "2026-09-16T17:10:51.593932+00:00",
"writerId": "eccb3e0018f640d19e578c271a207b5c" "writerId": "eccb3e0018f640d19e578c271a207b5c"
} }

View File

@@ -85,7 +85,6 @@
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-06T15:18:32.207677+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26119812080, "estimatedRequiredGiB": 29.233, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26156887523, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26156887523}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T07:17:44.185498+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4663094", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-06T15:18:32.207677+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26119812080, "estimatedRequiredGiB": 29.233, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26156887523, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26156887523}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T07:17:44.185498+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4663094", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-06T16:14:21.709544+00:00", "modelId": "empero-ai/Qwen3.8-4B-Distill", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5777, "estimatedRequiredGiB": 10.441, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9342747188, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4659865088, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:empero-ai", "custom_tag:qwen3.5", "custom_tag:qwen3.8", "custom_tag:distillation", "custom_tag:reasoning", "custom_tag:function-calling", "custom_tag:sft"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9342747188}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T08:12:14.777960+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4663715", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-06T16:14:21.709544+00:00", "modelId": "empero-ai/Qwen3.8-4B-Distill", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5777, "estimatedRequiredGiB": 10.441, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9342747188, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4659865088, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:empero-ai", "custom_tag:qwen3.5", "custom_tag:qwen3.8", "custom_tag:distillation", "custom_tag:reasoning", "custom_tag:function-calling", "custom_tag:sft"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9342747188}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T08:12:14.777960+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4663715", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-06T16:36:18.900122+00:00", "modelId": "empero-ai/Qwen3.8-4B-Distill", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5777, "estimatedRequiredGiB": 10.441, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9342747188, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4659865088, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:empero-ai", "custom_tag:qwen3.5", "custom_tag:qwen3.8", "custom_tag:distillation", "custom_tag:reasoning", "custom_tag:function-calling", "custom_tag:sft"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9342747188}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T08:34:09.456899+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4664143", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-06T16:36:18.900122+00:00", "modelId": "empero-ai/Qwen3.8-4B-Distill", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5777, "estimatedRequiredGiB": 10.441, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9342747188, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4659865088, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:empero-ai", "custom_tag:qwen3.5", "custom_tag:qwen3.8", "custom_tag:distillation", "custom_tag:reasoning", "custom_tag:function-calling", "custom_tag:sft"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9342747188}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T08:34:09.456899+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4664143", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-06T17:08:48.403518+00:00", "modelId": "empero-ai/Qwen3.8-4B-Distill", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5777, "estimatedRequiredGiB": 10.441, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9342747188, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4659865088, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:empero-ai", "custom_tag:qwen3.5", "custom_tag:qwen3.8", "custom_tag:distillation", "custom_tag:reasoning", "custom_tag:function-calling", "custom_tag:sft"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9342747188}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T09:06:12.932838+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4664478", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-06T17:46:47.144783+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26119812080, "estimatedRequiredGiB": 29.233, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:50"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26156887523, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26156887523}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T09:46:33.145376+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4664976", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-06T17:46:47.144783+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26119812080, "estimatedRequiredGiB": 29.233, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:50"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26156887523, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26156887523}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T09:46:33.145376+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4664976", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-06T18:28:08.305046+00:00", "modelId": "empero-ai/Qwen3.8-4B-Distill", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5777, "estimatedRequiredGiB": 10.441, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:50"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9342747188, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4659865088, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:empero-ai", "custom_tag:qwen3.5", "custom_tag:qwen3.8", "custom_tag:distillation", "custom_tag:reasoning", "custom_tag:function-calling", "custom_tag:sft"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9342747188}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T10:24:50.144356+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4665470", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-06T18:28:08.305046+00:00", "modelId": "empero-ai/Qwen3.8-4B-Distill", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5777, "estimatedRequiredGiB": 10.441, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:50"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9342747188, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4659865088, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:empero-ai", "custom_tag:qwen3.5", "custom_tag:qwen3.8", "custom_tag:distillation", "custom_tag:reasoning", "custom_tag:function-calling", "custom_tag:sft"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9342747188}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T10:24:50.144356+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4665470", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T20:39:29.797774+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26119812080, "estimatedRequiredGiB": 29.233, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26156887523, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26156887523}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T12:37:22.515030+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4667091", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T20:39:29.797774+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26119812080, "estimatedRequiredGiB": 29.233, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26156887523, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26156887523}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T12:37:22.515030+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4667091", "taskType": "text-generation", "verifyResult": null}
@@ -621,7 +620,7 @@
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": "2026-09-16T06:24:42.305201+00:00", "modelId": "mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 13249080076, "estimatedRequiredGiB": 14.83, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 13269529926, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3825799024, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:expert-pruning", "custom_tag:reap", "custom_tag:moe", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13269529926}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-15T22:23:03.691857+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4881076", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm-customized", "lastSyncTime": "2026-09-16T06:24:42.305201+00:00", "modelId": "mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 13249080076, "estimatedRequiredGiB": 14.83, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 13269529926, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3825799024, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:expert-pruning", "custom_tag:reap", "custom_tag:moe", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13269529926}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-15T22:23:03.691857+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4881076", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-16T07:23:50.803385+00:00", "modelId": "hcnote/SparkMuse-4B-INT8", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4453695184, "estimatedRequiredGiB": 4.989, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 4463853848, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4114366080, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:creative-writing", "custom_tag:novel-generation", "custom_tag:roleplay", "custom_tag:nsfw", "custom_tag:quantized", "custom_tag:torchao", "custom_tag:long-context", "custom_tag:spark"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "torchao", "repositoryOnDiskBytes": 4463853848}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-15T23:23:24.114544+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4881705", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-16T07:23:50.803385+00:00", "modelId": "hcnote/SparkMuse-4B-INT8", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4453695184, "estimatedRequiredGiB": 4.989, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 4463853848, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4114366080, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:creative-writing", "custom_tag:novel-generation", "custom_tag:roleplay", "custom_tag:nsfw", "custom_tag:quantized", "custom_tag:torchao", "custom_tag:long-context", "custom_tag:spark"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "torchao", "repositoryOnDiskBytes": 4463853848}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-15T23:23:24.114544+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4881705", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-16T16:51:05.849377+00:00", "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-16T08:44:19.185430+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4889469", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-16T16:51:05.849377+00:00", "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-16T08:44:19.185430+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4889469", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T09:03:09.666793+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4889677", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-16T17:09:11.209752+00:00", "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-16T09:03:09.666793+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4889677", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T09:20:28.174669+00:00", "targetGpu": "MetaX_c-500", "taskId": "4889839", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T09:20:28.174669+00:00", "targetGpu": "MetaX_c-500", "taskId": "4889839", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T09:37:59.405066+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4890080", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm-mlu", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T09:37:59.405066+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4890080", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T09:49:03.662873+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4890217", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T09:49:03.662873+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4890217", "taskType": "text-generation", "verifyResult": null}