state: generation 9333 (intent)

This commit is contained in:
2026-09-18 15:43:06 +00:00
parent 35c936c470
commit e44debf056
7 changed files with 916 additions and 823 deletions

View File

@@ -1743,7 +1743,7 @@
"taskType": "text-generation" "taskType": "text-generation"
} }
}, },
"generatedAt": "2026-09-18T15:40:27.987699+00:00", "generatedAt": "2026-09-18T15:41:29.055443+00:00",
"summary": { "summary": {
"activeBlockCount": 84, "activeBlockCount": 84,
"byGpuFramework": { "byGpuFramework": {

View File

@@ -416,7 +416,7 @@
} }
}, },
"frameworkUpdatedAt": null, "frameworkUpdatedAt": null,
"generatedAt": "2026-09-18T15:38:42.432531+00:00", "generatedAt": "2026-09-18T15:41:36.141349+00:00",
"gpuStats": { "gpuStats": {
"Ascend_910-b3": { "Ascend_910-b3": {
"available": true, "available": true,

File diff suppressed because it is too large Load Diff

View File

@@ -1,6 +1,6 @@
{ {
"generatedAt": "2026-09-18T15:24:48.950945+00:00", "generatedAt": "2026-09-18T15:41:29.020363+00:00",
"lastSyncTime": "2026-09-18T15:24:48.682231+00:00", "lastSyncTime": "2026-09-18T15:41:28.787881+00:00",
"recentLimit": 300, "recentLimit": 300,
"report": { "report": {
"architectureCompatibilityBlocks": { "architectureCompatibilityBlocks": {
@@ -2634,26 +2634,26 @@
"unresolvedFailureCount": 145 "unresolvedFailureCount": 145
}, },
"Vastai_va16|vllm_fix_tokenizer|text-generation": { "Vastai_va16|vllm_fix_tokenizer|text-generation": {
"attributableFailureCount": 3, "attributableFailureCount": 4,
"decisionFailureRate": 0.75, "decisionFailureRate": 0.8,
"decisionSuccessRate": 0.25, "decisionSuccessRate": 0.2,
"decisionTotal": 4, "decisionTotal": 5,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 5, "ambiguous_runtime": 5,
"framework_architecture_unsupported": 2, "framework_architecture_unsupported": 3,
"model_load": 1 "model_load": 1
}, },
"failureCount": 8, "failureCount": 9,
"failureRate": 0.8889, "failureRate": 0.9,
"framework": "vllm_fix_tokenizer", "framework": "vllm_fix_tokenizer",
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 1, "successCount": 1,
"successRate": 0.1111, "successRate": 0.1,
"targetGpu": "Vastai_va16", "targetGpu": "Vastai_va16",
"taskType": "text-generation", "taskType": "text-generation",
"total": 9, "total": 10,
"unresolvedFailureCount": 5 "unresolvedFailureCount": 5
}, },
"Vastai_va16|vllm|text-generation": { "Vastai_va16|vllm|text-generation": {
@@ -2724,18 +2724,18 @@
"unresolvedFailureCount": 29 "unresolvedFailureCount": 29
}, },
"hygon_k100-ai|vllm-patch-tokenizer|text-generation": { "hygon_k100-ai|vllm-patch-tokenizer|text-generation": {
"attributableFailureCount": 12, "attributableFailureCount": 13,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 12, "decisionTotal": 13,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 10, "ambiguous_runtime": 10,
"framework_architecture_unsupported": 8, "framework_architecture_unsupported": 9,
"model_load": 1, "model_load": 1,
"runtime_memory": 3, "runtime_memory": 3,
"参数/模板问题": 8 "参数/模板问题": 8
}, },
"failureCount": 30, "failureCount": 31,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm-patch-tokenizer", "framework": "vllm-patch-tokenizer",
"pendingCount": 0, "pendingCount": 0,
@@ -2745,7 +2745,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "hygon_k100-ai", "targetGpu": "hygon_k100-ai",
"taskType": "text-generation", "taskType": "text-generation",
"total": 30, "total": 31,
"unresolvedFailureCount": 18 "unresolvedFailureCount": 18
}, },
"hygon_k100-ai|vllm|text-generation": { "hygon_k100-ai|vllm|text-generation": {
@@ -2889,25 +2889,25 @@
"unresolvedFailureCount": 40 "unresolvedFailureCount": 40
}, },
"vllm-patch-tokenizer": { "vllm-patch-tokenizer": {
"attributableFailureCount": 12, "attributableFailureCount": 13,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 12, "decisionTotal": 13,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 10, "ambiguous_runtime": 10,
"framework_architecture_unsupported": 8, "framework_architecture_unsupported": 9,
"model_load": 1, "model_load": 1,
"runtime_memory": 3, "runtime_memory": 3,
"参数/模板问题": 8 "参数/模板问题": 8
}, },
"failureCount": 30, "failureCount": 31,
"failureRate": 1.0, "failureRate": 1.0,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 0, "successCount": 0,
"successRate": 0.0, "successRate": 0.0,
"total": 30, "total": 31,
"unresolvedFailureCount": 18 "unresolvedFailureCount": 18
}, },
"vllm_0_17_0_corex_4_4_0": { "vllm_0_17_0_corex_4_4_0": {
@@ -2934,27 +2934,27 @@
"unresolvedFailureCount": 13 "unresolvedFailureCount": 13
}, },
"vllm_fix_tokenizer": { "vllm_fix_tokenizer": {
"attributableFailureCount": 65, "attributableFailureCount": 66,
"decisionFailureRate": 0.9848, "decisionFailureRate": 0.9851,
"decisionSuccessRate": 0.0152, "decisionSuccessRate": 0.0149,
"decisionTotal": 66, "decisionTotal": 67,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 73, "ambiguous_runtime": 73,
"backend_operator": 7, "backend_operator": 7,
"framework_architecture_unsupported": 9, "framework_architecture_unsupported": 10,
"model_load": 1, "model_load": 1,
"platform_infrastructure": 8, "platform_infrastructure": 8,
"tokenizer_compatibility": 48, "tokenizer_compatibility": 48,
"参数/模板问题": 8 "参数/模板问题": 8
}, },
"failureCount": 154, "failureCount": 155,
"failureRate": 0.9935, "failureRate": 0.9936,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 8, "platformFailureCount": 8,
"successCount": 1, "successCount": 1,
"successRate": 0.0065, "successRate": 0.0064,
"total": 155, "total": 156,
"unresolvedFailureCount": 81 "unresolvedFailureCount": 81
}, },
"vllm_tokenizer_patch": { "vllm_tokenizer_patch": {
@@ -2977,7 +2977,7 @@
"unresolvedFailureCount": 18 "unresolvedFailureCount": 18
} }
}, },
"generatedAt": "2026-09-18T15:24:48.909424+00:00", "generatedAt": "2026-09-18T15:41:29.014648+00:00",
"gpuSummaries": { "gpuSummaries": {
"Ascend_910-b3": { "Ascend_910-b3": {
"attributableFailureCount": 58, "attributableFailureCount": 58,
@@ -3267,36 +3267,36 @@
"unresolvedFailureCount": 51 "unresolvedFailureCount": 51
}, },
"Vastai_va16": { "Vastai_va16": {
"attributableFailureCount": 61, "attributableFailureCount": 62,
"decisionFailureRate": 0.7625, "decisionFailureRate": 0.7654,
"decisionSuccessRate": 0.2375, "decisionSuccessRate": 0.2346,
"decisionTotal": 80, "decisionTotal": 81,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 79, "ambiguous_runtime": 79,
"framework_architecture_unsupported": 54, "framework_architecture_unsupported": 55,
"memory_capacity": 1, "memory_capacity": 1,
"model_load": 6, "model_load": 6,
"参数/模板问题": 15, "参数/模板问题": 15,
"验证失败": 130 "验证失败": 130
}, },
"failureCount": 285, "failureCount": 286,
"failureRate": 0.9375, "failureRate": 0.9377,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 19, "successCount": 19,
"successRate": 0.0625, "successRate": 0.0623,
"total": 304, "total": 305,
"unresolvedFailureCount": 224 "unresolvedFailureCount": 224
}, },
"hygon_k100-ai": { "hygon_k100-ai": {
"attributableFailureCount": 68, "attributableFailureCount": 69,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 68, "decisionTotal": 69,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 37, "ambiguous_runtime": 37,
"framework_architecture_unsupported": 52, "framework_architecture_unsupported": 53,
"memory_capacity": 1, "memory_capacity": 1,
"model_load": 6, "model_load": 6,
"repository_structure": 4, "repository_structure": 4,
@@ -3304,14 +3304,14 @@
"参数/模板问题": 10, "参数/模板问题": 10,
"验证失败": 27 "验证失败": 27
}, },
"failureCount": 142, "failureCount": 143,
"failureRate": 1.0, "failureRate": 1.0,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 0, "successCount": 0,
"successRate": 0.0, "successRate": 0.0,
"total": 142, "total": 143,
"unresolvedFailureCount": 74 "unresolvedFailureCount": 74
} }
}, },
@@ -7183,6 +7183,29 @@
"total": 1, "total": 1,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Vastai_va16|vllm_fix_tokenizer|text-generation|qwen3_5|modelopt": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"modelType": "qwen3_5",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "modelopt",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Vastai_va16",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Vastai_va16|vllm_fix_tokenizer|text-generation|qwen3|none": { "Vastai_va16|vllm_fix_tokenizer|text-generation|qwen3|none": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
"decisionFailureRate": 0.0, "decisionFailureRate": 0.0,
@@ -7344,6 +7367,29 @@
"total": 1, "total": 1,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"hygon_k100-ai|vllm-patch-tokenizer|text-generation|qwen3_5_text|none": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm-patch-tokenizer",
"modelType": "qwen3_5_text",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "hygon_k100-ai",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"hygon_k100-ai|vllm-patch-tokenizer|text-generation|qwen3_5|none": { "hygon_k100-ai|vllm-patch-tokenizer|text-generation|qwen3_5|none": {
"attributableFailureCount": 2, "attributableFailureCount": 2,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
@@ -14118,6 +14164,30 @@
"total": 1, "total": 1,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Vastai_va16|vllm_fix_tokenizer|text-generation|qwen3_5|modelopt|33": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"loadSizeLog2Bucket": 33,
"modelType": "qwen3_5",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "modelopt",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Vastai_va16",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Vastai_va16|vllm_fix_tokenizer|text-generation|qwen3|none|32": { "Vastai_va16|vllm_fix_tokenizer|text-generation|qwen3|none|32": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
"decisionFailureRate": 0.0, "decisionFailureRate": 0.0,
@@ -14309,6 +14379,30 @@
"total": 1, "total": 1,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"hygon_k100-ai|vllm-patch-tokenizer|text-generation|qwen3_5_text|none|32": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm-patch-tokenizer",
"loadSizeLog2Bucket": 32,
"modelType": "qwen3_5_text",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "hygon_k100-ai",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"hygon_k100-ai|vllm-patch-tokenizer|text-generation|qwen3_5|none|33": { "hygon_k100-ai|vllm-patch-tokenizer|text-generation|qwen3_5|none|33": {
"attributableFailureCount": 1, "attributableFailureCount": 1,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
@@ -14455,17 +14549,17 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
} }
}, },
"terminalRecords": 2257, "terminalRecords": 2259,
"totalRecords": 2360, "totalRecords": 2362,
"totals": { "totals": {
"attributableFailureCount": 734, "attributableFailureCount": 736,
"decisionFailureRate": 0.8408, "decisionFailureRate": 0.8411,
"decisionSuccessRate": 0.1592, "decisionSuccessRate": 0.1589,
"decisionTotal": 873, "decisionTotal": 875,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 567, "ambiguous_runtime": 567,
"backend_operator": 55, "backend_operator": 55,
"framework_architecture_unsupported": 485, "framework_architecture_unsupported": 487,
"memory_capacity": 12, "memory_capacity": 12,
"model_load": 88, "model_load": 88,
"platform_infrastructure": 11, "platform_infrastructure": 11,
@@ -14475,14 +14569,14 @@
"参数/模板问题": 133, "参数/模板问题": 133,
"验证失败": 673 "验证失败": 673
}, },
"failureCount": 2118, "failureCount": 2120,
"failureRate": 0.9384, "failureRate": 0.9385,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 11, "platformFailureCount": 11,
"successCount": 139, "successCount": 139,
"successRate": 0.0616, "successRate": 0.0615,
"total": 2257, "total": 2259,
"unresolvedFailureCount": 1373 "unresolvedFailureCount": 1373
}, },
"warnings": [ "warnings": [
@@ -14522,6 +14616,6 @@
] ]
}, },
"storageMode": "decision_state_only", "storageMode": "decision_state_only",
"summarizedRecords": 2360, "summarizedRecords": 2362,
"version": 1 "version": 1
} }

View File

@@ -914,6 +914,7 @@
{"batchId": "604b79b346d24c56a3a5aaac6cbbc7e2", "completedAt": "2026-09-18T15:02:37.237440+00:00", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-18T15:02:31.567843+00:00", "framework": "vllm_fix_tokenizer", "intentId": "06368add7bae4d378ca0dae55bb87edd", "lastModified": "2026-09-18T05:43:55+00:00", "modelAddress": "https://modelscope.cn/models/prism-ml/Ternary-Bonsai-2-27B-mlx-2bit", "reason": null, "reconciledAt": "2026-09-18T15:34:28.738779+00:00", "repoId": "prism-ml/Ternary-Bonsai-2-27B-mlx-2bit", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Kunlunxin_p-800", "taskId": "4936404", "taskType": "text-generation"} {"batchId": "604b79b346d24c56a3a5aaac6cbbc7e2", "completedAt": "2026-09-18T15:02:37.237440+00:00", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-18T15:02:31.567843+00:00", "framework": "vllm_fix_tokenizer", "intentId": "06368add7bae4d378ca0dae55bb87edd", "lastModified": "2026-09-18T05:43:55+00:00", "modelAddress": "https://modelscope.cn/models/prism-ml/Ternary-Bonsai-2-27B-mlx-2bit", "reason": null, "reconciledAt": "2026-09-18T15:34:28.738779+00:00", "repoId": "prism-ml/Ternary-Bonsai-2-27B-mlx-2bit", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Kunlunxin_p-800", "taskId": "4936404", "taskType": "text-generation"}
{"batchId": "33210cc7776e4d1f80266991e4a6686f", "completedAt": "2026-09-18T15:07:38.945993+00:00", "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configSource": "modelhub_live", "createdAt": "2026-09-18T15:07:32.755874+00:00", "framework": "vllm", "intentId": "2847d1563509462f90fb104319e6d307", "lastModified": "2026-09-16T16:32:52+00:00", "modelAddress": "https://modelscope.cn/models/IFM/K2-Horizon-7B-FP8", "reason": null, "reconciledAt": "2026-09-18T15:34:28.738641+00:00", "repoId": "IFM/K2-Horizon-7B-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "MetaX_c-500", "taskId": "4936437", "taskType": "text-generation"} {"batchId": "33210cc7776e4d1f80266991e4a6686f", "completedAt": "2026-09-18T15:07:38.945993+00:00", "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configSource": "modelhub_live", "createdAt": "2026-09-18T15:07:32.755874+00:00", "framework": "vllm", "intentId": "2847d1563509462f90fb104319e6d307", "lastModified": "2026-09-16T16:32:52+00:00", "modelAddress": "https://modelscope.cn/models/IFM/K2-Horizon-7B-FP8", "reason": null, "reconciledAt": "2026-09-18T15:34:28.738641+00:00", "repoId": "IFM/K2-Horizon-7B-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "MetaX_c-500", "taskId": "4936437", "taskType": "text-generation"}
{"batchId": "135ba641c64640bb86f256fda8e253d6", "completedAt": "2026-09-18T15:26:37.206310+00:00", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-18T15:26:30.207963+00:00", "framework": "vllm_tokenizer_patch", "intentId": "f750850f52b4424585d77b2897874a03", "lastModified": "2026-09-16T16:32:52+00:00", "modelAddress": "https://modelscope.cn/models/IFM/K2-Horizon-7B-FP8", "reason": null, "reconciledAt": "2026-09-18T15:34:28.739041+00:00", "repoId": "IFM/K2-Horizon-7B-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Ascend_910-b3", "taskId": "4936617", "taskType": "text-generation"} {"batchId": "135ba641c64640bb86f256fda8e253d6", "completedAt": "2026-09-18T15:26:37.206310+00:00", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-18T15:26:30.207963+00:00", "framework": "vllm_tokenizer_patch", "intentId": "f750850f52b4424585d77b2897874a03", "lastModified": "2026-09-16T16:32:52+00:00", "modelAddress": "https://modelscope.cn/models/IFM/K2-Horizon-7B-FP8", "reason": null, "reconciledAt": "2026-09-18T15:34:28.739041+00:00", "repoId": "IFM/K2-Horizon-7B-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Ascend_910-b3", "taskId": "4936617", "taskType": "text-generation"}
{"batchId": "609a2e994d8e43b49ac79f542d7cbdea", "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configSource": "modelhub_live", "createdAt": "2026-09-18T15:43:06.335172+00:00", "framework": "vllm-mlu", "intentId": "f8f544e4fa4e4cdfafa021536b5c8cb5", "lastModified": "2026-09-16T16:32:52+00:00", "modelAddress": "https://modelscope.cn/models/IFM/K2-Horizon-7B-FP8", "repoId": "IFM/K2-Horizon-7B-FP8", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x4", "taskType": "text-generation"}
{"batchId": "8d27d6d5a99748efbdd97223e640c598", "completedAt": "2026-09-18T14:35:12.136557+00:00", "configFingerprint": "359725c5cf0821e26681bc436dad944747db4ac72a1bf9e70b5f6f8e63a023b4", "configSource": "modelhub_live", "createdAt": "2026-09-18T14:35:06.015304+00:00", "framework": "llamacpp", "intentId": "10303336f4f24e2f98aa86764cc2f0fc", "repoId": "voconly/parakeet-unified-en-0.6b-gguf", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "uniqueness_rejected", "targetGpu": "Mthreads_s4000", "taskId": null, "taskType": "text-generation"} {"batchId": "8d27d6d5a99748efbdd97223e640c598", "completedAt": "2026-09-18T14:35:12.136557+00:00", "configFingerprint": "359725c5cf0821e26681bc436dad944747db4ac72a1bf9e70b5f6f8e63a023b4", "configSource": "modelhub_live", "createdAt": "2026-09-18T14:35:06.015304+00:00", "framework": "llamacpp", "intentId": "10303336f4f24e2f98aa86764cc2f0fc", "repoId": "voconly/parakeet-unified-en-0.6b-gguf", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "uniqueness_rejected", "targetGpu": "Mthreads_s4000", "taskId": null, "taskType": "text-generation"}
{"batchId": "bfd8069a25c5437f87c1eb1f45b313b1", "completedAt": "2026-09-18T14:33:31.434713+00:00", "configFingerprint": "e1276bf11e548198418262ddca64a6c4514301bb26e65ecc7e04f8d8f0442b6b", "configSource": "modelhub_live", "createdAt": "2026-09-18T14:33:23.815503+00:00", "framework": "llamacpp", "intentId": "3c35cff71fae4e2997f1d1fa492cdda8", "repoId": "voconly/parakeet-unified-en-0.6b-gguf", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "uniqueness_rejected", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"} {"batchId": "bfd8069a25c5437f87c1eb1f45b313b1", "completedAt": "2026-09-18T14:33:31.434713+00:00", "configFingerprint": "e1276bf11e548198418262ddca64a6c4514301bb26e65ecc7e04f8d8f0442b6b", "configSource": "modelhub_live", "createdAt": "2026-09-18T14:33:23.815503+00:00", "framework": "llamacpp", "intentId": "3c35cff71fae4e2997f1d1fa492cdda8", "repoId": "voconly/parakeet-unified-en-0.6b-gguf", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "uniqueness_rejected", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"}
{"batchId": "18470fd259dd4f84ad6911206103ac42", "completedAt": "2026-09-18T14:31:05.798368+00:00", "configFingerprint": "680fcdbc0952bbae15ed7a5a114782d3eaa8806d60778df64b3e1a99d41a066b", "configSource": "modelhub_live", "createdAt": "2026-09-18T14:30:53.123908+00:00", "framework": "llamacpp", "intentId": "c7daff34b5a4470ab91e00ee1da68653", "repoId": "voconly/parakeet-unified-en-0.6b-gguf", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "uniqueness_rejected", "targetGpu": "hygon_k100-ai", "taskId": null, "taskType": "text-generation"} {"batchId": "18470fd259dd4f84ad6911206103ac42", "completedAt": "2026-09-18T14:31:05.798368+00:00", "configFingerprint": "680fcdbc0952bbae15ed7a5a114782d3eaa8806d60778df64b3e1a99d41a066b", "configSource": "modelhub_live", "createdAt": "2026-09-18T14:30:53.123908+00:00", "framework": "llamacpp", "intentId": "c7daff34b5a4470ab91e00ee1da68653", "repoId": "voconly/parakeet-unified-en-0.6b-gguf", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "uniqueness_rejected", "targetGpu": "hygon_k100-ai", "taskId": null, "taskType": "text-generation"}

View File

@@ -1,24 +1,24 @@
{ {
"agentVersion": "2026.09.04.4", "agentVersion": "2026.09.04.4",
"checksums": { "checksums": {
".modelhub_state/architecture_compatibility_blacklist.json": "11f3bdef29f22e0dc3a24d5d6450fa550331192291488c8fccd0fc4c0920dc89", ".modelhub_state/architecture_compatibility_blacklist.json": "ec6ba98fd88cf520b31f591802304ff31e11f4d080c671536ef747db87981042",
".modelhub_state/architecture_history_backfill.json": "821e04077ec7de2ac6605b85ee681e1c11c772d887ca2aa403c2105ae8849d90", ".modelhub_state/architecture_history_backfill.json": "821e04077ec7de2ac6605b85ee681e1c11c772d887ca2aa403c2105ae8849d90",
".modelhub_state/market_intelligence.json": "28b63199bcf11ff198072131815626278eda65630817211ba8ff43c13c40f4a7", ".modelhub_state/market_intelligence.json": "cfe19cd5f49608de8611707e7f05f6bc5baf10fc1c538bcc86eaf6b8ce802802",
".modelhub_state/official_capabilities.json": "3bfcba3766080431a9336839f78b7a9258f4eb008cd0b65a2de13d2ac8e17585", ".modelhub_state/official_capabilities.json": "50cc8fc808144c7d57d26a9de2e329cf177b96b26a14822eed2c22960a0480ba",
".modelhub_state/outcome_checkpoint.json": "373119f9fbb148e30c9fffd621a4d73d5cdae3d141124de71d3b77783891c192", ".modelhub_state/outcome_checkpoint.json": "207c01d06e253bec7628f56745b2bff92ecc3cac15abebdca7d5a18033f3497b",
".modelhub_state/queue_cleanup_latest.json": "3c142558258883a3e45602d09e9392c46bfde3deae83aabc51d01f3b1c2d91d9", ".modelhub_state/queue_cleanup_latest.json": "3c142558258883a3e45602d09e9392c46bfde3deae83aabc51d01f3b1c2d91d9",
".modelhub_state/recent_outcomes.jsonl": "2d5bc67a36c72a2c3a6906f9d6fc23b5e73c0a4175033a1552d7b0978d7da2c4", ".modelhub_state/recent_outcomes.jsonl": "2d5bc67a36c72a2c3a6906f9d6fc23b5e73c0a4175033a1552d7b0978d7da2c4",
".modelhub_state/recovery_active_tasks.jsonl": "53dbe015f17d9e0655e877ba697b10b499e64bce097f8fa77932bbfe7c7146a7", ".modelhub_state/recovery_active_tasks.jsonl": "53dbe015f17d9e0655e877ba697b10b499e64bce097f8fa77932bbfe7c7146a7",
".modelhub_state/recovery_intents.jsonl": "9855d324e2176cc7b70342bb78fff8ce8daefdb95c8d18998341679c4e9819ed", ".modelhub_state/recovery_intents.jsonl": "1781fba523a806fe20ca532c676162013b179e2f67b2245fd80cee3ab55e3b81",
".modelhub_state/routing_intelligence.json": "29b74d2e65ccbfb42955842460c122572b3e0ba489361a87b71b0de63896d926", ".modelhub_state/routing_intelligence.json": "29b74d2e65ccbfb42955842460c122572b3e0ba489361a87b71b0de63896d926",
".modelhub_state/submission_exclusions.jsonl": "cd362b9c480f30eb8106e264359b19b37b894120db54586e63b0a0d1d41a73b0", ".modelhub_state/submission_exclusions.jsonl": "cd362b9c480f30eb8106e264359b19b37b894120db54586e63b0a0d1d41a73b0",
".modelhub_state/worker_crashes.jsonl": "a494412048eca6dce83e7284303a6b4b56a445c1274ac201ab8723c739a14f23", ".modelhub_state/worker_crashes.jsonl": "a494412048eca6dce83e7284303a6b4b56a445c1274ac201ab8723c739a14f23",
"ledger/submissions.jsonl": "cd8fbff656eafffcbb08597c77d4960cc3e1121cdf75a655d8a2f3345a5b052a", "ledger/submissions.jsonl": "cd8fbff656eafffcbb08597c77d4960cc3e1121cdf75a655d8a2f3345a5b052a",
"outcomes/submissions.jsonl": "bd5e1f0384d929aa62023086993ac1b7f2c27f58735782afc83ca2f2110ac22d" "outcomes/submissions.jsonl": "d3075fd3846f38c550b7b0b703d4b8f9028e0b5b151893d1f63233e9ae743cfc"
}, },
"generation": 9332, "generation": 9333,
"phase": "cycle", "phase": "intent",
"schemaVersion": 1, "schemaVersion": 1,
"updatedAt": "2026-09-18T15:40:28.109992+00:00", "updatedAt": "2026-09-18T15:43:06.414027+00:00",
"writerId": "217f1806b49c401894bed7ce18be9095" "writerId": "217f1806b49c401894bed7ce18be9095"
} }

View File

@@ -24,7 +24,6 @@
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-05T22:44:15.950234+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "modelProfile": {"architectures": [], "configFingerprint": "05c5a771e8ede0f61c61c1696314e8cb295d87d07704eedc68e43f7483931e45", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 356491104, "estimatedRequiredGiB": 1.363, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 1219879026, "modelscopeLicense": "other", "modelscopeParams": 327707521, "modelscopeTags": ["license:other", "library:gguf", "library:pytorch", "task:text-generation", "custom_tag:speculative-decoding", "custom_tag:dspark", "custom_tag:lfm2", "custom_tag:draft-model", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1219879026}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T14:39:36.525018+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4650654", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-05T22:44:15.950234+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "modelProfile": {"architectures": [], "configFingerprint": "05c5a771e8ede0f61c61c1696314e8cb295d87d07704eedc68e43f7483931e45", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 356491104, "estimatedRequiredGiB": 1.363, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 1219879026, "modelscopeLicense": "other", "modelscopeParams": 327707521, "modelscopeTags": ["license:other", "library:gguf", "library:pytorch", "task:text-generation", "custom_tag:speculative-decoding", "custom_tag:dspark", "custom_tag:lfm2", "custom_tag:draft-model", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1219879026}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T14:39:36.525018+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4650654", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113688+00:00", "modelId": "Mia-AiLab/Qwen3.8-27B-EXL3-3.5bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15338408461, "estimatedRequiredGiB": 17.168, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 15362064483, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7669052656, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:exl3", "custom_tag:exllamav3", "custom_tag:quantization", "custom_tag:long-context", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "exl3", "repositoryOnDiskBytes": 15362064483}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:38.391817+00:00", "targetGpu": "Vastai_va16", "taskId": "4651794", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113688+00:00", "modelId": "Mia-AiLab/Qwen3.8-27B-EXL3-3.5bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15338408461, "estimatedRequiredGiB": 17.168, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 15362064483, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7669052656, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:exl3", "custom_tag:exllamav3", "custom_tag:quantization", "custom_tag:long-context", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "exl3", "repositoryOnDiskBytes": 15362064483}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:38.391817+00:00", "targetGpu": "Vastai_va16", "taskId": "4651794", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113645+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:43.655577+00:00", "targetGpu": "Vastai_va16", "taskId": "4651783", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113645+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:43.655577+00:00", "targetGpu": "Vastai_va16", "taskId": "4651783", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:33:16.913577+00:00", "modelId": "ornith-ai/Ornith-1.5-9B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8818103892, "estimatedRequiredGiB": 9.883, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8842796317, "modelscopeLicense": "mit", "modelscopeParams": 6728625904, "modelscopeTags": ["license:mit", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 8842796317}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:30:08.472711+00:00", "targetGpu": "Vastai_va16", "taskId": "4651931", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T08:35:17.503426+00:00", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W8A8-FP8-Channelwise-compressed-tensors", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084014480, "estimatedRequiredGiB": 10.162, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093198689, "modelscopeLicense": null, "modelscopeParams": 8030261248, "modelscopeTags": ["model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9093198689}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T00:30:35.389266+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4657903", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T08:35:17.503426+00:00", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W8A8-FP8-Channelwise-compressed-tensors", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084014480, "estimatedRequiredGiB": 10.162, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093198689, "modelscopeLicense": null, "modelscopeParams": 8030261248, "modelscopeTags": ["model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9093198689}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T00:30:35.389266+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4657903", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T08:35:17.503433+00:00", "modelId": "ornith-ai/Ornith-1.5-35B-A3B-FP8", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 39365175520, "estimatedRequiredGiB": 44.029, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 39396646235, "modelscopeLicense": "mit", "modelscopeParams": 35951822704, "modelscopeTags": ["license:mit", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 39396646235}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T00:30:35.385964+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4657900", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T08:35:17.503433+00:00", "modelId": "ornith-ai/Ornith-1.5-35B-A3B-FP8", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 39365175520, "estimatedRequiredGiB": 44.029, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 39396646235, "modelscopeLicense": "mit", "modelscopeParams": 35951822704, "modelscopeTags": ["license:mit", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 39396646235}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T00:30:35.385964+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4657900", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T08:35:17.503369+00:00", "modelId": "zstack/qwen3-4b-fake", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 25165824, "estimatedRequiredGiB": 0.028, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 25169485, "modelscopeLicense": null, "modelscopeParams": 25165435, "modelscopeTags": ["model_type:qwen3", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 25169485}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T00:30:35.392602+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4657901", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T08:35:17.503369+00:00", "modelId": "zstack/qwen3-4b-fake", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 25165824, "estimatedRequiredGiB": 0.028, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 25169485, "modelscopeLicense": null, "modelscopeParams": 25165435, "modelscopeTags": ["model_type:qwen3", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 25169485}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T00:30:35.392602+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4657901", "taskType": "text-generation", "verifyResult": null}
@@ -129,7 +128,6 @@
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-09T05:32:40.605323+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T21:29:34.930719+00:00", "targetGpu": "Vastai_va16", "taskId": "4727973", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-09T05:32:40.605323+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T21:29:34.930719+00:00", "targetGpu": "Vastai_va16", "taskId": "4727973", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-09T05:49:44.122625+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T21:48:30.702694+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4728207", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-09T05:49:44.122625+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T21:48:30.702694+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4728207", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-09T06:06:56.823871+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T22:05:28.693168+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4728506", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-09T06:06:56.823871+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T22:05:28.693168+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4728506", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-09T06:23:50.502374+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T22:23:01.734631+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4728907", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-09T06:27:15.307637+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T22:24:50.419230+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4728926", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-09T06:27:15.307637+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T22:24:50.419230+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4728926", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-09T06:43:52.702366+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T22:42:29.518873+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4729173", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-09T06:43:52.702366+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T22:42:29.518873+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4729173", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-09T06:50:35.321512+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T22:49:56.925801+00:00", "targetGpu": "Biren_166m", "taskId": "4729275", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-09T06:50:35.321512+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794148, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794148}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T22:49:56.925801+00:00", "targetGpu": "Biren_166m", "taskId": "4729275", "taskType": "text-generation", "verifyResult": null}