state: generation 11291 (intent)
This commit is contained in:
@@ -765,6 +765,25 @@
|
|||||||
"targetGpu": "Cambricon_mlu-370-x4",
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
"taskType": "text-generation"
|
"taskType": "text-generation"
|
||||||
},
|
},
|
||||||
|
"cambricon_mlu-370-x4|vllm|text-generation|model_type:qwen3_5_text": {
|
||||||
|
"architectureSignature": "model_type:qwen3_5_text",
|
||||||
|
"architectures": [],
|
||||||
|
"evidenceCount": 1,
|
||||||
|
"expiresAt": "2026-10-20T02:52:52.089383+00:00",
|
||||||
|
"framework": "vllm",
|
||||||
|
"latestFailureAt": "2026-09-20T02:52:52.089383+00:00",
|
||||||
|
"latestSuccessfulAt": null,
|
||||||
|
"matchType": "model_type",
|
||||||
|
"modelType": "qwen3_5_text",
|
||||||
|
"sourceModelIds": [
|
||||||
|
"TokenRhythm/NeoHorse-1-9B"
|
||||||
|
],
|
||||||
|
"sourceTaskIds": [
|
||||||
|
"4969078"
|
||||||
|
],
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation"
|
||||||
|
},
|
||||||
"cambricon_mlu-370-x8|vllm-customized|text-generation|model_type:qwen3_5": {
|
"cambricon_mlu-370-x8|vllm-customized|text-generation|model_type:qwen3_5": {
|
||||||
"architectureSignature": "model_type:qwen3_5",
|
"architectureSignature": "model_type:qwen3_5",
|
||||||
"architectures": [],
|
"architectures": [],
|
||||||
@@ -2599,16 +2618,16 @@
|
|||||||
"taskType": "text-generation"
|
"taskType": "text-generation"
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"generatedAt": "2026-09-21T07:13:18.333709+00:00",
|
"generatedAt": "2026-09-21T07:17:51.359324+00:00",
|
||||||
"summary": {
|
"summary": {
|
||||||
"activeBlockCount": 130,
|
"activeBlockCount": 131,
|
||||||
"byGpuFramework": {
|
"byGpuFramework": {
|
||||||
"Ascend_910-b3|vllm": 13,
|
"Ascend_910-b3|vllm": 13,
|
||||||
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
||||||
"Ascend_910-b4|vllm": 12,
|
"Ascend_910-b4|vllm": 12,
|
||||||
"Ascend_910-b4|vllm_tokenizer_patch": 5,
|
"Ascend_910-b4|vllm_tokenizer_patch": 5,
|
||||||
"Biren_166m|vllm": 3,
|
"Biren_166m|vllm": 3,
|
||||||
"Cambricon_mlu-370-x4|vllm": 1,
|
"Cambricon_mlu-370-x4|vllm": 2,
|
||||||
"Cambricon_mlu-370-x8|vllm": 12,
|
"Cambricon_mlu-370-x8|vllm": 12,
|
||||||
"Cambricon_mlu-370-x8|vllm-customized": 1,
|
"Cambricon_mlu-370-x8|vllm-customized": 1,
|
||||||
"Cambricon_mlu-370-x8|vllm-mlu": 5,
|
"Cambricon_mlu-370-x8|vllm-mlu": 5,
|
||||||
|
|||||||
@@ -425,7 +425,7 @@
|
|||||||
}
|
}
|
||||||
},
|
},
|
||||||
"frameworkUpdatedAt": null,
|
"frameworkUpdatedAt": null,
|
||||||
"generatedAt": "2026-09-21T07:16:47.460864+00:00",
|
"generatedAt": "2026-09-21T07:18:01.742848+00:00",
|
||||||
"gpuStats": {
|
"gpuStats": {
|
||||||
"Ascend_910-b3": {
|
"Ascend_910-b3": {
|
||||||
"available": true,
|
"available": true,
|
||||||
|
|||||||
File diff suppressed because it is too large
Load Diff
@@ -1,6 +1,6 @@
|
|||||||
{
|
{
|
||||||
"generatedAt": "2026-09-21T07:09:45.936241+00:00",
|
"generatedAt": "2026-09-21T07:18:01.667464+00:00",
|
||||||
"lastSyncTime": "2026-09-21T07:09:45.554098+00:00",
|
"lastSyncTime": "2026-09-21T07:18:01.568496+00:00",
|
||||||
"recentLimit": 300,
|
"recentLimit": 300,
|
||||||
"report": {
|
"report": {
|
||||||
"architectureCompatibilityBlocks": {
|
"architectureCompatibilityBlocks": {
|
||||||
@@ -769,6 +769,25 @@
|
|||||||
"targetGpu": "Cambricon_mlu-370-x4",
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
"taskType": "text-generation"
|
"taskType": "text-generation"
|
||||||
},
|
},
|
||||||
|
"cambricon_mlu-370-x4|vllm|text-generation|model_type:qwen3_5_text": {
|
||||||
|
"architectureSignature": "model_type:qwen3_5_text",
|
||||||
|
"architectures": [],
|
||||||
|
"evidenceCount": 1,
|
||||||
|
"expiresAt": "2026-10-20T02:52:52.089383+00:00",
|
||||||
|
"framework": "vllm",
|
||||||
|
"latestFailureAt": "2026-09-20T02:52:52.089383+00:00",
|
||||||
|
"latestSuccessfulAt": null,
|
||||||
|
"matchType": "model_type",
|
||||||
|
"modelType": "qwen3_5_text",
|
||||||
|
"sourceModelIds": [
|
||||||
|
"TokenRhythm/NeoHorse-1-9B"
|
||||||
|
],
|
||||||
|
"sourceTaskIds": [
|
||||||
|
"4969078"
|
||||||
|
],
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation"
|
||||||
|
},
|
||||||
"cambricon_mlu-370-x8|vllm-customized|text-generation|model_type:qwen3_5": {
|
"cambricon_mlu-370-x8|vllm-customized|text-generation|model_type:qwen3_5": {
|
||||||
"architectureSignature": "model_type:qwen3_5",
|
"architectureSignature": "model_type:qwen3_5",
|
||||||
"architectures": [],
|
"architectures": [],
|
||||||
@@ -2604,14 +2623,14 @@
|
|||||||
}
|
}
|
||||||
},
|
},
|
||||||
"architectureCompatibilitySummary": {
|
"architectureCompatibilitySummary": {
|
||||||
"activeBlockCount": 130,
|
"activeBlockCount": 131,
|
||||||
"byGpuFramework": {
|
"byGpuFramework": {
|
||||||
"Ascend_910-b3|vllm": 13,
|
"Ascend_910-b3|vllm": 13,
|
||||||
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
||||||
"Ascend_910-b4|vllm": 12,
|
"Ascend_910-b4|vllm": 12,
|
||||||
"Ascend_910-b4|vllm_tokenizer_patch": 5,
|
"Ascend_910-b4|vllm_tokenizer_patch": 5,
|
||||||
"Biren_166m|vllm": 3,
|
"Biren_166m|vllm": 3,
|
||||||
"Cambricon_mlu-370-x4|vllm": 1,
|
"Cambricon_mlu-370-x4|vllm": 2,
|
||||||
"Cambricon_mlu-370-x8|vllm": 12,
|
"Cambricon_mlu-370-x8|vllm": 12,
|
||||||
"Cambricon_mlu-370-x8|vllm-customized": 1,
|
"Cambricon_mlu-370-x8|vllm-customized": 1,
|
||||||
"Cambricon_mlu-370-x8|vllm-mlu": 5,
|
"Cambricon_mlu-370-x8|vllm-mlu": 5,
|
||||||
@@ -3135,32 +3154,33 @@
|
|||||||
"decisionSuccessRate": 0.5333,
|
"decisionSuccessRate": 0.5333,
|
||||||
"decisionTotal": 15,
|
"decisionTotal": 15,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 14,
|
"ambiguous_runtime": 15,
|
||||||
"framework_architecture_unsupported": 6,
|
"framework_architecture_unsupported": 6,
|
||||||
"model_load": 1
|
"model_load": 1
|
||||||
},
|
},
|
||||||
"failureCount": 21,
|
"failureCount": 22,
|
||||||
"failureRate": 0.7241,
|
"failureRate": 0.7333,
|
||||||
"framework": "vllm-mlu",
|
"framework": "vllm-mlu",
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 0,
|
"platformFailureCount": 0,
|
||||||
"successCount": 8,
|
"successCount": 8,
|
||||||
"successRate": 0.2759,
|
"successRate": 0.2667,
|
||||||
"targetGpu": "Cambricon_mlu-370-x4",
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 29,
|
"total": 30,
|
||||||
"unresolvedFailureCount": 14
|
"unresolvedFailureCount": 15
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x4|vllm|text-generation": {
|
"Cambricon_mlu-370-x4|vllm|text-generation": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 2,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 1,
|
"decisionTotal": 2,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"framework_architecture_unsupported": 1
|
"ambiguous_runtime": 1,
|
||||||
|
"framework_architecture_unsupported": 2
|
||||||
},
|
},
|
||||||
"failureCount": 1,
|
"failureCount": 3,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm",
|
"framework": "vllm",
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
@@ -3170,8 +3190,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Cambricon_mlu-370-x4",
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 1,
|
"total": 3,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 1
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8|unknown|feature_emb": {
|
"Cambricon_mlu-370-x8|unknown|feature_emb": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
@@ -5219,17 +5239,17 @@
|
|||||||
"unresolvedFailureCount": 6442
|
"unresolvedFailureCount": 6442
|
||||||
},
|
},
|
||||||
"vllm": {
|
"vllm": {
|
||||||
"attributableFailureCount": 3516,
|
"attributableFailureCount": 3517,
|
||||||
"decisionFailureRate": 0.9764,
|
"decisionFailureRate": 0.9764,
|
||||||
"decisionSuccessRate": 0.0236,
|
"decisionSuccessRate": 0.0236,
|
||||||
"decisionTotal": 3601,
|
"decisionTotal": 3602,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 1470,
|
"ambiguous_runtime": 1471,
|
||||||
"architecture_compatibility": 112,
|
"architecture_compatibility": 112,
|
||||||
"attention_backend": 1,
|
"attention_backend": 1,
|
||||||
"backend_operator": 84,
|
"backend_operator": 84,
|
||||||
"context_length": 161,
|
"context_length": 161,
|
||||||
"framework_architecture_unsupported": 1288,
|
"framework_architecture_unsupported": 1289,
|
||||||
"memory_capacity": 720,
|
"memory_capacity": 720,
|
||||||
"model_load": 206,
|
"model_load": 206,
|
||||||
"platform_infrastructure": 865,
|
"platform_infrastructure": 865,
|
||||||
@@ -5238,15 +5258,15 @@
|
|||||||
"tokenizer_compatibility": 410,
|
"tokenizer_compatibility": 410,
|
||||||
"参数/模板问题": 41
|
"参数/模板问题": 41
|
||||||
},
|
},
|
||||||
"failureCount": 5892,
|
"failureCount": 5894,
|
||||||
"failureRate": 0.9858,
|
"failureRate": 0.9858,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 865,
|
"platformFailureCount": 865,
|
||||||
"successCount": 85,
|
"successCount": 85,
|
||||||
"successRate": 0.0142,
|
"successRate": 0.0142,
|
||||||
"total": 5977,
|
"total": 5979,
|
||||||
"unresolvedFailureCount": 1511
|
"unresolvedFailureCount": 1512
|
||||||
},
|
},
|
||||||
"vllm-customized": {
|
"vllm-customized": {
|
||||||
"attributableFailureCount": 3,
|
"attributableFailureCount": 3,
|
||||||
@@ -5274,21 +5294,21 @@
|
|||||||
"decisionSuccessRate": 0.5,
|
"decisionSuccessRate": 0.5,
|
||||||
"decisionTotal": 40,
|
"decisionTotal": 40,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 50,
|
"ambiguous_runtime": 51,
|
||||||
"framework_architecture_unsupported": 17,
|
"framework_architecture_unsupported": 17,
|
||||||
"model_load": 2,
|
"model_load": 2,
|
||||||
"tokenizer_compatibility": 1,
|
"tokenizer_compatibility": 1,
|
||||||
"参数/模板问题": 3
|
"参数/模板问题": 3
|
||||||
},
|
},
|
||||||
"failureCount": 73,
|
"failureCount": 74,
|
||||||
"failureRate": 0.7849,
|
"failureRate": 0.7872,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 0,
|
"platformFailureCount": 0,
|
||||||
"successCount": 20,
|
"successCount": 20,
|
||||||
"successRate": 0.2151,
|
"successRate": 0.2128,
|
||||||
"total": 93,
|
"total": 94,
|
||||||
"unresolvedFailureCount": 53
|
"unresolvedFailureCount": 54
|
||||||
},
|
},
|
||||||
"vllm-patch-tokenizer": {
|
"vllm-patch-tokenizer": {
|
||||||
"attributableFailureCount": 27,
|
"attributableFailureCount": 27,
|
||||||
@@ -5387,7 +5407,7 @@
|
|||||||
"unresolvedFailureCount": 67
|
"unresolvedFailureCount": 67
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"generatedAt": "2026-09-21T07:09:45.924060+00:00",
|
"generatedAt": "2026-09-21T07:18:01.620340+00:00",
|
||||||
"gpuSummaries": {
|
"gpuSummaries": {
|
||||||
"Ascend_910-b3": {
|
"Ascend_910-b3": {
|
||||||
"attributableFailureCount": 91,
|
"attributableFailureCount": 91,
|
||||||
@@ -5474,15 +5494,15 @@
|
|||||||
"unresolvedFailureCount": 442
|
"unresolvedFailureCount": 442
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x4": {
|
"Cambricon_mlu-370-x4": {
|
||||||
"attributableFailureCount": 713,
|
"attributableFailureCount": 714,
|
||||||
"decisionFailureRate": 0.9071,
|
"decisionFailureRate": 0.9072,
|
||||||
"decisionSuccessRate": 0.0929,
|
"decisionSuccessRate": 0.0928,
|
||||||
"decisionTotal": 786,
|
"decisionTotal": 787,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 429,
|
"ambiguous_runtime": 431,
|
||||||
"architecture_compatibility": 53,
|
"architecture_compatibility": 53,
|
||||||
"context_length": 48,
|
"context_length": 48,
|
||||||
"framework_architecture_unsupported": 269,
|
"framework_architecture_unsupported": 270,
|
||||||
"memory_capacity": 218,
|
"memory_capacity": 218,
|
||||||
"model_load": 5,
|
"model_load": 5,
|
||||||
"platform_infrastructure": 15,
|
"platform_infrastructure": 15,
|
||||||
@@ -5492,15 +5512,15 @@
|
|||||||
"日志缺失": 13,
|
"日志缺失": 13,
|
||||||
"验证失败": 39
|
"验证失败": 39
|
||||||
},
|
},
|
||||||
"failureCount": 1493,
|
"failureCount": 1496,
|
||||||
"failureRate": 0.9534,
|
"failureRate": 0.9535,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 15,
|
"platformFailureCount": 15,
|
||||||
"successCount": 73,
|
"successCount": 73,
|
||||||
"successRate": 0.0466,
|
"successRate": 0.0465,
|
||||||
"total": 1566,
|
"total": 1569,
|
||||||
"unresolvedFailureCount": 765
|
"unresolvedFailureCount": 767
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8": {
|
"Cambricon_mlu-370-x8": {
|
||||||
"attributableFailureCount": 44,
|
"attributableFailureCount": 44,
|
||||||
@@ -5818,7 +5838,7 @@
|
|||||||
"hygon_k100-ai": 64.0
|
"hygon_k100-ai": 64.0
|
||||||
},
|
},
|
||||||
"pendingRecords": 0,
|
"pendingRecords": 0,
|
||||||
"policyCancelledRecords": 181,
|
"policyCancelledRecords": 182,
|
||||||
"profileCombinationStats": {
|
"profileCombinationStats": {
|
||||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|bailing_hybrid|none": {
|
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|bailing_hybrid|none": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 1,
|
||||||
@@ -8119,6 +8139,52 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 1
|
"unresolvedFailureCount": 1
|
||||||
},
|
},
|
||||||
|
"Cambricon_mlu-370-x4|vllm-mlu|text-generation|stablelm|none": {
|
||||||
|
"attributableFailureCount": 0,
|
||||||
|
"decisionFailureRate": 0.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 0,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"ambiguous_runtime": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm-mlu",
|
||||||
|
"modelType": "stablelm",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "none",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 1
|
||||||
|
},
|
||||||
|
"Cambricon_mlu-370-x4|vllm|text-generation|llama|awq": {
|
||||||
|
"attributableFailureCount": 0,
|
||||||
|
"decisionFailureRate": 0.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 0,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"ambiguous_runtime": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"modelType": "llama",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "awq",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 1
|
||||||
|
},
|
||||||
"Cambricon_mlu-370-x4|vllm|text-generation|nemotron_h|compressed-tensors": {
|
"Cambricon_mlu-370-x4|vllm|text-generation|nemotron_h|compressed-tensors": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 1,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
@@ -8142,6 +8208,29 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
|
"Cambricon_mlu-370-x4|vllm|text-generation|qwen3_5_text|none": {
|
||||||
|
"attributableFailureCount": 1,
|
||||||
|
"decisionFailureRate": 1.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 1,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"framework_architecture_unsupported": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"modelType": "qwen3_5_text",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "none",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 0
|
||||||
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm-customized|text-generation|granite|none": {
|
"Cambricon_mlu-370-x8|vllm-customized|text-generation|granite|none": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
"decisionFailureRate": 0.0,
|
"decisionFailureRate": 0.0,
|
||||||
@@ -14768,19 +14857,19 @@
|
|||||||
"unresolvedFailureCount": 11
|
"unresolvedFailureCount": 11
|
||||||
},
|
},
|
||||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation": {
|
"Ascend_910-b3|vllm_tokenizer_patch|text-generation": {
|
||||||
"attributableFailureCount": 4,
|
"attributableFailureCount": 3,
|
||||||
"consecutiveFailures": 4,
|
"consecutiveFailures": 3,
|
||||||
"consecutivePlatformFailures": 0,
|
"consecutivePlatformFailures": 0,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 4,
|
"decisionTotal": 3,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 4,
|
"ambiguous_runtime": 4,
|
||||||
"context_length": 1,
|
"context_length": 1,
|
||||||
"framework_architecture_unsupported": 3,
|
"framework_architecture_unsupported": 2,
|
||||||
"参数/模板问题": 5
|
"参数/模板问题": 5
|
||||||
},
|
},
|
||||||
"failureCount": 13,
|
"failureCount": 12,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm_tokenizer_patch",
|
"framework": "vllm_tokenizer_patch",
|
||||||
"lastPlatformFailureAt": null,
|
"lastPlatformFailureAt": null,
|
||||||
@@ -14792,7 +14881,7 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Ascend_910-b3",
|
"targetGpu": "Ascend_910-b3",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 13,
|
"total": 12,
|
||||||
"unresolvedFailureCount": 9
|
"unresolvedFailureCount": 9
|
||||||
},
|
},
|
||||||
"Ascend_910-b3|vllm|text-generation": {
|
"Ascend_910-b3|vllm|text-generation": {
|
||||||
@@ -14951,20 +15040,21 @@
|
|||||||
"unresolvedFailureCount": 8
|
"unresolvedFailureCount": 8
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x4|vllm|text-generation": {
|
"Cambricon_mlu-370-x4|vllm|text-generation": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 2,
|
||||||
"consecutiveFailures": 1,
|
"consecutiveFailures": 2,
|
||||||
"consecutivePlatformFailures": 0,
|
"consecutivePlatformFailures": 0,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 1,
|
"decisionTotal": 2,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"framework_architecture_unsupported": 1
|
"ambiguous_runtime": 1,
|
||||||
|
"framework_architecture_unsupported": 2
|
||||||
},
|
},
|
||||||
"failureCount": 1,
|
"failureCount": 3,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm",
|
"framework": "vllm",
|
||||||
"lastPlatformFailureAt": null,
|
"lastPlatformFailureAt": null,
|
||||||
"lastTerminalAt": "2026-09-21T01:35:45.556177+00:00",
|
"lastTerminalAt": "2026-09-21T07:17:50.956808+00:00",
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 0,
|
"platformFailureCount": 0,
|
||||||
@@ -14972,8 +15062,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Cambricon_mlu-370-x4",
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 1,
|
"total": 3,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 1
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8|unknown|text-generation": {
|
"Cambricon_mlu-370-x8|unknown|text-generation": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
@@ -15342,31 +15432,6 @@
|
|||||||
"total": 14,
|
"total": 14,
|
||||||
"unresolvedFailureCount": 10
|
"unresolvedFailureCount": 10
|
||||||
},
|
},
|
||||||
"Kunlunxin_p-800|vllm|text-generation": {
|
|
||||||
"attributableFailureCount": 1,
|
|
||||||
"consecutiveFailures": 1,
|
|
||||||
"consecutivePlatformFailures": 0,
|
|
||||||
"decisionFailureRate": 1.0,
|
|
||||||
"decisionSuccessRate": 0.0,
|
|
||||||
"decisionTotal": 1,
|
|
||||||
"failureBreakdown": {
|
|
||||||
"framework_architecture_unsupported": 1
|
|
||||||
},
|
|
||||||
"failureCount": 1,
|
|
||||||
"failureRate": 1.0,
|
|
||||||
"framework": "vllm",
|
|
||||||
"lastPlatformFailureAt": null,
|
|
||||||
"lastTerminalAt": "2026-09-20T20:49:19.683316+00:00",
|
|
||||||
"pendingCount": 0,
|
|
||||||
"pendingRate": 0.0,
|
|
||||||
"platformFailureCount": 0,
|
|
||||||
"successCount": 0,
|
|
||||||
"successRate": 0.0,
|
|
||||||
"targetGpu": "Kunlunxin_p-800",
|
|
||||||
"taskType": "text-generation",
|
|
||||||
"total": 1,
|
|
||||||
"unresolvedFailureCount": 0
|
|
||||||
},
|
|
||||||
"MetaX_c-500|unknown|text-generation": {
|
"MetaX_c-500|unknown|text-generation": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
"consecutiveFailures": 0,
|
"consecutiveFailures": 0,
|
||||||
@@ -15729,31 +15794,6 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|qwen3_5|compressed-tensors": {
|
|
||||||
"attributableFailureCount": 1,
|
|
||||||
"consecutiveFailures": 1,
|
|
||||||
"decisionFailureRate": 1.0,
|
|
||||||
"decisionSuccessRate": 0.0,
|
|
||||||
"decisionTotal": 1,
|
|
||||||
"failureBreakdown": {
|
|
||||||
"framework_architecture_unsupported": 1
|
|
||||||
},
|
|
||||||
"failureCount": 1,
|
|
||||||
"failureRate": 1.0,
|
|
||||||
"framework": "vllm_tokenizer_patch",
|
|
||||||
"lastTerminalAt": "2026-09-20T22:27:36.960854+00:00",
|
|
||||||
"modelType": "qwen3_5",
|
|
||||||
"pendingCount": 0,
|
|
||||||
"pendingRate": 0.0,
|
|
||||||
"platformFailureCount": 0,
|
|
||||||
"quantizationMethod": "compressed-tensors",
|
|
||||||
"successCount": 0,
|
|
||||||
"successRate": 0.0,
|
|
||||||
"targetGpu": "Ascend_910-b3",
|
|
||||||
"taskType": "text-generation",
|
|
||||||
"total": 1,
|
|
||||||
"unresolvedFailureCount": 0
|
|
||||||
},
|
|
||||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|stablelm|none": {
|
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|stablelm|none": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
"consecutiveFailures": 0,
|
"consecutiveFailures": 0,
|
||||||
@@ -16206,6 +16246,31 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
|
"Cambricon_mlu-370-x4|vllm|text-generation|llama|awq": {
|
||||||
|
"attributableFailureCount": 0,
|
||||||
|
"consecutiveFailures": 0,
|
||||||
|
"decisionFailureRate": 0.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 0,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"ambiguous_runtime": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"lastTerminalAt": "2026-09-21T07:17:50.956808+00:00",
|
||||||
|
"modelType": "llama",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "awq",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 1
|
||||||
|
},
|
||||||
"Cambricon_mlu-370-x4|vllm|text-generation|nemotron_h|compressed-tensors": {
|
"Cambricon_mlu-370-x4|vllm|text-generation|nemotron_h|compressed-tensors": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 1,
|
||||||
"consecutiveFailures": 1,
|
"consecutiveFailures": 1,
|
||||||
@@ -16231,6 +16296,31 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
|
"Cambricon_mlu-370-x4|vllm|text-generation|qwen3_5_text|none": {
|
||||||
|
"attributableFailureCount": 1,
|
||||||
|
"consecutiveFailures": 1,
|
||||||
|
"decisionFailureRate": 1.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 1,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"framework_architecture_unsupported": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"lastTerminalAt": "2026-09-21T07:17:50.956847+00:00",
|
||||||
|
"modelType": "qwen3_5_text",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "none",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 0
|
||||||
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm-customized|text-generation|granite|none": {
|
"Cambricon_mlu-370-x8|vllm-customized|text-generation|granite|none": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
"consecutiveFailures": 0,
|
"consecutiveFailures": 0,
|
||||||
@@ -18810,31 +18900,6 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
"Kunlunxin_p-800|vllm|text-generation|deepseek_v4|none": {
|
|
||||||
"attributableFailureCount": 1,
|
|
||||||
"consecutiveFailures": 1,
|
|
||||||
"decisionFailureRate": 1.0,
|
|
||||||
"decisionSuccessRate": 0.0,
|
|
||||||
"decisionTotal": 1,
|
|
||||||
"failureBreakdown": {
|
|
||||||
"framework_architecture_unsupported": 1
|
|
||||||
},
|
|
||||||
"failureCount": 1,
|
|
||||||
"failureRate": 1.0,
|
|
||||||
"framework": "vllm",
|
|
||||||
"lastTerminalAt": "2026-09-20T20:49:19.683316+00:00",
|
|
||||||
"modelType": "deepseek_v4",
|
|
||||||
"pendingCount": 0,
|
|
||||||
"pendingRate": 0.0,
|
|
||||||
"platformFailureCount": 0,
|
|
||||||
"quantizationMethod": "none",
|
|
||||||
"successCount": 0,
|
|
||||||
"successRate": 0.0,
|
|
||||||
"targetGpu": "Kunlunxin_p-800",
|
|
||||||
"taskType": "text-generation",
|
|
||||||
"total": 1,
|
|
||||||
"unresolvedFailureCount": 0
|
|
||||||
},
|
|
||||||
"MetaX_c-500|vllm|text-generation|gemma3|compressed-tensors": {
|
"MetaX_c-500|vllm|text-generation|gemma3|compressed-tensors": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 1,
|
||||||
"consecutiveFailures": 1,
|
"consecutiveFailures": 1,
|
||||||
@@ -22528,6 +22593,54 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 1
|
"unresolvedFailureCount": 1
|
||||||
},
|
},
|
||||||
|
"Cambricon_mlu-370-x4|vllm-mlu|text-generation|stablelm|none|31": {
|
||||||
|
"attributableFailureCount": 0,
|
||||||
|
"decisionFailureRate": 0.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 0,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"ambiguous_runtime": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm-mlu",
|
||||||
|
"loadSizeLog2Bucket": 31,
|
||||||
|
"modelType": "stablelm",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "none",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 1
|
||||||
|
},
|
||||||
|
"Cambricon_mlu-370-x4|vllm|text-generation|llama|awq|32": {
|
||||||
|
"attributableFailureCount": 0,
|
||||||
|
"decisionFailureRate": 0.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 0,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"ambiguous_runtime": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"loadSizeLog2Bucket": 32,
|
||||||
|
"modelType": "llama",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "awq",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 1
|
||||||
|
},
|
||||||
"Cambricon_mlu-370-x4|vllm|text-generation|nemotron_h|compressed-tensors|34": {
|
"Cambricon_mlu-370-x4|vllm|text-generation|nemotron_h|compressed-tensors|34": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 1,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
@@ -22552,6 +22665,30 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
|
"Cambricon_mlu-370-x4|vllm|text-generation|qwen3_5_text|none|34": {
|
||||||
|
"attributableFailureCount": 1,
|
||||||
|
"decisionFailureRate": 1.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 1,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"framework_architecture_unsupported": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"loadSizeLog2Bucket": 34,
|
||||||
|
"modelType": "qwen3_5_text",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "none",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x4",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 0
|
||||||
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm-customized|text-generation|granite|none|33": {
|
"Cambricon_mlu-370-x8|vllm-customized|text-generation|granite|none|33": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
"decisionFailureRate": 0.0,
|
"decisionFailureRate": 0.0,
|
||||||
@@ -32126,20 +32263,20 @@
|
|||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"terminalRecords": 16164,
|
"terminalRecords": 16167,
|
||||||
"totalRecords": 16345,
|
"totalRecords": 16349,
|
||||||
"totals": {
|
"totals": {
|
||||||
"attributableFailureCount": 5854,
|
"attributableFailureCount": 5855,
|
||||||
"decisionFailureRate": 0.8614,
|
"decisionFailureRate": 0.8614,
|
||||||
"decisionSuccessRate": 0.1386,
|
"decisionSuccessRate": 0.1386,
|
||||||
"decisionTotal": 6796,
|
"decisionTotal": 6797,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 3882,
|
"ambiguous_runtime": 3884,
|
||||||
"architecture_compatibility": 212,
|
"architecture_compatibility": 212,
|
||||||
"attention_backend": 1,
|
"attention_backend": 1,
|
||||||
"backend_operator": 102,
|
"backend_operator": 102,
|
||||||
"context_length": 319,
|
"context_length": 319,
|
||||||
"framework_architecture_unsupported": 2051,
|
"framework_architecture_unsupported": 2052,
|
||||||
"memory_capacity": 1197,
|
"memory_capacity": 1197,
|
||||||
"model_load": 494,
|
"model_load": 494,
|
||||||
"platform_infrastructure": 923,
|
"platform_infrastructure": 923,
|
||||||
@@ -32150,15 +32287,15 @@
|
|||||||
"日志缺失": 719,
|
"日志缺失": 719,
|
||||||
"验证失败": 673
|
"验证失败": 673
|
||||||
},
|
},
|
||||||
"failureCount": 15222,
|
"failureCount": 15225,
|
||||||
"failureRate": 0.9417,
|
"failureRate": 0.9417,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 923,
|
"platformFailureCount": 923,
|
||||||
"successCount": 942,
|
"successCount": 942,
|
||||||
"successRate": 0.0583,
|
"successRate": 0.0583,
|
||||||
"total": 16164,
|
"total": 16167,
|
||||||
"unresolvedFailureCount": 8445
|
"unresolvedFailureCount": 8447
|
||||||
},
|
},
|
||||||
"warnings": [
|
"warnings": [
|
||||||
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
@@ -32213,6 +32350,6 @@
|
|||||||
]
|
]
|
||||||
},
|
},
|
||||||
"storageMode": "decision_state_only",
|
"storageMode": "decision_state_only",
|
||||||
"summarizedRecords": 16345,
|
"summarizedRecords": 16349,
|
||||||
"version": 1
|
"version": 1
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -20,7 +20,7 @@
|
|||||||
"cleanupDisabled": true,
|
"cleanupDisabled": true,
|
||||||
"reason": "admission_only"
|
"reason": "admission_only"
|
||||||
},
|
},
|
||||||
"architectureBlockCount": 130,
|
"architectureBlockCount": 131,
|
||||||
"architectureFrameworkCatalog": {
|
"architectureFrameworkCatalog": {
|
||||||
"ascend_910-b3|text-generation": [
|
"ascend_910-b3|text-generation": [
|
||||||
"llamacpp",
|
"llamacpp",
|
||||||
@@ -38,11 +38,6 @@
|
|||||||
"vllm_fix_tokenizer",
|
"vllm_fix_tokenizer",
|
||||||
"vllm_tokenizer_patch"
|
"vllm_tokenizer_patch"
|
||||||
],
|
],
|
||||||
"cambricon_mlu-370-x4|text-generation": [
|
|
||||||
"vllm",
|
|
||||||
"vllm-customized",
|
|
||||||
"vllm-mlu"
|
|
||||||
],
|
|
||||||
"cambricon_mlu-370-x8|text-generation": [
|
"cambricon_mlu-370-x8|text-generation": [
|
||||||
"vllm",
|
"vllm",
|
||||||
"vllm-customized",
|
"vllm-customized",
|
||||||
@@ -53,12 +48,6 @@
|
|||||||
"vllm",
|
"vllm",
|
||||||
"vllm-patch-tokenizer"
|
"vllm-patch-tokenizer"
|
||||||
],
|
],
|
||||||
"iluvatar_bi-100|text-generation": [
|
|
||||||
"transformers",
|
|
||||||
"vllm",
|
|
||||||
"vllm-patch-tokenizer",
|
|
||||||
"vllm_fix_tokenizer"
|
|
||||||
],
|
|
||||||
"iluvatar_bi-150|text-generation": [
|
"iluvatar_bi-150|text-generation": [
|
||||||
"llamacpp",
|
"llamacpp",
|
||||||
"transformers",
|
"transformers",
|
||||||
@@ -67,19 +56,11 @@
|
|||||||
"vllm_fix_tokenizer",
|
"vllm_fix_tokenizer",
|
||||||
"vllm_tokenizer_patch"
|
"vllm_tokenizer_patch"
|
||||||
],
|
],
|
||||||
"iluvatar_mrv-100|text-generation": [
|
|
||||||
"transformers",
|
|
||||||
"vllm"
|
|
||||||
],
|
|
||||||
"kunlunxin_p-800|text-generation": [
|
"kunlunxin_p-800|text-generation": [
|
||||||
"vllm",
|
"vllm",
|
||||||
"vllm_fix_tokenizer",
|
"vllm_fix_tokenizer",
|
||||||
"vllm_tokenizer_patch"
|
"vllm_tokenizer_patch"
|
||||||
],
|
],
|
||||||
"metax_c-500|text-generation": [
|
|
||||||
"vllm",
|
|
||||||
"vllm_tokenizer_patch"
|
|
||||||
],
|
|
||||||
"mthreads_s4000|text-generation": [
|
"mthreads_s4000|text-generation": [
|
||||||
"llamacpp",
|
"llamacpp",
|
||||||
"vllm",
|
"vllm",
|
||||||
@@ -89,15 +70,33 @@
|
|||||||
"sglang",
|
"sglang",
|
||||||
"vllm",
|
"vllm",
|
||||||
"vllm_fix_tokenizer"
|
"vllm_fix_tokenizer"
|
||||||
],
|
|
||||||
"vastai_va16|text-generation": [
|
|
||||||
"vllm",
|
|
||||||
"vllm_fix_tokenizer"
|
|
||||||
]
|
]
|
||||||
},
|
},
|
||||||
"architectureFrameworkCatalogErrors": {},
|
"architectureFrameworkCatalogErrors": {},
|
||||||
"architectureIncompatibleCount": 0,
|
"architectureIncompatibleCount": 1,
|
||||||
"architectureIncompatibleTasks": [],
|
"architectureIncompatibleTasks": [
|
||||||
|
{
|
||||||
|
"accountIndex": 10,
|
||||||
|
"architectureBlockEvidenceCount": 1,
|
||||||
|
"architectureBlockExpiresAt": "2026-10-20T02:52:52.089383+00:00",
|
||||||
|
"architectureMatchScope": "exact_framework",
|
||||||
|
"architectureMatchType": "model_type",
|
||||||
|
"architectureSignature": "model_type:qwen3_5_text",
|
||||||
|
"architectureSignatures": [
|
||||||
|
"model_type:qwen3_5_text"
|
||||||
|
],
|
||||||
|
"evaluatedFrameworks": [
|
||||||
|
"vllm"
|
||||||
|
],
|
||||||
|
"framework": "vllm",
|
||||||
|
"gpuType": "Cambricon_mlu-370-x4",
|
||||||
|
"modelId": "TokenRhythm/NeoHorse-1-4B",
|
||||||
|
"reason": "known_framework_architecture_incompatible",
|
||||||
|
"status": "waiting",
|
||||||
|
"taskId": 4969086,
|
||||||
|
"taskType": "text-generation"
|
||||||
|
}
|
||||||
|
],
|
||||||
"architectureModelConfigErrors": {
|
"architectureModelConfigErrors": {
|
||||||
"GestaltLabs/Ornstein-3.5-9B-V2-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/GestaltLabs/Ornstein-3.5-9B-V2-GGUF/resolve/master/config.json (status=404)",
|
"GestaltLabs/Ornstein-3.5-9B-V2-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/GestaltLabs/Ornstein-3.5-9B-V2-GGUF/resolve/master/config.json (status=404)",
|
||||||
"LiquidAI/LFM2.5-2.6B-DSpark-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/LiquidAI/LFM2.5-2.6B-DSpark-GGUF/resolve/master/config.json (status=500)",
|
"LiquidAI/LFM2.5-2.6B-DSpark-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/LiquidAI/LFM2.5-2.6B-DSpark-GGUF/resolve/master/config.json (status=500)",
|
||||||
@@ -121,22 +120,47 @@
|
|||||||
"z-lab/Qwen3.8-27B-DFlash2-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/z-lab/Qwen3.8-27B-DFlash2-GGUF/resolve/master/config.json (status=404)"
|
"z-lab/Qwen3.8-27B-DFlash2-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/z-lab/Qwen3.8-27B-DFlash2-GGUF/resolve/master/config.json (status=404)"
|
||||||
},
|
},
|
||||||
"architectureModelConfigsComplete": 58,
|
"architectureModelConfigsComplete": 58,
|
||||||
"architectureOnly": false,
|
"architectureOnly": true,
|
||||||
"architecturePolicySkipped": {
|
"architecturePolicySkipped": {
|
||||||
"frameworkCatalogUnknown": 0,
|
"frameworkCatalogUnknown": 0,
|
||||||
"frameworkContextUnknown": 67,
|
"frameworkContextUnknown": 67,
|
||||||
"modelArchitectureUnknown": 133,
|
"modelArchitectureUnknown": 133,
|
||||||
"noMatchingBlock": 986,
|
"noMatchingBlock": 985,
|
||||||
"partiallyBlockedFrameworkSet": 22,
|
"partiallyBlockedFrameworkSet": 22,
|
||||||
"runningMatchedProtected": 0,
|
"runningMatchedProtected": 0,
|
||||||
"submissionContextMismatch": 0,
|
"submissionContextMismatch": 0,
|
||||||
"submissionContextUnknown": 0
|
"submissionContextUnknown": 0
|
||||||
},
|
},
|
||||||
"cancelledCount": 0,
|
"cancelledCount": 1,
|
||||||
"cancelledTasks": [],
|
"cancelledTasks": [
|
||||||
|
{
|
||||||
|
"accountIndex": 10,
|
||||||
|
"architectureBlockEvidenceCount": 1,
|
||||||
|
"architectureBlockExpiresAt": "2026-10-20T02:52:52.089383+00:00",
|
||||||
|
"architectureMatchScope": "exact_framework",
|
||||||
|
"architectureMatchType": "model_type",
|
||||||
|
"architectureSignature": "model_type:qwen3_5_text",
|
||||||
|
"architectureSignatures": [
|
||||||
|
"model_type:qwen3_5_text"
|
||||||
|
],
|
||||||
|
"cleanupReasons": [
|
||||||
|
"known_framework_architecture_incompatible"
|
||||||
|
],
|
||||||
|
"evaluatedFrameworks": [
|
||||||
|
"vllm"
|
||||||
|
],
|
||||||
|
"framework": "vllm",
|
||||||
|
"gpuType": "Cambricon_mlu-370-x4",
|
||||||
|
"modelId": "TokenRhythm/NeoHorse-1-4B",
|
||||||
|
"reason": "known_framework_architecture_incompatible",
|
||||||
|
"status": "waiting",
|
||||||
|
"taskId": 4969086,
|
||||||
|
"taskType": "text-generation"
|
||||||
|
}
|
||||||
|
],
|
||||||
"certainOomCount": 0,
|
"certainOomCount": 0,
|
||||||
"certainOomTasks": [],
|
"certainOomTasks": [],
|
||||||
"cleanupCandidateCount": 0,
|
"cleanupCandidateCount": 1,
|
||||||
"dryRun": false,
|
"dryRun": false,
|
||||||
"listingErrors": {},
|
"listingErrors": {},
|
||||||
"modelAgeErrors": {},
|
"modelAgeErrors": {},
|
||||||
@@ -161,19 +185,17 @@
|
|||||||
],
|
],
|
||||||
"oldOverflowCount": 0,
|
"oldOverflowCount": 0,
|
||||||
"oldOverflowTasks": [],
|
"oldOverflowTasks": [],
|
||||||
"policyCancelledRecorded": 0,
|
"policyCancelledRecorded": 1,
|
||||||
"policyNoLongerAppliesCount": 0,
|
"policyNoLongerAppliesCount": 0,
|
||||||
"policyNoLongerAppliesTasks": [],
|
"policyNoLongerAppliesTasks": [],
|
||||||
"recentModelDays": 7,
|
"recentModelDays": 7,
|
||||||
"recentModelReserveSlots": 5,
|
"recentModelReserveSlots": 5,
|
||||||
"repositorySizeErrors": {
|
"repositorySizeErrors": {},
|
||||||
"empero-ai/Qwen3.8-9B-GGUF": "recursive_repository_size_incomplete"
|
"repositorySizesComplete": 0,
|
||||||
},
|
|
||||||
"repositorySizesComplete": 339,
|
|
||||||
"skipped": {
|
"skipped": {
|
||||||
"fitsKnownCapacity": 1139,
|
"fitsKnownCapacity": 0,
|
||||||
"gpuCapacityUnknown": 0,
|
"gpuCapacityUnknown": 0,
|
||||||
"repositorySizeUnknown": 2
|
"repositorySizeUnknown": 0
|
||||||
},
|
},
|
||||||
"stopErrors": [],
|
"stopErrors": [],
|
||||||
"uniqueModels": 340
|
"uniqueModels": 340
|
||||||
|
|||||||
@@ -255,6 +255,7 @@
|
|||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-20T11:30:54.600018+00:00", "modelId": "primitive-ai/Nex-N2.5-mini-mixed-NVFP4-FP8", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "ac4b2845ddd4bdce478dd2b65134ee3ab3d259a119afe6841bf068b135fd4892", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26057989872, "estimatedRequiredGiB": 29.158, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26090095180, "modelscopeLicense": "apache-2.0", "modelscopeParams": 21416999792, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:nvfp4", "custom_tag:fp8", "custom_tag:mixed-precision", "custom_tag:quantized", "custom_tag:moe", "custom_tag:nex-n2.5", "custom_tag:agentic", "custom_tag:tool-calling", "custom_tag:single-gpu", "custom_tag:blackwell"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26090095180}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.989339+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4969386", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-20T11:30:54.600018+00:00", "modelId": "primitive-ai/Nex-N2.5-mini-mixed-NVFP4-FP8", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "ac4b2845ddd4bdce478dd2b65134ee3ab3d259a119afe6841bf068b135fd4892", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26057989872, "estimatedRequiredGiB": 29.158, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26090095180, "modelscopeLicense": "apache-2.0", "modelscopeParams": 21416999792, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:nvfp4", "custom_tag:fp8", "custom_tag:mixed-precision", "custom_tag:quantized", "custom_tag:moe", "custom_tag:nex-n2.5", "custom_tag:agentic", "custom_tag:tool-calling", "custom_tag:single-gpu", "custom_tag:blackwell"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26090095180}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.989339+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4969386", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T04:11:57.967192+00:00", "modelId": "aisingapore/Llama-SEA-LION-v2-8B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16060556376, "estimatedRequiredGiB": 17.961, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 16071447395, "modelscopeLicense": "llama3", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16071447395}, "outcome": "success", "status": "success", "submitTime": "2026-09-20T03:09:18.938123+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4969384", "taskType": "text-generation", "verifyResult": 1}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T04:11:57.967192+00:00", "modelId": "aisingapore/Llama-SEA-LION-v2-8B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16060556376, "estimatedRequiredGiB": 17.961, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 16071447395, "modelscopeLicense": "llama3", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16071447395}, "outcome": "success", "status": "success", "submitTime": "2026-09-20T03:09:18.938123+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4969384", "taskType": "text-generation", "verifyResult": 1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-20T22:57:49.771972+00:00", "modelId": "ibm-granite/granite-guardian-4.1-8b", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16761144464, "estimatedRequiredGiB": 18.743, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": 16770924251, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8380551168, "modelscopeTags": ["license:apache-2.0", "model_type:granite", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:granite", "custom_tag:guardian", "custom_tag:safety", "custom_tag:hallucination"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16770924251}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.880160+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4969382", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-20T22:57:49.771972+00:00", "modelId": "ibm-granite/granite-guardian-4.1-8b", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16761144464, "estimatedRequiredGiB": 18.743, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": 16770924251, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8380551168, "modelscopeTags": ["license:apache-2.0", "model_type:granite", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:granite", "custom_tag:guardian", "custom_tag:safety", "custom_tag:hallucination"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16770924251}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.880160+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4969382", "taskType": "text-generation", "verifyResult": -1}
|
||||||
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-21T07:17:50.956808+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.836648+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969379", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T18:53:01.473025+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.549276+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969363", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T18:53:01.473025+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.549276+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969363", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T20:31:51.861054+00:00", "modelId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34327989512, "estimatedRequiredGiB": 38.385, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 34346039951, "modelscopeLicense": "mit", "modelscopeParams": 32762123264, "modelscopeTags": ["license:mit", "model_type:qwen3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 34346039951}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.543315+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969362", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T20:31:51.861054+00:00", "modelId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34327989512, "estimatedRequiredGiB": 38.385, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 34346039951, "modelscopeLicense": "mit", "modelscopeParams": 32762123264, "modelscopeTags": ["license:mit", "model_type:qwen3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 34346039951}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.543315+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969362", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["zaya"], "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T11:30:54.600365+00:00", "modelId": "JANGQ-AI/Osaurus-AppleScript-8B-JANG_4M", "modelProfile": {"architectures": ["ZayaForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7915599888, "estimatedRequiredGiB": 8.885, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "zaya", "modelscopeFileSize": 7950463599, "modelscopeLicense": "other", "modelscopeParams": 2274126328, "modelscopeTags": ["license:other", "model_type:zaya", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:applescript", "custom_tag:macos", "custom_tag:computer-use", "custom_tag:agent", "custom_tag:tool-calling", "custom_tag:function-calling", "custom_tag:mlx", "custom_tag:moe", "custom_tag:osaurus"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7950463599}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.452113+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4969361", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["zaya"], "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T11:30:54.600365+00:00", "modelId": "JANGQ-AI/Osaurus-AppleScript-8B-JANG_4M", "modelProfile": {"architectures": ["ZayaForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7915599888, "estimatedRequiredGiB": 8.885, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "zaya", "modelscopeFileSize": 7950463599, "modelscopeLicense": "other", "modelscopeParams": 2274126328, "modelscopeTags": ["license:other", "model_type:zaya", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:applescript", "custom_tag:macos", "custom_tag:computer-use", "custom_tag:agent", "custom_tag:tool-calling", "custom_tag:function-calling", "custom_tag:mlx", "custom_tag:moe", "custom_tag:osaurus"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7950463599}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:18.452113+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4969361", "taskType": "text-generation", "verifyResult": -1}
|
||||||
@@ -292,9 +293,8 @@
|
|||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["nemotron_h"], "framework": "vllm", "lastSyncTime": "2026-09-21T01:35:45.556177+00:00", "modelId": "primitive-ai/Nemotron-3.5-Lightning-30B-A3B-mixed-INT4-INT8", "modelProfile": {"architectures": ["NemotronHForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20628596944, "estimatedRequiredGiB": 23.076, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "nemotron_h", "modelscopeFileSize": 20647674585, "modelscopeLicense": "other", "modelscopeParams": 33943909952, "modelscopeTags": ["license:other", "model_type:nemotron_h", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:int4", "custom_tag:int8", "custom_tag:mixed-precision", "custom_tag:quantized", "custom_tag:moe", "custom_tag:mamba", "custom_tag:blackwell"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 20647674585}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.152592+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969082", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["nemotron_h"], "framework": "vllm", "lastSyncTime": "2026-09-21T01:35:45.556177+00:00", "modelId": "primitive-ai/Nemotron-3.5-Lightning-30B-A3B-mixed-INT4-INT8", "modelProfile": {"architectures": ["NemotronHForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20628596944, "estimatedRequiredGiB": 23.076, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "nemotron_h", "modelscopeFileSize": 20647674585, "modelscopeLicense": "other", "modelscopeParams": 33943909952, "modelscopeTags": ["license:other", "model_type:nemotron_h", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:int4", "custom_tag:int8", "custom_tag:mixed-precision", "custom_tag:quantized", "custom_tag:moe", "custom_tag:mamba", "custom_tag:blackwell"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 20647674585}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.152592+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969082", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-21T05:07:16.468524+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955963132}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.136770+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969080", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-21T05:07:16.468524+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955963132}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.136770+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969080", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T19:19:40.458840+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-5bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 804816628, "estimatedRequiredGiB": 0.905, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 809686744, "modelscopeLicense": "other", "modelscopeParams": 219544320, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 809686744}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.134762+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969081", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T19:19:40.458840+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-5bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 804816628, "estimatedRequiredGiB": 0.905, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 809686744, "modelscopeLicense": "other", "modelscopeParams": 219544320, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 809686744}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.134762+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969081", "taskType": "text-generation", "verifyResult": -1}
|
||||||
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_text"], "framework": "vllm", "lastSyncTime": "2026-09-21T07:17:50.956847+00:00", "modelId": "TokenRhythm/NeoHorse-1-9B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17907657384, "estimatedRequiredGiB": 20.04, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 17931574767, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 8953803264, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17931574767}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.089383+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969078", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T18:53:01.473075+00:00", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 37184031394, "estimatedRequiredGiB": 41.579, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 37204398600, "modelscopeLicense": "apache-2.0", "modelscopeParams": 10061358960, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:mixed-precision", "custom_tag:qwen3.6", "custom_tag:agentic-search"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 37204398600}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.043508+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969075", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T18:53:01.473075+00:00", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 37184031394, "estimatedRequiredGiB": 41.579, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 37204398600, "modelscopeLicense": "apache-2.0", "modelscopeParams": 10061358960, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:mixed-precision", "custom_tag:qwen3.6", "custom_tag:agentic-search"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 37204398600}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:52.043508+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969075", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-21T02:15:17.867341+00:00", "modelId": "aisingapore/Llama-SEA-LION-v3-8B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16060556376, "estimatedRequiredGiB": 17.964, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 16073634582, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16073634582}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.948761+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4969074", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-21T02:15:17.867341+00:00", "modelId": "aisingapore/Llama-SEA-LION-v3-8B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16060556376, "estimatedRequiredGiB": 17.964, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 16073634582, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16073634582}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.948761+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4969074", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "runtime_memory", "failureAction": "lower_memory_risk_or_reject_combination", "failureCategory": "runtime_memory", "failureClassificationReason": "structured_device_oom", "failureCode": "DEVICE_OOM", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T19:56:29.962443+00:00", "modelId": "prithivMLmods/CEERS-2112-14B-Instruct", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29540133904, "estimatedRequiredGiB": 33.022, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 29547221374, "modelscopeLicense": "apache-2.0", "modelscopeParams": 14770033664, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:Code", "custom_tag:Math", "custom_tag:Reasoning", "custom_tag:text-generation-inference", "custom_tag:Reinforcement-learning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 29547221374}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.939142+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969070", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "runtime_memory", "failureAction": "lower_memory_risk_or_reject_combination", "failureCategory": "runtime_memory", "failureClassificationReason": "structured_device_oom", "failureCode": "DEVICE_OOM", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu", "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T19:56:29.962443+00:00", "modelId": "prithivMLmods/CEERS-2112-14B-Instruct", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29540133904, "estimatedRequiredGiB": 33.022, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 29547221374, "modelscopeLicense": "apache-2.0", "modelscopeParams": 14770033664, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:Code", "custom_tag:Math", "custom_tag:Reasoning", "custom_tag:text-generation-inference", "custom_tag:Reinforcement-learning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 29547221374}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.939142+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969070", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedArchitectures": ["Cohere2ForCausalLM"], "framework": "vllm", "lastSyncTime": "2026-09-20T17:09:45.266177+00:00", "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730845002, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730845002}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.937617+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4969071", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedArchitectures": ["Cohere2ForCausalLM"], "framework": "vllm", "lastSyncTime": "2026-09-20T17:09:45.266177+00:00", "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730845002, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730845002}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.937617+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4969071", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["deepseek_v4"], "framework": "vllm", "lastSyncTime": "2026-09-20T20:49:19.683316+00:00", "modelId": "JANGQ-AI/DeepSeek-V4-Flash-JANGTQ-K", "modelProfile": {"architectures": ["DeepseekV4ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 85870266317, "estimatedRequiredGiB": 95.98, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "deepseek_v4", "modelscopeFileSize": 85881510108, "modelscopeLicense": "mit", "modelscopeParams": 21758832856, "modelscopeTags": ["license:mit", "model_type:deepseek_v4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:deepseek", "custom_tag:deepseek-v4", "custom_tag:dsv4", "custom_tag:mixture-of-experts", "custom_tag:mla", "custom_tag:mhc", "custom_tag:sparse-indexer", "custom_tag:million-token-context", "custom_tag:reasoning", "custom_tag:tool-use", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:quantized", "custom_tag:jangtq", "custom_tag:jang"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 85881510108}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.889837+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969069", "taskType": "text-generation", "verifyResult": -1}
|
|
||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T22:27:36.960854+00:00", "modelId": "VerStella/Qwen3.8-27B-heretic-ara-W8A8-DCU", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 31228126952, "estimatedRequiredGiB": 34.934, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 31258377392, "modelscopeLicense": "apache-2.0", "modelscopeParams": 27781427952, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:w8a8", "custom_tag:int8", "custom_tag:quantized"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 31258377392}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T02:52:51.846373+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969063", "taskType": "text-generation", "verifyResult": -1}
|
|
||||||
|
|||||||
@@ -1989,6 +1989,66 @@
|
|||||||
{"batchId": "81a741f19ea448e7a4ee75fcdafba17d", "completedAt": "2026-09-21T07:03:22.631335+00:00", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T06:57:35.196542+00:00", "framework": "vllm", "intentId": "db350590ad414c5398d08e754092973c", "lastModified": "2026-08-24T20:29:58+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-quantized.w8a16", "reason": null, "reconciledAt": "2026-09-21T07:14:25.861449+00:00", "repoId": "RedHatAI/starcoder2-7b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4995632", "taskType": "text-generation"}
|
{"batchId": "81a741f19ea448e7a4ee75fcdafba17d", "completedAt": "2026-09-21T07:03:22.631335+00:00", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T06:57:35.196542+00:00", "framework": "vllm", "intentId": "db350590ad414c5398d08e754092973c", "lastModified": "2026-08-24T20:29:58+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-quantized.w8a16", "reason": null, "reconciledAt": "2026-09-21T07:14:25.861449+00:00", "repoId": "RedHatAI/starcoder2-7b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4995632", "taskType": "text-generation"}
|
||||||
{"batchId": "54eda77281ee43918fc0a587212e1565", "completedAt": "2026-09-21T07:04:01.303012+00:00", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:03:52.542652+00:00", "framework": "vllm-customized", "intentId": "7d5bb49ba2644090a14fe51334e5fe84", "lastModified": "2026-09-15T14:59:20+00:00", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-12B-it-qat-OptiQ-4bit", "reason": null, "reconciledAt": "2026-09-21T07:14:25.863151+00:00", "repoId": "mlx-community/gemma-4-12B-it-qat-OptiQ-4bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4995699", "taskType": "text-generation"}
|
{"batchId": "54eda77281ee43918fc0a587212e1565", "completedAt": "2026-09-21T07:04:01.303012+00:00", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:03:52.542652+00:00", "framework": "vllm-customized", "intentId": "7d5bb49ba2644090a14fe51334e5fe84", "lastModified": "2026-09-15T14:59:20+00:00", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-12B-it-qat-OptiQ-4bit", "reason": null, "reconciledAt": "2026-09-21T07:14:25.863151+00:00", "repoId": "mlx-community/gemma-4-12B-it-qat-OptiQ-4bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4995699", "taskType": "text-generation"}
|
||||||
{"batchId": "54eda77281ee43918fc0a587212e1565", "completedAt": "2026-09-21T07:04:01.303034+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:03:52.542823+00:00", "framework": "vllm", "intentId": "1568d9b63d8845ee9087425d228f3883", "lastModified": "2026-09-16T17:07:54+00:00", "modelAddress": "https://modelscope.cn/models/mlx-community/LFM2.5-1.2B-Thinking-6bit", "reason": null, "reconciledAt": "2026-09-21T07:14:25.863879+00:00", "repoId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Ascend_910-b4", "taskId": "4995698", "taskType": "text-generation"}
|
{"batchId": "54eda77281ee43918fc0a587212e1565", "completedAt": "2026-09-21T07:04:01.303034+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:03:52.542823+00:00", "framework": "vllm", "intentId": "1568d9b63d8845ee9087425d228f3883", "lastModified": "2026-09-16T17:07:54+00:00", "modelAddress": "https://modelscope.cn/models/mlx-community/LFM2.5-1.2B-Thinking-6bit", "reason": null, "reconciledAt": "2026-09-21T07:14:25.863879+00:00", "repoId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Ascend_910-b4", "taskId": "4995698", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787031+00:00", "framework": "vllm", "intentId": "4e55a933509846d09a74dff575f534ac", "lastModified": "2026-08-26T16:45:23+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/Gemma-SEA-LION-v4-27B-IT-NVFP4", "repoId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-NVFP4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787150+00:00", "framework": "vllm_tokenizer_patch", "intentId": "63743683bb5f4635927307a52b8b4c1d", "lastModified": "2026-08-24T19:39:27+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-3b-FP8", "repoId": "neuralmagic/starcoder2-3b-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787195+00:00", "framework": "vllm_tokenizer_patch", "intentId": "ed5555b13a784a35b6ea56be532c6df5", "lastModified": "2026-08-24T19:41:12+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-15b-quantized.w8a8", "repoId": "RedHatAI/starcoder2-15b-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787235+00:00", "framework": "vllm_tokenizer_patch", "intentId": "d0ae2ffd088243439a784c5a33dc312f", "lastModified": "2026-08-24T19:34:42+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-FP8", "repoId": "RedHatAI/Phi-3-mini-128k-instruct-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787273+00:00", "framework": "vllm_tokenizer_patch", "intentId": "bad56210275e443596aea54db268a4bb", "lastModified": "2026-08-24T20:05:18+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "repoId": "RedHatAI/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787310+00:00", "framework": "vllm_tokenizer_patch", "intentId": "69d516f810db4260b714e10ad05e1a4d", "lastModified": "2026-08-24T19:59:13+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-FP8", "repoId": "RedHatAI/starcoder2-7b-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787348+00:00", "framework": "vllm_tokenizer_patch", "intentId": "919cfd7a8fb84099b999721d31d5c1d6", "lastModified": "2026-08-24T19:38:03+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "repoId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787386+00:00", "framework": "vllm_tokenizer_patch", "intentId": "fc26ba33d92143b98c5a076d975a032e", "lastModified": "2026-08-24T19:51:52+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-small-128k-instruct-quantized.w8a16", "repoId": "RedHatAI/Phi-3-small-128k-instruct-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787423+00:00", "framework": "vllm_tokenizer_patch", "intentId": "aec457985806452bb891d96a84a483c3", "lastModified": "2026-08-24T19:36:07+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w4a16", "repoId": "RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w4a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "26bfee231695e882b96dee0191bc1dfec25bcad132af96a7ee5f0e26f3cb9fdc", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787461+00:00", "framework": "vllm_tokenizer_patch", "intentId": "c82b329e31ff45ed9ae760d83015b5a9", "lastModified": "2026-08-24T19:35:49+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/SmolLM-1.7B-Instruct-quantized.w8a16", "repoId": "RedHatAI/SmolLM-1.7B-Instruct-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787498+00:00", "framework": "vllm_tokenizer_patch", "intentId": "ea72285c8871488385f0dcf542f00de3", "lastModified": "2026-08-24T20:02:01+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-2b-quantized.w8a16", "repoId": "RedHatAI/gemma-2-2b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787536+00:00", "framework": "vllm_tokenizer_patch", "intentId": "3a99c1f5d7634b9487d97a7dfdf2d543", "lastModified": "2026-09-09T07:06:46+00:00", "modelAddress": "https://modelscope.cn/models/ddalcu/Muse-Glimmer-30B-MLX-Serve-8bit", "repoId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-8bit", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787588+00:00", "framework": "vllm_tokenizer_patch", "intentId": "0cf0c2a55e2444629ec0b181a74837a2", "lastModified": "2026-08-26T16:41:36+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "repoId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787638+00:00", "framework": "vllm", "intentId": "d96839efba24414aab8fbc81c3a14cbf", "lastModified": "2026-09-10T07:49:50+00:00", "modelAddress": "https://modelscope.cn/models/TokenRhythm/NeoHorse-1-9B", "repoId": "TokenRhythm/NeoHorse-1-9B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787708+00:00", "framework": "vllm", "intentId": "1620544f5b744ecd811c0b3c09522e64", "lastModified": "2026-09-08T13:43:47+00:00", "modelAddress": "https://modelscope.cn/models/JANGQ-AI/Osaurus-AppleScript-8B-JANG_6M", "repoId": "JANGQ-AI/Osaurus-AppleScript-8B-JANG_6M", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787766+00:00", "framework": "vllm", "intentId": "c50613781c5447138b16b4a51641a9f0", "lastModified": "2026-09-09T06:27:24+00:00", "modelAddress": "https://modelscope.cn/models/ddalcu/Muse-Glimmer-30B-MLX-Serve-4bit", "repoId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-4bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787823+00:00", "framework": "vllm_fix_tokenizer", "intentId": "c64c546897f64389bd820235ca61d06f", "lastModified": "2026-08-31T03:35:12+00:00", "modelAddress": "https://modelscope.cn/models/logic65/Qwen3.8-Whittle-w50-18.3B-unrepaired", "repoId": "logic65/Qwen3.8-Whittle-w50-18.3B-unrepaired", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787867+00:00", "framework": "vllm_fix_tokenizer", "intentId": "6475c549afc943ae8bd7a764f9c55242", "lastModified": "2026-09-10T14:58:25+00:00", "modelAddress": "https://modelscope.cn/models/webAI-Official/TwIL-LM3", "repoId": "webAI-Official/TwIL-LM3", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787910+00:00", "framework": "vllm", "intentId": "eb442b9b0b244d45b8a44314f22639dd", "lastModified": "2026-08-24T20:09:59+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Llama-2-7b-chat-quantized.w8a8", "repoId": "neuralmagic/Llama-2-7b-chat-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x4", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.787959+00:00", "framework": "vllm", "intentId": "b3e6711e4b554039a75c4c7c452c6da2", "lastModified": "2026-08-24T20:10:55+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-quantized.w8a8", "repoId": "RedHatAI/starcoder2-7b-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x4", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788007+00:00", "framework": "vllm", "intentId": "e49e4dd19d64485e8dc6ce26941686c9", "lastModified": "2026-09-08T15:45:26+00:00", "modelAddress": "https://modelscope.cn/models/JANGQ-AI/AppleScript-8B-JANG_4M", "repoId": "JANGQ-AI/AppleScript-8B-JANG_4M", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x4", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "f49713b9fe8e3e7683ab2722544c85a44993bcd72b9421c1599f634e68b3a2f7", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788054+00:00", "framework": "vllm", "intentId": "77a22ed13df14af0a13a07c6a2d51405", "lastModified": "2026-08-26T17:29:18+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/SEA-LION-v1-7B", "repoId": "aisingapore/SEA-LION-v1-7B", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x4", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788101+00:00", "framework": "vllm-customized", "intentId": "27f7260ca53847b1ad9ec976da265577", "lastModified": "2026-09-09T12:19:20+00:00", "modelAddress": "https://modelscope.cn/models/OpenBMB/BitCPM-CANN-8B", "repoId": "OpenBMB/BitCPM-CANN-8B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788160+00:00", "framework": "vllm-customized", "intentId": "613dbc868dc440418441a226876dee0b", "lastModified": "2026-08-24T20:22:00+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-2b-it-quantized.w8a16", "repoId": "neuralmagic/gemma-2-2b-it-quantized.w8a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788219+00:00", "framework": "vllm-customized", "intentId": "a829144c10fd4c2496a06d6c1f7913c6", "lastModified": "2026-08-26T17:18:42+00:00", "modelAddress": "https://modelscope.cn/models/nv-community/NVIDIA-Nemotron-3-Nano-30B-A3B-NVFP4", "repoId": "nv-community/NVIDIA-Nemotron-3-Nano-30B-A3B-NVFP4", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788277+00:00", "framework": "vllm-customized", "intentId": "f4f7b1db9f2640019c399e35f0d9630e", "lastModified": "2026-08-24T19:44:39+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-2b-it-quantized.w8a8", "repoId": "neuralmagic/gemma-2-2b-it-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788334+00:00", "framework": "vllm-customized", "intentId": "801e20f23fad413ea8e371fbda094a55", "lastModified": "2026-09-12T12:32:06+00:00", "modelAddress": "https://modelscope.cn/models/OpenBMB/MiniCPM5-2B-GPTQ", "repoId": "OpenBMB/MiniCPM5-2B-GPTQ", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788392+00:00", "framework": "vllm-customized", "intentId": "571ca8d3d4d74ce0af134d4edeb20d5d", "lastModified": "2026-09-10T07:52:22+00:00", "modelAddress": "https://modelscope.cn/models/TokenRhythm/NeoHorse-1-4B", "repoId": "TokenRhythm/NeoHorse-1-4B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788449+00:00", "framework": "vllm-customized", "intentId": "28d1afc397a34efa82edcf9082670c7e", "lastModified": "2026-08-24T20:30:03+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w4a16", "repoId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w4a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788506+00:00", "framework": "vllm-customized", "intentId": "832a93f9a3274cd687584d2d7e561305", "lastModified": "2026-08-24T20:09:49+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Llama-3.2-3B-Instruct-FP8", "repoId": "RedHatAI/Llama-3.2-3B-Instruct-FP8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788564+00:00", "framework": "vllm-customized", "intentId": "5b22fbdea03f444bb61f799e6e9ab5b6", "lastModified": "2026-08-24T20:33:18+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Qwen2-0.5B-Instruct-quantized.w8a8", "repoId": "RedHatAI/Qwen2-0.5B-Instruct-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788621+00:00", "framework": "vllm-customized", "intentId": "c9a56dfb5e684de9a76fbff3ca18c9eb", "lastModified": "2026-08-24T20:29:58+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-quantized.w8a16", "repoId": "RedHatAI/starcoder2-7b-quantized.w8a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788682+00:00", "framework": "vllm-customized", "intentId": "9217ba6e3da44074bb3453ac04f496ef", "lastModified": "2026-08-24T20:26:53+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-9b-it-quantized.w8a16", "repoId": "RedHatAI/gemma-2-9b-it-quantized.w8a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788740+00:00", "framework": "vllm-customized", "intentId": "1f051ce558134969ad7e7e2cd8a7a327", "lastModified": "2026-09-04T07:05:43+00:00", "modelAddress": "https://modelscope.cn/models/XHToken/Spark-X2.5-4B-FP8", "repoId": "XHToken/Spark-X2.5-4B-FP8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788797+00:00", "framework": "vllm-customized", "intentId": "429aa1b3867a4fd8b5efdfd677fdc636", "lastModified": "2026-09-01T03:52:21+00:00", "modelAddress": "https://modelscope.cn/models/primitive-ai/Nemotron-3.5-Lightning-30B-A3B-mixed-INT4-INT8", "repoId": "primitive-ai/Nemotron-3.5-Lightning-30B-A3B-mixed-INT4-INT8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788854+00:00", "framework": "vllm-customized", "intentId": "0738585ee92f46948c5eee59078ea58d", "lastModified": "2026-08-26T19:32:38+00:00", "modelAddress": "https://modelscope.cn/models/pfnet/plamo-3-nict-8b-base", "repoId": "pfnet/plamo-3-nict-8b-base", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788911+00:00", "framework": "vllm", "intentId": "20a963eb43de433789b48c6cfd87185b", "lastModified": "2026-08-24T19:46:26+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-15b-quantized.w8a8", "repoId": "neuralmagic/starcoder2-15b-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Mthreads_s4000", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.788967+00:00", "framework": "vllm", "intentId": "4852a2d576f34cd5b13220080c900764", "lastModified": "2026-08-24T19:48:09+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "repoId": "neuralmagic/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Mthreads_s4000", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789024+00:00", "framework": "vllm", "intentId": "43eee3503ed6480ca36b2b2611038277", "lastModified": "2026-08-24T20:31:44+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Meta-Llama-3-8B-Instruct-quantized.w8a8", "repoId": "neuralmagic/Meta-Llama-3-8B-Instruct-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789067+00:00", "framework": "vllm", "intentId": "8f95190836584586b706c1038fe5bdcf", "lastModified": "2026-08-24T19:53:06+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "repoId": "neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789109+00:00", "framework": "vllm", "intentId": "6f1f26c9b56742e18f25462f819063c6", "lastModified": "2026-08-24T19:36:40+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Qwen2-1.5B-Instruct-quantized.w8a8", "repoId": "neuralmagic/Qwen2-1.5B-Instruct-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789150+00:00", "framework": "vllm", "intentId": "b8df678dccb8412bbc42164a0c8a34c6", "lastModified": "2026-08-24T20:26:05+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Qwen2-7B-Instruct-quantized.w8a8", "repoId": "neuralmagic/Qwen2-7B-Instruct-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789191+00:00", "framework": "vllm", "intentId": "1bcdcf79ef7d467fa41bea261ad54cc3", "lastModified": "2026-08-24T20:01:20+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Llama-3.2-1B-Instruct-quantized.w8a8", "repoId": "neuralmagic/Llama-3.2-1B-Instruct-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789232+00:00", "framework": "vllm", "intentId": "04f9f7af5dfa4e0998b04096d47e9d90", "lastModified": "2026-08-24T20:10:58+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Meta-Llama-3.1-8B-quantized.w8a16", "repoId": "neuralmagic/Meta-Llama-3.1-8B-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789273+00:00", "framework": "vllm", "intentId": "a730917a5ae24e60bbf64072920d4aeb", "lastModified": "2026-08-24T20:22:06+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-15b-quantized.w8a16", "repoId": "neuralmagic/starcoder2-15b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789314+00:00", "framework": "vllm", "intentId": "1cde64e8172e45c68af0d743c4026a3c", "lastModified": "2026-08-24T20:14:53+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-9b-it-quantized.w4a16", "repoId": "neuralmagic/gemma-2-9b-it-quantized.w4a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789355+00:00", "framework": "vllm", "intentId": "8167594e93b2484b96377bf45cabcdbc", "lastModified": "2026-08-24T19:39:21+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-9b-it-quantized.w8a16", "repoId": "neuralmagic/gemma-2-9b-it-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789396+00:00", "framework": "vllm", "intentId": "74292320b4b44de6abb5b95d3993b171", "lastModified": "2026-08-24T19:42:52+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-7b-FP8", "repoId": "neuralmagic/starcoder2-7b-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789437+00:00", "framework": "vllm", "intentId": "d98bbab92e9f4fae83dd1cc110ed96b2", "lastModified": "2026-08-24T20:21:52+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-7b-quantized.w8a16", "repoId": "neuralmagic/starcoder2-7b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789478+00:00", "framework": "vllm", "intentId": "0eb5811de6664de38ee08402735fcff9", "lastModified": "2026-08-24T20:36:30+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Mistral-Nemo-Instruct-2407-quantized.w4a16", "repoId": "neuralmagic/Mistral-Nemo-Instruct-2407-quantized.w4a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789519+00:00", "framework": "vllm", "intentId": "cdd47c5706da4d3aada1738a7d13464c", "lastModified": "2026-08-24T19:39:20+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-3b-quantized.w8a16", "repoId": "neuralmagic/starcoder2-3b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789560+00:00", "framework": "vllm", "intentId": "439c9ad32f2a473083cdb3ed032fec85", "lastModified": "2026-08-24T19:40:10+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-7b-quantized.w8a8", "repoId": "neuralmagic/starcoder2-7b-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789602+00:00", "framework": "vllm", "intentId": "6185b7a217574b7daf76e2b34d9ae225", "lastModified": "2026-09-08T03:51:51+00:00", "modelAddress": "https://modelscope.cn/models/XHToken/Spark-X2.5-4B", "repoId": "XHToken/Spark-X2.5-4B", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789642+00:00", "framework": "vllm", "intentId": "d6261393978949ca8e7da2b6cfc33629", "lastModified": "2026-09-03T13:37:32+00:00", "modelAddress": "https://modelscope.cn/models/XHToken/Spark-X2.5-1.7B", "repoId": "XHToken/Spark-X2.5-1.7B", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789688+00:00", "framework": "vllm", "intentId": "41a624ce382642808ac968496ff838e8", "lastModified": "2026-08-24T19:39:57+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Meta-Llama-3.1-8B-quantized.w8a8", "repoId": "RedHatAI/Meta-Llama-3.1-8B-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789729+00:00", "framework": "vllm", "intentId": "4b592aab3463494b9c43781e6e5af239", "lastModified": "2026-08-24T20:15:22+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Llama-3.2-1B-Instruct-quantized.w8a8", "repoId": "RedHatAI/Llama-3.2-1B-Instruct-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789770+00:00", "framework": "vllm", "intentId": "a3d944aa0bba4c52add303dd341cb4b0", "lastModified": "2026-08-24T20:23:31+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Qwen2-1.5B-Instruct-quantized.w8a8", "repoId": "RedHatAI/Qwen2-1.5B-Instruct-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789811+00:00", "framework": "vllm", "intentId": "d13f7f442d38435ea5550d0c51962e43", "lastModified": "2026-08-24T20:08:44+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "repoId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789852+00:00", "framework": "vllm", "intentId": "768b06a238ec4d649b3564639f373e61", "lastModified": "2026-08-24T19:39:28+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-3b-quantized.w8a16", "repoId": "RedHatAI/starcoder2-3b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
|
{"batchId": "916dc36a449a401baff956fb471682a2", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-21T07:19:28.789894+00:00", "framework": "vllm", "intentId": "b13b574710b64370a7aa4ea45c0f61ee", "lastModified": "2026-08-24T20:22:22+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Meta-Llama-3.1-8B-quantized.w8a16", "repoId": "RedHatAI/Meta-Llama-3.1-8B-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Sunrise_pt-200-x1", "taskType": "text-generation"}
|
||||||
{"batchId": "81a741f19ea448e7a4ee75fcdafba17d", "completedAt": "2026-09-21T07:03:22.631445+00:00", "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configSource": "modelhub_live", "createdAt": "2026-09-21T06:57:35.198537+00:00", "framework": "vllm_fix_tokenizer", "intentId": "232b073aaa7f4a7ebada67ce20761c8d", "repoId": "LiquidAI/LFM2.5-1.2B-Instruct-MLX-bf16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Sunrise_pt-200-x1", "taskId": null, "taskType": "text-generation"}
|
{"batchId": "81a741f19ea448e7a4ee75fcdafba17d", "completedAt": "2026-09-21T07:03:22.631445+00:00", "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configSource": "modelhub_live", "createdAt": "2026-09-21T06:57:35.198537+00:00", "framework": "vllm_fix_tokenizer", "intentId": "232b073aaa7f4a7ebada67ce20761c8d", "repoId": "LiquidAI/LFM2.5-1.2B-Instruct-MLX-bf16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Sunrise_pt-200-x1", "taskId": null, "taskType": "text-generation"}
|
||||||
{"batchId": "81a741f19ea448e7a4ee75fcdafba17d", "completedAt": "2026-09-21T07:03:22.631441+00:00", "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configSource": "modelhub_live", "createdAt": "2026-09-21T06:57:35.198482+00:00", "framework": "vllm_fix_tokenizer", "intentId": "f5f73adefdb44115a538f3c29f95b6bd", "repoId": "inceptionai/Jais-2-8B-Chat", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Sunrise_pt-200-x1", "taskId": null, "taskType": "text-generation"}
|
{"batchId": "81a741f19ea448e7a4ee75fcdafba17d", "completedAt": "2026-09-21T07:03:22.631441+00:00", "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configSource": "modelhub_live", "createdAt": "2026-09-21T06:57:35.198482+00:00", "framework": "vllm_fix_tokenizer", "intentId": "f5f73adefdb44115a538f3c29f95b6bd", "repoId": "inceptionai/Jais-2-8B-Chat", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Sunrise_pt-200-x1", "taskId": null, "taskType": "text-generation"}
|
||||||
{"batchId": "81a741f19ea448e7a4ee75fcdafba17d", "completedAt": "2026-09-21T07:03:22.631438+00:00", "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configSource": "modelhub_live", "createdAt": "2026-09-21T06:57:35.198427+00:00", "framework": "vllm_fix_tokenizer", "intentId": "b60961dede1547f5a4788d8a56af2845", "repoId": "mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Sunrise_pt-200-x1", "taskId": null, "taskType": "text-generation"}
|
{"batchId": "81a741f19ea448e7a4ee75fcdafba17d", "completedAt": "2026-09-21T07:03:22.631438+00:00", "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configSource": "modelhub_live", "createdAt": "2026-09-21T06:57:35.198427+00:00", "framework": "vllm_fix_tokenizer", "intentId": "b60961dede1547f5a4788d8a56af2845", "repoId": "mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Sunrise_pt-200-x1", "taskId": null, "taskType": "text-generation"}
|
||||||
|
|||||||
@@ -2,24 +2,24 @@
|
|||||||
"agentVersion": "2026.09.20.2",
|
"agentVersion": "2026.09.20.2",
|
||||||
"checksums": {
|
"checksums": {
|
||||||
".modelhub_state/account_capacity.json": "68d2a8390ad022f2ddcd30841f3860c4535387fa64cb46dc7507eb16837e798f",
|
".modelhub_state/account_capacity.json": "68d2a8390ad022f2ddcd30841f3860c4535387fa64cb46dc7507eb16837e798f",
|
||||||
".modelhub_state/architecture_compatibility_blacklist.json": "72d2d47132274b9bd4ef1cf2b8337165f63eab93c077eb75191e282cd9c44a23",
|
".modelhub_state/architecture_compatibility_blacklist.json": "1cb364f1d04083d2c3812e95d3a6f1dc8dcd6625214a5af3b61e53d398c6baaf",
|
||||||
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
||||||
".modelhub_state/market_intelligence.json": "a0b8fa31187f7749b754808c1df462618f312c885513ec6d01402f6a5332d5fe",
|
".modelhub_state/market_intelligence.json": "c0759ef3d616a38a14f650056c542a8b8ac07e6bbda350f469ae66455689a610",
|
||||||
".modelhub_state/official_capabilities.json": "f2ee1443842af7db6ad57f8101a7812bab2be7473bf25e4684e08f6978f4a753",
|
".modelhub_state/official_capabilities.json": "9d2fca6762c33200d5f0e306fdbf1c37f7e95daf3ad4f6184a48879a9b66acd2",
|
||||||
".modelhub_state/outcome_checkpoint.json": "1de66e7cc10708965e4a18f953474f36c9bd16eeb2a2f2a7c27c88d3e31c7678",
|
".modelhub_state/outcome_checkpoint.json": "b21edacad3863feb727dd4a47c5e5001d80e58a7db1a5f6ce6b7fa8193bdfcdb",
|
||||||
".modelhub_state/queue_cleanup_latest.json": "40260fa59ed1cebe75a3619841d2c3c1c4bef51eaabb680e7c15ca1057f2e3b6",
|
".modelhub_state/queue_cleanup_latest.json": "56206916bd251aff8514079cb7cc354f9b75586442b701fa6494a14e681e7d86",
|
||||||
".modelhub_state/recent_outcomes.jsonl": "1fe20cd2a751ada37142cfa47eda97f8a94321f25ca341c76c7b7ba641fc8fdf",
|
".modelhub_state/recent_outcomes.jsonl": "3ced5f9f51f08c755ddf966e6b719a91eada07d4132fb64acc8f2503d06e2bd7",
|
||||||
".modelhub_state/recovery_active_tasks.jsonl": "dd4475c3a96beabdba2305b64f69d5c0441677806fb8a99c108c12f266162520",
|
".modelhub_state/recovery_active_tasks.jsonl": "dd4475c3a96beabdba2305b64f69d5c0441677806fb8a99c108c12f266162520",
|
||||||
".modelhub_state/recovery_intents.jsonl": "d0b840cada3a9bc8ea454810ee37dcea3578d8f5fe6938e1bcd10d9e2e8c2557",
|
".modelhub_state/recovery_intents.jsonl": "13022ed562a981520764bafe75ff5c85901a3fbab79d28174de9a87b37388842",
|
||||||
".modelhub_state/routing_intelligence.json": "fa356d271999851a6f272460edeadbc8da77a4f8d10aabc6368d0878af9863a3",
|
".modelhub_state/routing_intelligence.json": "fa356d271999851a6f272460edeadbc8da77a4f8d10aabc6368d0878af9863a3",
|
||||||
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
|
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
|
||||||
".modelhub_state/worker_crashes.jsonl": "07505b69e0b050da5f90f8b046a3069d3e1ca2574ca39896bc7d8a66293167f7",
|
".modelhub_state/worker_crashes.jsonl": "07505b69e0b050da5f90f8b046a3069d3e1ca2574ca39896bc7d8a66293167f7",
|
||||||
"ledger/submissions.jsonl": "d5273eaf6bcd9fc5cf4a1d95d1351b991ec397e0ff9fc9668e9c92d32d5d9585",
|
"ledger/submissions.jsonl": "d5273eaf6bcd9fc5cf4a1d95d1351b991ec397e0ff9fc9668e9c92d32d5d9585",
|
||||||
"outcomes/submissions.jsonl": "40d56326595aaafadb6910f3ac4fc3da41f52fe87ee0143f297fdaa4f9e60ed6"
|
"outcomes/submissions.jsonl": "482ac3d7b16937a07c686fc7ce9d83313b01ed00c0d87dbefaed54fe4d7e12e7"
|
||||||
},
|
},
|
||||||
"generation": 11290,
|
"generation": 11291,
|
||||||
"phase": "cycle",
|
"phase": "intent",
|
||||||
"schemaVersion": 1,
|
"schemaVersion": 1,
|
||||||
"updatedAt": "2026-09-21T07:16:47.944091+00:00",
|
"updatedAt": "2026-09-21T07:19:28.925639+00:00",
|
||||||
"writerId": "8b35139af6674067a339a670222d4b67"
|
"writerId": "8b35139af6674067a339a670222d4b67"
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -401,7 +401,6 @@
|
|||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T11:00:25.647557+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248409935, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248409935}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.040710+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969021", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T11:00:25.647557+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248409935, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248409935}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.040710+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969021", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T11:00:25.647466+00:00", "modelId": "inceptionai/Jais-2-8B-Chat", "modelProfile": {"architectures": ["Jais2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16180861064, "estimatedRequiredGiB": 18.097, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "jais2", "modelscopeFileSize": 16193187697, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8090401280, "modelscopeTags": ["license:apache-2.0", "model_type:jais2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16193187697}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.140009+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969028", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T11:00:25.647466+00:00", "modelId": "inceptionai/Jais-2-8B-Chat", "modelProfile": {"architectures": ["Jais2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16180861064, "estimatedRequiredGiB": 18.097, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "jais2", "modelscopeFileSize": 16193187697, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8090401280, "modelscopeTags": ["license:apache-2.0", "model_type:jais2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16193187697}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.140009+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969028", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:00:25.647787+00:00", "modelId": "OpenBMB/BitCPM-CANN-8B", "modelProfile": {"architectures": ["MiniCPMForCausalLM"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16370604914, "estimatedRequiredGiB": 18.305, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "minicpm", "modelscopeFileSize": 16378634956, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:minicpm", "library:pytorch", "library:transformer", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16378634956}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.100478+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969023", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:00:25.647787+00:00", "modelId": "OpenBMB/BitCPM-CANN-8B", "modelProfile": {"architectures": ["MiniCPMForCausalLM"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16370604914, "estimatedRequiredGiB": 18.305, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "minicpm", "modelscopeFileSize": 16378634956, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:minicpm", "library:pytorch", "library:transformer", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16378634956}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.100478+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969023", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:00:25.647366+00:00", "modelId": "IntervitensInc/kek_mk3", "modelProfile": {"architectures": ["StableLmForCausalLM"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3289069288, "estimatedRequiredGiB": 3.683, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "stablelm", "modelscopeFileSize": 3295875324, "modelscopeLicense": null, "modelscopeParams": 1644515328, "modelscopeTags": ["model_type:stablelm", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3295875324}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.141570+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969026", "taskType": "text-generation", "verifyResult": null}
|
|
||||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:00:25.647629+00:00", "modelId": "JANGQ-AI/AppleScript-8B-JANG_4M", "modelProfile": {"architectures": ["ZayaForCausalLM"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7915599888, "estimatedRequiredGiB": 8.885, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "zaya", "modelscopeFileSize": 7950463599, "modelscopeLicense": "other", "modelscopeParams": 2274126328, "modelscopeTags": ["license:other", "model_type:zaya", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:applescript", "custom_tag:macos", "custom_tag:computer-use", "custom_tag:agent", "custom_tag:tool-calling", "custom_tag:function-calling", "custom_tag:mlx", "custom_tag:moe", "custom_tag:osaurus"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7950463599}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.142896+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969027", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:00:25.647629+00:00", "modelId": "JANGQ-AI/AppleScript-8B-JANG_4M", "modelProfile": {"architectures": ["ZayaForCausalLM"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7915599888, "estimatedRequiredGiB": 8.885, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "zaya", "modelscopeFileSize": 7950463599, "modelscopeLicense": "other", "modelscopeParams": 2274126328, "modelscopeTags": ["license:other", "model_type:zaya", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:applescript", "custom_tag:macos", "custom_tag:computer-use", "custom_tag:agent", "custom_tag:tool-calling", "custom_tag:function-calling", "custom_tag:mlx", "custom_tag:moe", "custom_tag:osaurus"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7950463599}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.142896+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969027", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:00:25.647880+00:00", "modelId": "aisingapore/SEA-LION-v1-7B-IT", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "b3ef6bcf57e6bb4fc9ede989e67a765b0e40b19fbd407f1f8452833c12bd5c38", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361928, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008164431, "modelscopeLicense": "mit", "modelscopeParams": 7501651968, "modelscopeTags": ["license:mit", "model_type:mpt", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008164431}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.138593+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969025", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:00:25.647880+00:00", "modelId": "aisingapore/SEA-LION-v1-7B-IT", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "b3ef6bcf57e6bb4fc9ede989e67a765b0e40b19fbd407f1f8452833c12bd5c38", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361928, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008164431, "modelscopeLicense": "mit", "modelscopeParams": 7501651968, "modelscopeTags": ["license:mit", "model_type:mpt", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008164431}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.138593+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969025", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:00:25.647418+00:00", "modelId": "aisingapore/SEA-LION-v1-7B", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "b3ef6bcf57e6bb4fc9ede989e67a765b0e40b19fbd407f1f8452833c12bd5c38", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361968, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008115847, "modelscopeLicense": "mit", "modelscopeParams": 7501651968, "modelscopeTags": ["license:mit", "model_type:mpt", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008115847}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.244847+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969034", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:00:25.647418+00:00", "modelId": "aisingapore/SEA-LION-v1-7B", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "b3ef6bcf57e6bb4fc9ede989e67a765b0e40b19fbd407f1f8452833c12bd5c38", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361968, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008115847, "modelscopeLicense": "mit", "modelscopeParams": 7501651968, "modelscopeTags": ["license:mit", "model_type:mpt", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008115847}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:51.244847+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969034", "taskType": "text-generation", "verifyResult": null}
|
||||||
@@ -413,9 +412,7 @@
|
|||||||
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T11:00:25.647657+00:00", "modelId": "aisingapore/Llama-SEA-LION-v2-8B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16060556376, "estimatedRequiredGiB": 17.961, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 16071447395, "modelscopeLicense": "llama3", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16071447395}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.096766+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969079", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T11:00:25.647657+00:00", "modelId": "aisingapore/Llama-SEA-LION-v2-8B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16060556376, "estimatedRequiredGiB": 17.961, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 16071447395, "modelscopeLicense": "llama3", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16071447395}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.096766+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969079", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647811+00:00", "modelId": "ornith-ai/Ornith-1.5-9B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8818103892, "estimatedRequiredGiB": 9.883, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8842796317, "modelscopeLicense": "mit", "modelscopeParams": 6728625904, "modelscopeTags": ["license:mit", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 8842796317}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.078583+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969077", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647811+00:00", "modelId": "ornith-ai/Ornith-1.5-9B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8818103892, "estimatedRequiredGiB": 9.883, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8842796317, "modelscopeLicense": "mit", "modelscopeParams": 6728625904, "modelscopeTags": ["license:mit", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 8842796317}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.078583+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969077", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647611+00:00", "modelId": "XHToken/Spark-X2.5-4B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8224192408, "estimatedRequiredGiB": 9.209, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 8239718268, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4112079360, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8239718268}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.074544+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969076", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647611+00:00", "modelId": "XHToken/Spark-X2.5-4B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8224192408, "estimatedRequiredGiB": 9.209, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 8239718268, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4112079360, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8239718268}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.074544+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969076", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647581+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794426, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 4205751296, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794426}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.149341+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969086", "taskType": "text-generation", "verifyResult": null}
|
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647756+00:00", "modelId": "XHToken/Spark-X2.5-1.7B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3415340496, "estimatedRequiredGiB": 3.834, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 3430847031, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1707657216, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3430847031}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.140617+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969085", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647756+00:00", "modelId": "XHToken/Spark-X2.5-1.7B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3415340496, "estimatedRequiredGiB": 3.834, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 3430847031, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1707657216, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3430847031}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.140617+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969085", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647479+00:00", "modelId": "TokenRhythm/NeoHorse-1-9B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17907657384, "estimatedRequiredGiB": 20.04, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 17931574767, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 8953803264, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17931574767}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.089383+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969078", "taskType": "text-generation", "verifyResult": null}
|
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647863+00:00", "modelId": "XHToken/Spark-X2.5-4B-FP8", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4450260320, "estimatedRequiredGiB": 4.991, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 4465562659, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "fp8", "repositoryOnDiskBytes": 4465562659}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.150681+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969083", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647863+00:00", "modelId": "XHToken/Spark-X2.5-4B-FP8", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4450260320, "estimatedRequiredGiB": 4.991, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 4465562659, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "fp8", "repositoryOnDiskBytes": 4465562659}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.150681+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969083", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647728+00:00", "modelId": "Mia-AiLab/Qwen3.8-27B-EXL3-3.5bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15338408461, "estimatedRequiredGiB": 17.168, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 15362064483, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7669052656, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:exl3", "custom_tag:exllamav3", "custom_tag:quantization", "custom_tag:long-context", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "exl3", "repositoryOnDiskBytes": 15362064483}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.138616+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969084", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647728+00:00", "modelId": "Mia-AiLab/Qwen3.8-27B-EXL3-3.5bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15338408461, "estimatedRequiredGiB": 17.168, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 15362064483, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7669052656, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:exl3", "custom_tag:exllamav3", "custom_tag:quantization", "custom_tag:long-context", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "exl3", "repositoryOnDiskBytes": 15362064483}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.138616+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969084", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647805+00:00", "modelId": "mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 13736540899, "estimatedRequiredGiB": 15.375, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 13756990943, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3825799024, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:expert-pruning", "custom_tag:reap", "custom_tag:moe", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13756990943}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.247958+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969088", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:00:25.647805+00:00", "modelId": "mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 13736540899, "estimatedRequiredGiB": 15.375, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 13756990943, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3825799024, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:expert-pruning", "custom_tag:reap", "custom_tag:moe", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13756990943}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T02:52:52.247958+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969088", "taskType": "text-generation", "verifyResult": null}
|
||||||
@@ -470,7 +467,6 @@
|
|||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766189+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-DPO-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727954960, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737317601, "modelscopeLicense": "other", "modelscopeParams": 8030269440, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737317601}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.793172+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969378", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766189+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-DPO-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727954960, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737317601, "modelscopeLicense": "other", "modelscopeParams": 8030269440, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737317601}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.793172+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969378", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766496+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-v0.5-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727954960, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737318602, "modelscopeLicense": "other", "modelscopeParams": 8030269440, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737318602}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.791323+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969374", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766496+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-v0.5-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727954960, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737318602, "modelscopeLicense": "other", "modelscopeParams": 8030269440, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737318602}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.791323+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969374", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766081+00:00", "modelId": "solidrust/Llama-3-11B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7541230432, "estimatedRequiredGiB": 8.438, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 7550622334, "modelscopeLicense": "other", "modelscopeParams": 11520053248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 7550622334}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.838488+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969377", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766081+00:00", "modelId": "solidrust/Llama-3-11B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7541230432, "estimatedRequiredGiB": 8.438, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 7550622334, "modelscopeLicense": "other", "modelscopeParams": 11520053248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 7550622334}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.838488+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969377", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766136+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.836648+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969379", "taskType": "text-generation", "verifyResult": null}
|
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766455+00:00", "modelId": "mlx-community/Qwen3.6-35B-A3B-OptiQ-4bit-REAP-19B", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 14864536516, "estimatedRequiredGiB": 16.635, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 14884986381, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3818458992, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:expert-pruning", "custom_tag:reap", "custom_tag:moe", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 14884986381}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.840393+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969375", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766455+00:00", "modelId": "mlx-community/Qwen3.6-35B-A3B-OptiQ-4bit-REAP-19B", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 14864536516, "estimatedRequiredGiB": 16.635, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 14884986381, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3818458992, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:expert-pruning", "custom_tag:reap", "custom_tag:moe", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 14884986381}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.840393+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969375", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766025+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v3-9B", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 18483466448, "estimatedRequiredGiB": 20.702, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 18523972484, "modelscopeLicense": "gemma", "modelscopeParams": 9241705984, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 18523972484}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.835980+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969376", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766025+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v3-9B", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 18483466448, "estimatedRequiredGiB": 20.702, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 18523972484, "modelscopeLicense": "gemma", "modelscopeParams": 9241705984, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 18523972484}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.835980+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969376", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766413+00:00", "modelId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-4bit", "modelProfile": {"architectures": ["MuseGlimmerForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21407235494, "estimatedRequiredGiB": 23.956, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "muse_glimmer", "modelscopeFileSize": 21435726263, "modelscopeLicense": "apache-2.0", "modelscopeParams": 5792290816, "modelscopeTags": ["license:apache-2.0", "model_type:muse_glimmer", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:muse_glimmer"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21435726263}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.943995+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969385", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:10:38.766413+00:00", "modelId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-4bit", "modelProfile": {"architectures": ["MuseGlimmerForConditionalGeneration"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21407235494, "estimatedRequiredGiB": 23.956, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "muse_glimmer", "modelscopeFileSize": 21435726263, "modelscopeLicense": "apache-2.0", "modelscopeParams": 5792290816, "modelscopeTags": ["license:apache-2.0", "model_type:muse_glimmer", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:muse_glimmer"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21435726263}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:09:18.943995+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969385", "taskType": "text-generation", "verifyResult": null}
|
||||||
@@ -994,9 +990,9 @@
|
|||||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T06:41:38.663855+00:00", "modelId": "ornith-ai/Ornith-1.5-35B-A3B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23420419700, "estimatedRequiredGiB": 26.214, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 23455515266, "modelscopeLicense": "mit", "modelscopeParams": 19528501104, "modelscopeTags": ["license:mit", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 23455515266}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T22:29:12.298134+00:00", "targetGpu": "Biren_166m", "taskId": "4988756", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T06:41:38.663855+00:00", "modelId": "ornith-ai/Ornith-1.5-35B-A3B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23420419700, "estimatedRequiredGiB": 26.214, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 23455515266, "modelscopeLicense": "mit", "modelscopeParams": 19528501104, "modelscopeTags": ["license:mit", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 23455515266}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T22:29:12.298134+00:00", "targetGpu": "Biren_166m", "taskId": "4988756", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T06:41:38.663987+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B", "modelProfile": {"architectures": ["Lfm2MoeForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16936006912, "estimatedRequiredGiB": 18.948, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2_moe", "modelscopeFileSize": 16953948786, "modelscopeLicense": "other", "modelscopeParams": 8467856832, "modelscopeTags": ["license:other", "model_type:lfm2_moe", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16953948786}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T22:29:12.195398+00:00", "targetGpu": "Biren_166m", "taskId": "4988751", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T06:41:38.663987+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B", "modelProfile": {"architectures": ["Lfm2MoeForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16936006912, "estimatedRequiredGiB": 18.948, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2_moe", "modelscopeFileSize": 16953948786, "modelscopeLicense": "other", "modelscopeParams": 8467856832, "modelscopeTags": ["license:other", "model_type:lfm2_moe", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16953948786}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T22:29:12.195398+00:00", "targetGpu": "Biren_166m", "taskId": "4988751", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T07:13:18.258797+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663410366, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663410366}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T23:11:16.555645+00:00", "targetGpu": "Biren_166m", "taskId": "4989302", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T07:13:18.258797+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663410366, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663410366}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T23:11:16.555645+00:00", "targetGpu": "Biren_166m", "taskId": "4989302", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-6bit", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 22263760648, "estimatedRequiredGiB": 24.908, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 22286993822, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6019449856, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:qwen3_5", "custom_tag:apple-silicon", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 22286993822}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T23:11:57.335621+00:00", "targetGpu": "Biren_166m", "taskId": "4989327", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T07:17:50.956783+00:00", "modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-6bit", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 22263760648, "estimatedRequiredGiB": 24.908, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 22286993822, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6019449856, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:qwen3_5", "custom_tag:apple-silicon", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 22286993822}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T23:11:57.335621+00:00", "targetGpu": "Biren_166m", "taskId": "4989327", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": null, "modelId": "RedHatAI/Meta-Llama-3.1-8B-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084013360, "estimatedRequiredGiB": 10.162, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093203965, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm", "custom_tag:quantized", "custom_tag:8-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9093203965}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T23:16:02.274149+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4989370", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-21T07:17:50.956821+00:00", "modelId": "RedHatAI/Meta-Llama-3.1-8B-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084013360, "estimatedRequiredGiB": 10.162, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093203965, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm", "custom_tag:quantized", "custom_tag:8-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9093203965}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T23:16:02.274149+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4989370", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": null, "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w4a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7686303632, "estimatedRequiredGiB": 8.592, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 7688299781, "modelscopeLicense": "llama2", "modelscopeParams": 13960238080, "modelscopeTags": ["license:llama2", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7688299781}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T23:16:09.058789+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4989371", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-21T07:17:50.956839+00:00", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w4a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7686303632, "estimatedRequiredGiB": 8.592, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 7688299781, "modelscopeLicense": "llama2", "modelscopeParams": 13960238080, "modelscopeTags": ["license:llama2", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7688299781}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T23:16:09.058789+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4989371", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "solidrust/Llama-3-16B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 10261168632, "estimatedRequiredGiB": 11.478, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 10270608122, "modelscopeLicense": "other", "modelscopeParams": 16754741248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible - mergekit - merge - facebook - meta - pytorch - llama - llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 10270608122}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T23:23:58.501966+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4989457", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "solidrust/Llama-3-16B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 10261168632, "estimatedRequiredGiB": 11.478, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 10270608122, "modelscopeLicense": "other", "modelscopeParams": 16754741248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible - mergekit - merge - facebook - meta - pytorch - llama - llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 10270608122}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T23:23:58.501966+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4989457", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": null, "modelId": "IntervitensInc/kek_mk3", "modelProfile": {"architectures": ["StableLmForCausalLM"], "configFingerprint": "bb5d6992290b1b4fa8f4c37e884585a27bef87a1ebfb58e5ca8f098924550f05", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3289069288, "estimatedRequiredGiB": 3.683, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "stablelm", "modelscopeFileSize": 3295875324, "modelscopeLicense": null, "modelscopeParams": 1644515328, "modelscopeTags": ["model_type:stablelm", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3295875324}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T23:23:58.519510+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4989458", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": null, "modelId": "IntervitensInc/kek_mk3", "modelProfile": {"architectures": ["StableLmForCausalLM"], "configFingerprint": "bb5d6992290b1b4fa8f4c37e884585a27bef87a1ebfb58e5ca8f098924550f05", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3289069288, "estimatedRequiredGiB": 3.683, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "stablelm", "modelscopeFileSize": 3295875324, "modelscopeLicense": null, "modelscopeParams": 1644515328, "modelscopeTags": ["model_type:stablelm", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3295875324}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T23:23:58.519510+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4989458", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": null, "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "bb5d6992290b1b4fa8f4c37e884585a27bef87a1ebfb58e5ca8f098924550f05", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T23:23:58.507376+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4989456", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": null, "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "bb5d6992290b1b4fa8f4c37e884585a27bef87a1ebfb58e5ca8f098924550f05", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T23:23:58.507376+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4989456", "taskType": "text-generation", "verifyResult": null}
|
||||||
|
|||||||
Reference in New Issue
Block a user