state: generation 19040 (cycle)

This commit is contained in:
2026-10-01 03:02:04 +00:00
parent 5ecfd7e283
commit 5899baa814
8 changed files with 1408 additions and 1362 deletions

View File

@@ -3341,7 +3341,7 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-10-01T02:54:18.801291+00:00",
"generatedAt": "2026-10-01T03:00:54.812160+00:00",
"summary": {
"activeBlockCount": 168,
"byGpuFramework": {

View File

@@ -434,7 +434,7 @@
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-10-01T02:59:51.965940+00:00",
"generatedAt": "2026-10-01T03:01:57.764822+00:00",
"gpuStats": {
"Ascend_910-b3": {
"available": true,

View File

@@ -1,5 +1,5 @@
{
"catalogUpdatedAt": "2026-10-01T02:59:51.965940+00:00",
"catalogUpdatedAt": "2026-10-01T03:01:57.764822+00:00",
"configuredTaskTypes": [
"text-generation"
],
@@ -56,7 +56,7 @@
"time-series-forecasting"
],
"errors": [],
"generatedAt": "2026-10-01T02:59:51.965940+00:00",
"generatedAt": "2026-10-01T03:01:57.764822+00:00",
"gpuCatalog": {
"Ascend_910-b3": {
"canVerify": true,
@@ -6381,6 +6381,6 @@
"updateTime": "2025-12-22 08:59:53"
}
],
"taskTreeUpdatedAt": "2026-10-01T02:59:51.965940+00:00",
"taskTreeUpdatedAt": "2026-10-01T03:01:57.764822+00:00",
"version": 1
}

View File

@@ -1,6 +1,6 @@
{
"generatedAt": "2026-10-01T02:54:18.718305+00:00",
"lastSyncTime": "2026-10-01T02:54:18.466130+00:00",
"generatedAt": "2026-10-01T03:00:54.727019+00:00",
"lastSyncTime": "2026-10-01T03:00:54.465128+00:00",
"recentLimit": 300,
"report": {
"architectureCompatibilityBlocks": {
@@ -3490,28 +3490,28 @@
"unresolvedFailureCount": 87
},
"Ascend_910-b3|vllm|text-generation": {
"attributableFailureCount": 69,
"decisionFailureRate": 0.8625,
"decisionSuccessRate": 0.1375,
"decisionTotal": 80,
"attributableFailureCount": 70,
"decisionFailureRate": 0.8642,
"decisionSuccessRate": 0.1358,
"decisionTotal": 81,
"failureBreakdown": {
"ambiguous_runtime": 66,
"framework_architecture_unsupported": 64,
"framework_architecture_unsupported": 65,
"memory_capacity": 1,
"repository_structure": 1,
"tokenizer_compatibility": 3
},
"failureCount": 135,
"failureRate": 0.9247,
"failureCount": 136,
"failureRate": 0.9252,
"framework": "vllm",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"successCount": 11,
"successRate": 0.0753,
"successRate": 0.0748,
"targetGpu": "Ascend_910-b3",
"taskType": "text-generation",
"total": 146,
"total": 147,
"unresolvedFailureCount": 66
},
"Ascend_910-b3|vllm|visual-multi-modal": {
@@ -6040,17 +6040,17 @@
"unresolvedFailureCount": 6473
},
"vllm": {
"attributableFailureCount": 3635,
"attributableFailureCount": 3636,
"decisionFailureRate": 0.9727,
"decisionSuccessRate": 0.0273,
"decisionTotal": 3737,
"decisionTotal": 3738,
"failureBreakdown": {
"ambiguous_runtime": 1589,
"architecture_compatibility": 112,
"attention_backend": 3,
"backend_operator": 93,
"context_length": 161,
"framework_architecture_unsupported": 1365,
"framework_architecture_unsupported": 1366,
"memory_capacity": 725,
"model_load": 220,
"platform_infrastructure": 874,
@@ -6059,14 +6059,14 @@
"tokenizer_compatibility": 420,
"参数/模板问题": 50
},
"failureCount": 6148,
"failureCount": 6149,
"failureRate": 0.9837,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 874,
"successCount": 102,
"successRate": 0.0163,
"total": 6250,
"total": 6251,
"unresolvedFailureCount": 1639
},
"vllm-customized": {
@@ -6217,17 +6217,17 @@
"unresolvedFailureCount": 220
}
},
"generatedAt": "2026-10-01T02:54:18.705539+00:00",
"generatedAt": "2026-10-01T03:00:54.712873+00:00",
"gpuSummaries": {
"Ascend_910-b3": {
"attributableFailureCount": 111,
"decisionFailureRate": 0.8346,
"decisionSuccessRate": 0.1654,
"decisionTotal": 133,
"attributableFailureCount": 112,
"decisionFailureRate": 0.8358,
"decisionSuccessRate": 0.1642,
"decisionTotal": 134,
"failureBreakdown": {
"ambiguous_runtime": 147,
"context_length": 8,
"framework_architecture_unsupported": 97,
"framework_architecture_unsupported": 98,
"memory_capacity": 2,
"platform_infrastructure": 1,
"repository_structure": 1,
@@ -6236,14 +6236,14 @@
"日志缺失": 3,
"验证失败": 27
},
"failureCount": 328,
"failureRate": 0.9371,
"failureCount": 329,
"failureRate": 0.9373,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 1,
"successCount": 22,
"successRate": 0.0629,
"total": 350,
"successRate": 0.0627,
"total": 351,
"unresolvedFailureCount": 216
},
"Ascend_910-b4": {
@@ -7854,6 +7854,29 @@
"total": 2,
"unresolvedFailureCount": 0
},
"Ascend_910-b3|vllm|text-generation|param2moe|none": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm",
"modelType": "param2moe",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Ascend_910-b3",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Ascend_910-b3|vllm|text-generation|qwen2|compressed-tensors": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
@@ -29794,6 +29817,30 @@
"total": 1,
"unresolvedFailureCount": 0
},
"Ascend_910-b3|vllm|text-generation|param2moe|none|34": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm",
"loadSizeLog2Bucket": 34,
"modelType": "param2moe",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Ascend_910-b3",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Ascend_910-b3|vllm|text-generation|qwen2|compressed-tensors|29": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
@@ -54072,20 +54119,20 @@
"unresolvedFailureCount": 0
}
},
"terminalRecords": 17182,
"totalRecords": 17403,
"terminalRecords": 17183,
"totalRecords": 17404,
"totals": {
"attributableFailureCount": 6160,
"decisionFailureRate": 0.8633,
"decisionSuccessRate": 0.1367,
"decisionTotal": 7135,
"attributableFailureCount": 6161,
"decisionFailureRate": 0.8634,
"decisionSuccessRate": 0.1366,
"decisionTotal": 7136,
"failureBreakdown": {
"ambiguous_runtime": 4480,
"architecture_compatibility": 212,
"attention_backend": 3,
"backend_operator": 128,
"context_length": 326,
"framework_architecture_unsupported": 2205,
"framework_architecture_unsupported": 2206,
"memory_capacity": 1206,
"model_load": 551,
"platform_infrastructure": 944,
@@ -54096,30 +54143,30 @@
"日志缺失": 719,
"验证失败": 676
},
"failureCount": 16207,
"failureCount": 16208,
"failureRate": 0.9433,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 944,
"successCount": 975,
"successRate": 0.0567,
"total": 17182,
"total": 17183,
"unresolvedFailureCount": 9103
},
"warnings": [
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Iluvatar_bi-100 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU hygon_k100-ai 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Kunlunxin_p-800 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Sunrise_pt-200-x1 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Ascend_910-b3 本地统计失败率偏高(≥50%),建议重点关注。",
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Kunlunxin_p-800|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -54164,6 +54211,6 @@
]
},
"storageMode": "decision_state_only",
"summarizedRecords": 17403,
"summarizedRecords": 17404,
"version": 1
}

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -2,24 +2,24 @@
"agentVersion": "2026.09.22.1",
"checksums": {
".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1",
".modelhub_state/architecture_compatibility_blacklist.json": "5ba3776c189e26e5233a8905528da5af5ca0c4d26c2af46e7fc8a8b2cd441827",
".modelhub_state/architecture_compatibility_blacklist.json": "ec0d8a813b44b9ea56e4f0794765a361bac42def5fbc77fbc4aa02077a80f60a",
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
".modelhub_state/market_intelligence.json": "609215e4a7c07c92c86bb4b4e315d8d7e2924cadc42df8754664a6f3ea26e5d0",
".modelhub_state/official_capabilities.json": "e73e31ebab760aba5487e275cdf46e9a1692388151f334f7e473cc750492e517",
".modelhub_state/outcome_checkpoint.json": "b87feb8bba31bc238148b700eca2b3226cbfb6807ddb6caae35a67f13d8fde67",
".modelhub_state/market_intelligence.json": "4e3da9f97d69cfa4687cad45a3ec24f131c13dc5238ba6fc322fd88f9198b481",
".modelhub_state/official_capabilities.json": "af5da451f9ea2bec06e1d119efcff1024d80054ba4fbd6fdf99ca5a9b12fa18e",
".modelhub_state/outcome_checkpoint.json": "bd760ef60d0bba18f3b4dcb9089b6211924c731d2a03e654aef997a9c3eae23d",
".modelhub_state/queue_cleanup_latest.json": "2ab9713b25b7d4ba561d3c77969513a372db37a9de30ef70151fb0193b2bc7ce",
".modelhub_state/recent_outcomes.jsonl": "73f23566b5b4b8c0176d6e15dac2ca91200b082f28f05f39862b89485e9e092d",
".modelhub_state/recovery_active_tasks.jsonl": "3a708902b562db3aba7cfe4e335d8c18ad3dfaefcd4ff80917c2f97889afb24f",
".modelhub_state/recovery_intents.jsonl": "322aeb59a0d1e616d3002abc8e6b0211707597d6ca455310ed4810d506022638",
".modelhub_state/recovery_active_tasks.jsonl": "fd90e051bec2d3c5c0f90312bad065d0f77e0777bdd9d9d541cce15bb954a9d8",
".modelhub_state/recovery_intents.jsonl": "ad2bd8e40d06db05209352cd532efdae5c45ba302e55991d4be0c16e9356a180",
".modelhub_state/routing_intelligence.json": "09cc69960b345ef40104091dc4737ddc82df87d637c93fb6b2c12b4e9f5ead9e",
".modelhub_state/submission_exclusions.jsonl": "1e62f6ba5bd2ba5f2f360bc134b44f7b69190230dd0e274919a87a73ea66761c",
".modelhub_state/worker_crashes.jsonl": "21a892d638884250c04f19e07a433baae065285327484e361c9e5067b0bc17dd",
"ledger/submissions.jsonl": "ac0cf722099f2eda29a035bd4f2090406738ec924a0305895c0e8dfaf5acbd5f",
"outcomes/submissions.jsonl": "02a97bec31b2c0dea3677d3300af47c0b0774b718a9318ee3ac209c27c7a2550"
"outcomes/submissions.jsonl": "7fe75902b9f85bfe3010e6d9e26561f729c486f6cc6f7191161e49257af2e8fe"
},
"generation": 19039,
"generation": 19040,
"phase": "cycle",
"schemaVersion": 1,
"updatedAt": "2026-10-01T02:59:52.959265+00:00",
"updatedAt": "2026-10-01T03:02:04.224316+00:00",
"writerId": "a812d0e025c3474eb16b828a4d987334"
}

View File

@@ -289,7 +289,6 @@
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T13:18:18.977292+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T05:16:13.324611+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4971311", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T14:18:03.258951+00:00", "modelId": "RedHatAI/gemma-2-27b-it-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 30775446656, "estimatedRequiredGiB": 34.419, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 30797371689, "modelscopeLicense": "gemma", "modelscopeParams": 28406776320, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 30797371689}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T06:16:07.269550+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4972172", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T16:51:24.761939+00:00", "modelId": "inceptionai/Jais-2-8B-Chat", "modelProfile": {"architectures": ["Jais2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16180861064, "estimatedRequiredGiB": 18.097, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "jais2", "modelscopeFileSize": 16193187697, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8090401280, "modelscopeTags": ["license:apache-2.0", "model_type:jais2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16193187697}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T08:49:13.744928+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4974217", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T17:54:13.170240+00:00", "modelId": "bharatgenai/Param2-17B-A2.4B-Thinking", "modelProfile": {"architectures": ["Param2MoEForCausalLM"], "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34302758832, "estimatedRequiredGiB": 38.353, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "param2moe", "modelscopeFileSize": 34317809830, "modelscopeLicense": null, "modelscopeParams": 17151125376, "modelscopeTags": ["model_type:param2moe", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:mixture-of-experts", "custom_tag:multilingual", "custom_tag:indian-languages"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 34317809830}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T09:41:34.108703+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976623", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T21:06:45.850274+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T13:01:42.205703+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4979280", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T22:43:47.754909+00:00", "modelId": "neuralmagic/starcoder2-3b-quantized.w8a8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4093158536, "estimatedRequiredGiB": 4.578, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 4096457730, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 3181366272, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4096457730}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T14:38:07.337494+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4980575", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T23:08:35.573457+00:00", "modelId": "neuralmagic/Qwen2-1.5B-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2245331488, "estimatedRequiredGiB": 2.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 2256941144, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1777088000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 2256941144}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T15:06:16.550878+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4980890", "taskType": "text-generation", "verifyResult": null}
@@ -749,9 +748,9 @@
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T02:54:18.466070+00:00", "modelId": "webAI-Official/TwIL-LM3-Pro", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7319517272, "estimatedRequiredGiB": 29.512, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:granite-4.2", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 26407018955}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T18:43:37.348780+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200237", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T02:54:18.466091+00:00", "modelId": "webAI-Official/TwIL-LM2", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5067121544, "estimatedRequiredGiB": 20.413, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:twil-lm", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 18264867356}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T18:43:37.346955+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200239", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T02:54:18.466100+00:00", "modelId": "RWKV/RWKV7-G1k-2.9B-20260930", "modelProfile": {"architectures": ["Rwkv7ForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5896236664, "estimatedRequiredGiB": 6.592, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "rwkv7", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:rwkv7", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:rwkv", "custom_tag:rwkv7", "custom_tag:recurrent", "custom_tag:causal-lm", "custom_tag:conversational"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5898355154}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T18:43:37.247818+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200236", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "mlx-community/K2-Horizon-3.7B-8bit", "modelProfile": {"architectures": ["K2HorizonForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5374667178, "estimatedRequiredGiB": 6.03, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "k2_horizon", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:k2_horizon", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:dense", "custom_tag:k2-horizon"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5395463319}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:16.876280+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200515", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "webAI-Official/TwIL-LM3-Pro", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "645727a8c34425a2c3e191b5407d247f231e0d4e88f0bd3deb19333f39d09ba4", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3892651456, "estimatedRequiredGiB": 29.512, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:granite-4.2", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 26407018955}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:30.049878+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5200516", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "webAI-Official/TwIL-LM2", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "8120661075bb80389dd6d7b11ccbb58a0586dc072989d5570471de11994360d9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2694119616, "estimatedRequiredGiB": 20.413, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:twil-lm", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 18264867356}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:30.051413+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5200518", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "bench-labs/pulvis-v1", "modelProfile": {"architectures": ["PulvisForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 11853976, "estimatedRequiredGiB": 0.123, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "pulvis", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8886720, "modelscopeTags": ["license:apache-2.0", "model_type:pulvis", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:causal-lm", "custom_tag:language-model", "custom_tag:base-model", "custom_tag:pretrained-from-scratch", "custom_tag:small-language-model"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 109820987}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:30.054397+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200517", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "RWKV/RWKV7-G1k-2.9B-20260930", "modelProfile": {"architectures": ["Rwkv7ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5896236664, "estimatedRequiredGiB": 6.592, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "rwkv7", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:rwkv7", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:rwkv", "custom_tag:rwkv7", "custom_tag:recurrent", "custom_tag:causal-lm", "custom_tag:conversational"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5898355154}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:30.056861+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200519", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T03:00:54.465099+00:00", "modelId": "mlx-community/K2-Horizon-3.7B-8bit", "modelProfile": {"architectures": ["K2HorizonForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5374667178, "estimatedRequiredGiB": 6.03, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "k2_horizon", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:k2_horizon", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:dense", "custom_tag:k2-horizon"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5395463319}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:16.876280+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200515", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-10-01T03:00:54.465090+00:00", "modelId": "webAI-Official/TwIL-LM3-Pro", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "645727a8c34425a2c3e191b5407d247f231e0d4e88f0bd3deb19333f39d09ba4", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3892651456, "estimatedRequiredGiB": 29.512, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:granite-4.2", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 26407018955}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:30.049878+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5200516", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-10-01T03:00:54.465128+00:00", "modelId": "webAI-Official/TwIL-LM2", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "8120661075bb80389dd6d7b11ccbb58a0586dc072989d5570471de11994360d9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2694119616, "estimatedRequiredGiB": 20.413, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:twil-lm", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 18264867356}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:30.051413+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5200518", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T03:00:54.465107+00:00", "modelId": "bench-labs/pulvis-v1", "modelProfile": {"architectures": ["PulvisForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 11853976, "estimatedRequiredGiB": 0.123, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "pulvis", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8886720, "modelscopeTags": ["license:apache-2.0", "model_type:pulvis", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:causal-lm", "custom_tag:language-model", "custom_tag:base-model", "custom_tag:pretrained-from-scratch", "custom_tag:small-language-model"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 109820987}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:30.054397+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200517", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T03:00:54.465067+00:00", "modelId": "RWKV/RWKV7-G1k-2.9B-20260930", "modelProfile": {"architectures": ["Rwkv7ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5896236664, "estimatedRequiredGiB": 6.592, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "rwkv7", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:rwkv7", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:rwkv", "custom_tag:rwkv7", "custom_tag:recurrent", "custom_tag:causal-lm", "custom_tag:conversational"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5898355154}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:30.056861+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200519", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "mlx-community/K2-Horizon-3.7B-8bit", "modelProfile": {"architectures": ["K2HorizonForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5374667178, "estimatedRequiredGiB": 6.03, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "k2_horizon", "modelscopeFileSize": 5395463319, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:k2_horizon", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:dense", "custom_tag:k2-horizon"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5395463319}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:16:58.857252+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200707", "taskType": "text-generation", "verifyResult": null}