state: generation 7055 (result)

This commit is contained in:
2026-09-14 06:48:14 +00:00
parent b2a352453e
commit 207595a00a
5 changed files with 15 additions and 13 deletions

View File

@@ -622,3 +622,4 @@
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-13T22:28:30.242556+00:00", "modelId": "XHToken/Spark-X2.5-4B-INT8", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4454282952, "estimatedRequiredGiB": 4.995, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 4469587176, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4469587176}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-13T14:24:08.911334+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4829350", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-14T03:29:33.013591+00:00", "modelId": "XHToken/Spark-X2.5-4B-INT8", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4454282952, "estimatedRequiredGiB": 4.995, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 4469587176, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4469587176}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-13T19:22:45.622511+00:00", "targetGpu": "Biren_166m", "taskId": "4834315", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-14T06:17:54.916608+00:00", "modelId": "XHToken/Spark-X2.5-4B-INT8", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4454282952, "estimatedRequiredGiB": 4.995, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 4469587176, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4469587176}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-13T22:17:43.666489+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4836986", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "Edge0/Edge0-35B-A3B-preview", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 19689863445, "estimatedRequiredGiB": 22.059, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 19738195355, "modelscopeLicense": "apache-2.0", "modelscopeParams": 5419330688, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 19738195355}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-14T06:48:00.583349+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4844610", "taskType": "text-generation", "verifyResult": null}