state: generation 8658 (result)
This commit is contained in:
@@ -595,3 +595,4 @@
|
||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": null, "modelId": "rapid-mlx/Qwen3.8-27B-4bit-MTP-MLX", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293475486, "estimatedRequiredGiB": 18.239, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16320427112, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:rapid-mlx", "custom_tag:qwen3.8", "custom_tag:qwen3.8-27b", "custom_tag:4-bit", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16320427112}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-17T02:02:00.949464+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4906948", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": null, "modelId": "ornith-ai/Ornith-1.5-9B-MLX-6bit", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7276346365, "estimatedRequiredGiB": 8.156, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 7298158256, "modelscopeLicense": null, "modelscopeParams": 1959473664, "modelscopeTags": ["model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7298158256}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-17T02:02:00.946894+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4906950", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "RentedNoodle/Qwen3.8-27B-OrcaRouter-GSQ-RCO-IQ3_XXS-GGUF", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "bb269a9a8b4c1eeb18e3ef78bedcf638c3d1821a581b2f17113cde2fcd8be6c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 629247648, "estimatedRequiredGiB": 25.142, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 22496243465, "modelscopeLicense": "apache-2.0", "modelscopeParams": 27320697856, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:gguf", "task:text-generation", "custom_tag:qwen3", "custom_tag:qwen3.8", "custom_tag:orcarouter", "custom_tag:gsq", "custom_tag:rco", "custom_tag:gguf", "custom_tag:quantized", "custom_tag:speculative-decoding", "custom_tag:mtp", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-use", "custom_tag:16gb-vram", "custom_tag:image-text-to-text", "custom_tag:vision", "custom_tag:multimodal", "custom_tag:mmproj", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 22496243465}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-17T09:19:02.317218+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4912678", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "RentedNoodle/Qwen3.8-27B-OrcaRouter-GSQ-RCO-IQ3_XXS-GGUF", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "7df61fc66119e362f4b324c308b44e18bd70c6e8415d3cf79ad822bc9e6377c3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 629247648, "estimatedRequiredGiB": 25.142, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 22496243465, "modelscopeLicense": "apache-2.0", "modelscopeParams": 27320697856, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:gguf", "task:text-generation", "custom_tag:qwen3", "custom_tag:qwen3.8", "custom_tag:orcarouter", "custom_tag:gsq", "custom_tag:rco", "custom_tag:gguf", "custom_tag:quantized", "custom_tag:speculative-decoding", "custom_tag:mtp", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-use", "custom_tag:16gb-vram", "custom_tag:image-text-to-text", "custom_tag:vision", "custom_tag:multimodal", "custom_tag:mmproj", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 22496243465}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-17T09:36:01.072109+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4912861", "taskType": "text-generation", "verifyResult": null}
|
||||
|
||||
Reference in New Issue
Block a user