state: generation 9426 (result)
This commit is contained in:
@@ -538,3 +538,4 @@
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "barozp/Qwen3.8-Whittle-MoE-27B-MLX-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15155811051, "estimatedRequiredGiB": 16.961, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 15176196372, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4210722304, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-lm", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:reasoning", "custom_tag:opus-distill", "custom_tag:whittle", "custom_tag:expert-pruning", "custom_tag:moe", "custom_tag:router-healing", "custom_tag:quantized"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15176196372}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-18T18:34:31.787488+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4938618", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": null, "modelId": "barozp/Qwen3.8-Whittle-MoE-27B-MLX-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForCausalLM"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15155811051, "estimatedRequiredGiB": 16.961, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 15176196372, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4210722304, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-lm", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:reasoning", "custom_tag:opus-distill", "custom_tag:whittle", "custom_tag:expert-pruning", "custom_tag:moe", "custom_tag:router-healing", "custom_tag:quantized"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15176196372}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-18T18:41:55.617980+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4938690", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "barozp/Qwen3.8-Whittle-MoE-27B-MLX-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15155811051, "estimatedRequiredGiB": 16.961, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 15176196372, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4210722304, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-lm", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:reasoning", "custom_tag:opus-distill", "custom_tag:whittle", "custom_tag:expert-pruning", "custom_tag:moe", "custom_tag:router-healing", "custom_tag:quantized"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15176196372}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-18T18:58:28.435554+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4938823", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": null, "modelId": "barozp/Qwen3.8-Whittle-MoE-27B-MLX-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15155811051, "estimatedRequiredGiB": 16.961, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 15176196372, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4210722304, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-lm", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:reasoning", "custom_tag:opus-distill", "custom_tag:whittle", "custom_tag:expert-pruning", "custom_tag:moe", "custom_tag:router-healing", "custom_tag:quantized"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15176196372}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-18T19:15:30.844293+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4939007", "taskType": "text-generation", "verifyResult": null}
|
||||
|
||||
Reference in New Issue
Block a user