state: generation 9381 (intent)

This commit is contained in:
2026-09-18 17:46:20 +00:00
parent d2ccde4d43
commit 8978e204e8
6 changed files with 1588 additions and 1824 deletions

View File

@@ -524,7 +524,7 @@
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-18T16:31:25.093486+00:00", "modelId": "prism-ml/Ternary-Bonsai-2-27B-mlx-2bit", "modelProfile": {"architectures": [], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8595477990, "estimatedRequiredGiB": 9.621, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "prism_hadamard_qwen35", "modelscopeFileSize": 8608734989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2589078768, "modelscopeTags": ["license:apache-2.0", "model_type:prism_hadamard_qwen35", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:ternary", "custom_tag:2-bit", "custom_tag:mlx", "custom_tag:cuda", "custom_tag:metal", "custom_tag:on-device", "custom_tag:hybrid-attention", "custom_tag:prismml", "custom_tag:bonsai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8608734989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-18T08:30:26.717917+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4931821", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-18T16:50:06.681047+00:00", "modelId": "prism-ml/Ternary-Bonsai-2-27B-mlx-2bit", "modelProfile": {"architectures": [], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8595477990, "estimatedRequiredGiB": 9.621, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "prism_hadamard_qwen35", "modelscopeFileSize": 8608734989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2589078768, "modelscopeTags": ["license:apache-2.0", "model_type:prism_hadamard_qwen35", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:ternary", "custom_tag:2-bit", "custom_tag:mlx", "custom_tag:cuda", "custom_tag:metal", "custom_tag:on-device", "custom_tag:hybrid-attention", "custom_tag:prismml", "custom_tag:bonsai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8608734989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-18T08:47:48.285664+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4932114", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-18T17:26:39.586895+00:00", "modelId": "rengensheng/Ternary-Bonsai-2-27B-gguf", "modelProfile": {"architectures": [], "configFingerprint": "ae25b7990c7dee096113ec65c0e54ee3ae00205524523882157b7ed10068496b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 629246976, "estimatedRequiredGiB": 8.757, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 7835440764, "modelscopeLicense": "apache-2.0", "modelscopeParams": 460730096, "modelscopeTags": ["license:apache-2.0", "library:gguf", "task:text-generation", "custom_tag:ternary", "custom_tag:2-bit", "custom_tag:gguf", "custom_tag:llama-cpp", "custom_tag:cuda", "custom_tag:metal", "custom_tag:on-device", "custom_tag:hybrid-attention", "custom_tag:prismml", "custom_tag:bonsai", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7835440764}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-18T09:17:39.257829+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4932473", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "rengensheng/Ternary-Bonsai-2-27B-gguf", "modelProfile": {"architectures": [], "configFingerprint": "96fb2eca023ffaf5f3a7f1befecfe8080483f3ff3ed3b3fce6364bcc1f40fd92", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 629246976, "estimatedRequiredGiB": 8.757, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 7835440764, "modelscopeLicense": "apache-2.0", "modelscopeParams": 460730096, "modelscopeTags": ["license:apache-2.0", "library:gguf", "task:text-generation", "custom_tag:ternary", "custom_tag:2-bit", "custom_tag:gguf", "custom_tag:llama-cpp", "custom_tag:cuda", "custom_tag:metal", "custom_tag:on-device", "custom_tag:hybrid-attention", "custom_tag:prismml", "custom_tag:bonsai", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7835440764}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-18T09:36:29.093079+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4932714", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-18T17:44:35.896312+00:00", "modelId": "rengensheng/Ternary-Bonsai-2-27B-gguf", "modelProfile": {"architectures": [], "configFingerprint": "96fb2eca023ffaf5f3a7f1befecfe8080483f3ff3ed3b3fce6364bcc1f40fd92", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 629246976, "estimatedRequiredGiB": 8.757, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 7835440764, "modelscopeLicense": "apache-2.0", "modelscopeParams": 460730096, "modelscopeTags": ["license:apache-2.0", "library:gguf", "task:text-generation", "custom_tag:ternary", "custom_tag:2-bit", "custom_tag:gguf", "custom_tag:llama-cpp", "custom_tag:cuda", "custom_tag:metal", "custom_tag:on-device", "custom_tag:hybrid-attention", "custom_tag:prismml", "custom_tag:bonsai", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7835440764}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-18T09:36:29.093079+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4932714", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "rengensheng/Ternary-Bonsai-2-27B-gguf", "modelProfile": {"architectures": [], "configFingerprint": "23334d2dcbe57c09d582b0146945adcb58473c0d3e0dc0da7ef6d516e24eb645", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 629246976, "estimatedRequiredGiB": 8.757, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 7835440764, "modelscopeLicense": "apache-2.0", "modelscopeParams": 460730096, "modelscopeTags": ["license:apache-2.0", "library:gguf", "task:text-generation", "custom_tag:ternary", "custom_tag:2-bit", "custom_tag:gguf", "custom_tag:llama-cpp", "custom_tag:cuda", "custom_tag:metal", "custom_tag:on-device", "custom_tag:hybrid-attention", "custom_tag:prismml", "custom_tag:bonsai", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7835440764}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-18T09:54:06.949570+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4932877", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "rengensheng/Ternary-Bonsai-2-27B-gguf", "modelProfile": {"architectures": [], "configFingerprint": "064f77ce69ad1df6d2506a353bfb02950549c28ea28db152c22e15f209e87be7", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 629246976, "estimatedRequiredGiB": 8.757, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 7835440764, "modelscopeLicense": "apache-2.0", "modelscopeParams": 460730096, "modelscopeTags": ["license:apache-2.0", "library:gguf", "task:text-generation", "custom_tag:ternary", "custom_tag:2-bit", "custom_tag:gguf", "custom_tag:llama-cpp", "custom_tag:cuda", "custom_tag:metal", "custom_tag:on-device", "custom_tag:hybrid-attention", "custom_tag:prismml", "custom_tag:bonsai", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7835440764}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-18T10:11:22.196791+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4933209", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "ProCreations/Ternary-Bonsai-2-27B-MTP", "modelProfile": {"architectures": [], "configFingerprint": "ea1658f2c8c61659f2c0f3d125b718e7445bb2df5854decf74f202aac3336048", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7657489728, "estimatedRequiredGiB": 9.567, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 8560440422, "modelscopeLicense": "apache-2.0", "modelscopeParams": 424699392, "modelscopeTags": ["license:apache-2.0", "library:safetensors", "library:gguf", "library:pytorch", "task:text-generation", "custom_tag:gguf", "custom_tag:qwen3_5", "custom_tag:mtp", "custom_tag:speculative-decoding", "custom_tag:bonsai", "custom_tag:ternary", "custom_tag:blackwell", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8560440422}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-18T13:15:39.090237+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4935339", "taskType": "text-generation", "verifyResult": null}