commit caaec696961d66841bec0d06252509d95c587fb7 Author: ModelHub XC Date: Fri Jul 17 23:06:15 2026 +0800 初始化项目,由ModelHub XC社区提供模型 Model: mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF Source: Original Platform diff --git a/.gitattributes b/.gitattributes new file mode 100644 index 0000000..5ccadd1 --- /dev/null +++ b/.gitattributes @@ -0,0 +1,50 @@ +*.7z filter=lfs diff=lfs merge=lfs -text +*.arrow filter=lfs diff=lfs merge=lfs -text +*.bin filter=lfs diff=lfs merge=lfs -text +*.bz2 filter=lfs diff=lfs merge=lfs -text +*.ckpt filter=lfs diff=lfs merge=lfs -text +*.ftz filter=lfs diff=lfs merge=lfs -text +*.gz filter=lfs diff=lfs merge=lfs -text +*.h5 filter=lfs diff=lfs merge=lfs -text +*.joblib filter=lfs diff=lfs merge=lfs -text +*.lfs.* filter=lfs diff=lfs merge=lfs -text +*.mlmodel filter=lfs diff=lfs merge=lfs -text +*.model filter=lfs diff=lfs merge=lfs -text +*.msgpack filter=lfs diff=lfs merge=lfs -text +*.npy filter=lfs diff=lfs merge=lfs -text +*.npz filter=lfs diff=lfs merge=lfs -text +*.onnx filter=lfs diff=lfs merge=lfs -text +*.ot filter=lfs diff=lfs merge=lfs -text +*.parquet filter=lfs diff=lfs merge=lfs -text +*.pb filter=lfs diff=lfs merge=lfs -text +*.pickle filter=lfs diff=lfs merge=lfs -text +*.pkl filter=lfs diff=lfs merge=lfs -text +*.pt filter=lfs diff=lfs merge=lfs -text +*.pth filter=lfs diff=lfs merge=lfs -text +*.rar filter=lfs diff=lfs merge=lfs -text +*.safetensors filter=lfs diff=lfs merge=lfs -text +saved_model/**/* filter=lfs diff=lfs merge=lfs -text +*.tar.* filter=lfs diff=lfs merge=lfs -text +*.tar filter=lfs diff=lfs merge=lfs -text +*.tflite filter=lfs diff=lfs merge=lfs -text +*.tgz filter=lfs diff=lfs merge=lfs -text +*.wasm filter=lfs diff=lfs merge=lfs -text +*.xz filter=lfs diff=lfs merge=lfs -text +*.zip filter=lfs diff=lfs merge=lfs -text +*.zst filter=lfs diff=lfs merge=lfs -text +*tfevents* filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q2_K.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q8_0.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_L.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q6_K.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.f16.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_XS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ4_XS.gguf filter=lfs diff=lfs merge=lfs -text diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_M.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_M.gguf new file mode 100644 index 0000000..1e84859 --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:50147f748dcf00c75b7822c00b64d7b2f64a2d89b496e6dd5ca1bf94fc8d14be +size 3572216224 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_S.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_S.gguf new file mode 100644 index 0000000..553dd2a --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:df64451c3fd3823a5d5eac0f68032e5b6649dc12acf252b585f47d60cce82af3 +size 3497396640 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_XS.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_XS.gguf new file mode 100644 index 0000000..617f673 --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:05224a6e9f141d86ab246d5d83e3d3edae2e7d491071faea95b7fe794f6d1d6d +size 3344460192 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ4_XS.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ4_XS.gguf new file mode 100644 index 0000000..344b228 --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ4_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:187e6450305874dd8bab4e2ad4573c0e7529f2c3c76196140683cad57beeb4ec +size 4248357440 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q2_K.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q2_K.gguf new file mode 100644 index 0000000..2e32450 --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q2_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:7533368ba2bb3faabe985dcf668679d63fb51effdf82b018f916cbde7ef2d410 +size 3014289632 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_L.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_L.gguf new file mode 100644 index 0000000..46add15 --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_L.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:dafe07f5e811f6340dc814cf1d85872f3b04a0531978123fb4f179edb6a35fe5 +size 4086663584 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_M.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_M.gguf new file mode 100644 index 0000000..47887cc --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a20d265cc0db5b7ed121d02bcfc4cae81f61509e33876e7d02fc83ac896263cd +size 3806595488 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_S.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_S.gguf new file mode 100644 index 0000000..a2e739c --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:5d44a058256d4dca42faa1d82f0bb5c50fe6c9b238637a370ffa78eadc6dfc46 +size 3490572704 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_M.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_M.gguf new file mode 100644 index 0000000..595c67a --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:ff1f40a5661ed21e1725a927c1e8d3ea24da563618fe3a51603e2bccb379e841 +size 4681087904 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_S.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_S.gguf new file mode 100644 index 0000000..d96a5f8 --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:86ed91e7e810e19cbf3be7f0205303447ae69a980f678e9b62969e55e25ccb96 +size 4455783328 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_M.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_M.gguf new file mode 100644 index 0000000..c1c8eae --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:6055e3d04133bc6e276180dae04e6be7eaf78d5862c076dda8394d514799f9ae +size 5442666848 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_S.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_S.gguf new file mode 100644 index 0000000..59e7a43 --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a423f20d4a8761cf91134a66034834f31bcc5293e05230eec62f0eec1a8843e3 +size 5313012064 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q6_K.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q6_K.gguf new file mode 100644 index 0000000..e7e90ca --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q6_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:2455f0efd12896e6d59dcb79fe1823a3096604609c5467fa2e1d8ad8a8dad8ff +size 6251844480 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q8_0.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q8_0.gguf new file mode 100644 index 0000000..11fd49f --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q8_0.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:d5177f9c56796b519c2676759e3f78d082ebef4b0b33f4904590be2e6367c1ac +size 8095478208 diff --git a/Qwen2.5-Coder-7B-Chat-Instruct-TIES.f16.gguf b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.f16.gguf new file mode 100644 index 0000000..50e26ea --- /dev/null +++ b/Qwen2.5-Coder-7B-Chat-Instruct-TIES.f16.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:55dfc41f92ddbabd50a3b1b9a12be0ae412d643bd8d62068411a9343c6c7fbc5 +size 15232124928 diff --git a/README.md b/README.md new file mode 100644 index 0000000..d0b7047 --- /dev/null +++ b/README.md @@ -0,0 +1,69 @@ +--- +base_model: BenevolenceMessiah/Qwen2.5-Coder-7B-Chat-Instruct-TIES +language: +- en +library_name: transformers +quantized_by: mradermacher +tags: +- mergekit +- merge +--- +## About + + + + + + +static quants of https://huggingface.co/BenevolenceMessiah/Qwen2.5-Coder-7B-Chat-Instruct-TIES + + +weighted/imatrix quants are available at https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-i1-GGUF +## Usage + +If you are unsure how to use GGUF files, refer to one of [TheBloke's +READMEs](https://huggingface.co/TheBloke/KafkaLM-70B-German-V0.1-GGUF) for +more details, including on how to concatenate multi-part files. + +## Provided Quants + +(sorted by size, not necessarily quality. IQ-quants are often preferable over similar sized non-IQ quants) + +| Link | Type | Size/GB | Notes | +|:-----|:-----|--------:|:------| +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q2_K.gguf) | Q2_K | 3.1 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_XS.gguf) | IQ3_XS | 3.4 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_S.gguf) | Q3_K_S | 3.6 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_S.gguf) | IQ3_S | 3.6 | beats Q3_K* | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ3_M.gguf) | IQ3_M | 3.7 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_M.gguf) | Q3_K_M | 3.9 | lower quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q3_K_L.gguf) | Q3_K_L | 4.2 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.IQ4_XS.gguf) | IQ4_XS | 4.3 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_S.gguf) | Q4_K_S | 4.6 | fast, recommended | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q4_K_M.gguf) | Q4_K_M | 4.8 | fast, recommended | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_S.gguf) | Q5_K_S | 5.4 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q5_K_M.gguf) | Q5_K_M | 5.5 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q6_K.gguf) | Q6_K | 6.4 | very good quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.Q8_0.gguf) | Q8_0 | 8.2 | fast, best quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Coder-7B-Chat-Instruct-TIES-GGUF/resolve/main/Qwen2.5-Coder-7B-Chat-Instruct-TIES.f16.gguf) | f16 | 15.3 | 16 bpw, overkill | + +Here is a handy graph by ikawrakow comparing some lower-quality quant +types (lower is better): + +![image.png](https://www.nethype.de/huggingface_embed/quantpplgraph.png) + +And here are Artefact2's thoughts on the matter: +https://gist.github.com/Artefact2/b5f810600771265fc1e39442288e8ec9 + +## FAQ / Model Request + +See https://huggingface.co/mradermacher/model_requests for some answers to +questions you might have and/or if you want some other model quantized. + +## Thanks + +I thank my company, [nethype GmbH](https://www.nethype.de/), for letting +me use its servers and providing upgrades to my workstation to enable +this work in my free time. + +