commit 59ebb7e71f49447458ea009c924a9f7e5aec3b30 Author: ModelHub XC Date: Wed Sep 30 13:07:19 2026 +0800 初始化项目,由ModelHub XC社区提供模型 Model: mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF Source: Original Platform diff --git a/.gitattributes b/.gitattributes new file mode 100644 index 0000000..39e1414 --- /dev/null +++ b/.gitattributes @@ -0,0 +1,60 @@ +*.7z filter=lfs diff=lfs merge=lfs -text +*.arrow filter=lfs diff=lfs merge=lfs -text +*.bin filter=lfs diff=lfs merge=lfs -text +*.bz2 filter=lfs diff=lfs merge=lfs -text +*.ckpt filter=lfs diff=lfs merge=lfs -text +*.ftz filter=lfs diff=lfs merge=lfs -text +*.gz filter=lfs diff=lfs merge=lfs -text +*.h5 filter=lfs diff=lfs merge=lfs -text +*.joblib filter=lfs diff=lfs merge=lfs -text +*.lfs.* filter=lfs diff=lfs merge=lfs -text +*.mlmodel filter=lfs diff=lfs merge=lfs -text +*.model filter=lfs diff=lfs merge=lfs -text +*.msgpack filter=lfs diff=lfs merge=lfs -text +*.npy filter=lfs diff=lfs merge=lfs -text +*.npz filter=lfs diff=lfs merge=lfs -text +*.onnx filter=lfs diff=lfs merge=lfs -text +*.ot filter=lfs diff=lfs merge=lfs -text +*.parquet filter=lfs diff=lfs merge=lfs -text +*.pb filter=lfs diff=lfs merge=lfs -text +*.pickle filter=lfs diff=lfs merge=lfs -text +*.pkl filter=lfs diff=lfs merge=lfs -text +*.pt filter=lfs diff=lfs merge=lfs -text +*.pth filter=lfs diff=lfs merge=lfs -text +*.rar filter=lfs diff=lfs merge=lfs -text +*.safetensors filter=lfs diff=lfs merge=lfs -text +saved_model/**/* filter=lfs diff=lfs merge=lfs -text +*.tar.* filter=lfs diff=lfs merge=lfs -text +*.tar filter=lfs diff=lfs merge=lfs -text +*.tflite filter=lfs diff=lfs merge=lfs -text +*.tgz filter=lfs diff=lfs merge=lfs -text +*.wasm filter=lfs diff=lfs merge=lfs -text +*.xz filter=lfs diff=lfs merge=lfs -text +*.zip filter=lfs diff=lfs merge=lfs -text +*.zst filter=lfs diff=lfs merge=lfs -text +*tfevents* filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.imatrix.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q2_K.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ3_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XXS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ4_NL.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ2_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q6_K.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q2_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ4_XS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ1_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XXS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_L.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ2_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ1_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q4_0.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-Q4_1.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5Math-IVON-SFT-7B.i1-IQ3_S.gguf filter=lfs diff=lfs merge=lfs -text diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ1_M.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ1_M.gguf new file mode 100644 index 0000000..d5b6774 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ1_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:0e4f7e6e075abbd6e43bc438adbfbaccc7f5b31955b6727501761089db816fc2 +size 2042197088 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ1_S.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ1_S.gguf new file mode 100644 index 0000000..a839cd8 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ1_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:250ed10dd8b633c02f14bb622bf3c39dc633ebfb160dbeeae88fbea64d55523e +size 1903668320 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_M.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_M.gguf new file mode 100644 index 0000000..7c30a41 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:2ba2b84769b7dfe718e38f1edff34ea312c823dc3338aa9c345bb27906d642c9 +size 2780343392 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_S.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_S.gguf new file mode 100644 index 0000000..3227d68 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:36f5ea1c044858d600c2a5d59807a39e24988c1b2d057e24da50d1fa1cb74473 +size 2595638368 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XS.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XS.gguf new file mode 100644 index 0000000..ffb0b4c --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:34fd1ffb8f3f25e9128765205a197bd543a8452565d9c26037f7e701cce7b4bc +size 2469022816 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XXS.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XXS.gguf new file mode 100644 index 0000000..b486429 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XXS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:39823bfa4cdaf6fa8fe33c18a77db4e7cac80c991bf228e73fedd527867fd857 +size 2273078368 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_M.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_M.gguf new file mode 100644 index 0000000..2cb77b0 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:fc311ac2038b81bf8b04c243b2ccd5ff69fd7ef8709e2b7a5c1d639f73a700ad +size 3574013024 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_S.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_S.gguf new file mode 100644 index 0000000..b87472f --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e69ae00fd9f3ed596d08b42ecd2a384d6bbca2cd5006a84e20f2d50704574bbe +size 3499193440 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XS.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XS.gguf new file mode 100644 index 0000000..0e626d9 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:bd19941492134777db43ab5f52b689c3c729f1e381dd83a5bb2ae21ed38a3229 +size 3346256992 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XXS.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XXS.gguf new file mode 100644 index 0000000..8b92851 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XXS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:4e5ba971fa4f377fdbd16f30f8308b98b7973b2aea318886c29b0d557532d264 +size 3114515552 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ4_NL.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ4_NL.gguf new file mode 100644 index 0000000..3d804fc --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ4_NL.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:55517d52c3a723be3433cb3688c532ef563211eba39971edd333dfcd468bba1f +size 4437814368 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-IQ4_XS.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-IQ4_XS.gguf new file mode 100644 index 0000000..ffb7dee --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-IQ4_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:5cd06ad55ea0ba03722edee617b657f7dd7ba8598b3ce22e22fd42ce9ddd7b2a +size 4218473568 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q2_K.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q2_K.gguf new file mode 100644 index 0000000..29b233b --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q2_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:2124a86a229fcc4aa08c2710fafb9b3e2c185fd954f901c6d90e0ca184a4c09c +size 3015941216 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q2_K_S.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q2_K_S.gguf new file mode 100644 index 0000000..07b03e7 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q2_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:957fc922cb274063f1fdd1b27f5e5e80097fff4027f479381d084035387ba3b5 +size 2834074720 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_L.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_L.gguf new file mode 100644 index 0000000..37d8434 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_L.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:238383cb7fcde011cdf1191ccf807908ddf540e28518f34a691e488020d2c556 +size 4088460384 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_M.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_M.gguf new file mode 100644 index 0000000..97c5e5c --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:d8b11d692375340528d1e94395f70936fcd6864030f0567b64eb32abf871ab55 +size 3808392288 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_S.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_S.gguf new file mode 100644 index 0000000..0537df2 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:f6f8a5558eb41299182ce3c2a2e328a62b8e6ee3fe102cc44006fb7642dbf622 +size 3492369504 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q4_0.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q4_0.gguf new file mode 100644 index 0000000..cb2b5ac --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q4_0.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:7a7edccbe507c911f3f3d2e0f82ea760e8316ffdc7b4f97aa547b1605c01a145 +size 4444122208 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q4_1.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q4_1.gguf new file mode 100644 index 0000000..ddd3946 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q4_1.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:119fd078b042756e81d679ec27eb22998473071573e2e3ae7dc8b086a4bb4fd1 +size 4873284704 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_M.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_M.gguf new file mode 100644 index 0000000..add413a --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:9173c4bd4bf728fbc0fc420245c8dcf862b5d894e0a93c03ead7510a8f91a067 +size 4683074656 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_S.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_S.gguf new file mode 100644 index 0000000..33b4247 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:24f00d14708028d2d2df49693ac97ddeb38949137e30e92fba9b119919bf2d22 +size 4457770080 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_M.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_M.gguf new file mode 100644 index 0000000..b8bba20 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:60a735cf18b0a067fe17b7d3724a84d94beb4c653246ef778e693884b9000b0b +size 5444832352 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_S.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_S.gguf new file mode 100644 index 0000000..bb13369 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:806867851821e59690a8c4a9d22c10d7c3568321d52d896d7da09c8298335086 +size 5315177568 diff --git a/Qwen2.5Math-IVON-SFT-7B.i1-Q6_K.gguf b/Qwen2.5Math-IVON-SFT-7B.i1-Q6_K.gguf new file mode 100644 index 0000000..1056adc --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.i1-Q6_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:f0d03ca5b757efb6fc37976c168b8a1c1d7cc6b625f4ff227c2c7b9a85bb3054 +size 6254199904 diff --git a/Qwen2.5Math-IVON-SFT-7B.imatrix.gguf b/Qwen2.5Math-IVON-SFT-7B.imatrix.gguf new file mode 100644 index 0000000..6ead8d2 --- /dev/null +++ b/Qwen2.5Math-IVON-SFT-7B.imatrix.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:ecabc82796103928f0e2ddc9e137b0ef7541a3be101dfab5dd8b0be511d26066 +size 4560352 diff --git a/README.md b/README.md new file mode 100644 index 0000000..c377007 --- /dev/null +++ b/README.md @@ -0,0 +1,92 @@ +--- +base_model: BayesRL/Qwen2.5Math-IVON-SFT-7B +language: +- en +library_name: transformers +license: apache-2.0 +mradermacher: + readme_rev: 1 +quantized_by: mradermacher +tags: +- ivon +- variational-learning +- sft +- 3po +- math +- reasoning +--- +## About + + + + + + + + + +weighted/imatrix quants of https://huggingface.co/BayesRL/Qwen2.5Math-IVON-SFT-7B + + + +***For a convenient overview and download list, visit our [model page for this model](https://hf.tst.eu/model#Qwen2.5Math-IVON-SFT-7B-i1-GGUF).*** + +static quants are available at https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-GGUF +## Usage + +If you are unsure how to use GGUF files, refer to one of [TheBloke's +READMEs](https://huggingface.co/TheBloke/KafkaLM-70B-German-V0.1-GGUF) for +more details, including on how to concatenate multi-part files. + +## Provided Quants + +(sorted by size, not necessarily quality. IQ-quants are often preferable over similar sized non-IQ quants) + +| Link | Type | Size/GB | Notes | +|:-----|:-----|--------:|:------| +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.imatrix.gguf) | imatrix | 0.1 | imatrix file (for creating your own quants) | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ1_S.gguf) | i1-IQ1_S | 2.0 | for the desperate | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ1_M.gguf) | i1-IQ1_M | 2.1 | mostly desperate | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XXS.gguf) | i1-IQ2_XXS | 2.4 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_XS.gguf) | i1-IQ2_XS | 2.6 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_S.gguf) | i1-IQ2_S | 2.7 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ2_M.gguf) | i1-IQ2_M | 2.9 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q2_K_S.gguf) | i1-Q2_K_S | 2.9 | very low quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q2_K.gguf) | i1-Q2_K | 3.1 | IQ3_XXS probably better | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XXS.gguf) | i1-IQ3_XXS | 3.2 | lower quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_XS.gguf) | i1-IQ3_XS | 3.4 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_S.gguf) | i1-Q3_K_S | 3.6 | IQ3_XS probably better | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_S.gguf) | i1-IQ3_S | 3.6 | beats Q3_K* | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ3_M.gguf) | i1-IQ3_M | 3.7 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_M.gguf) | i1-Q3_K_M | 3.9 | IQ3_S probably better | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q3_K_L.gguf) | i1-Q3_K_L | 4.2 | IQ3_M probably better | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ4_XS.gguf) | i1-IQ4_XS | 4.3 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-IQ4_NL.gguf) | i1-IQ4_NL | 4.5 | prefer IQ4_XS | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q4_0.gguf) | i1-Q4_0 | 4.5 | fast, low quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_S.gguf) | i1-Q4_K_S | 4.6 | optimal size/speed/quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q4_K_M.gguf) | i1-Q4_K_M | 4.8 | fast, recommended | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q4_1.gguf) | i1-Q4_1 | 5.0 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_S.gguf) | i1-Q5_K_S | 5.4 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q5_K_M.gguf) | i1-Q5_K_M | 5.5 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5Math-IVON-SFT-7B-i1-GGUF/resolve/main/Qwen2.5Math-IVON-SFT-7B.i1-Q6_K.gguf) | i1-Q6_K | 6.4 | practically like static Q6_K | + +Here is a handy graph by ikawrakow comparing some lower-quality quant +types (lower is better): + +![image.png](https://www.nethype.de/huggingface_embed/quantpplgraph.png) + +And here are Artefact2's thoughts on the matter: +https://gist.github.com/Artefact2/b5f810600771265fc1e39442288e8ec9 + +## FAQ / Model Request + +See https://huggingface.co/mradermacher/model_requests for some answers to +questions you might have and/or if you want some other model quantized. + +## Thanks + +I thank my company, [nethype GmbH](https://www.nethype.de/), for letting +me use its servers and providing upgrades to my workstation to enable +this work in my free time. Additional thanks to [@nicoboss](https://huggingface.co/nicoboss) for giving me access to his private supercomputer, enabling me to provide many more imatrix quants, at much higher quality, than I would otherwise be able to. + +