From 8d67d3e0b9e2491e4f755b10699edd7c4d118eaf Mon Sep 17 00:00:00 2001 From: ModelHub XC Date: Thu, 1 Oct 2026 21:39:22 +0800 Subject: [PATCH] =?UTF-8?q?=E5=88=9D=E5=A7=8B=E5=8C=96=E9=A1=B9=E7=9B=AE?= =?UTF-8?q?=EF=BC=8C=E7=94=B1ModelHub=20XC=E7=A4=BE=E5=8C=BA=E6=8F=90?= =?UTF-8?q?=E4=BE=9B=E6=A8=A1=E5=9E=8B?= MIME-Version: 1.0 Content-Type: text/plain; charset=UTF-8 Content-Transfer-Encoding: 8bit Model: mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF Source: Original Platform --- .gitattributes | 60 ++++++++++++++++ Qwen2.5-Math-1.5B-Instruct.i1-IQ1_M.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ1_S.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ2_M.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ2_S.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XS.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XXS.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ3_M.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ3_S.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XS.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XXS.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-IQ4_XS.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q2_K.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_L.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_M.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_S.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q4_0.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_4.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_8.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_8_8.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_M.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_S.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_M.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_S.gguf | 3 + Qwen2.5-Math-1.5B-Instruct.i1-Q6_K.gguf | 3 + README.md | 79 +++++++++++++++++++++ imatrix.dat | 3 + 27 files changed, 214 insertions(+) create mode 100644 .gitattributes create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ1_M.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ1_S.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ2_M.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ2_S.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XS.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XXS.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ3_M.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ3_S.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XS.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XXS.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-IQ4_XS.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q2_K.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_L.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_M.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_S.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q4_0.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_4.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_8.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_8_8.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_M.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_S.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_M.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_S.gguf create mode 100644 Qwen2.5-Math-1.5B-Instruct.i1-Q6_K.gguf create mode 100644 README.md create mode 100644 imatrix.dat diff --git a/.gitattributes b/.gitattributes new file mode 100644 index 0000000..6dfd0e6 --- /dev/null +++ b/.gitattributes @@ -0,0 +1,60 @@ +*.7z filter=lfs diff=lfs merge=lfs -text +*.arrow filter=lfs diff=lfs merge=lfs -text +*.bin filter=lfs diff=lfs merge=lfs -text +*.bz2 filter=lfs diff=lfs merge=lfs -text +*.ckpt filter=lfs diff=lfs merge=lfs -text +*.ftz filter=lfs diff=lfs merge=lfs -text +*.gz filter=lfs diff=lfs merge=lfs -text +*.h5 filter=lfs diff=lfs merge=lfs -text +*.joblib filter=lfs diff=lfs merge=lfs -text +*.lfs.* filter=lfs diff=lfs merge=lfs -text +*.mlmodel filter=lfs diff=lfs merge=lfs -text +*.model filter=lfs diff=lfs merge=lfs -text +*.msgpack filter=lfs diff=lfs merge=lfs -text +*.npy filter=lfs diff=lfs merge=lfs -text +*.npz filter=lfs diff=lfs merge=lfs -text +*.onnx filter=lfs diff=lfs merge=lfs -text +*.ot filter=lfs diff=lfs merge=lfs -text +*.parquet filter=lfs diff=lfs merge=lfs -text +*.pb filter=lfs diff=lfs merge=lfs -text +*.pickle filter=lfs diff=lfs merge=lfs -text +*.pkl filter=lfs diff=lfs merge=lfs -text +*.pt filter=lfs diff=lfs merge=lfs -text +*.pth filter=lfs diff=lfs merge=lfs -text +*.rar filter=lfs diff=lfs merge=lfs -text +*.safetensors filter=lfs diff=lfs merge=lfs -text +saved_model/**/* filter=lfs diff=lfs merge=lfs -text +*.tar.* filter=lfs diff=lfs merge=lfs -text +*.tar filter=lfs diff=lfs merge=lfs -text +*.tflite filter=lfs diff=lfs merge=lfs -text +*.tgz filter=lfs diff=lfs merge=lfs -text +*.wasm filter=lfs diff=lfs merge=lfs -text +*.xz filter=lfs diff=lfs merge=lfs -text +*.zip filter=lfs diff=lfs merge=lfs -text +*.zst filter=lfs diff=lfs merge=lfs -text +*tfevents* filter=lfs diff=lfs merge=lfs -text +imatrix.dat filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q2_K.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ3_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XXS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ2_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q6_K.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ1_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ4_XS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XXS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_L.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ2_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_4.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ1_S.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q4_0.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XS.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_8.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_8_8.gguf filter=lfs diff=lfs merge=lfs -text +Qwen2.5-Math-1.5B-Instruct.i1-IQ3_S.gguf filter=lfs diff=lfs merge=lfs -text diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ1_M.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ1_M.gguf new file mode 100644 index 0000000..cefa6a3 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ1_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a47481b06ac49795798a55e232798973b116950f92eb5e382a40f3d222587635 +size 464462048 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ1_S.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ1_S.gguf new file mode 100644 index 0000000..bf3de53 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ1_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:5c627e98616be4fcfe882f646e9084952cedba7056dad17dc69783a79faf5ebe +size 436528352 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_M.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_M.gguf new file mode 100644 index 0000000..a8719ce --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:1d21efdd84eecc87f34d4bd7304bc3e94bc7790073445fc7cae12d9e7291b35d +size 601055456 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_S.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_S.gguf new file mode 100644 index 0000000..1df948a --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b998c3c64b48a1e0ec72b2d76260ce4f1bb379a54a9d08af6265cad5ac805be4 +size 563810528 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XS.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XS.gguf new file mode 100644 index 0000000..cf94f92 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3b52ade6f82643bf4b491dbe88552c7e5bf62064a18542c967dec100ca806db7 +size 550327520 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XXS.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XXS.gguf new file mode 100644 index 0000000..347473c --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XXS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b343c4acbb65af9b051b7d42aba2e9ceeb9d5d205f546eef06b575b660a68493 +size 511018208 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_M.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_M.gguf new file mode 100644 index 0000000..b9fc747 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:4d7a3c57b357f955ac439cfcaec8cb34620f22bdef4160ab9dcf85e0bec7244f +size 776664800 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_S.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_S.gguf new file mode 100644 index 0000000..c393095 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:2264a3532bfbc6c81b57841a99cb7f0a4c3fad8115e380a0670ee51736f08651 +size 762407648 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XS.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XS.gguf new file mode 100644 index 0000000..6ef71c0 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:bec88b0d102d31c08331420c3b3872f35909120737c0e9c35cf053b78e416040 +size 731699936 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XXS.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XXS.gguf new file mode 100644 index 0000000..88e8b0e --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XXS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e995c1e83249e71f9dfe0e351468e96446104de0a6dc3dc056c6e86b93c7c6e1 +size 668793056 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-IQ4_XS.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-IQ4_XS.gguf new file mode 100644 index 0000000..ef9460c --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-IQ4_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:fd9a567b1f2b525e8afa16903241409e396724563ae012e084d39dad32d0632b +size 895732448 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q2_K.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q2_K.gguf new file mode 100644 index 0000000..63453c9 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q2_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:1a7ad9a20a0c5c5693b4cea77982b63e99d50caf2dd5a7199eed56714094cf83 +size 676305632 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_L.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_L.gguf new file mode 100644 index 0000000..6a09980 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_L.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:845f6a9d1f22a212bf6ce1fa682a39ed3b91a8154fa5bc050a0b233749749a6a +size 880163552 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_M.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_M.gguf new file mode 100644 index 0000000..6903cce --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:c5acd1cda1b3cf5726618b6549859894147dfe7b00bbc0365a288ec8b72d29a5 +size 824179424 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_S.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_S.gguf new file mode 100644 index 0000000..39ca090 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b866526da27851cbdefb8773d12c854e49fff8d16f89af0ee67b1f567bcfd6c3 +size 760945376 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0.gguf new file mode 100644 index 0000000..99d13c9 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:2f116b06ebe455162abce34e460933e696f28068092a66c71be7027e011c5742 +size 937536224 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_4.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_4.gguf new file mode 100644 index 0000000..e1c913d --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_4.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:023dbea5f0c03cd5f2fb0b61731b72409902cb944881cdfba37597603d0a1346 +size 934955744 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_8.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_8.gguf new file mode 100644 index 0000000..73e44e9 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_8.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:d4064fc4d23fe8909561178abc776762719df004f1c86b9b3c32c41b1532c511 +size 934955744 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_8_8.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_8_8.gguf new file mode 100644 index 0000000..4a2bf5e --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_8_8.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e93ca92b8c0a597947cb46c9458548b6bc0f1a675afedd3447ab612ef27a9f87 +size 934955744 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_M.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_M.gguf new file mode 100644 index 0000000..99118ba --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:d20226c7ab968a4193300d20773d15ad80bbddd79ad28c69f32eac2f4e29ccb9 +size 986049248 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_S.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_S.gguf new file mode 100644 index 0000000..815f12d --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:dde1c9336d54b970da78ded881660e1d73ee4dc30c19ddf9195b978be46c01b0 +size 940313312 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_M.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_M.gguf new file mode 100644 index 0000000..3fcaba2 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:528cf0a38995d40a45493e226d3095b485d1626b85ce14f94b745afea3de755a +size 1125051104 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_S.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_S.gguf new file mode 100644 index 0000000..b6500a9 --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:604fe835f2163a85962e19ff850b3000170dcfa1a842257d3afe35cfd3a11c36 +size 1098730208 diff --git a/Qwen2.5-Math-1.5B-Instruct.i1-Q6_K.gguf b/Qwen2.5-Math-1.5B-Instruct.i1-Q6_K.gguf new file mode 100644 index 0000000..1e6510f --- /dev/null +++ b/Qwen2.5-Math-1.5B-Instruct.i1-Q6_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:6fb7db6148835eb949752b158ea058def0254172cf2b9ad197a3698052dc1a03 +size 1272740576 diff --git a/README.md b/README.md new file mode 100644 index 0000000..5f6dac9 --- /dev/null +++ b/README.md @@ -0,0 +1,79 @@ +--- +base_model: Qwen/Qwen2.5-Math-1.5B-Instruct +language: +- en +library_name: transformers +license: apache-2.0 +license_link: https://huggingface.co/Qwen/Qwen2.5-Math-1.5B-Instruct/blob/main/LICENSE +quantized_by: mradermacher +tags: +- chat +--- +## About + + + + + + +weighted/imatrix quants of https://huggingface.co/Qwen/Qwen2.5-Math-1.5B-Instruct + + +static quants are available at https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-GGUF +## Usage + +If you are unsure how to use GGUF files, refer to one of [TheBloke's +READMEs](https://huggingface.co/TheBloke/KafkaLM-70B-German-V0.1-GGUF) for +more details, including on how to concatenate multi-part files. + +## Provided Quants + +(sorted by size, not necessarily quality. IQ-quants are often preferable over similar sized non-IQ quants) + +| Link | Type | Size/GB | Notes | +|:-----|:-----|--------:|:------| +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ1_S.gguf) | i1-IQ1_S | 0.5 | for the desperate | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ1_M.gguf) | i1-IQ1_M | 0.6 | mostly desperate | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XXS.gguf) | i1-IQ2_XXS | 0.6 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_XS.gguf) | i1-IQ2_XS | 0.7 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_S.gguf) | i1-IQ2_S | 0.7 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ2_M.gguf) | i1-IQ2_M | 0.7 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XXS.gguf) | i1-IQ3_XXS | 0.8 | lower quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q2_K.gguf) | i1-Q2_K | 0.8 | IQ3_XXS probably better | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_XS.gguf) | i1-IQ3_XS | 0.8 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_S.gguf) | i1-Q3_K_S | 0.9 | IQ3_XS probably better | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_S.gguf) | i1-IQ3_S | 0.9 | beats Q3_K* | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ3_M.gguf) | i1-IQ3_M | 0.9 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_M.gguf) | i1-Q3_K_M | 0.9 | IQ3_S probably better | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q3_K_L.gguf) | i1-Q3_K_L | 1.0 | IQ3_M probably better | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-IQ4_XS.gguf) | i1-IQ4_XS | 1.0 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_4.gguf) | i1-Q4_0_4_4 | 1.0 | fast on arm, low quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_4_8.gguf) | i1-Q4_0_4_8 | 1.0 | fast on arm+i8mm, low quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0_8_8.gguf) | i1-Q4_0_8_8 | 1.0 | fast on arm+sve, low quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q4_0.gguf) | i1-Q4_0 | 1.0 | fast, low quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_S.gguf) | i1-Q4_K_S | 1.0 | optimal size/speed/quality | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q4_K_M.gguf) | i1-Q4_K_M | 1.1 | fast, recommended | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_S.gguf) | i1-Q5_K_S | 1.2 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q5_K_M.gguf) | i1-Q5_K_M | 1.2 | | +| [GGUF](https://huggingface.co/mradermacher/Qwen2.5-Math-1.5B-Instruct-i1-GGUF/resolve/main/Qwen2.5-Math-1.5B-Instruct.i1-Q6_K.gguf) | i1-Q6_K | 1.4 | practically like static Q6_K | + +Here is a handy graph by ikawrakow comparing some lower-quality quant +types (lower is better): + +![image.png](https://www.nethype.de/huggingface_embed/quantpplgraph.png) + +And here are Artefact2's thoughts on the matter: +https://gist.github.com/Artefact2/b5f810600771265fc1e39442288e8ec9 + +## FAQ / Model Request + +See https://huggingface.co/mradermacher/model_requests for some answers to +questions you might have and/or if you want some other model quantized. + +## Thanks + +I thank my company, [nethype GmbH](https://www.nethype.de/), for letting +me use its servers and providing upgrades to my workstation to enable +this work in my free time. Additional thanks to [@nicoboss](https://huggingface.co/nicoboss) for giving me access to his private supercomputer, enabling me to provide many more imatrix quants, at much higher quality, than I would otherwise be able to. + + diff --git a/imatrix.dat b/imatrix.dat new file mode 100644 index 0000000..0e5f264 --- /dev/null +++ b/imatrix.dat @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:0a95413380fc5af75f88ff9278a932069782f27c67c372aa9ae603c44ad17d07 +size 2042201