From 7f513c2fea2fad62f6acfaf0116b2e7c95e80c01 Mon Sep 17 00:00:00 2001 From: ModelHub XC Date: Sat, 13 Jun 2026 20:54:22 +0800 Subject: [PATCH] =?UTF-8?q?=E5=88=9D=E5=A7=8B=E5=8C=96=E9=A1=B9=E7=9B=AE?= =?UTF-8?q?=EF=BC=8C=E7=94=B1ModelHub=20XC=E7=A4=BE=E5=8C=BA=E6=8F=90?= =?UTF-8?q?=E4=BE=9B=E6=A8=A1=E5=9E=8B?= MIME-Version: 1.0 Content-Type: text/plain; charset=UTF-8 Content-Transfer-Encoding: 8bit Model: mradermacher/conversation_llama_esl-GGUF Source: Original Platform --- .gitattributes | 47 ++++++++++++++++++++ README.md | 69 ++++++++++++++++++++++++++++++ conversation_llama_esl.IQ4_XS.gguf | 3 ++ conversation_llama_esl.Q2_K.gguf | 3 ++ conversation_llama_esl.Q3_K_L.gguf | 3 ++ conversation_llama_esl.Q3_K_M.gguf | 3 ++ conversation_llama_esl.Q3_K_S.gguf | 3 ++ conversation_llama_esl.Q4_K_M.gguf | 3 ++ conversation_llama_esl.Q4_K_S.gguf | 3 ++ conversation_llama_esl.Q5_K_M.gguf | 3 ++ conversation_llama_esl.Q5_K_S.gguf | 3 ++ conversation_llama_esl.Q6_K.gguf | 3 ++ conversation_llama_esl.Q8_0.gguf | 3 ++ conversation_llama_esl.f16.gguf | 3 ++ 14 files changed, 152 insertions(+) create mode 100644 .gitattributes create mode 100644 README.md create mode 100644 conversation_llama_esl.IQ4_XS.gguf create mode 100644 conversation_llama_esl.Q2_K.gguf create mode 100644 conversation_llama_esl.Q3_K_L.gguf create mode 100644 conversation_llama_esl.Q3_K_M.gguf create mode 100644 conversation_llama_esl.Q3_K_S.gguf create mode 100644 conversation_llama_esl.Q4_K_M.gguf create mode 100644 conversation_llama_esl.Q4_K_S.gguf create mode 100644 conversation_llama_esl.Q5_K_M.gguf create mode 100644 conversation_llama_esl.Q5_K_S.gguf create mode 100644 conversation_llama_esl.Q6_K.gguf create mode 100644 conversation_llama_esl.Q8_0.gguf create mode 100644 conversation_llama_esl.f16.gguf diff --git a/.gitattributes b/.gitattributes new file mode 100644 index 0000000..a50b469 --- /dev/null +++ b/.gitattributes @@ -0,0 +1,47 @@ +*.7z filter=lfs diff=lfs merge=lfs -text +*.arrow filter=lfs diff=lfs merge=lfs -text +*.bin filter=lfs diff=lfs merge=lfs -text +*.bz2 filter=lfs diff=lfs merge=lfs -text +*.ckpt filter=lfs diff=lfs merge=lfs -text +*.ftz filter=lfs diff=lfs merge=lfs -text +*.gz filter=lfs diff=lfs merge=lfs -text +*.h5 filter=lfs diff=lfs merge=lfs -text +*.joblib filter=lfs diff=lfs merge=lfs -text +*.lfs.* filter=lfs diff=lfs merge=lfs -text +*.mlmodel filter=lfs diff=lfs merge=lfs -text +*.model filter=lfs diff=lfs merge=lfs -text +*.msgpack filter=lfs diff=lfs merge=lfs -text +*.npy filter=lfs diff=lfs merge=lfs -text +*.npz filter=lfs diff=lfs merge=lfs -text +*.onnx filter=lfs diff=lfs merge=lfs -text +*.ot filter=lfs diff=lfs merge=lfs -text +*.parquet filter=lfs diff=lfs merge=lfs -text +*.pb filter=lfs diff=lfs merge=lfs -text +*.pickle filter=lfs diff=lfs merge=lfs -text +*.pkl filter=lfs diff=lfs merge=lfs -text +*.pt filter=lfs diff=lfs merge=lfs -text +*.pth filter=lfs diff=lfs merge=lfs -text +*.rar filter=lfs diff=lfs merge=lfs -text +*.safetensors filter=lfs diff=lfs merge=lfs -text +saved_model/**/* filter=lfs diff=lfs merge=lfs -text +*.tar.* filter=lfs diff=lfs merge=lfs -text +*.tar filter=lfs diff=lfs merge=lfs -text +*.tflite filter=lfs diff=lfs merge=lfs -text +*.tgz filter=lfs diff=lfs merge=lfs -text +*.wasm filter=lfs diff=lfs merge=lfs -text +*.xz filter=lfs diff=lfs merge=lfs -text +*.zip filter=lfs diff=lfs merge=lfs -text +*.zst filter=lfs diff=lfs merge=lfs -text +*tfevents* filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q4_K_S.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q2_K.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.f16.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q8_0.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q6_K.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q3_K_M.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q3_K_S.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q3_K_L.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q4_K_M.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q5_K_S.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.Q5_K_M.gguf filter=lfs diff=lfs merge=lfs -text +conversation_llama_esl.IQ4_XS.gguf filter=lfs diff=lfs merge=lfs -text diff --git a/README.md b/README.md new file mode 100644 index 0000000..526f81c --- /dev/null +++ b/README.md @@ -0,0 +1,69 @@ +--- +base_model: sylviali/conversation_llama_esl +language: +- en +library_name: transformers +mradermacher: + readme_rev: 1 +quantized_by: mradermacher +tags: [] +--- +## About + + + + + + +static quants of https://huggingface.co/sylviali/conversation_llama_esl + + + +***For a convenient overview and download list, visit our [model page for this model](https://hf.tst.eu/model#conversation_llama_esl-GGUF).*** + +weighted/imatrix quants seem not to be available (by me) at this time. If they do not show up a week or so after the static ones, I have probably not planned for them. Feel free to request them by opening a Community Discussion. +## Usage + +If you are unsure how to use GGUF files, refer to one of [TheBloke's +READMEs](https://huggingface.co/TheBloke/KafkaLM-70B-German-V0.1-GGUF) for +more details, including on how to concatenate multi-part files. + +## Provided Quants + +(sorted by size, not necessarily quality. IQ-quants are often preferable over similar sized non-IQ quants) + +| Link | Type | Size/GB | Notes | +|:-----|:-----|--------:|:------| +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q2_K.gguf) | Q2_K | 2.6 | | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q3_K_S.gguf) | Q3_K_S | 3.0 | | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q3_K_M.gguf) | Q3_K_M | 3.4 | lower quality | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q3_K_L.gguf) | Q3_K_L | 3.7 | | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.IQ4_XS.gguf) | IQ4_XS | 3.7 | | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q4_K_S.gguf) | Q4_K_S | 4.0 | fast, recommended | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q4_K_M.gguf) | Q4_K_M | 4.2 | fast, recommended | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q5_K_S.gguf) | Q5_K_S | 4.8 | | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q5_K_M.gguf) | Q5_K_M | 4.9 | | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q6_K.gguf) | Q6_K | 5.6 | very good quality | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.Q8_0.gguf) | Q8_0 | 7.3 | fast, best quality | +| [GGUF](https://huggingface.co/mradermacher/conversation_llama_esl-GGUF/resolve/main/conversation_llama_esl.f16.gguf) | f16 | 13.6 | 16 bpw, overkill | + +Here is a handy graph by ikawrakow comparing some lower-quality quant +types (lower is better): + +![image.png](https://www.nethype.de/huggingface_embed/quantpplgraph.png) + +And here are Artefact2's thoughts on the matter: +https://gist.github.com/Artefact2/b5f810600771265fc1e39442288e8ec9 + +## FAQ / Model Request + +See https://huggingface.co/mradermacher/model_requests for some answers to +questions you might have and/or if you want some other model quantized. + +## Thanks + +I thank my company, [nethype GmbH](https://www.nethype.de/), for letting +me use its servers and providing upgrades to my workstation to enable +this work in my free time. + + diff --git a/conversation_llama_esl.IQ4_XS.gguf b/conversation_llama_esl.IQ4_XS.gguf new file mode 100644 index 0000000..4dbe2fc --- /dev/null +++ b/conversation_llama_esl.IQ4_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:ed5dc47e1e0389cdc3e43d888806eab665f670816a0af2285f466ef93a3578f9 +size 3647517952 diff --git a/conversation_llama_esl.Q2_K.gguf b/conversation_llama_esl.Q2_K.gguf new file mode 100644 index 0000000..fa396e4 --- /dev/null +++ b/conversation_llama_esl.Q2_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b314533ef00c37bfef950f801f9c8bbd2b2e45f9973b4b95ae0ddebc9d5b3d12 +size 2532865280 diff --git a/conversation_llama_esl.Q3_K_L.gguf b/conversation_llama_esl.Q3_K_L.gguf new file mode 100644 index 0000000..eeb5200 --- /dev/null +++ b/conversation_llama_esl.Q3_K_L.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3ffa5180293f55013f96046c853f3c39e1a129290bfdfae03e3ff11d7c6db522 +size 3597112576 diff --git a/conversation_llama_esl.Q3_K_M.gguf b/conversation_llama_esl.Q3_K_M.gguf new file mode 100644 index 0000000..d143b09 --- /dev/null +++ b/conversation_llama_esl.Q3_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:db0fdca0b7f48d6458a021423329be61b154036df6a090632b557d39c14d5fd3 +size 3298006272 diff --git a/conversation_llama_esl.Q3_K_S.gguf b/conversation_llama_esl.Q3_K_S.gguf new file mode 100644 index 0000000..f117592 --- /dev/null +++ b/conversation_llama_esl.Q3_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:70964b7d6b53f110981e926d3687b44b962ebd65ee8ae8bb8e7e6f4ab25e1aad +size 2948306176 diff --git a/conversation_llama_esl.Q4_K_M.gguf b/conversation_llama_esl.Q4_K_M.gguf new file mode 100644 index 0000000..e1e7ac3 --- /dev/null +++ b/conversation_llama_esl.Q4_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:0063d7879bc273266d0ad2276699682211d477670a8c12338bb59728326e8683 +size 4081005824 diff --git a/conversation_llama_esl.Q4_K_S.gguf b/conversation_llama_esl.Q4_K_S.gguf new file mode 100644 index 0000000..b66255f --- /dev/null +++ b/conversation_llama_esl.Q4_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:32d57ab3cc039d6da4037271c2c863aee0705fbdd12e8c4a37297aede1ef6ebc +size 3856741632 diff --git a/conversation_llama_esl.Q5_K_M.gguf b/conversation_llama_esl.Q5_K_M.gguf new file mode 100644 index 0000000..11144de --- /dev/null +++ b/conversation_llama_esl.Q5_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:557ac975ff62b7cdb8d7bbd889d1cc24b282f9621cdf87b729e871e8dc10161b +size 4783158528 diff --git a/conversation_llama_esl.Q5_K_S.gguf b/conversation_llama_esl.Q5_K_S.gguf new file mode 100644 index 0000000..9ef2153 --- /dev/null +++ b/conversation_llama_esl.Q5_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3ed72dcf40969aff354290237f1ebc51ab1616ced1fac075003c8852edd04b12 +size 4651693312 diff --git a/conversation_llama_esl.Q6_K.gguf b/conversation_llama_esl.Q6_K.gguf new file mode 100644 index 0000000..1361624 --- /dev/null +++ b/conversation_llama_esl.Q6_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:77a09a6bd37955e2e957b4ceb9bdb34d3c98e79cc9b34e4c219e6bf4d418725a +size 5529195776 diff --git a/conversation_llama_esl.Q8_0.gguf b/conversation_llama_esl.Q8_0.gguf new file mode 100644 index 0000000..2c89b63 --- /dev/null +++ b/conversation_llama_esl.Q8_0.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:557756e4f25d1b9f7956948bb060f24ce186330f17e2362f1a2acdce6c3ef582 +size 7161091328 diff --git a/conversation_llama_esl.f16.gguf b/conversation_llama_esl.f16.gguf new file mode 100644 index 0000000..85204c5 --- /dev/null +++ b/conversation_llama_esl.f16.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:351b9eb901b7b1be07f31844ad4bcd173e2ee083e5b1286ac0d0c6d54561c8a1 +size 13478106368