From 4b3545d27ac5697e4456a72d3bb4cc3bf1df6c09 Mon Sep 17 00:00:00 2001 From: ModelHub XC Date: Mon, 20 Jul 2026 16:06:09 +0800 Subject: [PATCH] =?UTF-8?q?=E5=88=9D=E5=A7=8B=E5=8C=96=E9=A1=B9=E7=9B=AE?= =?UTF-8?q?=EF=BC=8C=E7=94=B1ModelHub=20XC=E7=A4=BE=E5=8C=BA=E6=8F=90?= =?UTF-8?q?=E4=BE=9B=E6=A8=A1=E5=9E=8B?= MIME-Version: 1.0 Content-Type: text/plain; charset=UTF-8 Content-Transfer-Encoding: 8bit Model: mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF Source: Original Platform --- .gitattributes | 47 ++++++++++++ ...umma-ThaiLLM-qwen3-8b-it-2.0.0.IQ4_XS.gguf | 3 + Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q2_K.gguf | 3 + ...umma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_L.gguf | 3 + ...umma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_M.gguf | 3 + ...umma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_S.gguf | 3 + ...umma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_M.gguf | 3 + ...umma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_S.gguf | 3 + ...umma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_M.gguf | 3 + ...umma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_S.gguf | 3 + Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q6_K.gguf | 3 + Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q8_0.gguf | 3 + Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.f16.gguf | 3 + README.md | 72 +++++++++++++++++++ 14 files changed, 155 insertions(+) create mode 100644 .gitattributes create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.IQ4_XS.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q2_K.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_L.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_M.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_S.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_M.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_S.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_M.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_S.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q6_K.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q8_0.gguf create mode 100644 Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.f16.gguf create mode 100644 README.md diff --git a/.gitattributes b/.gitattributes new file mode 100644 index 0000000..fd6b7e8 --- /dev/null +++ b/.gitattributes @@ -0,0 +1,47 @@ +*.7z filter=lfs diff=lfs merge=lfs -text +*.arrow filter=lfs diff=lfs merge=lfs -text +*.bin filter=lfs diff=lfs merge=lfs -text +*.bz2 filter=lfs diff=lfs merge=lfs -text +*.ckpt filter=lfs diff=lfs merge=lfs -text +*.ftz filter=lfs diff=lfs merge=lfs -text +*.gz filter=lfs diff=lfs merge=lfs -text +*.h5 filter=lfs diff=lfs merge=lfs -text +*.joblib filter=lfs diff=lfs merge=lfs -text +*.lfs.* filter=lfs diff=lfs merge=lfs -text +*.mlmodel filter=lfs diff=lfs merge=lfs -text +*.model filter=lfs diff=lfs merge=lfs -text +*.msgpack filter=lfs diff=lfs merge=lfs -text +*.npy filter=lfs diff=lfs merge=lfs -text +*.npz filter=lfs diff=lfs merge=lfs -text +*.onnx filter=lfs diff=lfs merge=lfs -text +*.ot filter=lfs diff=lfs merge=lfs -text +*.parquet filter=lfs diff=lfs merge=lfs -text +*.pb filter=lfs diff=lfs merge=lfs -text +*.pickle filter=lfs diff=lfs merge=lfs -text +*.pkl filter=lfs diff=lfs merge=lfs -text +*.pt filter=lfs diff=lfs merge=lfs -text +*.pth filter=lfs diff=lfs merge=lfs -text +*.rar filter=lfs diff=lfs merge=lfs -text +*.safetensors filter=lfs diff=lfs merge=lfs -text +saved_model/**/* filter=lfs diff=lfs merge=lfs -text +*.tar.* filter=lfs diff=lfs merge=lfs -text +*.tar filter=lfs diff=lfs merge=lfs -text +*.tflite filter=lfs diff=lfs merge=lfs -text +*.tgz filter=lfs diff=lfs merge=lfs -text +*.wasm filter=lfs diff=lfs merge=lfs -text +*.xz filter=lfs diff=lfs merge=lfs -text +*.zip filter=lfs diff=lfs merge=lfs -text +*.zst filter=lfs diff=lfs merge=lfs -text +*tfevents* filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q2_K.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.f16.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q8_0.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q6_K.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_L.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_S.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_M.gguf filter=lfs diff=lfs merge=lfs -text +Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.IQ4_XS.gguf filter=lfs diff=lfs merge=lfs -text diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.IQ4_XS.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.IQ4_XS.gguf new file mode 100644 index 0000000..b4a68a6 --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.IQ4_XS.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:9b6758f4009a9051791da3442f4ba489f57db0a2ac1e7258a4ab23691c4a963b +size 4593296768 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q2_K.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q2_K.gguf new file mode 100644 index 0000000..7f2f4c6 --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q2_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:7f67874849067bb19160c23d93c089bde8c695a0b14284f8aa11eebb5da81184 +size 3281732992 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_L.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_L.gguf new file mode 100644 index 0000000..c613cde --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_L.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:329512b2a4f184cc51eede948dee1a5cf3ebe6126c5c1d1e7b29947591d6d619 +size 4431394176 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_M.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_M.gguf new file mode 100644 index 0000000..cc7af93 --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:973bde4bcb4cd8f8d1dd86191f59f914c927506b3bee5085497791f5c3098ec5 +size 4124161408 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_S.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_S.gguf new file mode 100644 index 0000000..9c11b52 --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e8c37fa6030caaf9aaa8ffee2e68ecce2172989a390b64d3035c08aabd9f4131 +size 3769611648 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_M.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_M.gguf new file mode 100644 index 0000000..bfc2c73 --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:d3ce2cc927882ee9a6b7a44f89b9e28f6e705372f9d5ff24e34e1b409a82e779 +size 5027784064 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_S.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_S.gguf new file mode 100644 index 0000000..9f4b86b --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:ef76ed2dd1ab63824139883f29458b0e2d0d490ba67ecca49888587e85883ed2 +size 4802012544 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_M.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_M.gguf new file mode 100644 index 0000000..3414b9b --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_M.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:c154b2ec35057491484fba96683782ebc131d368e3ce81447d48a4b30303691f +size 5851112832 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_S.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_S.gguf new file mode 100644 index 0000000..8a2d069 --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_S.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:5f249895a9190bd04fae9ca851078b47aa6a3d89562b8b6930dc2a7338018aa5 +size 5720761728 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q6_K.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q6_K.gguf new file mode 100644 index 0000000..150b766 --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q6_K.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:52912cae5e933060c66b3254f902068f219c48d2bd6d1185aa11a10f830f7a56 +size 6725899648 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q8_0.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q8_0.gguf new file mode 100644 index 0000000..e3183cc --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q8_0.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3017523c86e5b05c550c9beb1f6ed8d2f9db758eb55a9ddf69cca0c376770a83 +size 8709518720 diff --git a/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.f16.gguf b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.f16.gguf new file mode 100644 index 0000000..fe948c8 --- /dev/null +++ b/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.f16.gguf @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:6bda80612751bfe86c38b6c0ce79cd2a9c9d33ce4b1175b9f7946be468fe38c4 +size 16388044160 diff --git a/README.md b/README.md new file mode 100644 index 0000000..86c5d03 --- /dev/null +++ b/README.md @@ -0,0 +1,72 @@ +--- +base_model: nectec/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0 +language: +- en +library_name: transformers +license: apache-2.0 +mradermacher: + readme_rev: 1 +quantized_by: mradermacher +--- +## About + + + + + + + + + +static quants of https://huggingface.co/nectec/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0 + + + +***For a convenient overview and download list, visit our [model page for this model](https://hf.tst.eu/model#Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF).*** + +weighted/imatrix quants seem not to be available (by me) at this time. If they do not show up a week or so after the static ones, I have probably not planned for them. Feel free to request them by opening a Community Discussion. +## Usage + +If you are unsure how to use GGUF files, refer to one of [TheBloke's +READMEs](https://huggingface.co/TheBloke/KafkaLM-70B-German-V0.1-GGUF) for +more details, including on how to concatenate multi-part files. + +## Provided Quants + +(sorted by size, not necessarily quality. IQ-quants are often preferable over similar sized non-IQ quants) + +| Link | Type | Size/GB | Notes | +|:-----|:-----|--------:|:------| +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q2_K.gguf) | Q2_K | 3.4 | | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_S.gguf) | Q3_K_S | 3.9 | | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_M.gguf) | Q3_K_M | 4.2 | lower quality | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q3_K_L.gguf) | Q3_K_L | 4.5 | | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.IQ4_XS.gguf) | IQ4_XS | 4.7 | | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_S.gguf) | Q4_K_S | 4.9 | fast, recommended | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q4_K_M.gguf) | Q4_K_M | 5.1 | fast, recommended | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_S.gguf) | Q5_K_S | 5.8 | | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q5_K_M.gguf) | Q5_K_M | 6.0 | | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q6_K.gguf) | Q6_K | 6.8 | very good quality | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.Q8_0.gguf) | Q8_0 | 8.8 | fast, best quality | +| [GGUF](https://huggingface.co/mradermacher/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0-GGUF/resolve/main/Pathumma-ThaiLLM-qwen3-8b-it-2.0.0.f16.gguf) | f16 | 16.5 | 16 bpw, overkill | + +Here is a handy graph by ikawrakow comparing some lower-quality quant +types (lower is better): + +![image.png](https://www.nethype.de/huggingface_embed/quantpplgraph.png) + +And here are Artefact2's thoughts on the matter: +https://gist.github.com/Artefact2/b5f810600771265fc1e39442288e8ec9 + +## FAQ / Model Request + +See https://huggingface.co/mradermacher/model_requests for some answers to +questions you might have and/or if you want some other model quantized. + +## Thanks + +I thank my company, [nethype GmbH](https://www.nethype.de/), for letting +me use its servers and providing upgrades to my workstation to enable +this work in my free time. + +