初始化项目,由ModelHub XC社区提供模型

Model: bartowski/SILMA-9B-Instruct-v1.0-GGUF
Source: Original Platform
This commit is contained in:
ModelHub XC
2026-09-12 05:24:12 +08:00
commit c8a97c5920
28 changed files with 350 additions and 0 deletions

60
.gitattributes vendored Normal file
View File

@@ -0,0 +1,60 @@
*.7z filter=lfs diff=lfs merge=lfs -text
*.arrow filter=lfs diff=lfs merge=lfs -text
*.bin filter=lfs diff=lfs merge=lfs -text
*.bz2 filter=lfs diff=lfs merge=lfs -text
*.ckpt filter=lfs diff=lfs merge=lfs -text
*.ftz filter=lfs diff=lfs merge=lfs -text
*.gz filter=lfs diff=lfs merge=lfs -text
*.h5 filter=lfs diff=lfs merge=lfs -text
*.joblib filter=lfs diff=lfs merge=lfs -text
*.lfs.* filter=lfs diff=lfs merge=lfs -text
*.mlmodel filter=lfs diff=lfs merge=lfs -text
*.model filter=lfs diff=lfs merge=lfs -text
*.msgpack filter=lfs diff=lfs merge=lfs -text
*.npy filter=lfs diff=lfs merge=lfs -text
*.npz filter=lfs diff=lfs merge=lfs -text
*.onnx filter=lfs diff=lfs merge=lfs -text
*.ot filter=lfs diff=lfs merge=lfs -text
*.parquet filter=lfs diff=lfs merge=lfs -text
*.pb filter=lfs diff=lfs merge=lfs -text
*.pickle filter=lfs diff=lfs merge=lfs -text
*.pkl filter=lfs diff=lfs merge=lfs -text
*.pt filter=lfs diff=lfs merge=lfs -text
*.pth filter=lfs diff=lfs merge=lfs -text
*.rar filter=lfs diff=lfs merge=lfs -text
*.safetensors filter=lfs diff=lfs merge=lfs -text
saved_model/**/* filter=lfs diff=lfs merge=lfs -text
*.tar.* filter=lfs diff=lfs merge=lfs -text
*.tar filter=lfs diff=lfs merge=lfs -text
*.tflite filter=lfs diff=lfs merge=lfs -text
*.tgz filter=lfs diff=lfs merge=lfs -text
*.wasm filter=lfs diff=lfs merge=lfs -text
*.xz filter=lfs diff=lfs merge=lfs -text
*.zip filter=lfs diff=lfs merge=lfs -text
*.zst filter=lfs diff=lfs merge=lfs -text
*tfevents* filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q8_0.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q6_K_L.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q6_K.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q5_K_L.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q5_K_M.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q5_K_S.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q4_K_L.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q4_K_M.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q4_K_S.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q4_0_8_8.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q4_0_4_8.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q4_0_4_4.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q4_0.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-IQ4_XS.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q3_K_XL.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q3_K_L.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q3_K_M.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-IQ3_M.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q3_K_S.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-IQ3_XS.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q2_K_L.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-Q2_K.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-IQ2_M.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0-f16.gguf filter=lfs diff=lfs merge=lfs -text
SILMA-9B-Instruct-v1.0.imatrix filter=lfs diff=lfs merge=lfs -text

214
README.md Normal file
View File

@@ -0,0 +1,214 @@
---
base_model: silma-ai/SILMA-9B-Instruct-v1.0
language:
- ar
- en
library_name: transformers
license: gemma
pipeline_tag: text-generation
tags:
- conversational
quantized_by: bartowski
extra_gated_button_content: Acknowledge license
model-index:
- name: SILMA-9B-Instruct-v1.0
results:
- task:
type: text-generation
dataset:
name: MMLU (Arabic)
type: OALL/Arabic_MMLU
metrics:
- type: loglikelihood_acc_norm
value: 52.55
name: acc_norm
source:
url: https://huggingface.co/spaces/OALL/Open-Arabic-LLM-Leaderboard
name: Open Arabic LLM Leaderboard
- task:
type: text-generation
dataset:
name: AlGhafa
type: OALL/AlGhafa-Arabic-LLM-Benchmark-Native
metrics:
- type: loglikelihood_acc_norm
value: 71.85
name: acc_norm
source:
url: https://huggingface.co/spaces/OALL/Open-Arabic-LLM-Leaderboard
name: Open Arabic LLM Leaderboard
- task:
type: text-generation
dataset:
name: ARC Challenge (Arabic)
type: OALL/AlGhafa-Arabic-LLM-Benchmark-Translated
metrics:
- type: loglikelihood_acc_norm
value: 78.19
name: acc_norm
- type: loglikelihood_acc_norm
value: 86
name: acc_norm
- type: loglikelihood_acc_norm
value: 64.05
name: acc_norm
- type: loglikelihood_acc_norm
value: 78.89
name: acc_norm
- type: loglikelihood_acc_norm
value: 47.64
name: acc_norm
- type: loglikelihood_acc_norm
value: 72.93
name: acc_norm
- type: loglikelihood_acc_norm
value: 71.96
name: acc_norm
- type: loglikelihood_acc_norm
value: 75.55
name: acc_norm
- type: loglikelihood_acc_norm
value: 91.26
name: acc_norm
- type: loglikelihood_acc_norm
value: 67.59
name: acc_norm
source:
url: https://huggingface.co/spaces/OALL/Open-Arabic-LLM-Leaderboard
name: Open Arabic LLM Leaderboard
- task:
type: text-generation
dataset:
name: ACVA
type: OALL/ACVA
metrics:
- type: loglikelihood_acc_norm
value: 78.89
name: acc_norm
source:
url: https://huggingface.co/spaces/OALL/Open-Arabic-LLM-Leaderboard
name: Open Arabic LLM Leaderboard
- task:
type: text-generation
dataset:
name: Arabic_EXAMS
type: OALL/Arabic_EXAMS
metrics:
- type: loglikelihood_acc_norm
value: 51.4
name: acc_norm
source:
url: https://huggingface.co/spaces/OALL/Open-Arabic-LLM-Leaderboard
name: Open Arabic LLM Leaderboard
---
## Llamacpp imatrix Quantizations of SILMA-9B-Instruct-v1.0
Using <a href="https://github.com/ggerganov/llama.cpp/">llama.cpp</a> release <a href="https://github.com/ggerganov/llama.cpp/releases/tag/b3634">b3634</a> for quantization.
Original model: https://huggingface.co/silma-ai/SILMA-9B-Instruct-v1.0
All quants made using imatrix option with dataset from [here](https://gist.github.com/bartowski1182/eb213dccb3571f863da82e99418f81e8)
Run them in [LM Studio](https://lmstudio.ai/)
## Prompt format
```
<bos>{system_prompt}<start_of_turn>user
{prompt}<end_of_turn>
<start_of_turn>model
<end_of_turn>
```
## Download a file (not the whole branch) from below:
| Filename | Quant type | File Size | Split | Description |
| -------- | ---------- | --------- | ----- | ----------- |
| [SILMA-9B-Instruct-v1.0-f16.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-f16.gguf) | f16 | 18.49GB | false | Full F16 weights. |
| [SILMA-9B-Instruct-v1.0-Q8_0.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q8_0.gguf) | Q8_0 | 9.83GB | false | Extremely high quality, generally unneeded but max available quant. |
| [SILMA-9B-Instruct-v1.0-Q6_K_L.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q6_K_L.gguf) | Q6_K_L | 7.81GB | false | Uses Q8_0 for embed and output weights. Very high quality, near perfect, *recommended*. |
| [SILMA-9B-Instruct-v1.0-Q6_K.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q6_K.gguf) | Q6_K | 7.59GB | false | Very high quality, near perfect, *recommended*. |
| [SILMA-9B-Instruct-v1.0-Q5_K_L.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q5_K_L.gguf) | Q5_K_L | 6.87GB | false | Uses Q8_0 for embed and output weights. High quality, *recommended*. |
| [SILMA-9B-Instruct-v1.0-Q5_K_M.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q5_K_M.gguf) | Q5_K_M | 6.65GB | false | High quality, *recommended*. |
| [SILMA-9B-Instruct-v1.0-Q5_K_S.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q5_K_S.gguf) | Q5_K_S | 6.48GB | false | High quality, *recommended*. |
| [SILMA-9B-Instruct-v1.0-Q4_K_L.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q4_K_L.gguf) | Q4_K_L | 5.98GB | false | Uses Q8_0 for embed and output weights. Good quality, *recommended*. |
| [SILMA-9B-Instruct-v1.0-Q4_K_M.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q4_K_M.gguf) | Q4_K_M | 5.76GB | false | Good quality, default size for must use cases, *recommended*. |
| [SILMA-9B-Instruct-v1.0-Q4_K_S.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q4_K_S.gguf) | Q4_K_S | 5.48GB | false | Slightly lower quality with more space savings, *recommended*. |
| [SILMA-9B-Instruct-v1.0-Q4_0.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q4_0.gguf) | Q4_0 | 5.46GB | false | Legacy format, generally not worth using over similarly sized formats |
| [SILMA-9B-Instruct-v1.0-Q4_0_8_8.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q4_0_8_8.gguf) | Q4_0_8_8 | 5.44GB | false | Optimized for ARM and CPU inference, much faster than Q4_0 at similar quality. |
| [SILMA-9B-Instruct-v1.0-Q4_0_4_8.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q4_0_4_8.gguf) | Q4_0_4_8 | 5.44GB | false | Optimized for ARM and CPU inference, much faster than Q4_0 at similar quality. |
| [SILMA-9B-Instruct-v1.0-Q4_0_4_4.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q4_0_4_4.gguf) | Q4_0_4_4 | 5.44GB | false | Optimized for ARM and CPU inference, much faster than Q4_0 at similar quality. |
| [SILMA-9B-Instruct-v1.0-Q3_K_XL.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q3_K_XL.gguf) | Q3_K_XL | 5.35GB | false | Uses Q8_0 for embed and output weights. Lower quality but usable, good for low RAM availability. |
| [SILMA-9B-Instruct-v1.0-IQ4_XS.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-IQ4_XS.gguf) | IQ4_XS | 5.18GB | false | Decent quality, smaller than Q4_K_S with similar performance, *recommended*. |
| [SILMA-9B-Instruct-v1.0-Q3_K_L.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q3_K_L.gguf) | Q3_K_L | 5.13GB | false | Lower quality but usable, good for low RAM availability. |
| [SILMA-9B-Instruct-v1.0-Q3_K_M.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q3_K_M.gguf) | Q3_K_M | 4.76GB | false | Low quality. |
| [SILMA-9B-Instruct-v1.0-IQ3_M.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-IQ3_M.gguf) | IQ3_M | 4.49GB | false | Medium-low quality, new method with decent performance comparable to Q3_K_M. |
| [SILMA-9B-Instruct-v1.0-Q3_K_S.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q3_K_S.gguf) | Q3_K_S | 4.34GB | false | Low quality, not recommended. |
| [SILMA-9B-Instruct-v1.0-IQ3_XS.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-IQ3_XS.gguf) | IQ3_XS | 4.14GB | false | Lower quality, new method with decent performance, slightly better than Q3_K_S. |
| [SILMA-9B-Instruct-v1.0-Q2_K_L.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q2_K_L.gguf) | Q2_K_L | 4.03GB | false | Uses Q8_0 for embed and output weights. Very low quality but surprisingly usable. |
| [SILMA-9B-Instruct-v1.0-Q2_K.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-Q2_K.gguf) | Q2_K | 3.81GB | false | Very low quality but surprisingly usable. |
| [SILMA-9B-Instruct-v1.0-IQ2_M.gguf](https://huggingface.co/bartowski/SILMA-9B-Instruct-v1.0-GGUF/blob/main/SILMA-9B-Instruct-v1.0-IQ2_M.gguf) | IQ2_M | 3.43GB | false | Relatively low quality, uses SOTA techniques to be surprisingly usable. |
## Embed/output weights
Some of these quants (Q3_K_XL, Q4_K_L etc) are the standard quantization method with the embeddings and output weights quantized to Q8_0 instead of what they would normally default to.
Some say that this improves the quality, others don't notice any difference. If you use these models PLEASE COMMENT with your findings. I would like feedback that these are actually used and useful so I don't keep uploading quants no one is using.
Thanks!
## Credits
Thank you kalomaze and Dampf for assistance in creating the imatrix calibration dataset
Thank you ZeroWw for the inspiration to experiment with embed/output
## Downloading using huggingface-cli
First, make sure you have hugginface-cli installed:
```
pip install -U "huggingface_hub[cli]"
```
Then, you can target the specific file you want:
```
huggingface-cli download bartowski/SILMA-9B-Instruct-v1.0-GGUF --include "SILMA-9B-Instruct-v1.0-Q4_K_M.gguf" --local-dir ./
```
If the model is bigger than 50GB, it will have been split into multiple files. In order to download them all to a local folder, run:
```
huggingface-cli download bartowski/SILMA-9B-Instruct-v1.0-GGUF --include "SILMA-9B-Instruct-v1.0-Q8_0/*" --local-dir ./
```
You can either specify a new local-dir (SILMA-9B-Instruct-v1.0-Q8_0) or download them all in place (./)
## Which file should I choose?
A great write up with charts showing various performances is provided by Artefact2 [here](https://gist.github.com/Artefact2/b5f810600771265fc1e39442288e8ec9)
The first thing to figure out is how big a model you can run. To do this, you'll need to figure out how much RAM and/or VRAM you have.
If you want your model running as FAST as possible, you'll want to fit the whole thing on your GPU's VRAM. Aim for a quant with a file size 1-2GB smaller than your GPU's total VRAM.
If you want the absolute maximum quality, add both your system RAM and your GPU's VRAM together, then similarly grab a quant with a file size 1-2GB Smaller than that total.
Next, you'll need to decide if you want to use an 'I-quant' or a 'K-quant'.
If you don't want to think too much, grab one of the K-quants. These are in format 'QX_K_X', like Q5_K_M.
If you want to get more into the weeds, you can check out this extremely useful feature chart:
[llama.cpp feature matrix](https://github.com/ggerganov/llama.cpp/wiki/Feature-matrix)
But basically, if you're aiming for below Q4, and you're running cuBLAS (Nvidia) or rocBLAS (AMD), you should look towards the I-quants. These are in format IQX_X, like IQ3_M. These are newer and offer better performance for their size.
These I-quants can also be used on CPU and Apple Metal, but will be slower than their K-quant equivalent, so speed vs performance is a tradeoff you'll have to decide.
The I-quants are *not* compatible with Vulcan, which is also AMD, so if you have an AMD card double check if you're using the rocBLAS build or the Vulcan build. At the time of writing this, LM Studio has a preview with ROCm support, and other inference engines have specific builds for ROCm.
Want to support my work? Visit my ko-fi page here: https://ko-fi.com/bartowski

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:b9ea964c4c1ea3bbeaab918a59dbfec3d098bd64a0bc43697ac60a492a35caac
size 3434669280

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:6d3cd4352ad57ee65b11dfc5be5858361eb9c844510ec17ff31834b6512b5dd5
size 4494615776

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:013ad93479ca69147d05b9891786305cbde3027a4880d323a2832b78e19dfc35
size 4144989408

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:4e1e18f470b4457dab7330e4db19e6cd671448a8f6849acdba0202fe5e9dc2cd
size 5183030496

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:c03b59cc5b6e532d6d69e5089969d027a91c2df1211f8188927d2d69e7ec642c
size 3805398240

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:731d08af59b309dca7fa4c3f8296340b6b60ff91273b8848b2c6aefcf23660ab
size 4027606240

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:ce8e979d9c86a006b02d6e713873712e3ead10bd09f7849c175787218f2aecda
size 5132453088

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:94f70d5fa6b03782771c563f802fd7beca074dc68eecb6955623d641db514c64
size 4761781472

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:2e9dfedbb6a8921bb834762293edc19c16b5ac45874d5cf1f125b2753265117a
size 4337665248

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:8ad93a39949f1878bee182b0cdeac51aeb347e4e759ae73d2ff1989ddaf30434
size 5354661088

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:150d3cc6a6c56acfd17f423dc4e2c2ed110c64dfd41b404a32f71f0cfe54f9ff
size 5459199200

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:715296ad19fecad63a332c8731cc810c8dfaad5516937f70118637de0b48c728
size 5443142880

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:9c495f445278267306c83976e3af117deaec3cc6d8c1698b41538adba746b203
size 5443142880

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:b59b8948e464a9080051a4fb87919420a5f660e95640dc1772668bc810a28279
size 5443142880

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:8ca8de92edec13c5cb56e706fe9b744b4bbf39961fbd61f19a0dae048574b9d1
size 5983266016

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:59a5a63a42ff27fbd3b3978014dba8711baeeeaf3dad8c5a31bcc8ff8e31943c
size 5761058016

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:0e008ae2c4f96b810c0bd86343e668c6a1440ccd034e2169a7a77b3207915850
size 5478925536

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:fe6148ee9ec2c1b8131e0ec4eb6682141af7438222634ec4873cb91b2963f138
size 6869574880

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:26324310a187ac698da24af7fc5ad49d37ed3d7b6f246dbdd348859e799093e7
size 6647366880

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:ebbe42ef0533b005bc837d0bef700ca2bb869f2c4ebc59841de80032794fdac3
size 6483592416

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:5d2e7b69bf63344546eb8c993fa83fc18bdaa0f5f30c2de21607284a2b39bad4
size 7589070048

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:bce246eede2cc0a97b3c66024453e881403a119860844b39b58cd721ff6d4e5f
size 7811278048

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:9f2cc4f3560bd79fbb542a21df940b930d97320b6b66006353cc90280c60231f
size 9827149024

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:9115197f9ee6d486bb3b29b32db7605c246489584a3720cbb5546e392a43af73
size 18490680256

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:f4aeb729e3db99ba7217a930601431d72261d8f9e916dbaa77678363f1da8613
size 6116900

1
configuration.json Normal file
View File

@@ -0,0 +1 @@
{"framework": "pytorch", "task": "text-generation", "allow_remote": true}