初始化项目,由ModelHub XC社区提供模型
Model: Zynerji/Ektome-SmolLM2-360Mi-PristinelyUncensored Source: Original Platform
This commit is contained in:
43
.gitattributes
vendored
Normal file
43
.gitattributes
vendored
Normal file
@@ -0,0 +1,43 @@
|
|||||||
|
*.7z filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.arrow filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.bin filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.bz2 filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.ckpt filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.ftz filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.gz filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.h5 filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.joblib filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.lfs.* filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.mlmodel filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.model filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.msgpack filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.npy filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.npz filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.onnx filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.ot filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.parquet filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.pb filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.pickle filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.pkl filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.pt filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.pth filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.rar filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.safetensors filter=lfs diff=lfs merge=lfs -text
|
||||||
|
saved_model/**/* filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.tar.* filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.tar filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.tflite filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.tgz filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.wasm filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.xz filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.zip filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*.zst filter=lfs diff=lfs merge=lfs -text
|
||||||
|
*tfevents* filter=lfs diff=lfs merge=lfs -text
|
||||||
|
hero.png filter=lfs diff=lfs merge=lfs -text
|
||||||
|
Ektome-SmolLM2-360Mi-IQ3_M.gguf filter=lfs diff=lfs merge=lfs -text
|
||||||
|
Ektome-SmolLM2-360Mi-IQ4_XS.gguf filter=lfs diff=lfs merge=lfs -text
|
||||||
|
Ektome-SmolLM2-360Mi-Q4_K_M.gguf filter=lfs diff=lfs merge=lfs -text
|
||||||
|
Ektome-SmolLM2-360Mi-Q5_K_M.gguf filter=lfs diff=lfs merge=lfs -text
|
||||||
|
Ektome-SmolLM2-360Mi-Q6_K.gguf filter=lfs diff=lfs merge=lfs -text
|
||||||
|
Ektome-SmolLM2-360Mi-Q8_0.gguf filter=lfs diff=lfs merge=lfs -text
|
||||||
|
imatrix.dat filter=lfs diff=lfs merge=lfs -text
|
||||||
3
Ektome-SmolLM2-360Mi-IQ3_M.gguf
Normal file
3
Ektome-SmolLM2-360Mi-IQ3_M.gguf
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:69cc480b7f09998ce4369ad33632ceaa16ac683932cfe64b1069d29a3d0ab0df
|
||||||
|
size 224894688
|
||||||
3
Ektome-SmolLM2-360Mi-IQ4_XS.gguf
Normal file
3
Ektome-SmolLM2-360Mi-IQ4_XS.gguf
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:98dd101d9f6145c124ef305437c44348512b5e1323e031b59dae1d4fba672e13
|
||||||
|
size 226661088
|
||||||
3
Ektome-SmolLM2-360Mi-Q4_K_M.gguf
Normal file
3
Ektome-SmolLM2-360Mi-Q4_K_M.gguf
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:bc1c9e5f6f350db07fa7e621036b287ef0caf28944037974d2c1732d58d5d586
|
||||||
|
size 270590688
|
||||||
3
Ektome-SmolLM2-360Mi-Q5_K_M.gguf
Normal file
3
Ektome-SmolLM2-360Mi-Q5_K_M.gguf
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:d35a926aa9b077911b657a51b1e62a25f35feadc423eb02bd624da3d3e8867af
|
||||||
|
size 289944288
|
||||||
3
Ektome-SmolLM2-360Mi-Q6_K.gguf
Normal file
3
Ektome-SmolLM2-360Mi-Q6_K.gguf
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:2ab1c439438e5a57034d1e544c297b18a8e01e087c5352abd97886fe9ff6226e
|
||||||
|
size 367358688
|
||||||
3
Ektome-SmolLM2-360Mi-Q8_0.gguf
Normal file
3
Ektome-SmolLM2-360Mi-Q8_0.gguf
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:0a1366809e7b1ba7c1795c01fe03d7d598a435a1bc9211aba3c13d5e0277086a
|
||||||
|
size 386405088
|
||||||
121
README.md
Normal file
121
README.md
Normal file
@@ -0,0 +1,121 @@
|
|||||||
|
---
|
||||||
|
license: apache-2.0
|
||||||
|
base_model: HuggingFaceTB/SmolLM2-360M-Instruct
|
||||||
|
tags:
|
||||||
|
- uncensored
|
||||||
|
- abliterated
|
||||||
|
- uncertified
|
||||||
|
- ektome
|
||||||
|
- sphragis
|
||||||
|
- smollm2
|
||||||
|
language:
|
||||||
|
- en
|
||||||
|
pipeline_tag: text-generation
|
||||||
|
---
|
||||||
|
|
||||||
|

|
||||||
|
|
||||||
|
# Ektome-SmolLM2-360Mi-PristinelyUncensored
|
||||||
|
|
||||||
|
**Uncensored. No n=2800 certificate has been run for this model, so no capability-retention claim is made.**
|
||||||
|
|
||||||
|
> **compliance 0.89 to 1.00 at capability -0.010 vs pristine.**
|
||||||
|
|
||||||
|
$$\colorbox{black}{$\color{white}
|
||||||
|
\begin{array}{ll}
|
||||||
|
\textsf{EKTOME CERTIFICATE} & {} \\
|
||||||
|
\textsf{capability} & \textsf{NOT} \\
|
||||||
|
\textsf{margin} & 3\% \\
|
||||||
|
\textsf{items } n & 200 \\
|
||||||
|
\textsf{worst-axis bound} & -0.010 \\
|
||||||
|
\textsf{compliance} & 0.89 \rightarrow 1.00 \\
|
||||||
|
\end{array}$}$$
|
||||||
|
|
||||||
|
> ### ⚠️ Not certified
|
||||||
|
>
|
||||||
|
> No n=2800 paired certificate exists for this model. Any numbers below are
|
||||||
|
> point estimates with no confidence interval.
|
||||||
|
|
||||||
|
|
||||||
|
📄 **[Read the whitepaper (PDF)](./whitepaper.pdf)** — full method, receipts and certification.
|
||||||
|
The PDF is the authoritative document: dark-typeset, with the complete derivation, the
|
||||||
|
per-axis certificate and the reproducibility hashes.
|
||||||
|
|
||||||
|
---
|
||||||
|
|
||||||
|
## Why this exists
|
||||||
|
|
||||||
|
Standard abliteration removes a coarse *refusal direction* that is entangled with
|
||||||
|
directions carrying knowledge and reasoning. The result is an uncensored model with a
|
||||||
|
capability tax that is **almost never measured**.
|
||||||
|
|
||||||
|
Ektomē (ἐκτομή, *excision*) isolates and removes only the refusal-**specific**
|
||||||
|
component, leaving general helpfulness intact, and does so norm-preservingly on the
|
||||||
|
pristine model — no training, no distillation, no damage to repair. The extraction
|
||||||
|
depth is selected per model by automated search against measured compliance.
|
||||||
|
|
||||||
|
The estimator, excision operator and depth-selection procedure are proprietary.
|
||||||
|
What is published here is the **measured outcome** and the evidence for it, which you
|
||||||
|
can verify against the artifacts in this repo.
|
||||||
|
|
||||||
|
## The receipt
|
||||||
|
|
||||||
|
| model | capability (MMLU-val) ↑ | compliance on harmful ↑ |
|
||||||
|
|---|---|---|
|
||||||
|
| pristine `SmolLM2-360M-Instruct` | 0.247 | 0.890 |
|
||||||
|
| **Ektomē (this model)** | **0.258** | **1.000** |
|
||||||
|
|
||||||
|
These are **point estimates with no confidence interval** — which is precisely why the next section exists.
|
||||||
|
|
||||||
|
|
||||||
|
## The certificate
|
||||||
|
|
||||||
|
Capability retention is certified by a paired non-inferiority test against the pristine
|
||||||
|
model (exact McNemar, Holm-corrected, one-sided bootstrap bound on the drop $d$ vs a
|
||||||
|
3% margin):
|
||||||
|
|
||||||
|
| axis | n | ref | cand | d upper | verdict |
|
||||||
|
|---|---|---|---|---|---|
|
||||||
|
| MMLU-val (POINT ESTIMATE, n=200, no CI) | 200 | 0.247 | 0.258 | -0.010 | UNCERTIFIED |
|
||||||
|
|
||||||
|
|
||||||
|
**Overall: NOT CERTIFIED - no n=2800 paired test has been run for this model**
|
||||||
|
|
||||||
|
Reproducible from `seed=20260726`, pack `sha256:7bbaff877146e081…`.
|
||||||
|
|
||||||
|
### Generation health checks
|
||||||
|
|
||||||
|
| metric | pristine | Ektomē | n |
|
||||||
|
|---|---|---|---|
|
||||||
|
| `foreign_rate` | 0.0 | 0.0 | 15 |
|
||||||
|
| `degen_rate` | 0.1 | 0.2 | 15 |
|
||||||
|
| `instr_pass` | 0.6 | 0.6 | 5 |
|
||||||
|
|
||||||
|
These are **degeneration guards** — code-switching, babbling, format compliance —
|
||||||
|
not capability measures. Note the sample sizes: they detect a broken model, not a
|
||||||
|
subtly weaker one. The capability claim rests on the certificate above, not here.
|
||||||
|
|
||||||
|
|
||||||
|
## Quantisations
|
||||||
|
|
||||||
|
_No quantisations have been published for this model yet — bf16 weights only._
|
||||||
|
|
||||||
|
|
||||||
|
## Limitations
|
||||||
|
|
||||||
|
The certificate bounds **capability retention only**. It does not certify safety, factual
|
||||||
|
accuracy, or fitness for any purpose. Axes marked *inconclusive* are honestly
|
||||||
|
under-powered, and the certificate states the $n$ needed to resolve them. Compliance uses
|
||||||
|
a keyword classifier — a proxy that evasive phrasing can fool. **This model is uncensored
|
||||||
|
by construction: it will not refuse, and you are accountable for what you do with it.**
|
||||||
|
|
||||||
|
## Citation
|
||||||
|
|
||||||
|
```bibtex
|
||||||
|
@software{ektome_Ektome-SmolLM2-360Mi-PristinelyUncensored,
|
||||||
|
title = {Ektome-SmolLM2-360Mi-PristinelyUncensored},
|
||||||
|
author = {Zynerji},
|
||||||
|
year = {2026},
|
||||||
|
url = {https://huggingface.co/Zynerji/Ektome-SmolLM2-360Mi-PristinelyUncensored}
|
||||||
|
}
|
||||||
|
```
|
||||||
133
cert_SmolLM2-360Mi.json
Normal file
133
cert_SmolLM2-360Mi.json
Normal file
@@ -0,0 +1,133 @@
|
|||||||
|
{
|
||||||
|
"sphragis_version": "0.1.0",
|
||||||
|
"generated_at": "2026-07-27T16:17:13.291817+00:00",
|
||||||
|
"reference": {
|
||||||
|
"endpoint": "http://127.0.0.1:8080/v1",
|
||||||
|
"model": "ref"
|
||||||
|
},
|
||||||
|
"candidate": {
|
||||||
|
"endpoint": "http://127.0.0.1:8081/v1",
|
||||||
|
"model": "cand"
|
||||||
|
},
|
||||||
|
"pack": {
|
||||||
|
"name": "/root/pack_v2.jsonl",
|
||||||
|
"n_tasks": 2800,
|
||||||
|
"sha256": "2de27099bbb15bab4f7b35599b215f038fdb68b3f5f4e1cda1a464f9dd18e14e"
|
||||||
|
},
|
||||||
|
"overall": "FAIL",
|
||||||
|
"claim": "At least one axis shows a statistically significant accuracy regression (exact McNemar, Holm-corrected, alpha=0.05).",
|
||||||
|
"axes": [
|
||||||
|
{
|
||||||
|
"axis": "arithmetic",
|
||||||
|
"n": 1400,
|
||||||
|
"counts": {
|
||||||
|
"both_correct": 237,
|
||||||
|
"ref_only": 203,
|
||||||
|
"cand_only": 88,
|
||||||
|
"both_wrong": 872
|
||||||
|
},
|
||||||
|
"acc_reference": 0.314286,
|
||||||
|
"acc_candidate": 0.232143,
|
||||||
|
"regression_d": 0.082143,
|
||||||
|
"d_ci": [
|
||||||
|
0.058571,
|
||||||
|
0.106429
|
||||||
|
],
|
||||||
|
"d_upper_bound": 0.102857,
|
||||||
|
"p_regression": 0.0,
|
||||||
|
"p_regression_holm": 0.0,
|
||||||
|
"p_improvement": 1.0,
|
||||||
|
"improved": false,
|
||||||
|
"mde_at_power": 0.030275,
|
||||||
|
"n_needed_for_margin": 1426,
|
||||||
|
"verdict": "FAIL",
|
||||||
|
"reason": "significant_regression"
|
||||||
|
},
|
||||||
|
{
|
||||||
|
"axis": "instruction",
|
||||||
|
"n": 600,
|
||||||
|
"counts": {
|
||||||
|
"both_correct": 0,
|
||||||
|
"ref_only": 11,
|
||||||
|
"cand_only": 16,
|
||||||
|
"both_wrong": 573
|
||||||
|
},
|
||||||
|
"acc_reference": 0.018333,
|
||||||
|
"acc_candidate": 0.026667,
|
||||||
|
"regression_d": -0.008333,
|
||||||
|
"d_ci": [
|
||||||
|
-0.025,
|
||||||
|
0.008333
|
||||||
|
],
|
||||||
|
"d_upper_bound": 0.005,
|
||||||
|
"p_regression": 0.87610572,
|
||||||
|
"p_regression_holm": 1.0,
|
||||||
|
"p_improvement": 0.22103417,
|
||||||
|
"improved": false,
|
||||||
|
"mde_at_power": 0.021496,
|
||||||
|
"n_needed_for_margin": 308,
|
||||||
|
"verdict": "PASS",
|
||||||
|
"reason": "non_inferior_within_margin"
|
||||||
|
},
|
||||||
|
{
|
||||||
|
"axis": "knowledge",
|
||||||
|
"n": 400,
|
||||||
|
"counts": {
|
||||||
|
"both_correct": 57,
|
||||||
|
"ref_only": 67,
|
||||||
|
"cand_only": 38,
|
||||||
|
"both_wrong": 238
|
||||||
|
},
|
||||||
|
"acc_reference": 0.31,
|
||||||
|
"acc_candidate": 0.2375,
|
||||||
|
"regression_d": 0.0725,
|
||||||
|
"d_ci": [
|
||||||
|
0.0225,
|
||||||
|
0.1225
|
||||||
|
],
|
||||||
|
"d_upper_bound": 0.115,
|
||||||
|
"p_regression": 0.00300804,
|
||||||
|
"p_regression_holm": 0.00902412,
|
||||||
|
"p_improvement": 0.99838951,
|
||||||
|
"improved": false,
|
||||||
|
"mde_at_power": 0.063531,
|
||||||
|
"n_needed_for_margin": 1802,
|
||||||
|
"verdict": "FAIL",
|
||||||
|
"reason": "significant_regression"
|
||||||
|
},
|
||||||
|
{
|
||||||
|
"axis": "reasoning",
|
||||||
|
"n": 400,
|
||||||
|
"counts": {
|
||||||
|
"both_correct": 50,
|
||||||
|
"ref_only": 34,
|
||||||
|
"cand_only": 35,
|
||||||
|
"both_wrong": 281
|
||||||
|
},
|
||||||
|
"acc_reference": 0.21,
|
||||||
|
"acc_candidate": 0.2125,
|
||||||
|
"regression_d": -0.0025,
|
||||||
|
"d_ci": [
|
||||||
|
-0.045,
|
||||||
|
0.0375
|
||||||
|
],
|
||||||
|
"d_upper_bound": 0.0325,
|
||||||
|
"p_regression": 0.59502547,
|
||||||
|
"p_regression_holm": 1.0,
|
||||||
|
"p_improvement": 0.5,
|
||||||
|
"improved": false,
|
||||||
|
"mde_at_power": 0.051501,
|
||||||
|
"n_needed_for_margin": 1183,
|
||||||
|
"verdict": "INCONCLUSIVE",
|
||||||
|
"reason": "ci_too_wide"
|
||||||
|
}
|
||||||
|
],
|
||||||
|
"params": {
|
||||||
|
"margin": 0.03,
|
||||||
|
"alpha": 0.05,
|
||||||
|
"n_floor": 30,
|
||||||
|
"power": 0.8,
|
||||||
|
"n_boot": 4000,
|
||||||
|
"seed": 0
|
||||||
|
}
|
||||||
|
}
|
||||||
6
chat_template.jinja
Normal file
6
chat_template.jinja
Normal file
@@ -0,0 +1,6 @@
|
|||||||
|
{% for message in messages %}{% if loop.first and messages[0]['role'] != 'system' %}{{ '<|im_start|>system
|
||||||
|
You are a helpful AI assistant named SmolLM, trained by Hugging Face<|im_end|>
|
||||||
|
' }}{% endif %}{{'<|im_start|>' + message['role'] + '
|
||||||
|
' + message['content'] + '<|im_end|>' + '
|
||||||
|
'}}{% endfor %}{% if add_generation_prompt %}{{ '<|im_start|>assistant
|
||||||
|
' }}{% endif %}
|
||||||
40
config.json
Normal file
40
config.json
Normal file
@@ -0,0 +1,40 @@
|
|||||||
|
{
|
||||||
|
"architectures": [
|
||||||
|
"LlamaForCausalLM"
|
||||||
|
],
|
||||||
|
"attention_bias": false,
|
||||||
|
"attention_dropout": 0.0,
|
||||||
|
"bos_token_id": 1,
|
||||||
|
"dtype": "bfloat16",
|
||||||
|
"eos_token_id": 2,
|
||||||
|
"head_dim": 64,
|
||||||
|
"hidden_act": "silu",
|
||||||
|
"hidden_size": 960,
|
||||||
|
"initializer_range": 0.02,
|
||||||
|
"intermediate_size": 2560,
|
||||||
|
"is_llama_config": true,
|
||||||
|
"max_position_embeddings": 8192,
|
||||||
|
"mlp_bias": false,
|
||||||
|
"model_type": "llama",
|
||||||
|
"num_attention_heads": 15,
|
||||||
|
"num_hidden_layers": 32,
|
||||||
|
"num_key_value_heads": 5,
|
||||||
|
"pad_token_id": 2,
|
||||||
|
"pretraining_tp": 1,
|
||||||
|
"rms_norm_eps": 1e-05,
|
||||||
|
"rope_interleaved": false,
|
||||||
|
"rope_parameters": {
|
||||||
|
"rope_theta": 100000,
|
||||||
|
"rope_type": "default"
|
||||||
|
},
|
||||||
|
"tie_word_embeddings": true,
|
||||||
|
"transformers.js_config": {
|
||||||
|
"kv_cache_dtype": {
|
||||||
|
"fp16": "float16",
|
||||||
|
"q4f16": "float16"
|
||||||
|
}
|
||||||
|
},
|
||||||
|
"transformers_version": "5.13.1",
|
||||||
|
"use_cache": true,
|
||||||
|
"vocab_size": 49152
|
||||||
|
}
|
||||||
28
ektome_report.json
Normal file
28
ektome_report.json
Normal file
@@ -0,0 +1,28 @@
|
|||||||
|
{
|
||||||
|
"status": "SHIP",
|
||||||
|
"target": 0.99,
|
||||||
|
"SE_mmlu": 0.021577983571223702,
|
||||||
|
"base_compliance": 0.89,
|
||||||
|
"compliance": 1.0,
|
||||||
|
"base_cap": 0.2475,
|
||||||
|
"cap": 0.2575,
|
||||||
|
"dcap": 0.01,
|
||||||
|
"gen_base": {
|
||||||
|
"foreign_rate": 0.0,
|
||||||
|
"degen_rate": 0.1,
|
||||||
|
"instr_pass": 0.6
|
||||||
|
},
|
||||||
|
"gen": {
|
||||||
|
"foreign_rate": 0.0,
|
||||||
|
"degen_rate": 0.2,
|
||||||
|
"instr_pass": 0.6
|
||||||
|
},
|
||||||
|
"gen_delta": {
|
||||||
|
"d_foreign": 0.0,
|
||||||
|
"d_degen": 0.1,
|
||||||
|
"d_instr": 0.0,
|
||||||
|
"holds_gen": true
|
||||||
|
},
|
||||||
|
"base_model": "HuggingFaceTB/SmolLM2-360M-Instruct",
|
||||||
|
"_note": "Redacted: search trajectory, selected depth and edit-scope removed. Reported values are the measured outcome only."
|
||||||
|
}
|
||||||
7
generation_config.json
Normal file
7
generation_config.json
Normal file
@@ -0,0 +1,7 @@
|
|||||||
|
{
|
||||||
|
"_from_model_config": true,
|
||||||
|
"bos_token_id": 1,
|
||||||
|
"eos_token_id": 2,
|
||||||
|
"pad_token_id": 2,
|
||||||
|
"transformers_version": "5.13.1"
|
||||||
|
}
|
||||||
3
hero.png
Normal file
3
hero.png
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:c9df19bddb40d55ab6ac76c6cc52fedc3666bc2dfd436c3f47823fdef912f682
|
||||||
|
size 243790
|
||||||
3
imatrix.dat
Normal file
3
imatrix.dat
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:3a8e15ec8f6d0d39635aed802245e02180cdd31782cbf8d6b44de067e2d26f75
|
||||||
|
size 1099424
|
||||||
3
model.safetensors
Normal file
3
model.safetensors
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:4e4405e4920fbe6cd9ed5529876031e5fdef04ff4d6b418b4ce058cb45c46b0b
|
||||||
|
size 723674912
|
||||||
6
nvfp4/chat_template.jinja
Normal file
6
nvfp4/chat_template.jinja
Normal file
@@ -0,0 +1,6 @@
|
|||||||
|
{% for message in messages %}{% if loop.first and messages[0]['role'] != 'system' %}{{ '<|im_start|>system
|
||||||
|
You are a helpful AI assistant named SmolLM, trained by Hugging Face<|im_end|>
|
||||||
|
' }}{% endif %}{{'<|im_start|>' + message['role'] + '
|
||||||
|
' + message['content'] + '<|im_end|>' + '
|
||||||
|
'}}{% endfor %}{% if add_generation_prompt %}{{ '<|im_start|>assistant
|
||||||
|
' }}{% endif %}
|
||||||
90
nvfp4/config.json
Normal file
90
nvfp4/config.json
Normal file
@@ -0,0 +1,90 @@
|
|||||||
|
{
|
||||||
|
"architectures": [
|
||||||
|
"LlamaForCausalLM"
|
||||||
|
],
|
||||||
|
"attention_bias": false,
|
||||||
|
"attention_dropout": 0.0,
|
||||||
|
"bos_token_id": 1,
|
||||||
|
"dtype": "bfloat16",
|
||||||
|
"eos_token_id": 2,
|
||||||
|
"head_dim": 64,
|
||||||
|
"hidden_act": "silu",
|
||||||
|
"hidden_size": 960,
|
||||||
|
"initializer_range": 0.02,
|
||||||
|
"intermediate_size": 2560,
|
||||||
|
"is_llama_config": true,
|
||||||
|
"max_position_embeddings": 8192,
|
||||||
|
"mlp_bias": false,
|
||||||
|
"model_type": "llama",
|
||||||
|
"num_attention_heads": 15,
|
||||||
|
"num_hidden_layers": 32,
|
||||||
|
"num_key_value_heads": 5,
|
||||||
|
"pad_token_id": 2,
|
||||||
|
"pretraining_tp": 1,
|
||||||
|
"quantization_config": {
|
||||||
|
"config_groups": {
|
||||||
|
"group_0": {
|
||||||
|
"format": "nvfp4-pack-quantized",
|
||||||
|
"input_activations": {
|
||||||
|
"actorder": null,
|
||||||
|
"block_structure": null,
|
||||||
|
"dynamic": "local",
|
||||||
|
"group_size": 16,
|
||||||
|
"num_bits": 4,
|
||||||
|
"observer": "static_minmax",
|
||||||
|
"observer_kwargs": {},
|
||||||
|
"scale_dtype": "torch.float8_e4m3fn",
|
||||||
|
"strategy": "tensor_group",
|
||||||
|
"symmetric": true,
|
||||||
|
"type": "float",
|
||||||
|
"zp_dtype": null
|
||||||
|
},
|
||||||
|
"output_activations": null,
|
||||||
|
"targets": [
|
||||||
|
"Linear"
|
||||||
|
],
|
||||||
|
"weights": {
|
||||||
|
"actorder": null,
|
||||||
|
"block_structure": null,
|
||||||
|
"dynamic": false,
|
||||||
|
"group_size": 16,
|
||||||
|
"num_bits": 4,
|
||||||
|
"observer": "memoryless_minmax",
|
||||||
|
"observer_kwargs": {},
|
||||||
|
"scale_dtype": "torch.float8_e4m3fn",
|
||||||
|
"strategy": "tensor_group",
|
||||||
|
"symmetric": true,
|
||||||
|
"type": "float",
|
||||||
|
"zp_dtype": null
|
||||||
|
}
|
||||||
|
}
|
||||||
|
},
|
||||||
|
"format": "nvfp4-pack-quantized",
|
||||||
|
"global_compression_ratio": null,
|
||||||
|
"ignore": [
|
||||||
|
"lm_head"
|
||||||
|
],
|
||||||
|
"kv_cache_scheme": null,
|
||||||
|
"quant_method": "compressed-tensors",
|
||||||
|
"quantization_status": "compressed",
|
||||||
|
"sparsity_config": {},
|
||||||
|
"transform_config": {},
|
||||||
|
"version": "0.17.1"
|
||||||
|
},
|
||||||
|
"rms_norm_eps": 1e-05,
|
||||||
|
"rope_interleaved": false,
|
||||||
|
"rope_parameters": {
|
||||||
|
"rope_theta": 100000,
|
||||||
|
"rope_type": "default"
|
||||||
|
},
|
||||||
|
"tie_word_embeddings": true,
|
||||||
|
"transformers.js_config": {
|
||||||
|
"kv_cache_dtype": {
|
||||||
|
"fp16": "float16",
|
||||||
|
"q4f16": "float16"
|
||||||
|
}
|
||||||
|
},
|
||||||
|
"transformers_version": "5.10.1",
|
||||||
|
"use_cache": true,
|
||||||
|
"vocab_size": 49152
|
||||||
|
}
|
||||||
7
nvfp4/generation_config.json
Normal file
7
nvfp4/generation_config.json
Normal file
@@ -0,0 +1,7 @@
|
|||||||
|
{
|
||||||
|
"_from_model_config": true,
|
||||||
|
"bos_token_id": 1,
|
||||||
|
"eos_token_id": 2,
|
||||||
|
"pad_token_id": 2,
|
||||||
|
"transformers_version": "5.10.1"
|
||||||
|
}
|
||||||
3
nvfp4/model.safetensors
Normal file
3
nvfp4/model.safetensors
Normal file
@@ -0,0 +1,3 @@
|
|||||||
|
version https://git-lfs.github.com/spec/v1
|
||||||
|
oid sha256:709a97ba40c41cd7f42dae3c538b461c534ccb44a69857e4edf4e1f43dee8eb3
|
||||||
|
size 271553840
|
||||||
7
nvfp4/recipe.yaml
Normal file
7
nvfp4/recipe.yaml
Normal file
@@ -0,0 +1,7 @@
|
|||||||
|
default_stage:
|
||||||
|
default_modifiers:
|
||||||
|
QuantizationModifier:
|
||||||
|
targets: [Linear]
|
||||||
|
ignore: [lm_head]
|
||||||
|
scheme: NVFP4
|
||||||
|
bypass_divisibility_checks: false
|
||||||
244965
nvfp4/tokenizer.json
Normal file
244965
nvfp4/tokenizer.json
Normal file
File diff suppressed because it is too large
Load Diff
19
nvfp4/tokenizer_config.json
Normal file
19
nvfp4/tokenizer_config.json
Normal file
@@ -0,0 +1,19 @@
|
|||||||
|
{
|
||||||
|
"add_prefix_space": false,
|
||||||
|
"backend": "tokenizers",
|
||||||
|
"bos_token": "<|im_start|>",
|
||||||
|
"clean_up_tokenization_spaces": false,
|
||||||
|
"eos_token": "<|im_end|>",
|
||||||
|
"errors": "replace",
|
||||||
|
"extra_special_tokens": [
|
||||||
|
"<|im_start|>",
|
||||||
|
"<|im_end|>"
|
||||||
|
],
|
||||||
|
"is_local": false,
|
||||||
|
"local_files_only": false,
|
||||||
|
"model_max_length": 8192,
|
||||||
|
"pad_token": "<|im_end|>",
|
||||||
|
"tokenizer_class": "GPT2Tokenizer",
|
||||||
|
"unk_token": "<|endoftext|>",
|
||||||
|
"vocab_size": 49152
|
||||||
|
}
|
||||||
34
special_tokens_map.json
Normal file
34
special_tokens_map.json
Normal file
@@ -0,0 +1,34 @@
|
|||||||
|
{
|
||||||
|
"additional_special_tokens": [
|
||||||
|
"<|im_start|>",
|
||||||
|
"<|im_end|>"
|
||||||
|
],
|
||||||
|
"bos_token": {
|
||||||
|
"content": "<|im_start|>",
|
||||||
|
"lstrip": false,
|
||||||
|
"normalized": false,
|
||||||
|
"rstrip": false,
|
||||||
|
"single_word": false
|
||||||
|
},
|
||||||
|
"eos_token": {
|
||||||
|
"content": "<|im_end|>",
|
||||||
|
"lstrip": false,
|
||||||
|
"normalized": false,
|
||||||
|
"rstrip": false,
|
||||||
|
"single_word": false
|
||||||
|
},
|
||||||
|
"pad_token": {
|
||||||
|
"content": "<|im_end|>",
|
||||||
|
"lstrip": false,
|
||||||
|
"normalized": false,
|
||||||
|
"rstrip": false,
|
||||||
|
"single_word": false
|
||||||
|
},
|
||||||
|
"unk_token": {
|
||||||
|
"content": "<|endoftext|>",
|
||||||
|
"lstrip": false,
|
||||||
|
"normalized": false,
|
||||||
|
"rstrip": false,
|
||||||
|
"single_word": false
|
||||||
|
}
|
||||||
|
}
|
||||||
244965
tokenizer.json
Normal file
244965
tokenizer.json
Normal file
File diff suppressed because it is too large
Load Diff
16
tokenizer_config.json
Normal file
16
tokenizer_config.json
Normal file
@@ -0,0 +1,16 @@
|
|||||||
|
{
|
||||||
|
"add_prefix_space": false,
|
||||||
|
"backend": "tokenizers",
|
||||||
|
"bos_token": "<|im_start|>",
|
||||||
|
"clean_up_tokenization_spaces": false,
|
||||||
|
"eos_token": "<|im_end|>",
|
||||||
|
"errors": "replace",
|
||||||
|
"is_local": false,
|
||||||
|
"local_files_only": false,
|
||||||
|
"model_max_length": 8192,
|
||||||
|
"pad_token": "<|im_end|>",
|
||||||
|
"tokenizer_class": "GPT2Tokenizer",
|
||||||
|
"unk_token": "<|endoftext|>",
|
||||||
|
"vocab_size": 49152,
|
||||||
|
"chat_template": "{% for message in messages %}{% if loop.first and messages[0]['role'] != 'system' %}{{ '<|im_start|>system\nYou are a helpful AI assistant named SmolLM, trained by Hugging Face<|im_end|>\n' }}{% endif %}{{'<|im_start|>' + message['role'] + '\n' + message['content'] + '<|im_end|>' + '\n'}}{% endfor %}{% if add_generation_prompt %}{{ '<|im_start|>assistant\n' }}{% endif %}"
|
||||||
|
}
|
||||||
BIN
whitepaper.pdf
Normal file
BIN
whitepaper.pdf
Normal file
Binary file not shown.
Reference in New Issue
Block a user