初始化项目,由ModelHub XC社区提供模型

Model: fpadovani/hin-deva-100mb-after-ppt-shuff-dyck-100mb-ckpt500_seed3407
Source: Original Platform
This commit is contained in:
ModelHub XC
2026-06-25 05:20:19 +08:00
commit ef0939370e
360 changed files with 4266295 additions and 0 deletions

File diff suppressed because it is too large Load Diff

View File

@@ -0,0 +1,34 @@
{
"activation_function": "gelu",
"architectures": [
"GPT2LMHeadModel"
],
"attn_pdrop": 0.1,
"bos_token_id": 50000,
"dtype": "bfloat16",
"embd_pdrop": 0.1,
"eos_token_id": 50001,
"initializer_range": 0.02,
"layer_norm_epsilon": 1e-05,
"model_type": "gpt2",
"n_ctx": 512,
"n_embd": 768,
"n_head": 12,
"n_inner": 3072,
"n_layer": 12,
"n_positions": 512,
"pad_token_id": 50002,
"prefix": "[CLS]",
"reorder_and_upcast_attn": false,
"resid_pdrop": 0.1,
"scale_attn_by_inverse_layer_idx": false,
"scale_attn_weights": true,
"summary_activation": null,
"summary_first_dropout": 0.1,
"summary_proj_to_labels": true,
"summary_type": "cls_index",
"summary_use_proj": true,
"transformers_version": "4.56.2",
"use_cache": true,
"vocab_size": 51200
}

View File

@@ -0,0 +1,9 @@
{
"_from_model_config": true,
"bos_token_id": 50000,
"eos_token_id": [
50001
],
"pad_token_id": 50002,
"transformers_version": "4.56.2"
}

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:ed27411a0d82c272eb3f6619e556f585a8619d2d0845d427257bc6cfe67a9911
size 249556672

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:7392b053a6723e743a15b340b897fb35425a0e25ba282cc5dc9565c280f743bd
size 14645

File diff suppressed because it is too large Load Diff

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:cfcecb41352908b2a034dc4ec11e511d1a7ed2d594a679147bc9b51bfe6bd14a
size 1500527

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:1107a6b3ce1f160767a5fc9c3059c21f9a18c070d8fac619479741fa9e41dd02
size 6353