初始化项目,由ModelHub XC社区提供模型

Model: fpadovani/swa-latn-10mb-ppt-Dp-10mb_seed3407
Source: Original Platform
This commit is contained in:
ModelHub XC
2026-06-24 11:16:17 +08:00
commit 51fd6c50df
89 changed files with 156332 additions and 0 deletions

File diff suppressed because it is too large Load Diff

View File

@@ -0,0 +1,34 @@
{
"activation_function": "gelu",
"architectures": [
"GPT2LMHeadModel"
],
"attn_pdrop": 0.1,
"bos_token_id": 50000,
"dtype": "bfloat16",
"embd_pdrop": 0.1,
"eos_token_id": 50001,
"initializer_range": 0.02,
"layer_norm_epsilon": 1e-05,
"model_type": "gpt2",
"n_ctx": 512,
"n_embd": 512,
"n_head": 8,
"n_inner": 2048,
"n_layer": 4,
"n_positions": 512,
"pad_token_id": 50002,
"prefix": "[CLS]",
"reorder_and_upcast_attn": false,
"resid_pdrop": 0.1,
"scale_attn_by_inverse_layer_idx": false,
"scale_attn_weights": true,
"summary_activation": null,
"summary_first_dropout": 0.1,
"summary_proj_to_labels": true,
"summary_type": "cls_index",
"summary_use_proj": true,
"transformers_version": "4.56.2",
"use_cache": true,
"vocab_size": 51200
}

View File

@@ -0,0 +1,9 @@
{
"_from_model_config": true,
"bos_token_id": 50000,
"eos_token_id": [
50001
],
"pad_token_id": 50002,
"transformers_version": "4.56.2"
}

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:71358a3dba5b6d5cb008bfec054f079c83baea7d2b548a8a898a71514418b741
size 78179408

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:4199b4168be793e99ccc26ec487e9fa9ff8f47bb2d7e18feaab9a5e5aad9bf29
size 14645

File diff suppressed because it is too large Load Diff

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:fb0bca6df279e2e17a03defe0e75bcb121ed6bbd1fd100100a4be7c26e8ec87c
size 1091784

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -0,0 +1,3 @@
version https://git-lfs.github.com/spec/v1
oid sha256:d7eaa14197970497a4e1da1904e0c5d0bc409a3cde063d3c581c74e4221f8ed5
size 6289