Upload Qwen3ForCausalLM
Browse files- config.json +1 -1
- generation_config.json +2 -8
- model-00001-of-00002.safetensors +2 -2
- model-00002-of-00002.safetensors +2 -2
- model.safetensors.index.json +0 -0
config.json
CHANGED
|
@@ -60,7 +60,7 @@
|
|
| 60 |
"rope_scaling": null,
|
| 61 |
"rope_theta": 1000000,
|
| 62 |
"sliding_window": null,
|
| 63 |
-
"tie_word_embeddings":
|
| 64 |
"transformers_version": "4.57.6",
|
| 65 |
"use_cache": true,
|
| 66 |
"use_sliding_window": false,
|
|
|
|
| 60 |
"rope_scaling": null,
|
| 61 |
"rope_theta": 1000000,
|
| 62 |
"sliding_window": null,
|
| 63 |
+
"tie_word_embeddings": false,
|
| 64 |
"transformers_version": "4.57.6",
|
| 65 |
"use_cache": true,
|
| 66 |
"use_sliding_window": false,
|
generation_config.json
CHANGED
|
@@ -1,12 +1,6 @@
|
|
| 1 |
{
|
| 2 |
-
"
|
| 3 |
-
"eos_token_id":
|
| 4 |
-
151645,
|
| 5 |
-
151643
|
| 6 |
-
],
|
| 7 |
"pad_token_id": 151643,
|
| 8 |
-
"temperature": 0.6,
|
| 9 |
-
"top_k": 20,
|
| 10 |
-
"top_p": 0.95,
|
| 11 |
"transformers_version": "4.57.6"
|
| 12 |
}
|
|
|
|
| 1 |
{
|
| 2 |
+
"_from_model_config": true,
|
| 3 |
+
"eos_token_id": 151645,
|
|
|
|
|
|
|
|
|
|
| 4 |
"pad_token_id": 151643,
|
|
|
|
|
|
|
|
|
|
| 5 |
"transformers_version": "4.57.6"
|
| 6 |
}
|
model-00001-of-00002.safetensors
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
-
size
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:513475ee51df67d6893aa3d48d9948a1837ba4440d0161ef7d97598e379b810a
|
| 3 |
+
size 4967215360
|
model-00002-of-00002.safetensors
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
-
size
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:3de0142bb7677eec43104aa4df6ed2635b62a1febfafd0e9e60c2fb5823d2d49
|
| 3 |
+
size 3855679144
|
model.safetensors.index.json
CHANGED
|
The diff for this file is too large to render.
See raw diff
|
|
|