Text Generation
Transformers
Safetensors
English
qwen2
conversational
text-generation-inference
Inference Endpoints
kz919 commited on
Commit
5916c2d
·
verified ·
1 Parent(s): adcac77

Upload Qwen2ForCausalLM

Browse files
Files changed (2) hide show
  1. config.json +2 -2
  2. model.safetensors +1 -1
config.json CHANGED
@@ -1,5 +1,5 @@
1
  {
2
- "_name_or_path": "/scratch/09979/kaizhaol/sft_model_8dp/checkpoint-702/",
3
  "architectures": [
4
  "Qwen2ForCausalLM"
5
  ],
@@ -11,7 +11,7 @@
11
  "initializer_range": 0.02,
12
  "intermediate_size": 4864,
13
  "max_position_embeddings": 32768,
14
- "max_window_layers": 24,
15
  "model_type": "qwen2",
16
  "num_attention_heads": 14,
17
  "num_hidden_layers": 24,
 
1
  {
2
+ "_name_or_path": "/scratch/09979/kaizhaol/sft_model_8dp_small_batch_v3/checkpoint-11646/",
3
  "architectures": [
4
  "Qwen2ForCausalLM"
5
  ],
 
11
  "initializer_range": 0.02,
12
  "intermediate_size": 4864,
13
  "max_position_embeddings": 32768,
14
+ "max_window_layers": 21,
15
  "model_type": "qwen2",
16
  "num_attention_heads": 14,
17
  "num_hidden_layers": 24,
model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:929c2eb5b004a5ce075984524b1013e37ad71d20a0470ea360b3d829088f8e4d
3
  size 988097824
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5b88a84538252528f570918443ac27765ded64fab86bcf63eb0dd43e42035934
3
  size 988097824