nm-autobot commited on
Commit
c46767c
·
verified ·
1 Parent(s): 665003f

Upload folder using huggingface_hub

Browse files
config.json CHANGED
@@ -9,6 +9,7 @@
9
  "AutoModelForCausalLM": "modeling_phi3.Phi3ForCausalLM"
10
  },
11
  "bos_token_id": 1,
 
12
  "embd_pdrop": 0.0,
13
  "eos_token_id": 32000,
14
  "hidden_act": "silu",
@@ -27,154 +28,26 @@
27
  "config_groups": {},
28
  "format": "dense",
29
  "global_compression_ratio": null,
30
- "ignore": [
31
- "model.layers.0.self_attn.o_proj",
32
- "model.layers.0.self_attn.qkv_proj",
33
- "model.layers.0.mlp.gate_up_proj",
34
- "model.layers.0.mlp.down_proj",
35
- "model.layers.1.self_attn.o_proj",
36
- "model.layers.1.self_attn.qkv_proj",
37
- "model.layers.1.mlp.gate_up_proj",
38
- "model.layers.1.mlp.down_proj",
39
- "model.layers.2.self_attn.o_proj",
40
- "model.layers.2.self_attn.qkv_proj",
41
- "model.layers.2.mlp.gate_up_proj",
42
- "model.layers.2.mlp.down_proj",
43
- "model.layers.3.self_attn.o_proj",
44
- "model.layers.3.self_attn.qkv_proj",
45
- "model.layers.3.mlp.gate_up_proj",
46
- "model.layers.3.mlp.down_proj",
47
- "model.layers.4.self_attn.o_proj",
48
- "model.layers.4.self_attn.qkv_proj",
49
- "model.layers.4.mlp.gate_up_proj",
50
- "model.layers.4.mlp.down_proj",
51
- "model.layers.5.self_attn.o_proj",
52
- "model.layers.5.self_attn.qkv_proj",
53
- "model.layers.5.mlp.gate_up_proj",
54
- "model.layers.5.mlp.down_proj",
55
- "model.layers.6.self_attn.o_proj",
56
- "model.layers.6.self_attn.qkv_proj",
57
- "model.layers.6.mlp.gate_up_proj",
58
- "model.layers.6.mlp.down_proj",
59
- "model.layers.7.self_attn.o_proj",
60
- "model.layers.7.self_attn.qkv_proj",
61
- "model.layers.7.mlp.gate_up_proj",
62
- "model.layers.7.mlp.down_proj",
63
- "model.layers.8.self_attn.o_proj",
64
- "model.layers.8.self_attn.qkv_proj",
65
- "model.layers.8.mlp.gate_up_proj",
66
- "model.layers.8.mlp.down_proj",
67
- "model.layers.9.self_attn.o_proj",
68
- "model.layers.9.self_attn.qkv_proj",
69
- "model.layers.9.mlp.gate_up_proj",
70
- "model.layers.9.mlp.down_proj",
71
- "model.layers.10.self_attn.o_proj",
72
- "model.layers.10.self_attn.qkv_proj",
73
- "model.layers.10.mlp.gate_up_proj",
74
- "model.layers.10.mlp.down_proj",
75
- "model.layers.11.self_attn.o_proj",
76
- "model.layers.11.self_attn.qkv_proj",
77
- "model.layers.11.mlp.gate_up_proj",
78
- "model.layers.11.mlp.down_proj",
79
- "model.layers.12.self_attn.o_proj",
80
- "model.layers.12.self_attn.qkv_proj",
81
- "model.layers.12.mlp.gate_up_proj",
82
- "model.layers.12.mlp.down_proj",
83
- "model.layers.13.self_attn.o_proj",
84
- "model.layers.13.self_attn.qkv_proj",
85
- "model.layers.13.mlp.gate_up_proj",
86
- "model.layers.13.mlp.down_proj",
87
- "model.layers.14.self_attn.o_proj",
88
- "model.layers.14.self_attn.qkv_proj",
89
- "model.layers.14.mlp.gate_up_proj",
90
- "model.layers.14.mlp.down_proj",
91
- "model.layers.15.self_attn.o_proj",
92
- "model.layers.15.self_attn.qkv_proj",
93
- "model.layers.15.mlp.gate_up_proj",
94
- "model.layers.15.mlp.down_proj",
95
- "model.layers.16.self_attn.o_proj",
96
- "model.layers.16.self_attn.qkv_proj",
97
- "model.layers.16.mlp.gate_up_proj",
98
- "model.layers.16.mlp.down_proj",
99
- "model.layers.17.self_attn.o_proj",
100
- "model.layers.17.self_attn.qkv_proj",
101
- "model.layers.17.mlp.gate_up_proj",
102
- "model.layers.17.mlp.down_proj",
103
- "model.layers.18.self_attn.o_proj",
104
- "model.layers.18.self_attn.qkv_proj",
105
- "model.layers.18.mlp.gate_up_proj",
106
- "model.layers.18.mlp.down_proj",
107
- "model.layers.19.self_attn.o_proj",
108
- "model.layers.19.self_attn.qkv_proj",
109
- "model.layers.19.mlp.gate_up_proj",
110
- "model.layers.19.mlp.down_proj",
111
- "model.layers.20.self_attn.o_proj",
112
- "model.layers.20.self_attn.qkv_proj",
113
- "model.layers.20.mlp.gate_up_proj",
114
- "model.layers.20.mlp.down_proj",
115
- "model.layers.21.self_attn.o_proj",
116
- "model.layers.21.self_attn.qkv_proj",
117
- "model.layers.21.mlp.gate_up_proj",
118
- "model.layers.21.mlp.down_proj",
119
- "model.layers.22.self_attn.o_proj",
120
- "model.layers.22.self_attn.qkv_proj",
121
- "model.layers.22.mlp.gate_up_proj",
122
- "model.layers.22.mlp.down_proj",
123
- "model.layers.23.self_attn.o_proj",
124
- "model.layers.23.self_attn.qkv_proj",
125
- "model.layers.23.mlp.gate_up_proj",
126
- "model.layers.23.mlp.down_proj",
127
- "model.layers.24.self_attn.o_proj",
128
- "model.layers.24.self_attn.qkv_proj",
129
- "model.layers.24.mlp.gate_up_proj",
130
- "model.layers.24.mlp.down_proj",
131
- "model.layers.25.self_attn.o_proj",
132
- "model.layers.25.self_attn.qkv_proj",
133
- "model.layers.25.mlp.gate_up_proj",
134
- "model.layers.25.mlp.down_proj",
135
- "model.layers.26.self_attn.o_proj",
136
- "model.layers.26.self_attn.qkv_proj",
137
- "model.layers.26.mlp.gate_up_proj",
138
- "model.layers.26.mlp.down_proj",
139
- "model.layers.27.self_attn.o_proj",
140
- "model.layers.27.self_attn.qkv_proj",
141
- "model.layers.27.mlp.gate_up_proj",
142
- "model.layers.27.mlp.down_proj",
143
- "model.layers.28.self_attn.o_proj",
144
- "model.layers.28.self_attn.qkv_proj",
145
- "model.layers.28.mlp.gate_up_proj",
146
- "model.layers.28.mlp.down_proj",
147
- "model.layers.29.self_attn.o_proj",
148
- "model.layers.29.self_attn.qkv_proj",
149
- "model.layers.29.mlp.gate_up_proj",
150
- "model.layers.29.mlp.down_proj",
151
- "model.layers.30.self_attn.o_proj",
152
- "model.layers.30.self_attn.qkv_proj",
153
- "model.layers.30.mlp.gate_up_proj",
154
- "model.layers.30.mlp.down_proj",
155
- "model.layers.31.self_attn.o_proj",
156
- "model.layers.31.self_attn.qkv_proj",
157
- "model.layers.31.mlp.gate_up_proj",
158
- "model.layers.31.mlp.down_proj",
159
- "lm_head"
160
- ],
161
  "kv_cache_scheme": {
162
  "actorder": null,
163
  "block_structure": null,
164
  "dynamic": false,
165
  "group_size": null,
166
  "num_bits": 8,
167
- "observer": "minmax",
168
  "observer_kwargs": {},
 
169
  "strategy": "tensor",
170
  "symmetric": true,
171
- "type": "float"
 
172
  },
173
  "quant_method": "compressed-tensors",
174
  "quantization_status": "frozen",
175
  "sparsity_config": {},
176
  "transform_config": {},
177
- "version": "0.11.0"
178
  },
179
  "resid_pdrop": 0.0,
180
  "rms_norm_eps": 1e-05,
@@ -182,8 +55,7 @@
182
  "rope_theta": 10000.0,
183
  "sliding_window": 2047,
184
  "tie_word_embeddings": false,
185
- "torch_dtype": "bfloat16",
186
- "transformers_version": "4.55.2",
187
  "use_cache": true,
188
  "vocab_size": 32064
189
  }
 
9
  "AutoModelForCausalLM": "modeling_phi3.Phi3ForCausalLM"
10
  },
11
  "bos_token_id": 1,
12
+ "dtype": "bfloat16",
13
  "embd_pdrop": 0.0,
14
  "eos_token_id": 32000,
15
  "hidden_act": "silu",
 
28
  "config_groups": {},
29
  "format": "dense",
30
  "global_compression_ratio": null,
31
+ "ignore": [],
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
32
  "kv_cache_scheme": {
33
  "actorder": null,
34
  "block_structure": null,
35
  "dynamic": false,
36
  "group_size": null,
37
  "num_bits": 8,
38
+ "observer": "memoryless_minmax",
39
  "observer_kwargs": {},
40
+ "scale_dtype": null,
41
  "strategy": "tensor",
42
  "symmetric": true,
43
+ "type": "float",
44
+ "zp_dtype": null
45
  },
46
  "quant_method": "compressed-tensors",
47
  "quantization_status": "frozen",
48
  "sparsity_config": {},
49
  "transform_config": {},
50
+ "version": "0.14.0.1"
51
  },
52
  "resid_pdrop": 0.0,
53
  "rms_norm_eps": 1e-05,
 
55
  "rope_theta": 10000.0,
56
  "sliding_window": 2047,
57
  "tie_word_embeddings": false,
58
+ "transformers_version": "4.57.6",
 
59
  "use_cache": true,
60
  "vocab_size": 32064
61
  }
generation_config.json CHANGED
@@ -7,5 +7,5 @@
7
  32007
8
  ],
9
  "pad_token_id": 32000,
10
- "transformers_version": "4.55.2"
11
  }
 
7
  32007
8
  ],
9
  "pad_token_id": 32000,
10
+ "transformers_version": "4.57.6"
11
  }
model-00001-of-00002.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:78a019bc9c9f03ee35b821f024704cdde9e9034ab5061e254e6092c8ee02d6b4
3
  size 4972493960
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c4e7401a2294820680d12d92693eca67123317ebe631720eb831a03dbd716175
3
  size 4972493960
model-00002-of-00002.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:f6fc59e16719e1415bfd49f557a26b6ecaf2ad1a1c05b23525b2b57a3dcc6bcd
3
  size 2669694664
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:495ff095f0e52fbadffc667850f95b1a232a8d5983cb5dbd7a4e02d68b3d12e6
3
  size 2669694664
recipe.yaml CHANGED
@@ -12,5 +12,8 @@ quant_stage:
12
  block_structure: null
13
  dynamic: false
14
  actorder: null
15
- observer: minmax
 
 
16
  observer_kwargs: {}
 
 
12
  block_structure: null
13
  dynamic: false
14
  actorder: null
15
+ scale_dtype: null
16
+ zp_dtype: null
17
+ observer: memoryless_minmax
18
  observer_kwargs: {}
19
+ bypass_divisibility_checks: false