gchauhan
/

ycchen-3

gchauhan commited on Nov 18, 2023

Commit

f444d93

1 Parent(s): 0e1aaa3

Upload folder using huggingface_hub

Files changed (2) hide show

config.json CHANGED Viewed

@@ -1,16 +1,16 @@
 {
-  "_name_or_path": "Qwen-14B-8bit-wikitext-ex512-len4096-hf-safetensors",
   "architectures": [
     "QWenLMHeadModel"
   ],
   "attn_dropout_prob": 0.0,
   "auto_map": {
-    "AutoConfig": "configuration_qwen.QWenConfig",
-    "AutoModelForCausalLM": "modeling_qwen.QWenLMHeadModel"
   },
-  "bf16": false,
   "emb_dropout_prob": 0.0,
-  "fp16": true,
   "fp32": false,
   "hidden_size": 5120,
   "initializer_range": 0.02,
@@ -23,24 +23,6 @@
   "num_attention_heads": 40,
   "num_hidden_layers": 40,
   "onnx_safe": null,
-  "quantization_config": {
-    "batch_size": 1,
-    "bits": 8,
-    "block_name_to_quantize": null,
-    "damp_percent": 0.01,
-    "dataset": null,
-    "desc_act": false,
-    "disable_exllama": false,
-    "group_size": 128,
-    "model_seqlen": null,
-    "module_name_preceding_first_block": null,
-    "pad_token_id": null,
-    "quant_method": "gptq",
-    "sym": true,
-    "tokenizer": null,
-    "true_sequential": true,
-    "use_cuda_fp16": false
-  },
   "rotary_emb_base": 10000,
   "rotary_pct": 1.0,
   "scale_attn_weights": true,

 {
+  "_name_or_path": "Qwen/Qwen-14B",
   "architectures": [
     "QWenLMHeadModel"
   ],
   "attn_dropout_prob": 0.0,
   "auto_map": {
+    "AutoConfig": "Qwen/Qwen-14B--configuration_qwen.QWenConfig",
+    "AutoModelForCausalLM": "Qwen/Qwen-14B--modeling_qwen.QWenLMHeadModel"
   },
+  "bf16": true,
   "emb_dropout_prob": 0.0,
+  "fp16": false,
   "fp32": false,
   "hidden_size": 5120,
   "initializer_range": 0.02,
   "num_attention_heads": 40,
   "num_hidden_layers": 40,
   "onnx_safe": null,
   "rotary_emb_base": 10000,
   "rotary_pct": 1.0,
   "scale_attn_weights": true,

pytorch_model.bin.bin ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:cd3f1e669764bd161c6eeebc0cc70659173e4bb7c011474efcd9f2d47d1c880d
+size 16029839536