yzhuang commited on
Commit
c3b28d1
·
verified ·
1 Parent(s): 8121428

End of training

Browse files
README.md CHANGED
@@ -15,6 +15,7 @@ model-index:
15
  <!-- This model card has been generated automatically according to the information the Trainer had access to. You
16
  should probably proofread and complete it, then remove this comment. -->
17
 
 
18
  # Meta-Llama-3-8B-Instruct_fictional_arc_challenge_Italian_v2
19
 
20
  This model is a fine-tuned version of [meta-llama/Meta-Llama-3-8B-Instruct](https://huggingface.co/meta-llama/Meta-Llama-3-8B-Instruct) on the generator dataset.
@@ -52,7 +53,7 @@ The following hyperparameters were used during training:
52
 
53
  ### Framework versions
54
 
55
- - Transformers 4.40.2
56
  - Pytorch 2.1.0a0+32f93b1
57
  - Datasets 2.19.1
58
  - Tokenizers 0.19.1
 
15
  <!-- This model card has been generated automatically according to the information the Trainer had access to. You
16
  should probably proofread and complete it, then remove this comment. -->
17
 
18
+ [<img src="https://raw.githubusercontent.com/wandb/assets/main/wandb-github-badge-28.svg" alt="Visualize in Weights & Biases" width="200" height="32"/>](https://wandb.ai/yufanz/autotree/runs/7283970144.51595-887226ef-9076-4284-993d-3e22f4763aa6)
19
  # Meta-Llama-3-8B-Instruct_fictional_arc_challenge_Italian_v2
20
 
21
  This model is a fine-tuned version of [meta-llama/Meta-Llama-3-8B-Instruct](https://huggingface.co/meta-llama/Meta-Llama-3-8B-Instruct) on the generator dataset.
 
53
 
54
  ### Framework versions
55
 
56
+ - Transformers 4.41.0
57
  - Pytorch 2.1.0a0+32f93b1
58
  - Datasets 2.19.1
59
  - Tokenizers 0.19.1
config.json CHANGED
@@ -12,6 +12,7 @@
12
  "initializer_range": 0.02,
13
  "intermediate_size": 14336,
14
  "max_position_embeddings": 8192,
 
15
  "model_type": "llama",
16
  "num_attention_heads": 32,
17
  "num_hidden_layers": 32,
@@ -22,7 +23,7 @@
22
  "rope_theta": 500000.0,
23
  "tie_word_embeddings": false,
24
  "torch_dtype": "bfloat16",
25
- "transformers_version": "4.40.2",
26
  "use_cache": true,
27
  "vocab_size": 128256
28
  }
 
12
  "initializer_range": 0.02,
13
  "intermediate_size": 14336,
14
  "max_position_embeddings": 8192,
15
+ "mlp_bias": false,
16
  "model_type": "llama",
17
  "num_attention_heads": 32,
18
  "num_hidden_layers": 32,
 
23
  "rope_theta": 500000.0,
24
  "tie_word_embeddings": false,
25
  "torch_dtype": "bfloat16",
26
+ "transformers_version": "4.41.0",
27
  "use_cache": true,
28
  "vocab_size": 128256
29
  }
generation_config.json CHANGED
@@ -8,5 +8,5 @@
8
  "max_length": 4096,
9
  "temperature": 0.6,
10
  "top_p": 0.9,
11
- "transformers_version": "4.40.2"
12
  }
 
8
  "max_length": 4096,
9
  "temperature": 0.6,
10
  "top_p": 0.9,
11
+ "transformers_version": "4.41.0"
12
  }
model-00001-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:6478e8d3392ecb7cfe8ff69662ef094542f8847ff3f3ff54c39241c53043678b
3
  size 4976698672
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:f200372ab9d4d9b7b338eaa3c631945f56193c8ae202484e93c7be41d506513a
3
  size 4976698672
model-00002-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:9a4b34db73db3bf5d71995af172a66b411881d484a442268e047f82e063be5eb
3
  size 4999802720
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2e7387d448c402f1ca610556c1fb3029285775c3f5137532bd6702e1768f9d60
3
  size 4999802720
model-00003-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:f56875d1b9deb12497831833ab26c128c28dbf472a0aa14b7972ea2f39c24ab0
3
  size 4915916176
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:f39800c153fdd2b8fb5d6b05c1ba66e7ca97f04bac0004120d55828987c63d05
3
  size 4915916176
model-00004-of-00004.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:025ac67de6f9d1d0b3df835e4c06c1bba5dc97a729eaddf676189e002c89a5d8
3
  size 1168138808
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:742388115e15cb657522dcfdb9ca2c42a71f07154b9c9ea2f8d37c313f8fe3c0
3
  size 1168138808
runs/May18_07-05-47_node-0/events.out.tfevents.1716015949.node-0.28736.0 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:57218ee9944b0ba82ffe75d8de9606ac76aeef87280d07d66bd02e5202607eac
3
+ size 5293
training_args.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:e03d6370fd44f7e1569bf814a77af4bc14bcb9112901b048fa31ef30921d9de9
3
- size 5048
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:faaf7f53e56d6e20fd96086f8191b4e1d4fe2ef644797a39d3eff73196a731f6
3
+ size 5176