Norod78
/

hebrew_lyrics-gemma2_2b

Text Generation

text-generation-inference

Model card Files Files and versions

Norod78 commited on Feb 5

Commit

cbb1862

·

verified ·

1 Parent(s): 266e595

Update config.json

Files changed (1) hide show

config.json +34 -35

config.json CHANGED Viewed

@@ -1,36 +1,35 @@
 {
-  "_name_or_path": "./hebrew_lyrics-gemma2_2b",
-  "architectures": [
-    "Gemma2ForCausalLM"
-  ],
-  "attention_bias": false,
-  "attention_dropout": 0.0,
-  "attn_logit_softcapping": 50.0,
-  "bos_token_id": 2,
-  "cache_implementation": "hybrid",
-  "eos_token_id": [
-    1,
-    107
-  ],
-  "final_logit_softcapping": 30.0,
-  "head_dim": 256,
-  "hidden_act": "gelu_pytorch_tanh",
-  "hidden_activation": "gelu_pytorch_tanh",
-  "hidden_size": 2304,
-  "initializer_range": 0.02,
-  "intermediate_size": 9216,
-  "max_position_embeddings": 8192,
-  "model_type": "gemma2",
-  "num_attention_heads": 8,
-  "num_hidden_layers": 26,
-  "num_key_value_heads": 4,
-  "pad_token_id": 0,
-  "query_pre_attn_scalar": 256,
-  "rms_norm_eps": 1e-06,
-  "rope_theta": 10000.0,
-  "sliding_window": 4096,
-  "torch_dtype": "bfloat16",
-  "transformers_version": "4.48.2",
-  "use_cache": true,
-  "vocab_size": 256000
-}

 {
+architectures: [
+"Gemma2ForCausalLM"
+],
+attention_bias: false,
+attention_dropout: 0,
+attn_logit_softcapping: 50,
+bos_token_id: 2,
+cache_implementation: "hybrid",
+eos_token_id: [
+1,
+107
+],
+final_logit_softcapping: 30,
+head_dim: 256,
+hidden_act: "gelu_pytorch_tanh",
+hidden_activation: "gelu_pytorch_tanh",
+hidden_size: 2304,
+initializer_range: 0.02,
+intermediate_size: 9216,
+max_position_embeddings: 8192,
+model_type: "gemma2",
+num_attention_heads: 8,
+num_hidden_layers: 26,
+num_key_value_heads: 4,
+pad_token_id: 0,
+query_pre_attn_scalar: 256,
+rms_norm_eps: 0.000001,
+rope_theta: 10000,
+sliding_window: 4096,
+torch_dtype: "bfloat16",
+transformers_version: "4.42.4",
+use_cache: true,
+vocab_size: 256000
+}