Norod78
/

hebrew_lyrics-gemma2_2b

Text Generation

text-generation-inference

Model card Files Files and versions

Norod78 commited on Feb 5

Commit

645e02e

·

verified ·

1 Parent(s): 60894af

Update config.json

Files changed (1) hide show

config.json +33 -33

config.json CHANGED Viewed

@@ -1,35 +1,35 @@
 {
-architectures: [
-"Gemma2ForCausalLM"
-],
-attention_bias: false,
-attention_dropout: 0,
-attn_logit_softcapping: 50,
-bos_token_id: 2,
-cache_implementation: "hybrid",
-eos_token_id: [
-1,
-107
-],
-final_logit_softcapping: 30,
-head_dim: 256,
-hidden_act: "gelu_pytorch_tanh",
-hidden_activation: "gelu_pytorch_tanh",
-hidden_size: 2304,
-initializer_range: 0.02,
-intermediate_size: 9216,
-max_position_embeddings: 8192,
-model_type: "gemma2",
-num_attention_heads: 8,
-num_hidden_layers: 26,
-num_key_value_heads: 4,
-pad_token_id: 0,
-query_pre_attn_scalar: 256,
-rms_norm_eps: 0.000001,
-rope_theta: 10000,
-sliding_window: 4096,
-torch_dtype: "bfloat16",
-transformers_version: "4.42.4",
-use_cache: true,
-vocab_size: 256000
 }

 {
+  "architectures": [
+    "Gemma2ForCausalLM"
+  ],
+  "attention_bias": false,
+  "attention_dropout": 0.0,
+  "attn_logit_softcapping": 50.0,
+  "bos_token_id": 2,
+  "cache_implementation": "hybrid",
+  "eos_token_id": [
+    1,
+    107
+  ],
+  "final_logit_softcapping": 30.0,
+  "head_dim": 256,
+  "hidden_act": "gelu_pytorch_tanh",
+  "hidden_activation": "gelu_pytorch_tanh",
+  "hidden_size": 2304,
+  "initializer_range": 0.02,
+  "intermediate_size": 9216,
+  "max_position_embeddings": 8192,
+  "model_type": "gemma2",
+  "num_attention_heads": 8,
+  "num_hidden_layers": 26,
+  "num_key_value_heads": 4,
+  "pad_token_id": 0,
+  "query_pre_attn_scalar": 256,
+  "rms_norm_eps": 1e-06,
+  "rope_theta": 10000.0,
+  "sliding_window": 4096,
+  "torch_dtype": "bfloat16",
+  "transformers_version": "4.42.4",
+  "use_cache": true,
+  "vocab_size": 256000
 }