NextTerm-47M / config.json
N8Programs's picture
Upload 7 files
ac593f9 verified
raw
history blame contribute delete
380 Bytes
{
"model_type": "qwen3",
"hidden_size": 512,
"num_hidden_layers": 12,
"intermediate_size": 2048,
"num_attention_heads": 8,
"rms_norm_eps": 1e-06,
"vocab_size": 15,
"num_key_value_heads": 4,
"max_position_embeddings": 2048,
"rope_theta": 10000,
"head_dim": 64,
"tie_word_embeddings": false,
"bos_token_id": 12,
"eos_token_id": 13,
"pad_token_id": 14
}