HenrySentinel
/

tinyMind

Text Generation

Model card Files Files and versions

tinyMind / configuration_tinymind.py

HenrySentinel's picture

Add configuration class

590734c verified 10 days ago

history blame contribute delete

671 Bytes

	"""TinyMind configuration."""
	from transformers import PretrainedConfig


	class TinyMindConfig(PretrainedConfig):
	model_type = "tiny_smart_llm"

	def __init__(self, vocab_size=50257, n_embd=256, n_heads=8, n_layers=6, max_seq_len=512, dropout=0.1, **kwargs):
	self.vocab_size = vocab_size
	self.n_embd = n_embd
	self.n_heads = n_heads
	self.n_layers = n_layers
	self.num_hidden_layers = n_layers
	self.hidden_size = n_embd
	self.num_attention_heads = n_heads
	self.max_seq_len = max_seq_len
	self.max_position_embeddings = max_seq_len
	self.dropout = dropout
	super().__init__(**kwargs)