{ "model_type": "gpt2", "architectures": ["GPT2LMHeadModel"], "n_layer": 36, "n_head": 20, "n_embd": 1280 }