Upload README.md with huggingface_hub
Browse files
README.md
CHANGED
|
@@ -7,12 +7,13 @@ license: mit
|
|
| 7 |
|
| 8 |
# dm_qwen4b_emulator
|
| 9 |
|
| 10 |
-
|
| 11 |
|
| 12 |
## Config
|
| 13 |
- `input_dim`: 6
|
| 14 |
-
- `hidden_dim`:
|
| 15 |
- `output_dim`: 3
|
|
|
|
| 16 |
|
| 17 |
## Usage
|
| 18 |
|
|
@@ -32,13 +33,16 @@ class MLP(nn.Module):
|
|
| 32 |
nn.Linear(hidden_dim, hidden_dim),
|
| 33 |
nn.LayerNorm(hidden_dim),
|
| 34 |
nn.ReLU(),
|
|
|
|
|
|
|
|
|
|
| 35 |
nn.Linear(hidden_dim, output_dim),
|
| 36 |
)
|
| 37 |
def forward(self, x):
|
| 38 |
return self.mlp(x)
|
| 39 |
|
| 40 |
path = hf_hub_download("anonom799/dm_qwen4b_emulator", "model.safetensors")
|
| 41 |
-
model = MLP(input_dim=6, hidden_dim=
|
| 42 |
model.load_state_dict(load_file(path))
|
| 43 |
model.eval()
|
| 44 |
|
|
|
|
| 7 |
|
| 8 |
# dm_qwen4b_emulator
|
| 9 |
|
| 10 |
+
MLP with 3 hidden Linear->LayerNorm->ReLU block(s) and a linear output head.
|
| 11 |
|
| 12 |
## Config
|
| 13 |
- `input_dim`: 6
|
| 14 |
+
- `hidden_dim`: 512
|
| 15 |
- `output_dim`: 3
|
| 16 |
+
- `n_layers`: 3
|
| 17 |
|
| 18 |
## Usage
|
| 19 |
|
|
|
|
| 33 |
nn.Linear(hidden_dim, hidden_dim),
|
| 34 |
nn.LayerNorm(hidden_dim),
|
| 35 |
nn.ReLU(),
|
| 36 |
+
nn.Linear(hidden_dim, hidden_dim),
|
| 37 |
+
nn.LayerNorm(hidden_dim),
|
| 38 |
+
nn.ReLU(),
|
| 39 |
nn.Linear(hidden_dim, output_dim),
|
| 40 |
)
|
| 41 |
def forward(self, x):
|
| 42 |
return self.mlp(x)
|
| 43 |
|
| 44 |
path = hf_hub_download("anonom799/dm_qwen4b_emulator", "model.safetensors")
|
| 45 |
+
model = MLP(input_dim=6, hidden_dim=512, output_dim=3)
|
| 46 |
model.load_state_dict(load_file(path))
|
| 47 |
model.eval()
|
| 48 |
|