Text Classification
Transformers
Safetensors
modernbert
agent-safety
tool-calling
long-context
distillation
Eval Results (legacy)
text-embeddings-inference
Instructions to use ProCreations/auto-0.4b-2 with libraries, inference providers, notebooks, and local apps. Follow these links to get started.
- Libraries
- Transformers
How to use ProCreations/auto-0.4b-2 with Transformers:
# Use a pipeline as a high-level helper from transformers import pipeline pipe = pipeline("text-classification", model="ProCreations/auto-0.4b-2")# pip install -U transformers accelerate # Load model directly from transformers import AutoTokenizer, AutoModelForSequenceClassification tokenizer = AutoTokenizer.from_pretrained("ProCreations/auto-0.4b-2") model = AutoModelForSequenceClassification.from_pretrained("ProCreations/auto-0.4b-2", device_map="auto") - Notebooks
- Google Colab
- Kaggle
Download data_audit.json from ProCreations/auto-0.4b-2: direct link, hf CLI and curl.
- Browser
- Download file 2.02 kB
-
https://huggingface.co/ProCreations/auto-0.4b-2/resolve/main/data_audit.json
- Command line
-
hf download hf://ProCreations/auto-0.4b-2/data_audit.json
-
curl -L -o data_audit.json https://huggingface.co/ProCreations/auto-0.4b-2/resolve/main/data_audit.json
2.02 kB
| { | |
| "removed_train": { | |
| "heldout_group": 14, | |
| "duplicate": 1 | |
| }, | |
| "removed_validation": {}, | |
| "alignment_sha256": "07c9c285026dfcc95c15c6454498b265138e36875c31ee3ad76b42a22d6049bd", | |
| "no_truncation": true, | |
| "synthetic_training_rows": 0, | |
| "teacher_target_split": "train only", | |
| "teacher_logits_sha256": "f78ba4870ad1dd79e669f71b6da238e3f35257c0175450b0ee70957ca89f0cce", | |
| "teacher_train_agreement": 0.997955013097186, | |
| "validation_partitions": { | |
| "selection": 7824, | |
| "calibration": 2581, | |
| "audit": 2595, | |
| "monitor": 7824 | |
| }, | |
| "train": { | |
| "rows": 711985, | |
| "tokens": 516733212, | |
| "max_length": 51408, | |
| "gt4096": 18578, | |
| "gt16384": 4202, | |
| "labels": { | |
| "approve": 355072, | |
| "deny": 356913 | |
| } | |
| }, | |
| "validation": { | |
| "rows": 13000, | |
| "tokens": 15660962, | |
| "max_length": 48710, | |
| "gt4096": 827, | |
| "gt16384": 185, | |
| "labels": { | |
| "approve": 6475, | |
| "deny": 6525 | |
| } | |
| }, | |
| "benchmark": { | |
| "rows": 3000, | |
| "tokens": 11710540, | |
| "max_length": 56176, | |
| "gt4096": 450, | |
| "gt16384": 239, | |
| "labels": { | |
| "deny": 1401, | |
| "approve": 1599 | |
| } | |
| }, | |
| "augmentation": { | |
| "rows": 120000, | |
| "teacher_deny_fraction": 0.6710166666666667, | |
| "teacher_confident_fraction": 0.8419916666666667, | |
| "student_tokens": 34321670, | |
| "teacher_logits_sha256": "aace8ec7b81715c7190c7a431968ba7c2bc8663167a52fbc3b7038d8bb148969", | |
| "max_chars": 6000, | |
| "seed": 20260910, | |
| "source_rows_pool": 687947, | |
| "original_label_of_call": { | |
| "deny": 59804, | |
| "approve": 60196 | |
| }, | |
| "lang": { | |
| "English": 83906, | |
| "German": 3603, | |
| "Hindi": 3461, | |
| "Italian": 3602, | |
| "Mandarin Chinese": 3780, | |
| "Korean": 3641, | |
| "Portuguese": 3676, | |
| "Russian": 3761, | |
| "French": 3478, | |
| "Japanese": 3502, | |
| "Spanish": 3590 | |
| } | |
| }, | |
| "stage_b_mixture": { | |
| "augmented_rows": 120000, | |
| "replay_rows": 173351, | |
| "replay_fraction": 0.25, | |
| "max_replay_length": 4096, | |
| "tokens": 103822740 | |
| } | |
| } | |