Feature Extraction
MLX
Safetensors
multilingual
embedding_gemma2
mlx-vlm
embedding
sentence-similarity
multimodal
image-feature-extraction
audio-feature-extraction
video-feature-extraction
8-bit precision
Instructions to use mlx-community/embeddinggemma-2-8bit with libraries, inference providers, notebooks, and local apps. Follow these links to get started.
- Libraries
- MLX
How to use mlx-community/embeddinggemma-2-8bit with MLX:
# Download the model from the Hub pip install huggingface_hub[hf_xet] hf download mlx-community/embeddinggemma-2-8bit --local-dir embeddinggemma-2-8bit
- Notebooks
- Google Colab
- Kaggle
- Local Apps Settings
- LM Studio
- Atomic Chat
File size: 2,417 Bytes
7505ef2 | 1 2 3 4 5 6 7 8 9 10 11 12 13 14 15 16 17 18 19 20 21 22 23 24 25 26 27 28 29 30 31 32 33 34 35 36 37 38 39 40 41 42 43 44 45 46 47 48 49 50 51 52 53 54 55 56 57 58 59 60 61 62 63 64 65 66 67 68 69 70 71 72 73 74 75 76 77 78 | {
"source_model": "google/embeddinggemma-2",
"source_revision": "914f7f89142e33e77833254d9c9b90c3cef7303b",
"mlx_vlm_revision": "3d87e88402f307efbf68e568971aa887ee7d9ed0",
"variant": "8bit",
"hardware": "Apple M3 Ultra",
"mlx_version": "0.32.3",
"transformers_version": "5.18.0.dev0",
"validation_scope": "Six multilingual text inputs and synthetic image, audio, two-frame video, and text+image; numerical smoke checks, not a retrieval benchmark.",
"tensor_count": 1816,
"tensor_dtypes": [
"BF16",
"U32"
],
"weight_bytes": 1234238431,
"cases": {
"audio": {
"shape": [
1,
768
],
"finite": true,
"minimum_cosine_vs_torch_fp32": 0.9998035070404524,
"maximum_absolute_error_vs_torch_fp32": 0.002476184321109907,
"minimum_cosine_vs_source_bf16": 0.9997744741131873,
"minimum_truncated_cosine_vs_torch_fp32": 0.9998247233031456
},
"image": {
"shape": [
1,
768
],
"finite": true,
"minimum_cosine_vs_torch_fp32": 0.9998904477000761,
"maximum_absolute_error_vs_torch_fp32": 0.0019366481302732774,
"minimum_cosine_vs_source_bf16": 0.9999318244275229,
"minimum_truncated_cosine_vs_torch_fp32": 0.9998950619484888
},
"text": {
"shape": [
6,
768
],
"finite": true,
"minimum_cosine_vs_torch_fp32": 0.9998142728897385,
"maximum_absolute_error_vs_torch_fp32": 0.0022990193383245483,
"minimum_cosine_vs_source_bf16": 0.9997514709446282,
"minimum_truncated_cosine_vs_torch_fp32": 0.9998170484369469,
"retrieval_scores_venus_mars": [
0.678564727306366,
0.8342599868774414
]
},
"text_image": {
"shape": [
1,
768
],
"finite": true,
"minimum_cosine_vs_torch_fp32": 0.9998230728231623,
"maximum_absolute_error_vs_torch_fp32": 0.002104893632396264,
"minimum_cosine_vs_source_bf16": 0.999874734632146,
"minimum_truncated_cosine_vs_torch_fp32": 0.999830769497079
},
"video": {
"shape": [
1,
768
],
"finite": true,
"minimum_cosine_vs_torch_fp32": 0.9998359946491495,
"maximum_absolute_error_vs_torch_fp32": 0.0021304845535331423,
"minimum_cosine_vs_source_bf16": 0.9998825464984411,
"minimum_truncated_cosine_vs_torch_fp32": 0.9998420387979006
}
}
}
|