Aleph-Alpha/Kolibri-1

hf Aleph-Alpha/Kolibri-1 2 sources, 6 claims · Watch

What each source says

PropertySourceSaidMeans here
Author
author
not compared
GGUF quantisationsHob-forge
receipt
Source
GGUF quantisations
Its words
Hob-forge
Read by
field:author
Said since
2026-10-04 06:13 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
What the source handed over
{
  "_id": "6ac181821b92ae3a1fdd3fbb",
  "author": "Hob-forge",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "de",
      "en"
    ],
    "library_name": "gguf",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "llama.cpp",
      "kolibri1",
      "mixture-of-experts",
      "reasoning",
      "tool-calling"
    ]
  },
  "createdAt": "2026-10-03T22:28:18.000Z",
  "downloads": 3268,
  "gated": false,
  "id": "Hob-forge/Kolibri-1-GGUF",
  "lastModified": "2026-10-04T05:12:11.000Z",
  "library_name": "gguf",
  "likes": 5,
  "modelId": "Hob-forge/Kolibri-1-GGUF",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "b08405e141706457126fa75b9f2f9e0729b30d0a",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "Kolibri-1-Q4_K_M.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00002-of-00002.gguf"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "kolibri1-llama.cpp.patch"
    }
  ],
  "tags": [
    "gguf",
    "llama.cpp",
    "kolibri1",
    "mixture-of-experts",
    "reasoning",
    "tool-calling",
    "text-generation",
    "de",
    "en",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "endpoints_compatible",
    "region:us",
    "conversational"
  ]
}
—
Author
author
not compared
Prompt48
receipt
Source
GGUF quantisations
Its words
Prompt48
Read by
field:author
Said since
2026-10-04 12:16 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
What the source handed over
{
  "_id": "6ac2408b2497d9f46d61cba6",
  "author": "Prompt48",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "en",
      "de"
    ],
    "library_name": "gguf",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "llama.cpp",
      "kolibri1",
      "mixture-of-experts"
    ]
  },
  "createdAt": "2026-10-04T12:03:23.000Z",
  "downloads": 333,
  "gated": false,
  "id": "Prompt48/Kolibri-1-GGUF",
  "lastModified": "2026-10-04T19:27:12.000Z",
  "library_name": "gguf",
  "likes": 2,
  "modelId": "Prompt48/Kolibri-1-GGUF",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "e4a85ea72e5eea480d01be4eb73284517e157f1c",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "Kolibri-1-Q4_K_M.gguf"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00002-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00002-of-00002.gguf"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "SHA256SUMS"
    },
    {
      "rfilename": "kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/.env.example"
    },
    {
      "rfilename": "tutorial/ATTRIBUTION.md"
    },
    {
      "rfilename": "tutorial/Kolibri-RunPod-Viewer-Code.zip"
    },
    {
      "rfilename": "tutorial/LICENSE"
    },
    {
      "rfilename": "tutorial/LLAMA-CPP-LICENSE"
    },
    {
      "rfilename": "tutorial/MODEL-CARD-EXAMPLE.md"
    },
    {
      "rfilename": "tutorial/README.md"
    },
    {
      "rfilename": "tutorial/evidence/fresh/download-source.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/executed-code.py"
    },
    {
      "rfilename": "tutorial/evidence/fresh/fresh-exercises.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/generated-summary.py"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/conversion.log"
    },
    {
      "rfilename": "tutorial/evidence/original/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/more-quants.log"
    },
    {
      "rfilename": "tutorial/evidence/original/remote-checksums-verified.json"
    },
    {
      "rfilename": "tutorial/evidence/original/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/source.json"
    },
    {
      "rfilename": "tutorial/historical/add-build-instructions.py"
    },
    {
      "rfilename": "tutorial/historical/convert-on-pod.sh"
    },
    {
      "rfilename": "tutorial/historical/more-quants.sh"
    },
    {
      "rfilename": "tutorial/historical/publish-more-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/publish-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/resume-conversion.sh"
    },
    {
      "rfilename": "tutorial/historical/validate-on-pod.py"
    },
    {
      "rfilename": "tutorial/kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/requirements.txt"
    },
    {
      "rfilename": "tutorial/runpod.py"
    },
    {
      "rfilename": "tutorial/steps/00-setup.sh"
    },
    {
      "rfilename": "tutorial/steps/10-download-source.py"
    },
    {
      "rfilename": "tutorial/steps/100-exercises.py"
    },
    {
      "rfilename": "tutorial/steps/20-convert-bf16.sh"
    },
    {
      "rfilename": "tutorial/steps/30-quantize.sh"
    },
    {
      "rfilename": "tutorial/steps/40-inspect.py"
    },
    {
      "rfilename": "tutorial/steps/50-serve.sh"
    },
    {
      "rfilename": "tutorial/steps/60-smoke-test.py"
    },
    {
      "rfilename": "tutorial/steps/70-split.py"
    },
    {
      "rfilename": "tutorial/steps/80-publish.py"
    },
    {
      "rfilename": "tutorial/steps/90-download-q4.py"
    },
    {
      "rfilename": "tutorial/video-production/README.md"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-comprehensive-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/build_full_video.py"
    },
    {
      "rfilename": "tutorial/video-production/build_video.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_exercises.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_full_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-SCRIPT.md"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-chapters.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-scenes.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-timing.json"
    },
    {
      "rfilename": "tutorial/video-production/finish_script.py"
    },
    {
      "rfilename": "tutorial/video-production/full_content.py"
    },
    {
      "rfilename": "tutorial/video-production/package_video.py"
    },
    {
      "rfilename": "tutorial/video-production/qa_video.py"
    },
    {
      "rfilename": "tutorial/video-production/render_cloud.py"
    },
    {
      "rfilename": "tutorial/video-production/tts_client.py"
    },
    {
      "rfilename": "tutorial/video-production/write_full_script.py"
    },
    {
      "rfilename": "tutorial/video-production/write_script.py"
    },
    {
      "rfilename": "validation/Q6_K/metadata.json"
    },
    {
      "rfilename": "validation/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "validation/Q8_0/metadata.json"
    },
    {
      "rfilename": "validation/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "validation/file-sizes.txt"
    },
    {
      "rfilename": "validation/llama-revision.txt"
    },
    {
      "rfilename": "validation/metadata.json"
    },
    {
      "rfilename": "validation/smoke-tests.json"
    },
    {
      "rfilename": "validation/source.json"
    }
  ],
  "tags": [
    "gguf",
    "llama.cpp",
    "kolibri1",
    "mixture-of-experts",
    "text-generation",
    "en",
    "de",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "endpoints_compatible",
    "region:us",
    "conversational"
  ]
}
—
Author
author
not compared
aparusel
receipt
Source
GGUF quantisations
Its words
aparusel
Read by
field:author
Said since
2026-10-05 18:25 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
What the source handed over
{
  "_id": "6ac3eb6ce74738c8515401e3",
  "author": "aparusel",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "de",
      "en"
    ],
    "library_name": "ds4",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "kolibri",
      "moe",
      "reasoning",
      "tool-calling",
      "dwarfstar"
    ]
  },
  "createdAt": "2026-10-05T18:24:44.000Z",
  "downloads": 0,
  "gated": false,
  "id": "aparusel/kolibri-1-gguf",
  "lastModified": "2026-10-05T18:24:47.000Z",
  "library_name": "ds4",
  "likes": 0,
  "modelId": "aparusel/kolibri-1-gguf",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "7acb9ebd3e1d75c4486c41d0a4141da1e9af5c48",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    }
  ],
  "tags": [
    "ds4",
    "gguf",
    "kolibri",
    "moe",
    "reasoning",
    "tool-calling",
    "dwarfstar",
    "text-generation",
    "de",
    "en",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "region:us"
  ]
}
—
Author
author
not compared
webmp3
receipt
Source
GGUF quantisations
Its words
webmp3
Read by
field:author
Said since
2026-10-05 18:25 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
What the source handed over
{
  "_id": "6ac3d55597599e961e9e1304",
  "author": "webmp3",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "library_name": "gguf",
    "license": "apache-2.0",
    "license_link": "LICENSE",
    "pipeline_tag": "text-generation",
    "tags": [
      "sakura",
      "sakura-micro",
      "gguf",
      "llama.cpp",
      "quantization",
      "imatrix",
      "mixed-quant",
      "kolibri1"
    ]
  },
  "createdAt": "2026-10-05T16:50:29.000Z",
  "downloads": 0,
  "gated": false,
  "id": "webmp3/Sakura-MicroQuality-Kolibri-1-365E-GGUF",
  "lastModified": "2026-10-05T17:04:43.000Z",
  "library_name": "gguf",
  "likes": 0,
  "modelId": "webmp3/Sakura-MicroQuality-Kolibri-1-365E-GGUF",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "b65a4f6360d484b2c142505c7869c3fc9df87df2",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "SHA256SUMS"
    },
    {
      "rfilename": "Sakura-Kolibri-1-365E-IQ2_XS-19.99GiB.gguf"
    },
    {
      "rfilename": "assets/sakura-logo.png"
    },
    {
      "rfilename": "assets/sakura-logo.svg"
    }
  ],
  "tags": [
    "gguf",
    "sakura",
    "sakura-micro",
    "llama.cpp",
    "quantization",
    "imatrix",
    "mixed-quant",
    "kolibri1",
    "text-generation",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "endpoints_compatible",
    "region:us",
    "conversational"
  ]
}
—
Author
author
not compared
Hugging Face modelsAleph-Alpha
receipt
Source
Hugging Face models
Its words
Aleph-Alpha
Read by
field:author
Said since
2026-10-04 06:13 UTC
Original
open at the source
What the source handed over
{
  "_asked": "Aleph-Alpha/Kolibri-1",
  "_id": "6abf46fa71befb841d4f4fff",
  "author": "Aleph-Alpha",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1-BF16",
    "language": [
      "de",
      "en"
    ],
    "library_name": "vllm",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "reasoning",
      "moe"
    ]
  },
  "config": {
    "architectures": [
      "Kolibri1ForCausalLM"
    ],
    "model_type": "kolibri1",
    "num_experts": 384,
    "num_experts_per_tok": 6,
    "quantization_config": {
      "modules_to_not_convert": [
        "model.layers.0.mlp.gate",
        "model.layers.1.mlp.gate",
        "model.layers.2.mlp.gate",
        "model.layers.3.mlp.gate",
        "model.layers.4.mlp.gate",
        "model.layers.5.mlp.gate",
        "model.layers.6.mlp.gate",
        "model.layers.7.mlp.gate",
        "model.layers.8.mlp.gate",
        "model.layers.9.mlp.gate",
        "model.layers.10.mlp.gate",
        "model.layers.11.mlp.gate",
        "model.layers.12.mlp.gate",
        "model.layers.13.mlp.gate",
        "model.layers.14.mlp.gate",
        "model.layers.15.mlp.gate",
        "model.layers.16.mlp.gate",
        "model.layers.17.mlp.gate",
        "model.layers.18.mlp.gate",
        "model.layers.19.mlp.gate",
        "model.layers.20.mlp.gate",
        "model.layers.21.mlp.gate",
        "model.layers.22.mlp.gate",
        "model.layers.23.mlp.gate",
        "model.layers.24.mlp.gate",
        "model.layers.25.mlp.gate",
        "model.layers.26.mlp.gate",
        "model.layers.27.mlp.gate",
        "model.layers.28.mlp.gate",
        "model.layers.29.mlp.gate",
        "model.layers.30.mlp.gate",
        "model.layers.31.mlp.gate",
        "model.layers.32.mlp.gate",
        "model.layers.33.mlp.gate",
        "model.layers.34.mlp.gate",
        "model.layers.35.mlp.gate",
        "model.layers.36.mlp.gate",
        "model.layers.37.mlp.gate",
        "model.layers.38.mlp.gate",
        "model.layers.39.mlp.gate",
        "model.layers.40.mlp.gate",
        "model.layers.41.mlp.gate",
        "model.layers.42.mlp.gate",
        "model.layers.43.mlp.gate",
        "model.layers.44.mlp.gate",
        "model.layers.45.mlp.gate",
        "model.layers.46.mlp.gate",
        "model.layers.47.mlp.gate",
        "model.layers.48.mlp.gate",
        "model.layers.49.mlp.gate"
      ],
      "quant_method": "fp8"
    },
    "tokenizer_config": {
      "bos_token": null,
      "chat_template": "{%- set _sv_reasoning_effort = reasoning_effort | default(none) -%}\n{%- set _sv_thinking_disabled = false -%}\n{%- if _sv_reasoning_effort is not none -%}\n    {%- set _sv_thinking_disabled = _sv_reasoning_effort == 'none' -%}\n{%- elif enable_thinking is defined and enable_thinking is false -%}\n    {%- set _sv_thinking_disabled = true -%}\n{%- endif -%}\n{%- set _sv_no_reasoning_sentence = \"Reasoning is disabled. Proceed straight to answering according to the user's instructions.\" -%}\n{%- set _sv_low_reasoning_sentence = \"Reasoning effort is set to low. Think briefly through only the essential steps in the user's language, then proceed directly to the answer.\" -%}\n{%- set _sv_medium_reasoning_sentence = \"Reasoning effort is set to medium. Think through the task methodically in the user's language, check key assumptions, and provide a well-supported answer.\" -%}\n{%- set _sv_high_reasoning_sentence = \"Reasoning effort is set to high. Think carefully through the task in the user's language, validate key assumptions, consider plausible alternatives, and prioritize correctness and clarity.\" -%}\n{%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- if _sv_thinking_disabled -%}\n    {%- set _sv_reasoning_sentence = _sv_no_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['minimal', 'low'] -%}\n    {%- set _sv_reasoning_sentence = _sv_low_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort == 'medium' -%}\n    {%- set _sv_reasoning_sentence = _sv_medium_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['high', 'xhigh', 'max'] or _sv_reasoning_effort is none -%}\n    {%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- endif -%}\n{%- set _sv_has_system = messages | length > 0 and messages[0].role == 'system' -%}\n{%- if tools or _sv_reasoning_sentence is not none or _sv_has_system %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if _sv_has_system %}\n        {{- messages[0].content }}\n    {%- endif %}\n    {%- if _sv_reasoning_sentence is not none %}\n        {%- if _sv_has_system %}{{- '\\n\\n' }}{%- endif %}\n        {{- '# Reasoning effort\\n\\n' + _sv_reasoning_sentence }}\n    {%- endif %}\n    {%- if tools %}\n        {%- if _sv_has_system or _sv_reasoning_sentence is not none %}{{- '\\n\\n' }}{%- endif %}\n        {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n        {%- for tool in tools %}\n            {{- \"\\n\" }}\n            {{- tool | tojson }}\n        {%- endfor %}\n        {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call>\" }}\n    {%- endif %}\n    {{- '<|im_end|>\\n' }}\n{%- endif %}\n{%- set ns = namespace(multi_step_tool=true, last_query_index=messages|length - 1) %}\n{%- for message in messages[::-1] %}\n    {%- set index = (messages|length - 1) - loop.index0 %}\n    {%- if ns.multi_step_tool and message.role == \"user\" and message.content is string and not(message.content.startswith('<tool_response>') and message.content.endswith('</tool_response>')) %}\n        {%- set ns.multi_step_tool = false %}\n        {%- set ns.last_query_index = index %}\n    {%- endif %}\n{%- endfor %}\n{%- for message in messages %}\n    {%- if (message.role == \"user\") or (message.role == \"system\" and not loop.first) %}\n        {{- '<|im_start|>' + message.role + '\\n' + message.content + '<|im_end|>' + '\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {%- set content = message.content if message.content is string else '' %}\n        {%- set reasoning = none %}\n        {%- if message.reasoning is defined and message.reasoning is not none %}\n            {%- set reasoning = message.reasoning %}\n        {%- elif message.reasoning_content is defined and message.reasoning_content is not none %}\n            {#- Deprecated vLLM compatibility. Prefer the `reasoning` field. -#}\n            {%- set reasoning = message.reasoning_content %}\n        {%- elif content is string and '</think>' in content %}\n            {%- set reasoning = content.split('</think>')[0].rstrip('\\n').split('<think>')[-1].lstrip('\\n') %}\n            {%- set content = content.split('</think>')[-1].lstrip('\\n') %}\n        {%- endif %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if (loop.index0 > ns.last_query_index) or (preserve_thinking is defined and preserve_thinking is true) %}\n            {{- '<think>\\n' }}\n            {%- if reasoning is not none and reasoning | trim %}\n                {{- reasoning.strip('\\n') }}\n            {%- endif %}\n            {{- '\\n</think>\\n\\n' }}\n        {%- endif %}\n        {{- content.lstrip('\\n') }}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {{- message.content }}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n    {%- if _sv_thinking_disabled %}\n        {{- '<think>\\n\\n</think>\\n\\n' }}\n    {%- endif %}\n{%- endif %}",
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-10-02T05:54:02.000Z",
  "disabled": false,
  "downloads": 2453,
  "gated": false,
  "id": "Aleph-Alpha/Kolibri-1",
  "lastModified": "2026-10-03T08:43:53.000Z",
  "library_name": "vllm",
  "likes": 523,
  "model-index": null,
  "modelId": "Aleph-Alpha/Kolibri-1",
  "pipeline_tag": "text-generation",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 705058560,
      "F8_E4M3": 77398016000
    },
    "total": 78103074560
  },
  "sha": "e52eb4627d11516b0c01de49210ab5a4e4061444",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "generation_config.json"
    },
    {
      "rfilename": "model-00001-of-00032.safetensors"
    },
    {
      "rfilename": "model-00002-of-00032.safetensors"
    },
    {
      "rfilename": "model-00003-of-00032.safetensors"
    },
    {
      "rfilename": "model-00004-of-00032.safetensors"
    },
    {
      "rfilename": "model-00005-of-00032.safetensors"
    },
    {
      "rfilename": "model-00006-of-00032.safetensors"
    },
    {
      "rfilename": "model-00007-of-00032.safetensors"
    },
    {
      "rfilename": "model-00008-of-00032.safetensors"
    },
    {
      "rfilename": "model-00009-of-00032.safetensors"
    },
    {
      "rfilename": "model-00010-of-00032.safetensors"
    },
    {
      "rfilename": "model-00011-of-00032.safetensors"
    },
    {
      "rfilename": "model-00012-of-00032.safetensors"
    },
    {
      "rfilename": "model-00013-of-00032.safetensors"
    },
    {
      "rfilename": "model-00014-of-00032.safetensors"
    },
    {
      "rfilename": "model-00015-of-00032.safetensors"
    },
    {
      "rfilename": "model-00016-of-00032.safetensors"
    },
    {
      "rfilename": "model-00017-of-00032.safetensors"
    },
    {
      "rfilename": "model-00018-of-00032.safetensors"
    },
    {
      "rfilename": "model-00019-of-00032.safetensors"
    },
    {
      "rfilename": "model-00020-of-00032.safetensors"
    },
    {
      "rfilename": "model-00021-of-00032.safetensors"
    },
    {
      "rfilename": "model-00022-of-00032.safetensors"
    },
    {
      "rfilename": "model-00023-of-00032.safetensors"
    },
    {
      "rfilename": "model-00024-of-00032.safetensors"
    },
    {
      "rfilename": "model-00025-of-00032.safetensors"
    },
    {
      "rfilename": "model-00026-of-00032.safetensors"
    },
    {
      "rfilename": "model-00027-of-00032.safetensors"
    },
    {
      "rfilename": "model-00028-of-00032.safetensors"
    },
    {
      "rfilename": "model-00029-of-00032.safetensors"
    },
    {
      "rfilename": "model-00030-of-00032.safetensors"
    },
    {
      "rfilename": "model-00031-of-00032.safetensors"
    },
    {
      "rfilename": "model-00032-of-00032.safetensors"
    },
    {
      "rfilename": "model.safetensors.index.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "tardellirs/model-pulse",
    "mrfakename/kolibri-1"
  ],
  "tags": [
    "vllm",
    "safetensors",
    "kolibri1",
    "reasoning",
    "moe",
    "text-generation",
    "conversational",
    "de",
    "en",
    "arxiv:2512.11614",
    "arxiv:2601.17858",
    "base_model:Aleph-Alpha/Kolibri-1-BF16",
    "base_model:quantized:Aleph-Alpha/Kolibri-1-BF16",
    "license:apache-2.0",
    "eval-results",
    "fp8",
    "region:us"
  ],
  "usedStorage": 78852606661
}
—
Base
base
GGUF quantisationsAleph-Alpha/Kolibri-1
receipt
Source
GGUF quantisations
Its words
Aleph-Alpha/Kolibri-1
Read by
field:cardData.base_model[]
Said since
2026-10-04 12:16 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
What the source handed over
{
  "_id": "6ac2408b2497d9f46d61cba6",
  "author": "Prompt48",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "en",
      "de"
    ],
    "library_name": "gguf",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "llama.cpp",
      "kolibri1",
      "mixture-of-experts"
    ]
  },
  "createdAt": "2026-10-04T12:03:23.000Z",
  "downloads": 333,
  "gated": false,
  "id": "Prompt48/Kolibri-1-GGUF",
  "lastModified": "2026-10-04T19:27:12.000Z",
  "library_name": "gguf",
  "likes": 2,
  "modelId": "Prompt48/Kolibri-1-GGUF",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "e4a85ea72e5eea480d01be4eb73284517e157f1c",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "Kolibri-1-Q4_K_M.gguf"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00002-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00002-of-00002.gguf"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "SHA256SUMS"
    },
    {
      "rfilename": "kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/.env.example"
    },
    {
      "rfilename": "tutorial/ATTRIBUTION.md"
    },
    {
      "rfilename": "tutorial/Kolibri-RunPod-Viewer-Code.zip"
    },
    {
      "rfilename": "tutorial/LICENSE"
    },
    {
      "rfilename": "tutorial/LLAMA-CPP-LICENSE"
    },
    {
      "rfilename": "tutorial/MODEL-CARD-EXAMPLE.md"
    },
    {
      "rfilename": "tutorial/README.md"
    },
    {
      "rfilename": "tutorial/evidence/fresh/download-source.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/executed-code.py"
    },
    {
      "rfilename": "tutorial/evidence/fresh/fresh-exercises.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/generated-summary.py"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/conversion.log"
    },
    {
      "rfilename": "tutorial/evidence/original/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/more-quants.log"
    },
    {
      "rfilename": "tutorial/evidence/original/remote-checksums-verified.json"
    },
    {
      "rfilename": "tutorial/evidence/original/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/source.json"
    },
    {
      "rfilename": "tutorial/historical/add-build-instructions.py"
    },
    {
      "rfilename": "tutorial/historical/convert-on-pod.sh"
    },
    {
      "rfilename": "tutorial/historical/more-quants.sh"
    },
    {
      "rfilename": "tutorial/historical/publish-more-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/publish-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/resume-conversion.sh"
    },
    {
      "rfilename": "tutorial/historical/validate-on-pod.py"
    },
    {
      "rfilename": "tutorial/kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/requirements.txt"
    },
    {
      "rfilename": "tutorial/runpod.py"
    },
    {
      "rfilename": "tutorial/steps/00-setup.sh"
    },
    {
      "rfilename": "tutorial/steps/10-download-source.py"
    },
    {
      "rfilename": "tutorial/steps/100-exercises.py"
    },
    {
      "rfilename": "tutorial/steps/20-convert-bf16.sh"
    },
    {
      "rfilename": "tutorial/steps/30-quantize.sh"
    },
    {
      "rfilename": "tutorial/steps/40-inspect.py"
    },
    {
      "rfilename": "tutorial/steps/50-serve.sh"
    },
    {
      "rfilename": "tutorial/steps/60-smoke-test.py"
    },
    {
      "rfilename": "tutorial/steps/70-split.py"
    },
    {
      "rfilename": "tutorial/steps/80-publish.py"
    },
    {
      "rfilename": "tutorial/steps/90-download-q4.py"
    },
    {
      "rfilename": "tutorial/video-production/README.md"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-comprehensive-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/build_full_video.py"
    },
    {
      "rfilename": "tutorial/video-production/build_video.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_exercises.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_full_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-SCRIPT.md"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-chapters.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-scenes.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-timing.json"
    },
    {
      "rfilename": "tutorial/video-production/finish_script.py"
    },
    {
      "rfilename": "tutorial/video-production/full_content.py"
    },
    {
      "rfilename": "tutorial/video-production/package_video.py"
    },
    {
      "rfilename": "tutorial/video-production/qa_video.py"
    },
    {
      "rfilename": "tutorial/video-production/render_cloud.py"
    },
    {
      "rfilename": "tutorial/video-production/tts_client.py"
    },
    {
      "rfilename": "tutorial/video-production/write_full_script.py"
    },
    {
      "rfilename": "tutorial/video-production/write_script.py"
    },
    {
      "rfilename": "validation/Q6_K/metadata.json"
    },
    {
      "rfilename": "validation/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "validation/Q8_0/metadata.json"
    },
    {
      "rfilename": "validation/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "validation/file-sizes.txt"
    },
    {
      "rfilename": "validation/llama-revision.txt"
    },
    {
      "rfilename": "validation/metadata.json"
    },
    {
      "rfilename": "validation/smoke-tests.json"
    },
    {
      "rfilename": "validation/source.json"
    }
  ],
  "tags": [
    "gguf",
    "llama.cpp",
    "kolibri1",
    "mixture-of-experts",
    "text-generation",
    "en",
    "de",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "endpoints_compatible",
    "region:us",
    "conversational"
  ]
}
—
Downloads
downloads
not compared
GGUF quantisations0
receipt
Source
GGUF quantisations
Its words
0
Read by
field:downloads
Said since
2026-10-05 18:25 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
What the source handed over
{
  "_id": "6ac3eb6ce74738c8515401e3",
  "author": "aparusel",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "de",
      "en"
    ],
    "library_name": "ds4",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "kolibri",
      "moe",
      "reasoning",
      "tool-calling",
      "dwarfstar"
    ]
  },
  "createdAt": "2026-10-05T18:24:44.000Z",
  "downloads": 0,
  "gated": false,
  "id": "aparusel/kolibri-1-gguf",
  "lastModified": "2026-10-05T18:24:47.000Z",
  "library_name": "ds4",
  "likes": 0,
  "modelId": "aparusel/kolibri-1-gguf",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "7acb9ebd3e1d75c4486c41d0a4141da1e9af5c48",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    }
  ],
  "tags": [
    "ds4",
    "gguf",
    "kolibri",
    "moe",
    "reasoning",
    "tool-calling",
    "dwarfstar",
    "text-generation",
    "de",
    "en",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "region:us"
  ]
}
—
Downloads
downloads
not compared
3268
receipt
Source
GGUF quantisations
Its words
3268
Read by
field:downloads
Said since
2026-10-05 12:24 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
2026-10-05 12:24 UTC3268
2026-10-04 12:16 UTC108
2026-10-04 06:13 UTC0
What the source handed over
{
  "_id": "6ac181821b92ae3a1fdd3fbb",
  "author": "Hob-forge",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "de",
      "en"
    ],
    "library_name": "gguf",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "llama.cpp",
      "kolibri1",
      "mixture-of-experts",
      "reasoning",
      "tool-calling"
    ]
  },
  "createdAt": "2026-10-03T22:28:18.000Z",
  "downloads": 3268,
  "gated": false,
  "id": "Hob-forge/Kolibri-1-GGUF",
  "lastModified": "2026-10-04T05:12:11.000Z",
  "library_name": "gguf",
  "likes": 5,
  "modelId": "Hob-forge/Kolibri-1-GGUF",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "b08405e141706457126fa75b9f2f9e0729b30d0a",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "Kolibri-1-Q4_K_M.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00002-of-00002.gguf"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "kolibri1-llama.cpp.patch"
    }
  ],
  "tags": [
    "gguf",
    "llama.cpp",
    "kolibri1",
    "mixture-of-experts",
    "reasoning",
    "tool-calling",
    "text-generation",
    "de",
    "en",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "endpoints_compatible",
    "region:us",
    "conversational"
  ]
}
—
Downloads
downloads
not compared
333
receipt
Source
GGUF quantisations
Its words
333
Read by
field:downloads
Said since
2026-10-05 12:24 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
2026-10-05 12:24 UTC333
2026-10-04 12:16 UTC0
What the source handed over
{
  "_id": "6ac2408b2497d9f46d61cba6",
  "author": "Prompt48",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "en",
      "de"
    ],
    "library_name": "gguf",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "llama.cpp",
      "kolibri1",
      "mixture-of-experts"
    ]
  },
  "createdAt": "2026-10-04T12:03:23.000Z",
  "downloads": 333,
  "gated": false,
  "id": "Prompt48/Kolibri-1-GGUF",
  "lastModified": "2026-10-04T19:27:12.000Z",
  "library_name": "gguf",
  "likes": 2,
  "modelId": "Prompt48/Kolibri-1-GGUF",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "e4a85ea72e5eea480d01be4eb73284517e157f1c",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "Kolibri-1-Q4_K_M.gguf"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00002-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00002-of-00002.gguf"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "SHA256SUMS"
    },
    {
      "rfilename": "kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/.env.example"
    },
    {
      "rfilename": "tutorial/ATTRIBUTION.md"
    },
    {
      "rfilename": "tutorial/Kolibri-RunPod-Viewer-Code.zip"
    },
    {
      "rfilename": "tutorial/LICENSE"
    },
    {
      "rfilename": "tutorial/LLAMA-CPP-LICENSE"
    },
    {
      "rfilename": "tutorial/MODEL-CARD-EXAMPLE.md"
    },
    {
      "rfilename": "tutorial/README.md"
    },
    {
      "rfilename": "tutorial/evidence/fresh/download-source.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/executed-code.py"
    },
    {
      "rfilename": "tutorial/evidence/fresh/fresh-exercises.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/generated-summary.py"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/conversion.log"
    },
    {
      "rfilename": "tutorial/evidence/original/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/more-quants.log"
    },
    {
      "rfilename": "tutorial/evidence/original/remote-checksums-verified.json"
    },
    {
      "rfilename": "tutorial/evidence/original/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/source.json"
    },
    {
      "rfilename": "tutorial/historical/add-build-instructions.py"
    },
    {
      "rfilename": "tutorial/historical/convert-on-pod.sh"
    },
    {
      "rfilename": "tutorial/historical/more-quants.sh"
    },
    {
      "rfilename": "tutorial/historical/publish-more-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/publish-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/resume-conversion.sh"
    },
    {
      "rfilename": "tutorial/historical/validate-on-pod.py"
    },
    {
      "rfilename": "tutorial/kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/requirements.txt"
    },
    {
      "rfilename": "tutorial/runpod.py"
    },
    {
      "rfilename": "tutorial/steps/00-setup.sh"
    },
    {
      "rfilename": "tutorial/steps/10-download-source.py"
    },
    {
      "rfilename": "tutorial/steps/100-exercises.py"
    },
    {
      "rfilename": "tutorial/steps/20-convert-bf16.sh"
    },
    {
      "rfilename": "tutorial/steps/30-quantize.sh"
    },
    {
      "rfilename": "tutorial/steps/40-inspect.py"
    },
    {
      "rfilename": "tutorial/steps/50-serve.sh"
    },
    {
      "rfilename": "tutorial/steps/60-smoke-test.py"
    },
    {
      "rfilename": "tutorial/steps/70-split.py"
    },
    {
      "rfilename": "tutorial/steps/80-publish.py"
    },
    {
      "rfilename": "tutorial/steps/90-download-q4.py"
    },
    {
      "rfilename": "tutorial/video-production/README.md"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-comprehensive-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/build_full_video.py"
    },
    {
      "rfilename": "tutorial/video-production/build_video.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_exercises.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_full_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-SCRIPT.md"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-chapters.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-scenes.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-timing.json"
    },
    {
      "rfilename": "tutorial/video-production/finish_script.py"
    },
    {
      "rfilename": "tutorial/video-production/full_content.py"
    },
    {
      "rfilename": "tutorial/video-production/package_video.py"
    },
    {
      "rfilename": "tutorial/video-production/qa_video.py"
    },
    {
      "rfilename": "tutorial/video-production/render_cloud.py"
    },
    {
      "rfilename": "tutorial/video-production/tts_client.py"
    },
    {
      "rfilename": "tutorial/video-production/write_full_script.py"
    },
    {
      "rfilename": "tutorial/video-production/write_script.py"
    },
    {
      "rfilename": "validation/Q6_K/metadata.json"
    },
    {
      "rfilename": "validation/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "validation/Q8_0/metadata.json"
    },
    {
      "rfilename": "validation/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "validation/file-sizes.txt"
    },
    {
      "rfilename": "validation/llama-revision.txt"
    },
    {
      "rfilename": "validation/metadata.json"
    },
    {
      "rfilename": "validation/smoke-tests.json"
    },
    {
      "rfilename": "validation/source.json"
    }
  ],
  "tags": [
    "gguf",
    "llama.cpp",
    "kolibri1",
    "mixture-of-experts",
    "text-generation",
    "en",
    "de",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "endpoints_compatible",
    "region:us",
    "conversational"
  ]
}
—
Downloads
downloads
not compared
Hugging Face models2453
receipt
Source
Hugging Face models
Its words
2453
Read by
field:downloads
Said since
2026-10-05 12:24 UTC
Original
open at the source
2026-10-05 12:24 UTC2453
2026-10-04 12:16 UTC1135
2026-10-04 06:13 UTC0
What the source handed over
{
  "_asked": "Aleph-Alpha/Kolibri-1",
  "_id": "6abf46fa71befb841d4f4fff",
  "author": "Aleph-Alpha",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1-BF16",
    "language": [
      "de",
      "en"
    ],
    "library_name": "vllm",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "reasoning",
      "moe"
    ]
  },
  "config": {
    "architectures": [
      "Kolibri1ForCausalLM"
    ],
    "model_type": "kolibri1",
    "num_experts": 384,
    "num_experts_per_tok": 6,
    "quantization_config": {
      "modules_to_not_convert": [
        "model.layers.0.mlp.gate",
        "model.layers.1.mlp.gate",
        "model.layers.2.mlp.gate",
        "model.layers.3.mlp.gate",
        "model.layers.4.mlp.gate",
        "model.layers.5.mlp.gate",
        "model.layers.6.mlp.gate",
        "model.layers.7.mlp.gate",
        "model.layers.8.mlp.gate",
        "model.layers.9.mlp.gate",
        "model.layers.10.mlp.gate",
        "model.layers.11.mlp.gate",
        "model.layers.12.mlp.gate",
        "model.layers.13.mlp.gate",
        "model.layers.14.mlp.gate",
        "model.layers.15.mlp.gate",
        "model.layers.16.mlp.gate",
        "model.layers.17.mlp.gate",
        "model.layers.18.mlp.gate",
        "model.layers.19.mlp.gate",
        "model.layers.20.mlp.gate",
        "model.layers.21.mlp.gate",
        "model.layers.22.mlp.gate",
        "model.layers.23.mlp.gate",
        "model.layers.24.mlp.gate",
        "model.layers.25.mlp.gate",
        "model.layers.26.mlp.gate",
        "model.layers.27.mlp.gate",
        "model.layers.28.mlp.gate",
        "model.layers.29.mlp.gate",
        "model.layers.30.mlp.gate",
        "model.layers.31.mlp.gate",
        "model.layers.32.mlp.gate",
        "model.layers.33.mlp.gate",
        "model.layers.34.mlp.gate",
        "model.layers.35.mlp.gate",
        "model.layers.36.mlp.gate",
        "model.layers.37.mlp.gate",
        "model.layers.38.mlp.gate",
        "model.layers.39.mlp.gate",
        "model.layers.40.mlp.gate",
        "model.layers.41.mlp.gate",
        "model.layers.42.mlp.gate",
        "model.layers.43.mlp.gate",
        "model.layers.44.mlp.gate",
        "model.layers.45.mlp.gate",
        "model.layers.46.mlp.gate",
        "model.layers.47.mlp.gate",
        "model.layers.48.mlp.gate",
        "model.layers.49.mlp.gate"
      ],
      "quant_method": "fp8"
    },
    "tokenizer_config": {
      "bos_token": null,
      "chat_template": "{%- set _sv_reasoning_effort = reasoning_effort | default(none) -%}\n{%- set _sv_thinking_disabled = false -%}\n{%- if _sv_reasoning_effort is not none -%}\n    {%- set _sv_thinking_disabled = _sv_reasoning_effort == 'none' -%}\n{%- elif enable_thinking is defined and enable_thinking is false -%}\n    {%- set _sv_thinking_disabled = true -%}\n{%- endif -%}\n{%- set _sv_no_reasoning_sentence = \"Reasoning is disabled. Proceed straight to answering according to the user's instructions.\" -%}\n{%- set _sv_low_reasoning_sentence = \"Reasoning effort is set to low. Think briefly through only the essential steps in the user's language, then proceed directly to the answer.\" -%}\n{%- set _sv_medium_reasoning_sentence = \"Reasoning effort is set to medium. Think through the task methodically in the user's language, check key assumptions, and provide a well-supported answer.\" -%}\n{%- set _sv_high_reasoning_sentence = \"Reasoning effort is set to high. Think carefully through the task in the user's language, validate key assumptions, consider plausible alternatives, and prioritize correctness and clarity.\" -%}\n{%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- if _sv_thinking_disabled -%}\n    {%- set _sv_reasoning_sentence = _sv_no_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['minimal', 'low'] -%}\n    {%- set _sv_reasoning_sentence = _sv_low_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort == 'medium' -%}\n    {%- set _sv_reasoning_sentence = _sv_medium_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['high', 'xhigh', 'max'] or _sv_reasoning_effort is none -%}\n    {%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- endif -%}\n{%- set _sv_has_system = messages | length > 0 and messages[0].role == 'system' -%}\n{%- if tools or _sv_reasoning_sentence is not none or _sv_has_system %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if _sv_has_system %}\n        {{- messages[0].content }}\n    {%- endif %}\n    {%- if _sv_reasoning_sentence is not none %}\n        {%- if _sv_has_system %}{{- '\\n\\n' }}{%- endif %}\n        {{- '# Reasoning effort\\n\\n' + _sv_reasoning_sentence }}\n    {%- endif %}\n    {%- if tools %}\n        {%- if _sv_has_system or _sv_reasoning_sentence is not none %}{{- '\\n\\n' }}{%- endif %}\n        {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n        {%- for tool in tools %}\n            {{- \"\\n\" }}\n            {{- tool | tojson }}\n        {%- endfor %}\n        {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call>\" }}\n    {%- endif %}\n    {{- '<|im_end|>\\n' }}\n{%- endif %}\n{%- set ns = namespace(multi_step_tool=true, last_query_index=messages|length - 1) %}\n{%- for message in messages[::-1] %}\n    {%- set index = (messages|length - 1) - loop.index0 %}\n    {%- if ns.multi_step_tool and message.role == \"user\" and message.content is string and not(message.content.startswith('<tool_response>') and message.content.endswith('</tool_response>')) %}\n        {%- set ns.multi_step_tool = false %}\n        {%- set ns.last_query_index = index %}\n    {%- endif %}\n{%- endfor %}\n{%- for message in messages %}\n    {%- if (message.role == \"user\") or (message.role == \"system\" and not loop.first) %}\n        {{- '<|im_start|>' + message.role + '\\n' + message.content + '<|im_end|>' + '\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {%- set content = message.content if message.content is string else '' %}\n        {%- set reasoning = none %}\n        {%- if message.reasoning is defined and message.reasoning is not none %}\n            {%- set reasoning = message.reasoning %}\n        {%- elif message.reasoning_content is defined and message.reasoning_content is not none %}\n            {#- Deprecated vLLM compatibility. Prefer the `reasoning` field. -#}\n            {%- set reasoning = message.reasoning_content %}\n        {%- elif content is string and '</think>' in content %}\n            {%- set reasoning = content.split('</think>')[0].rstrip('\\n').split('<think>')[-1].lstrip('\\n') %}\n            {%- set content = content.split('</think>')[-1].lstrip('\\n') %}\n        {%- endif %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if (loop.index0 > ns.last_query_index) or (preserve_thinking is defined and preserve_thinking is true) %}\n            {{- '<think>\\n' }}\n            {%- if reasoning is not none and reasoning | trim %}\n                {{- reasoning.strip('\\n') }}\n            {%- endif %}\n            {{- '\\n</think>\\n\\n' }}\n        {%- endif %}\n        {{- content.lstrip('\\n') }}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {{- message.content }}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n    {%- if _sv_thinking_disabled %}\n        {{- '<think>\\n\\n</think>\\n\\n' }}\n    {%- endif %}\n{%- endif %}",
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-10-02T05:54:02.000Z",
  "disabled": false,
  "downloads": 2453,
  "gated": false,
  "id": "Aleph-Alpha/Kolibri-1",
  "lastModified": "2026-10-03T08:43:53.000Z",
  "library_name": "vllm",
  "likes": 523,
  "model-index": null,
  "modelId": "Aleph-Alpha/Kolibri-1",
  "pipeline_tag": "text-generation",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 705058560,
      "F8_E4M3": 77398016000
    },
    "total": 78103074560
  },
  "sha": "e52eb4627d11516b0c01de49210ab5a4e4061444",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "generation_config.json"
    },
    {
      "rfilename": "model-00001-of-00032.safetensors"
    },
    {
      "rfilename": "model-00002-of-00032.safetensors"
    },
    {
      "rfilename": "model-00003-of-00032.safetensors"
    },
    {
      "rfilename": "model-00004-of-00032.safetensors"
    },
    {
      "rfilename": "model-00005-of-00032.safetensors"
    },
    {
      "rfilename": "model-00006-of-00032.safetensors"
    },
    {
      "rfilename": "model-00007-of-00032.safetensors"
    },
    {
      "rfilename": "model-00008-of-00032.safetensors"
    },
    {
      "rfilename": "model-00009-of-00032.safetensors"
    },
    {
      "rfilename": "model-00010-of-00032.safetensors"
    },
    {
      "rfilename": "model-00011-of-00032.safetensors"
    },
    {
      "rfilename": "model-00012-of-00032.safetensors"
    },
    {
      "rfilename": "model-00013-of-00032.safetensors"
    },
    {
      "rfilename": "model-00014-of-00032.safetensors"
    },
    {
      "rfilename": "model-00015-of-00032.safetensors"
    },
    {
      "rfilename": "model-00016-of-00032.safetensors"
    },
    {
      "rfilename": "model-00017-of-00032.safetensors"
    },
    {
      "rfilename": "model-00018-of-00032.safetensors"
    },
    {
      "rfilename": "model-00019-of-00032.safetensors"
    },
    {
      "rfilename": "model-00020-of-00032.safetensors"
    },
    {
      "rfilename": "model-00021-of-00032.safetensors"
    },
    {
      "rfilename": "model-00022-of-00032.safetensors"
    },
    {
      "rfilename": "model-00023-of-00032.safetensors"
    },
    {
      "rfilename": "model-00024-of-00032.safetensors"
    },
    {
      "rfilename": "model-00025-of-00032.safetensors"
    },
    {
      "rfilename": "model-00026-of-00032.safetensors"
    },
    {
      "rfilename": "model-00027-of-00032.safetensors"
    },
    {
      "rfilename": "model-00028-of-00032.safetensors"
    },
    {
      "rfilename": "model-00029-of-00032.safetensors"
    },
    {
      "rfilename": "model-00030-of-00032.safetensors"
    },
    {
      "rfilename": "model-00031-of-00032.safetensors"
    },
    {
      "rfilename": "model-00032-of-00032.safetensors"
    },
    {
      "rfilename": "model.safetensors.index.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "tardellirs/model-pulse",
    "mrfakename/kolibri-1"
  ],
  "tags": [
    "vllm",
    "safetensors",
    "kolibri1",
    "reasoning",
    "moe",
    "text-generation",
    "conversational",
    "de",
    "en",
    "arxiv:2512.11614",
    "arxiv:2601.17858",
    "base_model:Aleph-Alpha/Kolibri-1-BF16",
    "base_model:quantized:Aleph-Alpha/Kolibri-1-BF16",
    "license:apache-2.0",
    "eval-results",
    "fp8",
    "region:us"
  ],
  "usedStorage": 78852606661
}
—
Gated
gated
Hugging Face modelsfalse
receipt
Source
Hugging Face models
Its words
false
Read by
field:gated
Said since
2026-10-04 06:13 UTC
Original
open at the source
What the source handed over
{
  "_asked": "Aleph-Alpha/Kolibri-1",
  "_id": "6abf46fa71befb841d4f4fff",
  "author": "Aleph-Alpha",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1-BF16",
    "language": [
      "de",
      "en"
    ],
    "library_name": "vllm",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "reasoning",
      "moe"
    ]
  },
  "config": {
    "architectures": [
      "Kolibri1ForCausalLM"
    ],
    "model_type": "kolibri1",
    "num_experts": 384,
    "num_experts_per_tok": 6,
    "quantization_config": {
      "modules_to_not_convert": [
        "model.layers.0.mlp.gate",
        "model.layers.1.mlp.gate",
        "model.layers.2.mlp.gate",
        "model.layers.3.mlp.gate",
        "model.layers.4.mlp.gate",
        "model.layers.5.mlp.gate",
        "model.layers.6.mlp.gate",
        "model.layers.7.mlp.gate",
        "model.layers.8.mlp.gate",
        "model.layers.9.mlp.gate",
        "model.layers.10.mlp.gate",
        "model.layers.11.mlp.gate",
        "model.layers.12.mlp.gate",
        "model.layers.13.mlp.gate",
        "model.layers.14.mlp.gate",
        "model.layers.15.mlp.gate",
        "model.layers.16.mlp.gate",
        "model.layers.17.mlp.gate",
        "model.layers.18.mlp.gate",
        "model.layers.19.mlp.gate",
        "model.layers.20.mlp.gate",
        "model.layers.21.mlp.gate",
        "model.layers.22.mlp.gate",
        "model.layers.23.mlp.gate",
        "model.layers.24.mlp.gate",
        "model.layers.25.mlp.gate",
        "model.layers.26.mlp.gate",
        "model.layers.27.mlp.gate",
        "model.layers.28.mlp.gate",
        "model.layers.29.mlp.gate",
        "model.layers.30.mlp.gate",
        "model.layers.31.mlp.gate",
        "model.layers.32.mlp.gate",
        "model.layers.33.mlp.gate",
        "model.layers.34.mlp.gate",
        "model.layers.35.mlp.gate",
        "model.layers.36.mlp.gate",
        "model.layers.37.mlp.gate",
        "model.layers.38.mlp.gate",
        "model.layers.39.mlp.gate",
        "model.layers.40.mlp.gate",
        "model.layers.41.mlp.gate",
        "model.layers.42.mlp.gate",
        "model.layers.43.mlp.gate",
        "model.layers.44.mlp.gate",
        "model.layers.45.mlp.gate",
        "model.layers.46.mlp.gate",
        "model.layers.47.mlp.gate",
        "model.layers.48.mlp.gate",
        "model.layers.49.mlp.gate"
      ],
      "quant_method": "fp8"
    },
    "tokenizer_config": {
      "bos_token": null,
      "chat_template": "{%- set _sv_reasoning_effort = reasoning_effort | default(none) -%}\n{%- set _sv_thinking_disabled = false -%}\n{%- if _sv_reasoning_effort is not none -%}\n    {%- set _sv_thinking_disabled = _sv_reasoning_effort == 'none' -%}\n{%- elif enable_thinking is defined and enable_thinking is false -%}\n    {%- set _sv_thinking_disabled = true -%}\n{%- endif -%}\n{%- set _sv_no_reasoning_sentence = \"Reasoning is disabled. Proceed straight to answering according to the user's instructions.\" -%}\n{%- set _sv_low_reasoning_sentence = \"Reasoning effort is set to low. Think briefly through only the essential steps in the user's language, then proceed directly to the answer.\" -%}\n{%- set _sv_medium_reasoning_sentence = \"Reasoning effort is set to medium. Think through the task methodically in the user's language, check key assumptions, and provide a well-supported answer.\" -%}\n{%- set _sv_high_reasoning_sentence = \"Reasoning effort is set to high. Think carefully through the task in the user's language, validate key assumptions, consider plausible alternatives, and prioritize correctness and clarity.\" -%}\n{%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- if _sv_thinking_disabled -%}\n    {%- set _sv_reasoning_sentence = _sv_no_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['minimal', 'low'] -%}\n    {%- set _sv_reasoning_sentence = _sv_low_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort == 'medium' -%}\n    {%- set _sv_reasoning_sentence = _sv_medium_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['high', 'xhigh', 'max'] or _sv_reasoning_effort is none -%}\n    {%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- endif -%}\n{%- set _sv_has_system = messages | length > 0 and messages[0].role == 'system' -%}\n{%- if tools or _sv_reasoning_sentence is not none or _sv_has_system %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if _sv_has_system %}\n        {{- messages[0].content }}\n    {%- endif %}\n    {%- if _sv_reasoning_sentence is not none %}\n        {%- if _sv_has_system %}{{- '\\n\\n' }}{%- endif %}\n        {{- '# Reasoning effort\\n\\n' + _sv_reasoning_sentence }}\n    {%- endif %}\n    {%- if tools %}\n        {%- if _sv_has_system or _sv_reasoning_sentence is not none %}{{- '\\n\\n' }}{%- endif %}\n        {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n        {%- for tool in tools %}\n            {{- \"\\n\" }}\n            {{- tool | tojson }}\n        {%- endfor %}\n        {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call>\" }}\n    {%- endif %}\n    {{- '<|im_end|>\\n' }}\n{%- endif %}\n{%- set ns = namespace(multi_step_tool=true, last_query_index=messages|length - 1) %}\n{%- for message in messages[::-1] %}\n    {%- set index = (messages|length - 1) - loop.index0 %}\n    {%- if ns.multi_step_tool and message.role == \"user\" and message.content is string and not(message.content.startswith('<tool_response>') and message.content.endswith('</tool_response>')) %}\n        {%- set ns.multi_step_tool = false %}\n        {%- set ns.last_query_index = index %}\n    {%- endif %}\n{%- endfor %}\n{%- for message in messages %}\n    {%- if (message.role == \"user\") or (message.role == \"system\" and not loop.first) %}\n        {{- '<|im_start|>' + message.role + '\\n' + message.content + '<|im_end|>' + '\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {%- set content = message.content if message.content is string else '' %}\n        {%- set reasoning = none %}\n        {%- if message.reasoning is defined and message.reasoning is not none %}\n            {%- set reasoning = message.reasoning %}\n        {%- elif message.reasoning_content is defined and message.reasoning_content is not none %}\n            {#- Deprecated vLLM compatibility. Prefer the `reasoning` field. -#}\n            {%- set reasoning = message.reasoning_content %}\n        {%- elif content is string and '</think>' in content %}\n            {%- set reasoning = content.split('</think>')[0].rstrip('\\n').split('<think>')[-1].lstrip('\\n') %}\n            {%- set content = content.split('</think>')[-1].lstrip('\\n') %}\n        {%- endif %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if (loop.index0 > ns.last_query_index) or (preserve_thinking is defined and preserve_thinking is true) %}\n            {{- '<think>\\n' }}\n            {%- if reasoning is not none and reasoning | trim %}\n                {{- reasoning.strip('\\n') }}\n            {%- endif %}\n            {{- '\\n</think>\\n\\n' }}\n        {%- endif %}\n        {{- content.lstrip('\\n') }}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {{- message.content }}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n    {%- if _sv_thinking_disabled %}\n        {{- '<think>\\n\\n</think>\\n\\n' }}\n    {%- endif %}\n{%- endif %}",
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-10-02T05:54:02.000Z",
  "disabled": false,
  "downloads": 2453,
  "gated": false,
  "id": "Aleph-Alpha/Kolibri-1",
  "lastModified": "2026-10-03T08:43:53.000Z",
  "library_name": "vllm",
  "likes": 523,
  "model-index": null,
  "modelId": "Aleph-Alpha/Kolibri-1",
  "pipeline_tag": "text-generation",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 705058560,
      "F8_E4M3": 77398016000
    },
    "total": 78103074560
  },
  "sha": "e52eb4627d11516b0c01de49210ab5a4e4061444",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "generation_config.json"
    },
    {
      "rfilename": "model-00001-of-00032.safetensors"
    },
    {
      "rfilename": "model-00002-of-00032.safetensors"
    },
    {
      "rfilename": "model-00003-of-00032.safetensors"
    },
    {
      "rfilename": "model-00004-of-00032.safetensors"
    },
    {
      "rfilename": "model-00005-of-00032.safetensors"
    },
    {
      "rfilename": "model-00006-of-00032.safetensors"
    },
    {
      "rfilename": "model-00007-of-00032.safetensors"
    },
    {
      "rfilename": "model-00008-of-00032.safetensors"
    },
    {
      "rfilename": "model-00009-of-00032.safetensors"
    },
    {
      "rfilename": "model-00010-of-00032.safetensors"
    },
    {
      "rfilename": "model-00011-of-00032.safetensors"
    },
    {
      "rfilename": "model-00012-of-00032.safetensors"
    },
    {
      "rfilename": "model-00013-of-00032.safetensors"
    },
    {
      "rfilename": "model-00014-of-00032.safetensors"
    },
    {
      "rfilename": "model-00015-of-00032.safetensors"
    },
    {
      "rfilename": "model-00016-of-00032.safetensors"
    },
    {
      "rfilename": "model-00017-of-00032.safetensors"
    },
    {
      "rfilename": "model-00018-of-00032.safetensors"
    },
    {
      "rfilename": "model-00019-of-00032.safetensors"
    },
    {
      "rfilename": "model-00020-of-00032.safetensors"
    },
    {
      "rfilename": "model-00021-of-00032.safetensors"
    },
    {
      "rfilename": "model-00022-of-00032.safetensors"
    },
    {
      "rfilename": "model-00023-of-00032.safetensors"
    },
    {
      "rfilename": "model-00024-of-00032.safetensors"
    },
    {
      "rfilename": "model-00025-of-00032.safetensors"
    },
    {
      "rfilename": "model-00026-of-00032.safetensors"
    },
    {
      "rfilename": "model-00027-of-00032.safetensors"
    },
    {
      "rfilename": "model-00028-of-00032.safetensors"
    },
    {
      "rfilename": "model-00029-of-00032.safetensors"
    },
    {
      "rfilename": "model-00030-of-00032.safetensors"
    },
    {
      "rfilename": "model-00031-of-00032.safetensors"
    },
    {
      "rfilename": "model-00032-of-00032.safetensors"
    },
    {
      "rfilename": "model.safetensors.index.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "tardellirs/model-pulse",
    "mrfakename/kolibri-1"
  ],
  "tags": [
    "vllm",
    "safetensors",
    "kolibri1",
    "reasoning",
    "moe",
    "text-generation",
    "conversational",
    "de",
    "en",
    "arxiv:2512.11614",
    "arxiv:2601.17858",
    "base_model:Aleph-Alpha/Kolibri-1-BF16",
    "base_model:quantized:Aleph-Alpha/Kolibri-1-BF16",
    "license:apache-2.0",
    "eval-results",
    "fp8",
    "region:us"
  ],
  "usedStorage": 78852606661
}
—
Licence
licence
GGUF quantisationsapache-2.0
receipt
Source
GGUF quantisations
Its words
apache-2.0
Read by
field:cardData.license
Said since
2026-10-04 12:16 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
What the source handed over
{
  "_id": "6ac2408b2497d9f46d61cba6",
  "author": "Prompt48",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "en",
      "de"
    ],
    "library_name": "gguf",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "llama.cpp",
      "kolibri1",
      "mixture-of-experts"
    ]
  },
  "createdAt": "2026-10-04T12:03:23.000Z",
  "downloads": 333,
  "gated": false,
  "id": "Prompt48/Kolibri-1-GGUF",
  "lastModified": "2026-10-04T19:27:12.000Z",
  "library_name": "gguf",
  "likes": 2,
  "modelId": "Prompt48/Kolibri-1-GGUF",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "e4a85ea72e5eea480d01be4eb73284517e157f1c",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "Kolibri-1-Q4_K_M.gguf"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00002-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00002-of-00002.gguf"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "SHA256SUMS"
    },
    {
      "rfilename": "kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/.env.example"
    },
    {
      "rfilename": "tutorial/ATTRIBUTION.md"
    },
    {
      "rfilename": "tutorial/Kolibri-RunPod-Viewer-Code.zip"
    },
    {
      "rfilename": "tutorial/LICENSE"
    },
    {
      "rfilename": "tutorial/LLAMA-CPP-LICENSE"
    },
    {
      "rfilename": "tutorial/MODEL-CARD-EXAMPLE.md"
    },
    {
      "rfilename": "tutorial/README.md"
    },
    {
      "rfilename": "tutorial/evidence/fresh/download-source.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/executed-code.py"
    },
    {
      "rfilename": "tutorial/evidence/fresh/fresh-exercises.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/generated-summary.py"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/conversion.log"
    },
    {
      "rfilename": "tutorial/evidence/original/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/more-quants.log"
    },
    {
      "rfilename": "tutorial/evidence/original/remote-checksums-verified.json"
    },
    {
      "rfilename": "tutorial/evidence/original/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/source.json"
    },
    {
      "rfilename": "tutorial/historical/add-build-instructions.py"
    },
    {
      "rfilename": "tutorial/historical/convert-on-pod.sh"
    },
    {
      "rfilename": "tutorial/historical/more-quants.sh"
    },
    {
      "rfilename": "tutorial/historical/publish-more-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/publish-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/resume-conversion.sh"
    },
    {
      "rfilename": "tutorial/historical/validate-on-pod.py"
    },
    {
      "rfilename": "tutorial/kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/requirements.txt"
    },
    {
      "rfilename": "tutorial/runpod.py"
    },
    {
      "rfilename": "tutorial/steps/00-setup.sh"
    },
    {
      "rfilename": "tutorial/steps/10-download-source.py"
    },
    {
      "rfilename": "tutorial/steps/100-exercises.py"
    },
    {
      "rfilename": "tutorial/steps/20-convert-bf16.sh"
    },
    {
      "rfilename": "tutorial/steps/30-quantize.sh"
    },
    {
      "rfilename": "tutorial/steps/40-inspect.py"
    },
    {
      "rfilename": "tutorial/steps/50-serve.sh"
    },
    {
      "rfilename": "tutorial/steps/60-smoke-test.py"
    },
    {
      "rfilename": "tutorial/steps/70-split.py"
    },
    {
      "rfilename": "tutorial/steps/80-publish.py"
    },
    {
      "rfilename": "tutorial/steps/90-download-q4.py"
    },
    {
      "rfilename": "tutorial/video-production/README.md"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-comprehensive-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/build_full_video.py"
    },
    {
      "rfilename": "tutorial/video-production/build_video.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_exercises.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_full_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-SCRIPT.md"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-chapters.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-scenes.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-timing.json"
    },
    {
      "rfilename": "tutorial/video-production/finish_script.py"
    },
    {
      "rfilename": "tutorial/video-production/full_content.py"
    },
    {
      "rfilename": "tutorial/video-production/package_video.py"
    },
    {
      "rfilename": "tutorial/video-production/qa_video.py"
    },
    {
      "rfilename": "tutorial/video-production/render_cloud.py"
    },
    {
      "rfilename": "tutorial/video-production/tts_client.py"
    },
    {
      "rfilename": "tutorial/video-production/write_full_script.py"
    },
    {
      "rfilename": "tutorial/video-production/write_script.py"
    },
    {
      "rfilename": "validation/Q6_K/metadata.json"
    },
    {
      "rfilename": "validation/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "validation/Q8_0/metadata.json"
    },
    {
      "rfilename": "validation/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "validation/file-sizes.txt"
    },
    {
      "rfilename": "validation/llama-revision.txt"
    },
    {
      "rfilename": "validation/metadata.json"
    },
    {
      "rfilename": "validation/smoke-tests.json"
    },
    {
      "rfilename": "validation/source.json"
    }
  ],
  "tags": [
    "gguf",
    "llama.cpp",
    "kolibri1",
    "mixture-of-experts",
    "text-generation",
    "en",
    "de",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "endpoints_compatible",
    "region:us",
    "conversational"
  ]
}
—
Licence
licence
Hugging Face modelsapache-2.0
receipt
Source
Hugging Face models
Its words
apache-2.0
Read by
field:cardData.license
Said since
2026-10-04 06:13 UTC
Original
open at the source
What the source handed over
{
  "_asked": "Aleph-Alpha/Kolibri-1",
  "_id": "6abf46fa71befb841d4f4fff",
  "author": "Aleph-Alpha",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1-BF16",
    "language": [
      "de",
      "en"
    ],
    "library_name": "vllm",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "reasoning",
      "moe"
    ]
  },
  "config": {
    "architectures": [
      "Kolibri1ForCausalLM"
    ],
    "model_type": "kolibri1",
    "num_experts": 384,
    "num_experts_per_tok": 6,
    "quantization_config": {
      "modules_to_not_convert": [
        "model.layers.0.mlp.gate",
        "model.layers.1.mlp.gate",
        "model.layers.2.mlp.gate",
        "model.layers.3.mlp.gate",
        "model.layers.4.mlp.gate",
        "model.layers.5.mlp.gate",
        "model.layers.6.mlp.gate",
        "model.layers.7.mlp.gate",
        "model.layers.8.mlp.gate",
        "model.layers.9.mlp.gate",
        "model.layers.10.mlp.gate",
        "model.layers.11.mlp.gate",
        "model.layers.12.mlp.gate",
        "model.layers.13.mlp.gate",
        "model.layers.14.mlp.gate",
        "model.layers.15.mlp.gate",
        "model.layers.16.mlp.gate",
        "model.layers.17.mlp.gate",
        "model.layers.18.mlp.gate",
        "model.layers.19.mlp.gate",
        "model.layers.20.mlp.gate",
        "model.layers.21.mlp.gate",
        "model.layers.22.mlp.gate",
        "model.layers.23.mlp.gate",
        "model.layers.24.mlp.gate",
        "model.layers.25.mlp.gate",
        "model.layers.26.mlp.gate",
        "model.layers.27.mlp.gate",
        "model.layers.28.mlp.gate",
        "model.layers.29.mlp.gate",
        "model.layers.30.mlp.gate",
        "model.layers.31.mlp.gate",
        "model.layers.32.mlp.gate",
        "model.layers.33.mlp.gate",
        "model.layers.34.mlp.gate",
        "model.layers.35.mlp.gate",
        "model.layers.36.mlp.gate",
        "model.layers.37.mlp.gate",
        "model.layers.38.mlp.gate",
        "model.layers.39.mlp.gate",
        "model.layers.40.mlp.gate",
        "model.layers.41.mlp.gate",
        "model.layers.42.mlp.gate",
        "model.layers.43.mlp.gate",
        "model.layers.44.mlp.gate",
        "model.layers.45.mlp.gate",
        "model.layers.46.mlp.gate",
        "model.layers.47.mlp.gate",
        "model.layers.48.mlp.gate",
        "model.layers.49.mlp.gate"
      ],
      "quant_method": "fp8"
    },
    "tokenizer_config": {
      "bos_token": null,
      "chat_template": "{%- set _sv_reasoning_effort = reasoning_effort | default(none) -%}\n{%- set _sv_thinking_disabled = false -%}\n{%- if _sv_reasoning_effort is not none -%}\n    {%- set _sv_thinking_disabled = _sv_reasoning_effort == 'none' -%}\n{%- elif enable_thinking is defined and enable_thinking is false -%}\n    {%- set _sv_thinking_disabled = true -%}\n{%- endif -%}\n{%- set _sv_no_reasoning_sentence = \"Reasoning is disabled. Proceed straight to answering according to the user's instructions.\" -%}\n{%- set _sv_low_reasoning_sentence = \"Reasoning effort is set to low. Think briefly through only the essential steps in the user's language, then proceed directly to the answer.\" -%}\n{%- set _sv_medium_reasoning_sentence = \"Reasoning effort is set to medium. Think through the task methodically in the user's language, check key assumptions, and provide a well-supported answer.\" -%}\n{%- set _sv_high_reasoning_sentence = \"Reasoning effort is set to high. Think carefully through the task in the user's language, validate key assumptions, consider plausible alternatives, and prioritize correctness and clarity.\" -%}\n{%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- if _sv_thinking_disabled -%}\n    {%- set _sv_reasoning_sentence = _sv_no_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['minimal', 'low'] -%}\n    {%- set _sv_reasoning_sentence = _sv_low_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort == 'medium' -%}\n    {%- set _sv_reasoning_sentence = _sv_medium_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['high', 'xhigh', 'max'] or _sv_reasoning_effort is none -%}\n    {%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- endif -%}\n{%- set _sv_has_system = messages | length > 0 and messages[0].role == 'system' -%}\n{%- if tools or _sv_reasoning_sentence is not none or _sv_has_system %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if _sv_has_system %}\n        {{- messages[0].content }}\n    {%- endif %}\n    {%- if _sv_reasoning_sentence is not none %}\n        {%- if _sv_has_system %}{{- '\\n\\n' }}{%- endif %}\n        {{- '# Reasoning effort\\n\\n' + _sv_reasoning_sentence }}\n    {%- endif %}\n    {%- if tools %}\n        {%- if _sv_has_system or _sv_reasoning_sentence is not none %}{{- '\\n\\n' }}{%- endif %}\n        {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n        {%- for tool in tools %}\n            {{- \"\\n\" }}\n            {{- tool | tojson }}\n        {%- endfor %}\n        {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call>\" }}\n    {%- endif %}\n    {{- '<|im_end|>\\n' }}\n{%- endif %}\n{%- set ns = namespace(multi_step_tool=true, last_query_index=messages|length - 1) %}\n{%- for message in messages[::-1] %}\n    {%- set index = (messages|length - 1) - loop.index0 %}\n    {%- if ns.multi_step_tool and message.role == \"user\" and message.content is string and not(message.content.startswith('<tool_response>') and message.content.endswith('</tool_response>')) %}\n        {%- set ns.multi_step_tool = false %}\n        {%- set ns.last_query_index = index %}\n    {%- endif %}\n{%- endfor %}\n{%- for message in messages %}\n    {%- if (message.role == \"user\") or (message.role == \"system\" and not loop.first) %}\n        {{- '<|im_start|>' + message.role + '\\n' + message.content + '<|im_end|>' + '\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {%- set content = message.content if message.content is string else '' %}\n        {%- set reasoning = none %}\n        {%- if message.reasoning is defined and message.reasoning is not none %}\n            {%- set reasoning = message.reasoning %}\n        {%- elif message.reasoning_content is defined and message.reasoning_content is not none %}\n            {#- Deprecated vLLM compatibility. Prefer the `reasoning` field. -#}\n            {%- set reasoning = message.reasoning_content %}\n        {%- elif content is string and '</think>' in content %}\n            {%- set reasoning = content.split('</think>')[0].rstrip('\\n').split('<think>')[-1].lstrip('\\n') %}\n            {%- set content = content.split('</think>')[-1].lstrip('\\n') %}\n        {%- endif %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if (loop.index0 > ns.last_query_index) or (preserve_thinking is defined and preserve_thinking is true) %}\n            {{- '<think>\\n' }}\n            {%- if reasoning is not none and reasoning | trim %}\n                {{- reasoning.strip('\\n') }}\n            {%- endif %}\n            {{- '\\n</think>\\n\\n' }}\n        {%- endif %}\n        {{- content.lstrip('\\n') }}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {{- message.content }}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n    {%- if _sv_thinking_disabled %}\n        {{- '<think>\\n\\n</think>\\n\\n' }}\n    {%- endif %}\n{%- endif %}",
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-10-02T05:54:02.000Z",
  "disabled": false,
  "downloads": 2453,
  "gated": false,
  "id": "Aleph-Alpha/Kolibri-1",
  "lastModified": "2026-10-03T08:43:53.000Z",
  "library_name": "vllm",
  "likes": 523,
  "model-index": null,
  "modelId": "Aleph-Alpha/Kolibri-1",
  "pipeline_tag": "text-generation",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 705058560,
      "F8_E4M3": 77398016000
    },
    "total": 78103074560
  },
  "sha": "e52eb4627d11516b0c01de49210ab5a4e4061444",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "generation_config.json"
    },
    {
      "rfilename": "model-00001-of-00032.safetensors"
    },
    {
      "rfilename": "model-00002-of-00032.safetensors"
    },
    {
      "rfilename": "model-00003-of-00032.safetensors"
    },
    {
      "rfilename": "model-00004-of-00032.safetensors"
    },
    {
      "rfilename": "model-00005-of-00032.safetensors"
    },
    {
      "rfilename": "model-00006-of-00032.safetensors"
    },
    {
      "rfilename": "model-00007-of-00032.safetensors"
    },
    {
      "rfilename": "model-00008-of-00032.safetensors"
    },
    {
      "rfilename": "model-00009-of-00032.safetensors"
    },
    {
      "rfilename": "model-00010-of-00032.safetensors"
    },
    {
      "rfilename": "model-00011-of-00032.safetensors"
    },
    {
      "rfilename": "model-00012-of-00032.safetensors"
    },
    {
      "rfilename": "model-00013-of-00032.safetensors"
    },
    {
      "rfilename": "model-00014-of-00032.safetensors"
    },
    {
      "rfilename": "model-00015-of-00032.safetensors"
    },
    {
      "rfilename": "model-00016-of-00032.safetensors"
    },
    {
      "rfilename": "model-00017-of-00032.safetensors"
    },
    {
      "rfilename": "model-00018-of-00032.safetensors"
    },
    {
      "rfilename": "model-00019-of-00032.safetensors"
    },
    {
      "rfilename": "model-00020-of-00032.safetensors"
    },
    {
      "rfilename": "model-00021-of-00032.safetensors"
    },
    {
      "rfilename": "model-00022-of-00032.safetensors"
    },
    {
      "rfilename": "model-00023-of-00032.safetensors"
    },
    {
      "rfilename": "model-00024-of-00032.safetensors"
    },
    {
      "rfilename": "model-00025-of-00032.safetensors"
    },
    {
      "rfilename": "model-00026-of-00032.safetensors"
    },
    {
      "rfilename": "model-00027-of-00032.safetensors"
    },
    {
      "rfilename": "model-00028-of-00032.safetensors"
    },
    {
      "rfilename": "model-00029-of-00032.safetensors"
    },
    {
      "rfilename": "model-00030-of-00032.safetensors"
    },
    {
      "rfilename": "model-00031-of-00032.safetensors"
    },
    {
      "rfilename": "model-00032-of-00032.safetensors"
    },
    {
      "rfilename": "model.safetensors.index.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "tardellirs/model-pulse",
    "mrfakename/kolibri-1"
  ],
  "tags": [
    "vllm",
    "safetensors",
    "kolibri1",
    "reasoning",
    "moe",
    "text-generation",
    "conversational",
    "de",
    "en",
    "arxiv:2512.11614",
    "arxiv:2601.17858",
    "base_model:Aleph-Alpha/Kolibri-1-BF16",
    "base_model:quantized:Aleph-Alpha/Kolibri-1-BF16",
    "license:apache-2.0",
    "eval-results",
    "fp8",
    "region:us"
  ],
  "usedStorage": 78852606661
}
—
Likes
likes
not compared
GGUF quantisations0
receipt
Source
GGUF quantisations
Its words
0
Read by
field:likes
Said since
2026-10-05 18:25 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
What the source handed over
{
  "_id": "6ac3eb6ce74738c8515401e3",
  "author": "aparusel",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "de",
      "en"
    ],
    "library_name": "ds4",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "kolibri",
      "moe",
      "reasoning",
      "tool-calling",
      "dwarfstar"
    ]
  },
  "createdAt": "2026-10-05T18:24:44.000Z",
  "downloads": 0,
  "gated": false,
  "id": "aparusel/kolibri-1-gguf",
  "lastModified": "2026-10-05T18:24:47.000Z",
  "library_name": "ds4",
  "likes": 0,
  "modelId": "aparusel/kolibri-1-gguf",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "7acb9ebd3e1d75c4486c41d0a4141da1e9af5c48",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    }
  ],
  "tags": [
    "ds4",
    "gguf",
    "kolibri",
    "moe",
    "reasoning",
    "tool-calling",
    "dwarfstar",
    "text-generation",
    "de",
    "en",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "region:us"
  ]
}
—
Likes
likes
not compared
2
receipt
Source
GGUF quantisations
Its words
2
Read by
field:likes
Said since
2026-10-05 12:24 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
2026-10-05 12:24 UTC2
2026-10-04 18:16 UTC1
2026-10-04 12:16 UTC0
What the source handed over
{
  "_id": "6ac2408b2497d9f46d61cba6",
  "author": "Prompt48",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "en",
      "de"
    ],
    "library_name": "gguf",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "llama.cpp",
      "kolibri1",
      "mixture-of-experts"
    ]
  },
  "createdAt": "2026-10-04T12:03:23.000Z",
  "downloads": 333,
  "gated": false,
  "id": "Prompt48/Kolibri-1-GGUF",
  "lastModified": "2026-10-04T19:27:12.000Z",
  "library_name": "gguf",
  "likes": 2,
  "modelId": "Prompt48/Kolibri-1-GGUF",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "e4a85ea72e5eea480d01be4eb73284517e157f1c",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "Kolibri-1-Q4_K_M.gguf"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q6_K/Kolibri-1-Q6_K-00002-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00002-of-00002.gguf"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "SHA256SUMS"
    },
    {
      "rfilename": "kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/.env.example"
    },
    {
      "rfilename": "tutorial/ATTRIBUTION.md"
    },
    {
      "rfilename": "tutorial/Kolibri-RunPod-Viewer-Code.zip"
    },
    {
      "rfilename": "tutorial/LICENSE"
    },
    {
      "rfilename": "tutorial/LLAMA-CPP-LICENSE"
    },
    {
      "rfilename": "tutorial/MODEL-CARD-EXAMPLE.md"
    },
    {
      "rfilename": "tutorial/README.md"
    },
    {
      "rfilename": "tutorial/evidence/fresh/download-source.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/executed-code.py"
    },
    {
      "rfilename": "tutorial/evidence/fresh/fresh-exercises.json"
    },
    {
      "rfilename": "tutorial/evidence/fresh/generated-summary.py"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/conversion.log"
    },
    {
      "rfilename": "tutorial/evidence/original/metadata.json"
    },
    {
      "rfilename": "tutorial/evidence/original/more-quants.log"
    },
    {
      "rfilename": "tutorial/evidence/original/remote-checksums-verified.json"
    },
    {
      "rfilename": "tutorial/evidence/original/server.log"
    },
    {
      "rfilename": "tutorial/evidence/original/smoke-tests.json"
    },
    {
      "rfilename": "tutorial/evidence/original/source.json"
    },
    {
      "rfilename": "tutorial/historical/add-build-instructions.py"
    },
    {
      "rfilename": "tutorial/historical/convert-on-pod.sh"
    },
    {
      "rfilename": "tutorial/historical/more-quants.sh"
    },
    {
      "rfilename": "tutorial/historical/publish-more-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/publish-on-pod.py"
    },
    {
      "rfilename": "tutorial/historical/resume-conversion.sh"
    },
    {
      "rfilename": "tutorial/historical/validate-on-pod.py"
    },
    {
      "rfilename": "tutorial/kolibri1-llama.cpp.patch"
    },
    {
      "rfilename": "tutorial/requirements.txt"
    },
    {
      "rfilename": "tutorial/runpod.py"
    },
    {
      "rfilename": "tutorial/steps/00-setup.sh"
    },
    {
      "rfilename": "tutorial/steps/10-download-source.py"
    },
    {
      "rfilename": "tutorial/steps/100-exercises.py"
    },
    {
      "rfilename": "tutorial/steps/20-convert-bf16.sh"
    },
    {
      "rfilename": "tutorial/steps/30-quantize.sh"
    },
    {
      "rfilename": "tutorial/steps/40-inspect.py"
    },
    {
      "rfilename": "tutorial/steps/50-serve.sh"
    },
    {
      "rfilename": "tutorial/steps/60-smoke-test.py"
    },
    {
      "rfilename": "tutorial/steps/70-split.py"
    },
    {
      "rfilename": "tutorial/steps/80-publish.py"
    },
    {
      "rfilename": "tutorial/steps/90-download-q4.py"
    },
    {
      "rfilename": "tutorial/video-production/README.md"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-comprehensive-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/bootstrap-render.sh"
    },
    {
      "rfilename": "tutorial/video-production/build_full_video.py"
    },
    {
      "rfilename": "tutorial/video-production/build_video.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_exercises.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_full_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/capture_sources.py"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-SCRIPT.md"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-chapters.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-scenes.json"
    },
    {
      "rfilename": "tutorial/video-production/comprehensive-timing.json"
    },
    {
      "rfilename": "tutorial/video-production/finish_script.py"
    },
    {
      "rfilename": "tutorial/video-production/full_content.py"
    },
    {
      "rfilename": "tutorial/video-production/package_video.py"
    },
    {
      "rfilename": "tutorial/video-production/qa_video.py"
    },
    {
      "rfilename": "tutorial/video-production/render_cloud.py"
    },
    {
      "rfilename": "tutorial/video-production/tts_client.py"
    },
    {
      "rfilename": "tutorial/video-production/write_full_script.py"
    },
    {
      "rfilename": "tutorial/video-production/write_script.py"
    },
    {
      "rfilename": "validation/Q6_K/metadata.json"
    },
    {
      "rfilename": "validation/Q6_K/smoke-tests.json"
    },
    {
      "rfilename": "validation/Q8_0/metadata.json"
    },
    {
      "rfilename": "validation/Q8_0/smoke-tests.json"
    },
    {
      "rfilename": "validation/file-sizes.txt"
    },
    {
      "rfilename": "validation/llama-revision.txt"
    },
    {
      "rfilename": "validation/metadata.json"
    },
    {
      "rfilename": "validation/smoke-tests.json"
    },
    {
      "rfilename": "validation/source.json"
    }
  ],
  "tags": [
    "gguf",
    "llama.cpp",
    "kolibri1",
    "mixture-of-experts",
    "text-generation",
    "en",
    "de",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "endpoints_compatible",
    "region:us",
    "conversational"
  ]
}
—
Likes
likes
not compared
5
receipt
Source
GGUF quantisations
Its words
5
Read by
field:likes
Said since
2026-10-05 12:24 UTC
Last answered
2026-10-05 18:25 UTC
Original
open at the source
2026-10-05 12:24 UTC5
2026-10-05 06:21 UTC3
2026-10-04 12:16 UTC2
2026-10-04 06:13 UTC0
What the source handed over
{
  "_id": "6ac181821b92ae3a1fdd3fbb",
  "author": "Hob-forge",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1",
    "base_model_relation": "quantized",
    "language": [
      "de",
      "en"
    ],
    "library_name": "gguf",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "gguf",
      "llama.cpp",
      "kolibri1",
      "mixture-of-experts",
      "reasoning",
      "tool-calling"
    ]
  },
  "createdAt": "2026-10-03T22:28:18.000Z",
  "downloads": 3268,
  "gated": false,
  "id": "Hob-forge/Kolibri-1-GGUF",
  "lastModified": "2026-10-04T05:12:11.000Z",
  "library_name": "gguf",
  "likes": 5,
  "modelId": "Hob-forge/Kolibri-1-GGUF",
  "pipeline_tag": "text-generation",
  "private": false,
  "sha": "b08405e141706457126fa75b9f2f9e0729b30d0a",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "Kolibri-1-Q4_K_M.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00001-of-00002.gguf"
    },
    {
      "rfilename": "Q8_0/Kolibri-1-Q8_0-00002-of-00002.gguf"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "kolibri1-llama.cpp.patch"
    }
  ],
  "tags": [
    "gguf",
    "llama.cpp",
    "kolibri1",
    "mixture-of-experts",
    "reasoning",
    "tool-calling",
    "text-generation",
    "de",
    "en",
    "base_model:Aleph-Alpha/Kolibri-1",
    "base_model:quantized:Aleph-Alpha/Kolibri-1",
    "license:apache-2.0",
    "endpoints_compatible",
    "region:us",
    "conversational"
  ]
}
—
Likes
likes
not compared
Hugging Face models523
receipt
Source
Hugging Face models
Its words
523
Read by
field:likes
Said since
2026-10-05 12:24 UTC
Original
open at the source
2026-10-05 12:24 UTC523
2026-10-05 06:21 UTC434
2026-10-05 00:20 UTC389
2026-10-04 18:16 UTC359
2026-10-04 12:16 UTC319
2026-10-04 06:13 UTC270
What the source handed over
{
  "_asked": "Aleph-Alpha/Kolibri-1",
  "_id": "6abf46fa71befb841d4f4fff",
  "author": "Aleph-Alpha",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1-BF16",
    "language": [
      "de",
      "en"
    ],
    "library_name": "vllm",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "reasoning",
      "moe"
    ]
  },
  "config": {
    "architectures": [
      "Kolibri1ForCausalLM"
    ],
    "model_type": "kolibri1",
    "num_experts": 384,
    "num_experts_per_tok": 6,
    "quantization_config": {
      "modules_to_not_convert": [
        "model.layers.0.mlp.gate",
        "model.layers.1.mlp.gate",
        "model.layers.2.mlp.gate",
        "model.layers.3.mlp.gate",
        "model.layers.4.mlp.gate",
        "model.layers.5.mlp.gate",
        "model.layers.6.mlp.gate",
        "model.layers.7.mlp.gate",
        "model.layers.8.mlp.gate",
        "model.layers.9.mlp.gate",
        "model.layers.10.mlp.gate",
        "model.layers.11.mlp.gate",
        "model.layers.12.mlp.gate",
        "model.layers.13.mlp.gate",
        "model.layers.14.mlp.gate",
        "model.layers.15.mlp.gate",
        "model.layers.16.mlp.gate",
        "model.layers.17.mlp.gate",
        "model.layers.18.mlp.gate",
        "model.layers.19.mlp.gate",
        "model.layers.20.mlp.gate",
        "model.layers.21.mlp.gate",
        "model.layers.22.mlp.gate",
        "model.layers.23.mlp.gate",
        "model.layers.24.mlp.gate",
        "model.layers.25.mlp.gate",
        "model.layers.26.mlp.gate",
        "model.layers.27.mlp.gate",
        "model.layers.28.mlp.gate",
        "model.layers.29.mlp.gate",
        "model.layers.30.mlp.gate",
        "model.layers.31.mlp.gate",
        "model.layers.32.mlp.gate",
        "model.layers.33.mlp.gate",
        "model.layers.34.mlp.gate",
        "model.layers.35.mlp.gate",
        "model.layers.36.mlp.gate",
        "model.layers.37.mlp.gate",
        "model.layers.38.mlp.gate",
        "model.layers.39.mlp.gate",
        "model.layers.40.mlp.gate",
        "model.layers.41.mlp.gate",
        "model.layers.42.mlp.gate",
        "model.layers.43.mlp.gate",
        "model.layers.44.mlp.gate",
        "model.layers.45.mlp.gate",
        "model.layers.46.mlp.gate",
        "model.layers.47.mlp.gate",
        "model.layers.48.mlp.gate",
        "model.layers.49.mlp.gate"
      ],
      "quant_method": "fp8"
    },
    "tokenizer_config": {
      "bos_token": null,
      "chat_template": "{%- set _sv_reasoning_effort = reasoning_effort | default(none) -%}\n{%- set _sv_thinking_disabled = false -%}\n{%- if _sv_reasoning_effort is not none -%}\n    {%- set _sv_thinking_disabled = _sv_reasoning_effort == 'none' -%}\n{%- elif enable_thinking is defined and enable_thinking is false -%}\n    {%- set _sv_thinking_disabled = true -%}\n{%- endif -%}\n{%- set _sv_no_reasoning_sentence = \"Reasoning is disabled. Proceed straight to answering according to the user's instructions.\" -%}\n{%- set _sv_low_reasoning_sentence = \"Reasoning effort is set to low. Think briefly through only the essential steps in the user's language, then proceed directly to the answer.\" -%}\n{%- set _sv_medium_reasoning_sentence = \"Reasoning effort is set to medium. Think through the task methodically in the user's language, check key assumptions, and provide a well-supported answer.\" -%}\n{%- set _sv_high_reasoning_sentence = \"Reasoning effort is set to high. Think carefully through the task in the user's language, validate key assumptions, consider plausible alternatives, and prioritize correctness and clarity.\" -%}\n{%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- if _sv_thinking_disabled -%}\n    {%- set _sv_reasoning_sentence = _sv_no_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['minimal', 'low'] -%}\n    {%- set _sv_reasoning_sentence = _sv_low_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort == 'medium' -%}\n    {%- set _sv_reasoning_sentence = _sv_medium_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['high', 'xhigh', 'max'] or _sv_reasoning_effort is none -%}\n    {%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- endif -%}\n{%- set _sv_has_system = messages | length > 0 and messages[0].role == 'system' -%}\n{%- if tools or _sv_reasoning_sentence is not none or _sv_has_system %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if _sv_has_system %}\n        {{- messages[0].content }}\n    {%- endif %}\n    {%- if _sv_reasoning_sentence is not none %}\n        {%- if _sv_has_system %}{{- '\\n\\n' }}{%- endif %}\n        {{- '# Reasoning effort\\n\\n' + _sv_reasoning_sentence }}\n    {%- endif %}\n    {%- if tools %}\n        {%- if _sv_has_system or _sv_reasoning_sentence is not none %}{{- '\\n\\n' }}{%- endif %}\n        {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n        {%- for tool in tools %}\n            {{- \"\\n\" }}\n            {{- tool | tojson }}\n        {%- endfor %}\n        {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call>\" }}\n    {%- endif %}\n    {{- '<|im_end|>\\n' }}\n{%- endif %}\n{%- set ns = namespace(multi_step_tool=true, last_query_index=messages|length - 1) %}\n{%- for message in messages[::-1] %}\n    {%- set index = (messages|length - 1) - loop.index0 %}\n    {%- if ns.multi_step_tool and message.role == \"user\" and message.content is string and not(message.content.startswith('<tool_response>') and message.content.endswith('</tool_response>')) %}\n        {%- set ns.multi_step_tool = false %}\n        {%- set ns.last_query_index = index %}\n    {%- endif %}\n{%- endfor %}\n{%- for message in messages %}\n    {%- if (message.role == \"user\") or (message.role == \"system\" and not loop.first) %}\n        {{- '<|im_start|>' + message.role + '\\n' + message.content + '<|im_end|>' + '\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {%- set content = message.content if message.content is string else '' %}\n        {%- set reasoning = none %}\n        {%- if message.reasoning is defined and message.reasoning is not none %}\n            {%- set reasoning = message.reasoning %}\n        {%- elif message.reasoning_content is defined and message.reasoning_content is not none %}\n            {#- Deprecated vLLM compatibility. Prefer the `reasoning` field. -#}\n            {%- set reasoning = message.reasoning_content %}\n        {%- elif content is string and '</think>' in content %}\n            {%- set reasoning = content.split('</think>')[0].rstrip('\\n').split('<think>')[-1].lstrip('\\n') %}\n            {%- set content = content.split('</think>')[-1].lstrip('\\n') %}\n        {%- endif %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if (loop.index0 > ns.last_query_index) or (preserve_thinking is defined and preserve_thinking is true) %}\n            {{- '<think>\\n' }}\n            {%- if reasoning is not none and reasoning | trim %}\n                {{- reasoning.strip('\\n') }}\n            {%- endif %}\n            {{- '\\n</think>\\n\\n' }}\n        {%- endif %}\n        {{- content.lstrip('\\n') }}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {{- message.content }}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n    {%- if _sv_thinking_disabled %}\n        {{- '<think>\\n\\n</think>\\n\\n' }}\n    {%- endif %}\n{%- endif %}",
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-10-02T05:54:02.000Z",
  "disabled": false,
  "downloads": 2453,
  "gated": false,
  "id": "Aleph-Alpha/Kolibri-1",
  "lastModified": "2026-10-03T08:43:53.000Z",
  "library_name": "vllm",
  "likes": 523,
  "model-index": null,
  "modelId": "Aleph-Alpha/Kolibri-1",
  "pipeline_tag": "text-generation",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 705058560,
      "F8_E4M3": 77398016000
    },
    "total": 78103074560
  },
  "sha": "e52eb4627d11516b0c01de49210ab5a4e4061444",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "generation_config.json"
    },
    {
      "rfilename": "model-00001-of-00032.safetensors"
    },
    {
      "rfilename": "model-00002-of-00032.safetensors"
    },
    {
      "rfilename": "model-00003-of-00032.safetensors"
    },
    {
      "rfilename": "model-00004-of-00032.safetensors"
    },
    {
      "rfilename": "model-00005-of-00032.safetensors"
    },
    {
      "rfilename": "model-00006-of-00032.safetensors"
    },
    {
      "rfilename": "model-00007-of-00032.safetensors"
    },
    {
      "rfilename": "model-00008-of-00032.safetensors"
    },
    {
      "rfilename": "model-00009-of-00032.safetensors"
    },
    {
      "rfilename": "model-00010-of-00032.safetensors"
    },
    {
      "rfilename": "model-00011-of-00032.safetensors"
    },
    {
      "rfilename": "model-00012-of-00032.safetensors"
    },
    {
      "rfilename": "model-00013-of-00032.safetensors"
    },
    {
      "rfilename": "model-00014-of-00032.safetensors"
    },
    {
      "rfilename": "model-00015-of-00032.safetensors"
    },
    {
      "rfilename": "model-00016-of-00032.safetensors"
    },
    {
      "rfilename": "model-00017-of-00032.safetensors"
    },
    {
      "rfilename": "model-00018-of-00032.safetensors"
    },
    {
      "rfilename": "model-00019-of-00032.safetensors"
    },
    {
      "rfilename": "model-00020-of-00032.safetensors"
    },
    {
      "rfilename": "model-00021-of-00032.safetensors"
    },
    {
      "rfilename": "model-00022-of-00032.safetensors"
    },
    {
      "rfilename": "model-00023-of-00032.safetensors"
    },
    {
      "rfilename": "model-00024-of-00032.safetensors"
    },
    {
      "rfilename": "model-00025-of-00032.safetensors"
    },
    {
      "rfilename": "model-00026-of-00032.safetensors"
    },
    {
      "rfilename": "model-00027-of-00032.safetensors"
    },
    {
      "rfilename": "model-00028-of-00032.safetensors"
    },
    {
      "rfilename": "model-00029-of-00032.safetensors"
    },
    {
      "rfilename": "model-00030-of-00032.safetensors"
    },
    {
      "rfilename": "model-00031-of-00032.safetensors"
    },
    {
      "rfilename": "model-00032-of-00032.safetensors"
    },
    {
      "rfilename": "model.safetensors.index.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "tardellirs/model-pulse",
    "mrfakename/kolibri-1"
  ],
  "tags": [
    "vllm",
    "safetensors",
    "kolibri1",
    "reasoning",
    "moe",
    "text-generation",
    "conversational",
    "de",
    "en",
    "arxiv:2512.11614",
    "arxiv:2601.17858",
    "base_model:Aleph-Alpha/Kolibri-1-BF16",
    "base_model:quantized:Aleph-Alpha/Kolibri-1-BF16",
    "license:apache-2.0",
    "eval-results",
    "fp8",
    "region:us"
  ],
  "usedStorage": 78852606661
}
—
Task
task
Hugging Face modelstext-generation
receipt
Source
Hugging Face models
Its words
text-generation
Read by
field:pipeline_tag
Said since
2026-10-04 06:13 UTC
Original
open at the source
What the source handed over
{
  "_asked": "Aleph-Alpha/Kolibri-1",
  "_id": "6abf46fa71befb841d4f4fff",
  "author": "Aleph-Alpha",
  "cardData": {
    "base_model": "Aleph-Alpha/Kolibri-1-BF16",
    "language": [
      "de",
      "en"
    ],
    "library_name": "vllm",
    "license": "apache-2.0",
    "pipeline_tag": "text-generation",
    "tags": [
      "reasoning",
      "moe"
    ]
  },
  "config": {
    "architectures": [
      "Kolibri1ForCausalLM"
    ],
    "model_type": "kolibri1",
    "num_experts": 384,
    "num_experts_per_tok": 6,
    "quantization_config": {
      "modules_to_not_convert": [
        "model.layers.0.mlp.gate",
        "model.layers.1.mlp.gate",
        "model.layers.2.mlp.gate",
        "model.layers.3.mlp.gate",
        "model.layers.4.mlp.gate",
        "model.layers.5.mlp.gate",
        "model.layers.6.mlp.gate",
        "model.layers.7.mlp.gate",
        "model.layers.8.mlp.gate",
        "model.layers.9.mlp.gate",
        "model.layers.10.mlp.gate",
        "model.layers.11.mlp.gate",
        "model.layers.12.mlp.gate",
        "model.layers.13.mlp.gate",
        "model.layers.14.mlp.gate",
        "model.layers.15.mlp.gate",
        "model.layers.16.mlp.gate",
        "model.layers.17.mlp.gate",
        "model.layers.18.mlp.gate",
        "model.layers.19.mlp.gate",
        "model.layers.20.mlp.gate",
        "model.layers.21.mlp.gate",
        "model.layers.22.mlp.gate",
        "model.layers.23.mlp.gate",
        "model.layers.24.mlp.gate",
        "model.layers.25.mlp.gate",
        "model.layers.26.mlp.gate",
        "model.layers.27.mlp.gate",
        "model.layers.28.mlp.gate",
        "model.layers.29.mlp.gate",
        "model.layers.30.mlp.gate",
        "model.layers.31.mlp.gate",
        "model.layers.32.mlp.gate",
        "model.layers.33.mlp.gate",
        "model.layers.34.mlp.gate",
        "model.layers.35.mlp.gate",
        "model.layers.36.mlp.gate",
        "model.layers.37.mlp.gate",
        "model.layers.38.mlp.gate",
        "model.layers.39.mlp.gate",
        "model.layers.40.mlp.gate",
        "model.layers.41.mlp.gate",
        "model.layers.42.mlp.gate",
        "model.layers.43.mlp.gate",
        "model.layers.44.mlp.gate",
        "model.layers.45.mlp.gate",
        "model.layers.46.mlp.gate",
        "model.layers.47.mlp.gate",
        "model.layers.48.mlp.gate",
        "model.layers.49.mlp.gate"
      ],
      "quant_method": "fp8"
    },
    "tokenizer_config": {
      "bos_token": null,
      "chat_template": "{%- set _sv_reasoning_effort = reasoning_effort | default(none) -%}\n{%- set _sv_thinking_disabled = false -%}\n{%- if _sv_reasoning_effort is not none -%}\n    {%- set _sv_thinking_disabled = _sv_reasoning_effort == 'none' -%}\n{%- elif enable_thinking is defined and enable_thinking is false -%}\n    {%- set _sv_thinking_disabled = true -%}\n{%- endif -%}\n{%- set _sv_no_reasoning_sentence = \"Reasoning is disabled. Proceed straight to answering according to the user's instructions.\" -%}\n{%- set _sv_low_reasoning_sentence = \"Reasoning effort is set to low. Think briefly through only the essential steps in the user's language, then proceed directly to the answer.\" -%}\n{%- set _sv_medium_reasoning_sentence = \"Reasoning effort is set to medium. Think through the task methodically in the user's language, check key assumptions, and provide a well-supported answer.\" -%}\n{%- set _sv_high_reasoning_sentence = \"Reasoning effort is set to high. Think carefully through the task in the user's language, validate key assumptions, consider plausible alternatives, and prioritize correctness and clarity.\" -%}\n{%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- if _sv_thinking_disabled -%}\n    {%- set _sv_reasoning_sentence = _sv_no_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['minimal', 'low'] -%}\n    {%- set _sv_reasoning_sentence = _sv_low_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort == 'medium' -%}\n    {%- set _sv_reasoning_sentence = _sv_medium_reasoning_sentence -%}\n{%- elif _sv_reasoning_effort in ['high', 'xhigh', 'max'] or _sv_reasoning_effort is none -%}\n    {%- set _sv_reasoning_sentence = _sv_high_reasoning_sentence -%}\n{%- endif -%}\n{%- set _sv_has_system = messages | length > 0 and messages[0].role == 'system' -%}\n{%- if tools or _sv_reasoning_sentence is not none or _sv_has_system %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if _sv_has_system %}\n        {{- messages[0].content }}\n    {%- endif %}\n    {%- if _sv_reasoning_sentence is not none %}\n        {%- if _sv_has_system %}{{- '\\n\\n' }}{%- endif %}\n        {{- '# Reasoning effort\\n\\n' + _sv_reasoning_sentence }}\n    {%- endif %}\n    {%- if tools %}\n        {%- if _sv_has_system or _sv_reasoning_sentence is not none %}{{- '\\n\\n' }}{%- endif %}\n        {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n        {%- for tool in tools %}\n            {{- \"\\n\" }}\n            {{- tool | tojson }}\n        {%- endfor %}\n        {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call>\" }}\n    {%- endif %}\n    {{- '<|im_end|>\\n' }}\n{%- endif %}\n{%- set ns = namespace(multi_step_tool=true, last_query_index=messages|length - 1) %}\n{%- for message in messages[::-1] %}\n    {%- set index = (messages|length - 1) - loop.index0 %}\n    {%- if ns.multi_step_tool and message.role == \"user\" and message.content is string and not(message.content.startswith('<tool_response>') and message.content.endswith('</tool_response>')) %}\n        {%- set ns.multi_step_tool = false %}\n        {%- set ns.last_query_index = index %}\n    {%- endif %}\n{%- endfor %}\n{%- for message in messages %}\n    {%- if (message.role == \"user\") or (message.role == \"system\" and not loop.first) %}\n        {{- '<|im_start|>' + message.role + '\\n' + message.content + '<|im_end|>' + '\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {%- set content = message.content if message.content is string else '' %}\n        {%- set reasoning = none %}\n        {%- if message.reasoning is defined and message.reasoning is not none %}\n            {%- set reasoning = message.reasoning %}\n        {%- elif message.reasoning_content is defined and message.reasoning_content is not none %}\n            {#- Deprecated vLLM compatibility. Prefer the `reasoning` field. -#}\n            {%- set reasoning = message.reasoning_content %}\n        {%- elif content is string and '</think>' in content %}\n            {%- set reasoning = content.split('</think>')[0].rstrip('\\n').split('<think>')[-1].lstrip('\\n') %}\n            {%- set content = content.split('</think>')[-1].lstrip('\\n') %}\n        {%- endif %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if (loop.index0 > ns.last_query_index) or (preserve_thinking is defined and preserve_thinking is true) %}\n            {{- '<think>\\n' }}\n            {%- if reasoning is not none and reasoning | trim %}\n                {{- reasoning.strip('\\n') }}\n            {%- endif %}\n            {{- '\\n</think>\\n\\n' }}\n        {%- endif %}\n        {{- content.lstrip('\\n') }}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {{- message.content }}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n    {%- if _sv_thinking_disabled %}\n        {{- '<think>\\n\\n</think>\\n\\n' }}\n    {%- endif %}\n{%- endif %}",
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-10-02T05:54:02.000Z",
  "disabled": false,
  "downloads": 2453,
  "gated": false,
  "id": "Aleph-Alpha/Kolibri-1",
  "lastModified": "2026-10-03T08:43:53.000Z",
  "library_name": "vllm",
  "likes": 523,
  "model-index": null,
  "modelId": "Aleph-Alpha/Kolibri-1",
  "pipeline_tag": "text-generation",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 705058560,
      "F8_E4M3": 77398016000
    },
    "total": 78103074560
  },
  "sha": "e52eb4627d11516b0c01de49210ab5a4e4061444",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "LICENSE"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "generation_config.json"
    },
    {
      "rfilename": "model-00001-of-00032.safetensors"
    },
    {
      "rfilename": "model-00002-of-00032.safetensors"
    },
    {
      "rfilename": "model-00003-of-00032.safetensors"
    },
    {
      "rfilename": "model-00004-of-00032.safetensors"
    },
    {
      "rfilename": "model-00005-of-00032.safetensors"
    },
    {
      "rfilename": "model-00006-of-00032.safetensors"
    },
    {
      "rfilename": "model-00007-of-00032.safetensors"
    },
    {
      "rfilename": "model-00008-of-00032.safetensors"
    },
    {
      "rfilename": "model-00009-of-00032.safetensors"
    },
    {
      "rfilename": "model-00010-of-00032.safetensors"
    },
    {
      "rfilename": "model-00011-of-00032.safetensors"
    },
    {
      "rfilename": "model-00012-of-00032.safetensors"
    },
    {
      "rfilename": "model-00013-of-00032.safetensors"
    },
    {
      "rfilename": "model-00014-of-00032.safetensors"
    },
    {
      "rfilename": "model-00015-of-00032.safetensors"
    },
    {
      "rfilename": "model-00016-of-00032.safetensors"
    },
    {
      "rfilename": "model-00017-of-00032.safetensors"
    },
    {
      "rfilename": "model-00018-of-00032.safetensors"
    },
    {
      "rfilename": "model-00019-of-00032.safetensors"
    },
    {
      "rfilename": "model-00020-of-00032.safetensors"
    },
    {
      "rfilename": "model-00021-of-00032.safetensors"
    },
    {
      "rfilename": "model-00022-of-00032.safetensors"
    },
    {
      "rfilename": "model-00023-of-00032.safetensors"
    },
    {
      "rfilename": "model-00024-of-00032.safetensors"
    },
    {
      "rfilename": "model-00025-of-00032.safetensors"
    },
    {
      "rfilename": "model-00026-of-00032.safetensors"
    },
    {
      "rfilename": "model-00027-of-00032.safetensors"
    },
    {
      "rfilename": "model-00028-of-00032.safetensors"
    },
    {
      "rfilename": "model-00029-of-00032.safetensors"
    },
    {
      "rfilename": "model-00030-of-00032.safetensors"
    },
    {
      "rfilename": "model-00031-of-00032.safetensors"
    },
    {
      "rfilename": "model-00032-of-00032.safetensors"
    },
    {
      "rfilename": "model.safetensors.index.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "tardellirs/model-pulse",
    "mrfakename/kolibri-1"
  ],
  "tags": [
    "vllm",
    "safetensors",
    "kolibri1",
    "reasoning",
    "moe",
    "text-generation",
    "conversational",
    "de",
    "en",
    "arxiv:2512.11614",
    "arxiv:2601.17858",
    "base_model:Aleph-Alpha/Kolibri-1-BF16",
    "base_model:quantized:Aleph-Alpha/Kolibri-1-BF16",
    "license:apache-2.0",
    "eval-results",
    "fp8",
    "region:us"
  ],
  "usedStorage": 78852606661
}
—

model

Aleph-Alpha/Kolibri-1
zetlyn/models-hf · 2026-10-02
author Aleph-Alpha downloads 2453 gated false licence apache-2.0 likes 523 task text-generation source

quantisation

Prompt48/Kolibri-1-GGUF
zetlyn/models-gguf · 2026-10-04
author Prompt48 base Aleph-Alpha/Kolibri-1 downloads 333 licence apache-2.0 likes 2 source
Hob-forge/Kolibri-1-GGUF
zetlyn/models-gguf · 2026-10-03
author Hob-forge base Aleph-Alpha/Kolibri-1 downloads 3268 licence apache-2.0 likes 5 source
aparusel/kolibri-1-gguf
zetlyn/models-gguf · 2026-10-05
author aparusel base Aleph-Alpha/Kolibri-1 downloads 0 licence apache-2.0 likes 0 source
webmp3/Sakura-MicroQuality-Kolibri-1-365E-GGUF
zetlyn/models-gguf · 2026-10-05
author webmp3 base Aleph-Alpha/Kolibri-1 downloads 0 licence apache-2.0 likes 0 source
webmp3/Sakura-MicroQuality-Kolibri-1-GGUF
zetlyn/models-gguf · 2026-10-05
author webmp3 base Aleph-Alpha/Kolibri-1 downloads 0 licence apache-2.0 likes 0 source