DeepCybo/PhysBrain1.5-2B

zetlyn/models-hf model hf DeepCybo/PhysBrain1.5-2B known 2026-09-08

https://huggingface.co/DeepCybo/PhysBrain1.5-2B

Properties

authorDeepCybo
receipt
Source
Hugging Face models
Its words
DeepCybo
Read by
field:author
Said since
2026-10-02 12:02 UTC
Last answered
2026-10-04 18:21 UTC
Original
open at the source
What the source handed over
{
  "_asked": "DeepCybo/PhysBrain1.5-2B",
  "_id": "6aa04a07febd14f899e01a42",
  "author": "DeepCybo",
  "cardData": {
    "base_model": [
      "Qwen/Qwen3-VL-2B-Instruct"
    ],
    "language": [
      "en",
      "zh"
    ],
    "library_name": "transformers",
    "pipeline_tag": "any-to-any",
    "tags": [
      "embodied-ai",
      "vision-language-model",
      "robotics",
      "spatial-reasoning",
      "multimodal"
    ]
  },
  "config": {
    "architectures": [
      "Qwen3VLForConditionalGeneration"
    ],
    "chat_template_jinja": "{%- if tools %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if messages[0].role == 'system' %}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n\\n' }}\n    {%- endif %}\n    {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n    {%- for tool in tools %}\n        {{- \"\\n\" }}\n        {{- tool | tojson }}\n    {%- endfor %}\n    {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call><|im_end|>\\n\" }}\n{%- else %}\n    {%- if messages[0].role == 'system' %}\n        {{- '<|im_start|>system\\n' }}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- endif %}\n{%- endif %}\n{%- set image_count = namespace(value=0) %}\n{%- set video_count = namespace(value=0) %}\n{%- for message in messages %}\n    {%- if message.role == \"user\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content_item in message.content %}\n                {%- if 'text' in content_item %}\n                    {{- content_item.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and message.content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n{%- endif %}\n",
    "model_type": "qwen3_vl",
    "tokenizer_config": {
      "bos_token": null,
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-09-08T17:46:47.000Z",
  "disabled": false,
  "downloads": 541,
  "gated": false,
  "id": "DeepCybo/PhysBrain1.5-2B",
  "lastModified": "2026-09-16T03:44:04.000Z",
  "library_name": "transformers",
  "likes": 13,
  "model-index": null,
  "modelId": "DeepCybo/PhysBrain1.5-2B",
  "pipeline_tag": "any-to-any",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 2161610752
    },
    "total": 2161610752
  },
  "sha": "21af7b27ec9ce8f3ade336f682d1450455ccec52",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "assets/.DS_Store"
    },
    {
      "rfilename": "assets/case-action-trajectory.png"
    },
    {
      "rfilename": "assets/case-future-prediction-v2.png"
    },
    {
      "rfilename": "assets/case-future-prediction.png"
    },
    {
      "rfilename": "assets/embodied_benchmarks.png"
    },
    {
      "rfilename": "assets/evaluation-cases-v6.png"
    },
    {
      "rfilename": "assets/leaderboard.svg"
    },
    {
      "rfilename": "assets/logo.png"
    },
    {
      "rfilename": "assets/model-architecture-v2.png"
    },
    {
      "rfilename": "chat_template.jinja"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "model.safetensors"
    },
    {
      "rfilename": "preprocessor_config.json"
    },
    {
      "rfilename": "processor_config.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "hugging-apps/physbrain1-5-2b-demo"
  ],
  "tags": [
    "transformers",
    "safetensors",
    "qwen3_vl",
    "image-text-to-text",
    "embodied-ai",
    "vision-language-model",
    "robotics",
    "spatial-reasoning",
    "multimodal",
    "any-to-any",
    "en",
    "zh",
    "arxiv:2609.14973",
    "base_model:Qwen/Qwen3-VL-2B-Instruct",
    "base_model:finetune:Qwen/Qwen3-VL-2B-Instruct",
    "endpoints_compatible",
    "region:us"
  ],
  "transformersInfo": {
    "auto_model": "AutoModelForMultimodalLM",
    "pipeline_tag": "image-text-to-text",
    "processor": "AutoProcessor"
  },
  "usedStorage": 4347865690
}
downloads541
receipt
Source
Hugging Face models
Its words
541
Read by
field:downloads
Said since
2026-10-04 12:16 UTC
Last answered
2026-10-04 18:21 UTC
Original
open at the source
2026-10-04 12:16 UTC541
2026-10-03 18:10 UTC513
2026-10-02 12:02 UTC496
What the source handed over
{
  "_asked": "DeepCybo/PhysBrain1.5-2B",
  "_id": "6aa04a07febd14f899e01a42",
  "author": "DeepCybo",
  "cardData": {
    "base_model": [
      "Qwen/Qwen3-VL-2B-Instruct"
    ],
    "language": [
      "en",
      "zh"
    ],
    "library_name": "transformers",
    "pipeline_tag": "any-to-any",
    "tags": [
      "embodied-ai",
      "vision-language-model",
      "robotics",
      "spatial-reasoning",
      "multimodal"
    ]
  },
  "config": {
    "architectures": [
      "Qwen3VLForConditionalGeneration"
    ],
    "chat_template_jinja": "{%- if tools %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if messages[0].role == 'system' %}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n\\n' }}\n    {%- endif %}\n    {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n    {%- for tool in tools %}\n        {{- \"\\n\" }}\n        {{- tool | tojson }}\n    {%- endfor %}\n    {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call><|im_end|>\\n\" }}\n{%- else %}\n    {%- if messages[0].role == 'system' %}\n        {{- '<|im_start|>system\\n' }}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- endif %}\n{%- endif %}\n{%- set image_count = namespace(value=0) %}\n{%- set video_count = namespace(value=0) %}\n{%- for message in messages %}\n    {%- if message.role == \"user\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content_item in message.content %}\n                {%- if 'text' in content_item %}\n                    {{- content_item.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and message.content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n{%- endif %}\n",
    "model_type": "qwen3_vl",
    "tokenizer_config": {
      "bos_token": null,
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-09-08T17:46:47.000Z",
  "disabled": false,
  "downloads": 541,
  "gated": false,
  "id": "DeepCybo/PhysBrain1.5-2B",
  "lastModified": "2026-09-16T03:44:04.000Z",
  "library_name": "transformers",
  "likes": 13,
  "model-index": null,
  "modelId": "DeepCybo/PhysBrain1.5-2B",
  "pipeline_tag": "any-to-any",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 2161610752
    },
    "total": 2161610752
  },
  "sha": "21af7b27ec9ce8f3ade336f682d1450455ccec52",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "assets/.DS_Store"
    },
    {
      "rfilename": "assets/case-action-trajectory.png"
    },
    {
      "rfilename": "assets/case-future-prediction-v2.png"
    },
    {
      "rfilename": "assets/case-future-prediction.png"
    },
    {
      "rfilename": "assets/embodied_benchmarks.png"
    },
    {
      "rfilename": "assets/evaluation-cases-v6.png"
    },
    {
      "rfilename": "assets/leaderboard.svg"
    },
    {
      "rfilename": "assets/logo.png"
    },
    {
      "rfilename": "assets/model-architecture-v2.png"
    },
    {
      "rfilename": "chat_template.jinja"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "model.safetensors"
    },
    {
      "rfilename": "preprocessor_config.json"
    },
    {
      "rfilename": "processor_config.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "hugging-apps/physbrain1-5-2b-demo"
  ],
  "tags": [
    "transformers",
    "safetensors",
    "qwen3_vl",
    "image-text-to-text",
    "embodied-ai",
    "vision-language-model",
    "robotics",
    "spatial-reasoning",
    "multimodal",
    "any-to-any",
    "en",
    "zh",
    "arxiv:2609.14973",
    "base_model:Qwen/Qwen3-VL-2B-Instruct",
    "base_model:finetune:Qwen/Qwen3-VL-2B-Instruct",
    "endpoints_compatible",
    "region:us"
  ],
  "transformersInfo": {
    "auto_model": "AutoModelForMultimodalLM",
    "pipeline_tag": "image-text-to-text",
    "processor": "AutoProcessor"
  },
  "usedStorage": 4347865690
}
gatedfalse
receipt
Source
Hugging Face models
Its words
false
Read by
field:gated
Said since
2026-10-02 12:02 UTC
Last answered
2026-10-04 18:21 UTC
Original
open at the source
What the source handed over
{
  "_asked": "DeepCybo/PhysBrain1.5-2B",
  "_id": "6aa04a07febd14f899e01a42",
  "author": "DeepCybo",
  "cardData": {
    "base_model": [
      "Qwen/Qwen3-VL-2B-Instruct"
    ],
    "language": [
      "en",
      "zh"
    ],
    "library_name": "transformers",
    "pipeline_tag": "any-to-any",
    "tags": [
      "embodied-ai",
      "vision-language-model",
      "robotics",
      "spatial-reasoning",
      "multimodal"
    ]
  },
  "config": {
    "architectures": [
      "Qwen3VLForConditionalGeneration"
    ],
    "chat_template_jinja": "{%- if tools %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if messages[0].role == 'system' %}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n\\n' }}\n    {%- endif %}\n    {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n    {%- for tool in tools %}\n        {{- \"\\n\" }}\n        {{- tool | tojson }}\n    {%- endfor %}\n    {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call><|im_end|>\\n\" }}\n{%- else %}\n    {%- if messages[0].role == 'system' %}\n        {{- '<|im_start|>system\\n' }}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- endif %}\n{%- endif %}\n{%- set image_count = namespace(value=0) %}\n{%- set video_count = namespace(value=0) %}\n{%- for message in messages %}\n    {%- if message.role == \"user\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content_item in message.content %}\n                {%- if 'text' in content_item %}\n                    {{- content_item.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and message.content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n{%- endif %}\n",
    "model_type": "qwen3_vl",
    "tokenizer_config": {
      "bos_token": null,
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-09-08T17:46:47.000Z",
  "disabled": false,
  "downloads": 541,
  "gated": false,
  "id": "DeepCybo/PhysBrain1.5-2B",
  "lastModified": "2026-09-16T03:44:04.000Z",
  "library_name": "transformers",
  "likes": 13,
  "model-index": null,
  "modelId": "DeepCybo/PhysBrain1.5-2B",
  "pipeline_tag": "any-to-any",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 2161610752
    },
    "total": 2161610752
  },
  "sha": "21af7b27ec9ce8f3ade336f682d1450455ccec52",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "assets/.DS_Store"
    },
    {
      "rfilename": "assets/case-action-trajectory.png"
    },
    {
      "rfilename": "assets/case-future-prediction-v2.png"
    },
    {
      "rfilename": "assets/case-future-prediction.png"
    },
    {
      "rfilename": "assets/embodied_benchmarks.png"
    },
    {
      "rfilename": "assets/evaluation-cases-v6.png"
    },
    {
      "rfilename": "assets/leaderboard.svg"
    },
    {
      "rfilename": "assets/logo.png"
    },
    {
      "rfilename": "assets/model-architecture-v2.png"
    },
    {
      "rfilename": "chat_template.jinja"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "model.safetensors"
    },
    {
      "rfilename": "preprocessor_config.json"
    },
    {
      "rfilename": "processor_config.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "hugging-apps/physbrain1-5-2b-demo"
  ],
  "tags": [
    "transformers",
    "safetensors",
    "qwen3_vl",
    "image-text-to-text",
    "embodied-ai",
    "vision-language-model",
    "robotics",
    "spatial-reasoning",
    "multimodal",
    "any-to-any",
    "en",
    "zh",
    "arxiv:2609.14973",
    "base_model:Qwen/Qwen3-VL-2B-Instruct",
    "base_model:finetune:Qwen/Qwen3-VL-2B-Instruct",
    "endpoints_compatible",
    "region:us"
  ],
  "transformersInfo": {
    "auto_model": "AutoModelForMultimodalLM",
    "pipeline_tag": "image-text-to-text",
    "processor": "AutoProcessor"
  },
  "usedStorage": 4347865690
}
likes13
receipt
Source
Hugging Face models
Its words
13
Read by
field:likes
Said since
2026-10-02 12:02 UTC
Last answered
2026-10-04 18:21 UTC
Original
open at the source
What the source handed over
{
  "_asked": "DeepCybo/PhysBrain1.5-2B",
  "_id": "6aa04a07febd14f899e01a42",
  "author": "DeepCybo",
  "cardData": {
    "base_model": [
      "Qwen/Qwen3-VL-2B-Instruct"
    ],
    "language": [
      "en",
      "zh"
    ],
    "library_name": "transformers",
    "pipeline_tag": "any-to-any",
    "tags": [
      "embodied-ai",
      "vision-language-model",
      "robotics",
      "spatial-reasoning",
      "multimodal"
    ]
  },
  "config": {
    "architectures": [
      "Qwen3VLForConditionalGeneration"
    ],
    "chat_template_jinja": "{%- if tools %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if messages[0].role == 'system' %}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n\\n' }}\n    {%- endif %}\n    {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n    {%- for tool in tools %}\n        {{- \"\\n\" }}\n        {{- tool | tojson }}\n    {%- endfor %}\n    {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call><|im_end|>\\n\" }}\n{%- else %}\n    {%- if messages[0].role == 'system' %}\n        {{- '<|im_start|>system\\n' }}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- endif %}\n{%- endif %}\n{%- set image_count = namespace(value=0) %}\n{%- set video_count = namespace(value=0) %}\n{%- for message in messages %}\n    {%- if message.role == \"user\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content_item in message.content %}\n                {%- if 'text' in content_item %}\n                    {{- content_item.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and message.content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n{%- endif %}\n",
    "model_type": "qwen3_vl",
    "tokenizer_config": {
      "bos_token": null,
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-09-08T17:46:47.000Z",
  "disabled": false,
  "downloads": 541,
  "gated": false,
  "id": "DeepCybo/PhysBrain1.5-2B",
  "lastModified": "2026-09-16T03:44:04.000Z",
  "library_name": "transformers",
  "likes": 13,
  "model-index": null,
  "modelId": "DeepCybo/PhysBrain1.5-2B",
  "pipeline_tag": "any-to-any",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 2161610752
    },
    "total": 2161610752
  },
  "sha": "21af7b27ec9ce8f3ade336f682d1450455ccec52",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "assets/.DS_Store"
    },
    {
      "rfilename": "assets/case-action-trajectory.png"
    },
    {
      "rfilename": "assets/case-future-prediction-v2.png"
    },
    {
      "rfilename": "assets/case-future-prediction.png"
    },
    {
      "rfilename": "assets/embodied_benchmarks.png"
    },
    {
      "rfilename": "assets/evaluation-cases-v6.png"
    },
    {
      "rfilename": "assets/leaderboard.svg"
    },
    {
      "rfilename": "assets/logo.png"
    },
    {
      "rfilename": "assets/model-architecture-v2.png"
    },
    {
      "rfilename": "chat_template.jinja"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "model.safetensors"
    },
    {
      "rfilename": "preprocessor_config.json"
    },
    {
      "rfilename": "processor_config.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "hugging-apps/physbrain1-5-2b-demo"
  ],
  "tags": [
    "transformers",
    "safetensors",
    "qwen3_vl",
    "image-text-to-text",
    "embodied-ai",
    "vision-language-model",
    "robotics",
    "spatial-reasoning",
    "multimodal",
    "any-to-any",
    "en",
    "zh",
    "arxiv:2609.14973",
    "base_model:Qwen/Qwen3-VL-2B-Instruct",
    "base_model:finetune:Qwen/Qwen3-VL-2B-Instruct",
    "endpoints_compatible",
    "region:us"
  ],
  "transformersInfo": {
    "auto_model": "AutoModelForMultimodalLM",
    "pipeline_tag": "image-text-to-text",
    "processor": "AutoProcessor"
  },
  "usedStorage": 4347865690
}
taskany-to-any
receipt
Source
Hugging Face models
Its words
any-to-any
Read by
field:pipeline_tag
Said since
2026-10-02 12:02 UTC
Last answered
2026-10-04 18:21 UTC
Original
open at the source
What the source handed over
{
  "_asked": "DeepCybo/PhysBrain1.5-2B",
  "_id": "6aa04a07febd14f899e01a42",
  "author": "DeepCybo",
  "cardData": {
    "base_model": [
      "Qwen/Qwen3-VL-2B-Instruct"
    ],
    "language": [
      "en",
      "zh"
    ],
    "library_name": "transformers",
    "pipeline_tag": "any-to-any",
    "tags": [
      "embodied-ai",
      "vision-language-model",
      "robotics",
      "spatial-reasoning",
      "multimodal"
    ]
  },
  "config": {
    "architectures": [
      "Qwen3VLForConditionalGeneration"
    ],
    "chat_template_jinja": "{%- if tools %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if messages[0].role == 'system' %}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n\\n' }}\n    {%- endif %}\n    {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n    {%- for tool in tools %}\n        {{- \"\\n\" }}\n        {{- tool | tojson }}\n    {%- endfor %}\n    {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call><|im_end|>\\n\" }}\n{%- else %}\n    {%- if messages[0].role == 'system' %}\n        {{- '<|im_start|>system\\n' }}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- endif %}\n{%- endif %}\n{%- set image_count = namespace(value=0) %}\n{%- set video_count = namespace(value=0) %}\n{%- for message in messages %}\n    {%- if message.role == \"user\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content_item in message.content %}\n                {%- if 'text' in content_item %}\n                    {{- content_item.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and message.content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n{%- endif %}\n",
    "model_type": "qwen3_vl",
    "tokenizer_config": {
      "bos_token": null,
      "eos_token": "<|im_end|>",
      "pad_token": "<|endoftext|>",
      "unk_token": null
    }
  },
  "createdAt": "2026-09-08T17:46:47.000Z",
  "disabled": false,
  "downloads": 541,
  "gated": false,
  "id": "DeepCybo/PhysBrain1.5-2B",
  "lastModified": "2026-09-16T03:44:04.000Z",
  "library_name": "transformers",
  "likes": 13,
  "model-index": null,
  "modelId": "DeepCybo/PhysBrain1.5-2B",
  "pipeline_tag": "any-to-any",
  "private": false,
  "safetensors": {
    "parameters": {
      "BF16": 2161610752
    },
    "total": 2161610752
  },
  "sha": "21af7b27ec9ce8f3ade336f682d1450455ccec52",
  "siblings": [
    {
      "rfilename": ".gitattributes"
    },
    {
      "rfilename": "README.md"
    },
    {
      "rfilename": "assets/.DS_Store"
    },
    {
      "rfilename": "assets/case-action-trajectory.png"
    },
    {
      "rfilename": "assets/case-future-prediction-v2.png"
    },
    {
      "rfilename": "assets/case-future-prediction.png"
    },
    {
      "rfilename": "assets/embodied_benchmarks.png"
    },
    {
      "rfilename": "assets/evaluation-cases-v6.png"
    },
    {
      "rfilename": "assets/leaderboard.svg"
    },
    {
      "rfilename": "assets/logo.png"
    },
    {
      "rfilename": "assets/model-architecture-v2.png"
    },
    {
      "rfilename": "chat_template.jinja"
    },
    {
      "rfilename": "config.json"
    },
    {
      "rfilename": "model.safetensors"
    },
    {
      "rfilename": "preprocessor_config.json"
    },
    {
      "rfilename": "processor_config.json"
    },
    {
      "rfilename": "tokenizer.json"
    },
    {
      "rfilename": "tokenizer_config.json"
    }
  ],
  "spaces": [
    "hugging-apps/physbrain1-5-2b-demo"
  ],
  "tags": [
    "transformers",
    "safetensors",
    "qwen3_vl",
    "image-text-to-text",
    "embodied-ai",
    "vision-language-model",
    "robotics",
    "spatial-reasoning",
    "multimodal",
    "any-to-any",
    "en",
    "zh",
    "arxiv:2609.14973",
    "base_model:Qwen/Qwen3-VL-2B-Instruct",
    "base_model:finetune:Qwen/Qwen3-VL-2B-Instruct",
    "endpoints_compatible",
    "region:us"
  ],
  "transformersInfo": {
    "auto_model": "AutoModelForMultimodalLM",
    "pipeline_tag": "image-text-to-text",
    "processor": "AutoProcessor"
  },
  "usedStorage": 4347865690
}

Text

This source's terms allow its title, its values and a link here, not its text. It is at https://huggingface.co/DeepCybo/PhysBrain1.5-2B.