Models

All models available

deepseek-ai/DeepSeek-V4-Flash-0731

deepseek-ai/DeepSeek-V4-Flash-0731

-

moonshotai/Kimi-K3

moonshotai/Kimi-K3

-

thinkingmachines/Inkling-Small

thinkingmachines/Inkling-Small

-

zai-org/GLM-5.2

zai-org/GLM-5.2

-

deepseek-ai/DeepSeek-V4-Flash

deepseek-ai/DeepSeek-V4-Flash

-

meta-llama/Llama-3.1-8B-Instruct

meta-llama/Llama-3.1-8B-Instruct

-

Qwen/Qwen2.5-7B-Instruct

Qwen/Qwen2.5-7B-Instruct

-

Qwen/Qwen3.6-27B

Qwen/Qwen3.6-27B

-

prism-ml/Ternary-Bonsai-27B-gguf

prism-ml/Ternary-Bonsai-27B-gguf

-

deepseek-ai/DeepSeek-V3

deepseek-ai/DeepSeek-V3

-

MiniMaxAI/MiniMax-M3

MiniMaxAI/MiniMax-M3

-

thinkingmachines/Inkling

thinkingmachines/Inkling

-

Qwen/Qwen3.6-35B-A3B

Qwen/Qwen3.6-35B-A3B

-

google/gemma-4-31B-it

google/gemma-4-31B-it

-

deepseek-ai/DeepSeek-V4-Pro

deepseek-ai/DeepSeek-V4-Pro

-

openai/gpt-oss-120b

openai/gpt-oss-120b

-

openai/gpt-oss-20b

openai/gpt-oss-20b

-

moonshotai/Kimi-K2.7-Code

moonshotai/Kimi-K2.7-Code

-

XiaomiMiMo/MiMo-V2.5

XiaomiMiMo/MiMo-V2.5

-

Qwen/Qwen3.5-9B

Qwen/Qwen3.5-9B

-

tencent/Hy3

tencent/Hy3

-

deepseek-ai/DeepSeek-R1

deepseek-ai/DeepSeek-R1

-

swiss-ai/Apertus-v1.5-70B

swiss-ai/Apertus-v1.5-70B

-

Qwen/Qwen3-8B

Qwen/Qwen3-8B

-

google/gemma-4-26B-A4B-it

google/gemma-4-26B-A4B-it

-

XiaomiMiMo/MiMo-V2.5-Pro

XiaomiMiMo/MiMo-V2.5-Pro

-

nvidia/NVIDIA-Nemotron-3-Ultra-550B-A55B-NVFP4

nvidia/NVIDIA-Nemotron-3-Ultra-550B-A55B-NVFP4

-

swiss-ai/Apertus-v1.5-8B

swiss-ai/Apertus-v1.5-8B

-

Qwen/Qwen3.5-122B-A10B

Qwen/Qwen3.5-122B-A10B

-

Qwen/Qwen3-Coder-Next

Qwen/Qwen3-Coder-Next

-

meta-llama/Llama-3.3-70B-Instruct

meta-llama/Llama-3.3-70B-Instruct

-

zai-org/GLM-4.7-Flash

zai-org/GLM-4.7-Flash

-

moonshotai/Kimi-K2.6

moonshotai/Kimi-K2.6

-

zai-org/GLM-4.5-Air

zai-org/GLM-4.5-Air

-

Qwen/Qwen3.5-27B

Qwen/Qwen3.5-27B

-

Qwen/Qwen3-4B-Instruct-2507

Qwen/Qwen3-4B-Instruct-2507

-

MiniMaxAI/MiniMax-M2.7

MiniMaxAI/MiniMax-M2.7

-

zai-org/GLM-5.2-FP8

zai-org/GLM-5.2-FP8

-

stepfun-ai/Step-3.7-Flash

stepfun-ai/Step-3.7-Flash

-

speakleash/Bielik-11B-v3.0-Instruct

speakleash/Bielik-11B-v3.0-Instruct

-

zai-org/GLM-4.6V-Flash

zai-org/GLM-4.6V-Flash

-

Qwen/Qwen3-235B-A22B

Qwen/Qwen3-235B-A22B

-

nvidia/NVIDIA-Nemotron-3-Ultra-550B-A55B-BF16

nvidia/NVIDIA-Nemotron-3-Ultra-550B-A55B-BF16

-

Qwen/Qwen3.5-397B-A17B

Qwen/Qwen3.5-397B-A17B

-

Qwen/Qwen2.5-Coder-32B-Instruct

Qwen/Qwen2.5-Coder-32B-Instruct

-

Qwen/Qwen3.5-35B-A3B

Qwen/Qwen3.5-35B-A3B

-

google/gemma-3-4b-it

google/gemma-3-4b-it

-

openai/gpt-oss-safeguard-20b

openai/gpt-oss-safeguard-20b

-

Qwen/Qwen3-Coder-30B-A3B-Instruct

Qwen/Qwen3-Coder-30B-A3B-Instruct

-

Qwen/Qwen2.5-72B-Instruct

Qwen/Qwen2.5-72B-Instruct

-

Qwen/Qwen3-30B-A3B

Qwen/Qwen3-30B-A3B

-

deepseek-ai/DeepSeek-R1-Distill-Llama-70B

deepseek-ai/DeepSeek-R1-Distill-Llama-70B

-

Qwen/Qwen3-Next-80B-A3B-Instruct

Qwen/Qwen3-Next-80B-A3B-Instruct

-

Qwen/Qwen3-14B

Qwen/Qwen3-14B

-

google/gemma-3-12b-it

google/gemma-3-12b-it

-

microsoft/phi-4

microsoft/phi-4

-

swiss-ai/Apertus-8B-Instruct-2509

swiss-ai/Apertus-8B-Instruct-2509

-

prism-ml/Ternary-Bonsai-27B-AWQ-4bit

prism-ml/Ternary-Bonsai-27B-AWQ-4bit

-

MiniMaxAI/MiniMax-M2.1

MiniMaxAI/MiniMax-M2.1

-

swiss-ai/Apertus-70B-Instruct-2509

swiss-ai/Apertus-70B-Instruct-2509

-

zai-org/GLM-4.5

zai-org/GLM-4.5

-

baidu/ERNIE-4.5-VL-424B-A47B-Base-PT

baidu/ERNIE-4.5-VL-424B-A47B-Base-PT

-

Qwen/Qwen2.5-Coder-3B-Instruct

Qwen/Qwen2.5-Coder-3B-Instruct

-

moonshotai/Kimi-K2-Instruct-0905

moonshotai/Kimi-K2-Instruct-0905

-

utter-project/EuroLLM-22B-Instruct-2512

utter-project/EuroLLM-22B-Instruct-2512

-

zai-org/GLM-5.1

zai-org/GLM-5.1

-

Qwen/Qwen3-VL-235B-A22B-Instruct

Qwen/Qwen3-VL-235B-A22B-Instruct

-

deepseek-ai/DeepSeek-R1-Distill-Qwen-7B

deepseek-ai/DeepSeek-R1-Distill-Qwen-7B

-

deepseek-ai/DeepSeek-V3.2-Exp

deepseek-ai/DeepSeek-V3.2-Exp

-

moonshotai/Kimi-K2.5

moonshotai/Kimi-K2.5

-

deepseek-ai/DeepSeek-R1-Distill-Qwen-14B

deepseek-ai/DeepSeek-R1-Distill-Qwen-14B

-

MiniMaxAI/MiniMax-M2.5

MiniMaxAI/MiniMax-M2.5

-

deepcogito/cogito-671b-v2.1

deepcogito/cogito-671b-v2.1

-

meta-llama/Llama-Guard-4-12B

meta-llama/Llama-Guard-4-12B

-

NousResearch/Hermes-3-Llama-3.1-70B

NousResearch/Hermes-3-Llama-3.1-70B

-

Qwen/Qwen3-32B

Qwen/Qwen3-32B

-

stepfun-ai/Step-3.5-Flash

stepfun-ai/Step-3.5-Flash

-

CohereLabs/c4ai-command-r-08-2024

CohereLabs/c4ai-command-r-08-2024

-

Qwen/Qwen2.5-VL-72B-Instruct

Qwen/Qwen2.5-VL-72B-Instruct

-

meta-llama/Llama-4-Maverick-17B-128E-Instruct-FP8

meta-llama/Llama-4-Maverick-17B-128E-Instruct-FP8

-

Qwen/Qwen2.5-Coder-7B-Instruct

Qwen/Qwen2.5-Coder-7B-Instruct

-

deepseek-ai/DeepSeek-R1-0528

deepseek-ai/DeepSeek-R1-0528

-

meta-llama/Llama-4-Scout-17B-16E-Instruct

meta-llama/Llama-4-Scout-17B-16E-Instruct

-

CohereLabs/aya-expanse-32b

CohereLabs/aya-expanse-32b

-

zai-org/GLM-4.5V

zai-org/GLM-4.5V

-

deepseek-ai/DeepSeek-R1-Distill-Llama-8B

deepseek-ai/DeepSeek-R1-Distill-Llama-8B

-

Qwen/Qwen3-235B-A22B-Thinking-2507

Qwen/Qwen3-235B-A22B-Thinking-2507

-

CohereLabs/c4ai-command-r7b-12-2024

CohereLabs/c4ai-command-r7b-12-2024

-

Qwen/Qwen3-VL-235B-A22B-Thinking

Qwen/Qwen3-VL-235B-A22B-Thinking

-

Qwen/Qwen3-4B-Thinking-2507

Qwen/Qwen3-4B-Thinking-2507

-

deepseek-ai/DeepSeek-V3.1

deepseek-ai/DeepSeek-V3.1

-

CohereLabs/c4ai-command-r7b-arabic-02-2025

CohereLabs/c4ai-command-r7b-arabic-02-2025

-

Qwen/Qwen3-VL-30B-A3B-Instruct

Qwen/Qwen3-VL-30B-A3B-Instruct

-

zai-org/GLM-4.6-FP8

zai-org/GLM-4.6-FP8

-

deepseek-ai/DeepSeek-V3.1-Terminus

deepseek-ai/DeepSeek-V3.1-Terminus

-

CohereLabs/c4ai-command-a-03-2025

CohereLabs/c4ai-command-a-03-2025

-

zai-org/AutoGLM-Phone-9B-Multilingual

zai-org/AutoGLM-Phone-9B-Multilingual

-

zai-org/GLM-4.6

zai-org/GLM-4.6

-

deepseek-ai/DeepSeek-V3.2

deepseek-ai/DeepSeek-V3.2

-

CohereLabs/command-a-reasoning-08-2025

CohereLabs/command-a-reasoning-08-2025

-

aisingapore/Gemma-SEA-LION-v4-27B-IT

aisingapore/Gemma-SEA-LION-v4-27B-IT

-

zai-org/GLM-4.7-FP8

zai-org/GLM-4.7-FP8

-

zai-org/GLM-5

zai-org/GLM-5

-

CohereLabs/command-a-translate-08-2025

CohereLabs/command-a-translate-08-2025

-

aisingapore/Qwen-SEA-LION-v4-32B-IT

aisingapore/Qwen-SEA-LION-v4-32B-IT

-

zai-org/GLM-5.1-FP8

zai-org/GLM-5.1-FP8

-

alpindale/WizardLM-2-8x22B

alpindale/WizardLM-2-8x22B

-

CohereLabs/tiny-aya-global

CohereLabs/tiny-aya-global

-

allenai/Olmo-3-7B-Instruct

allenai/Olmo-3-7B-Instruct

-

zai-org/GLM-4.7

zai-org/GLM-4.7

-

Active
Sao10K/L3-8B-Stheno-v3.2

Sao10K/L3-8B-Stheno-v3.2

-

CohereLabs/tiny-aya-water

CohereLabs/tiny-aya-water

-

deepcogito/cogito-671b-v2.1-FP8

deepcogito/cogito-671b-v2.1-FP8

-

CohereLabs/aya-vision-32b

CohereLabs/aya-vision-32b

-

Sao10K/L3-8B-Lunaris-v1

Sao10K/L3-8B-Lunaris-v1

-

zai-org/GLM-4.5V-FP8

zai-org/GLM-4.5V-FP8

-

CohereLabs/tiny-aya-earth

CohereLabs/tiny-aya-earth

-

deepseek-ai/DeepSeek-V3-0324

deepseek-ai/DeepSeek-V3-0324

-

zai-org/GLM-4.6V

zai-org/GLM-4.6V

-

CohereLabs/tiny-aya-fire

CohereLabs/tiny-aya-fire

-

zai-org/GLM-4-32B-0414

zai-org/GLM-4-32B-0414

-

zai-org/GLM-4.6V-FP8

zai-org/GLM-4.6V-FP8

-

google/gemma-3-27b-it

google/gemma-3-27b-it

-

MiniMaxAI/MiniMax-M1-80k

MiniMaxAI/MiniMax-M1-80k

-

google/gemma-3n-E4B-it

google/gemma-3n-E4B-it

-

Qwen/Qwen3-235B-A22B-Instruct-2507

Qwen/Qwen3-235B-A22B-Instruct-2507

-

moonshotai/Kimi-K2-Instruct

moonshotai/Kimi-K2-Instruct

-

pearl-ai/Gemma-4-31B-it-pearl

pearl-ai/Gemma-4-31B-it-pearl

-

Qwen/Qwen3-Coder-480B-A35B-Instruct

Qwen/Qwen3-Coder-480B-A35B-Instruct

-

MiniMaxAI/MiniMax-M2

MiniMaxAI/MiniMax-M2

-

inclusionAI/Ling-2.6-1T

inclusionAI/Ling-2.6-1T

-