Skip to content
Closed
Show file tree
Hide file tree
Changes from all commits
Commits
Show all changes
27 commits
Select commit Hold shift + click to select a range
e63bf80
feat(sync): add Novita AI model catalog sync
Alex-yang00 Sep 14, 2026
b3eb982
fix(sync): accept Novita model list response
Alex-yang00 Sep 15, 2026
ffde31b
feat(sync): map Novita catalog metadata
Alex-yang00 Sep 17, 2026
010bff0
feat(sync): automatically sync Novita model metadata
Alex-yang00 Sep 17, 2026
a43f094
feat(sync): enable full Novita catalog lifecycle
Alex-yang00 Sep 17, 2026
5fbd74a
fix(sync): avoid guessing Novita pricing and model capabilities
Alex-yang00 Sep 17, 2026
89c2339
fix(sync): recognize Novita free models and verified controls
Alex-yang00 Sep 17, 2026
ed34f5b
fix(sync): harden Novita catalog automation and preserve authored prices
Alex-yang00 Sep 17, 2026
3dc8eeb
feat(novita-ai): sync DeepSeek V4.1 Flash with verified reasoning toggle
Alex-yang00 Sep 17, 2026
d9941f7
feat(novita-ai): sync models with verified thinking controls
Alex-yang00 Sep 17, 2026
f136248
merge dev and preserve Novita sync hardening
Alex-yang00 Sep 17, 2026
0ab93c0
fix(novita-ai): tolerate zero context catalog entries
Alex-yang00 Sep 18, 2026
318642d
fix(novita-ai): address automated review findings
Alex-yang00 Sep 18, 2026
a6acdf7
fix(novita-ai): resolve reasoning review findings
Alex-yang00 Sep 18, 2026
ffaa945
fix(novita-ai): address latest review findings
Alex-yang00 Sep 18, 2026
53d8701
fix(novita-ai): align sync controls with live API behavior
Alex-yang00 Sep 18, 2026
6552543
fix(novita-ai): preserve verified reasoning controls
Alex-yang00 Sep 18, 2026
8b0d622
fix(novita-ai): complete verified reasoning metadata
Alex-yang00 Sep 18, 2026
ed97138
fix(novita-ai): resolve remaining catalog review findings
Alex-yang00 Sep 20, 2026
9587a67
fix(novita-ai): factor restored models through lab metadata
Alex-yang00 Sep 20, 2026
ec72d3a
fix(novita-ai): align stale catalog entries with live routes
Alex-yang00 Sep 20, 2026
5c0b73d
fix(novita-ai): enforce catalog sync boundaries
Alex-yang00 Sep 20, 2026
1c4f353
chore: retry PR review
Alex-yang00 Sep 20, 2026
bd4ed0d
fix(novita-ai): verify retained routes and factor lab identities
Alex-yang00 Sep 20, 2026
b960c61
fix(novita-ai): factor remaining routes and correct host capabilities
Alex-yang00 Sep 20, 2026
2071974
fix(novita-ai): verify remaining reasoning review findings
Alex-yang00 Sep 20, 2026
867b777
Merge dev into feat/novita-ai-model-sync
Alex-yang00 Oct 9, 2026
File filter

Filter by extension

Filter by extension


Conversations
Failed to load comments.
Loading
Jump to
Jump to file
Failed to load files.
Loading
Diff view
Diff view
1 change: 1 addition & 0 deletions .github/workflows/sync-models.yml
Original file line number Diff line number Diff line change
Expand Up @@ -84,6 +84,7 @@ jobs:
VENICE_API_KEY: ${{ secrets.VENICE_API_KEY }}
LLMGATEWAY_API_KEY: ${{ secrets.LLMGATEWAY_API_KEY }}
MERGE_GATEWAY_API_KEY: ${{ secrets.MERGE_GATEWAY_API_KEY }}
NOVITA_API_KEY: ${{ secrets.NOVITA_API_KEY }}
MISTRAL_API_KEY: ${{ secrets.MISTRAL_API_KEY }}
KILO_API_KEY: ${{ secrets.KILO_API_KEY }}
GOOGLE_API_KEY: ${{ secrets.GOOGLE_API_KEY }}
Expand Down
19 changes: 19 additions & 0 deletions models/alibaba/qwen-mt-plus.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,19 @@
name = "Qwen MT Plus"
description = "Qwen translation model for multilingual conversion and localization"
family = "qwen"
release_date = "2025-09-03"
last_updated = "2025-09-03"
attachment = false
reasoning = false
temperature = true
tool_call = false
structured_output = false
open_weights = true

[limit]
context = 16_384
output = 8_192

[modalities]
input = ["text"]
output = ["text"]
20 changes: 20 additions & 0 deletions models/alibaba/qwen2.5-72b-instruct.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,20 @@
name = "Qwen2.5 72B Instruct"
description = "Open-weight Qwen2.5 instruction model for multilingual chat and coding"
family = "qwen"
release_date = "2024-10-15"
last_updated = "2024-10-15"
attachment = false
reasoning = false
temperature = true
tool_call = true
structured_output = true
knowledge = "2024-04"
open_weights = true

[limit]
context = 32_000
output = 8_192

[modalities]
input = ["text"]
output = ["text"]
25 changes: 25 additions & 0 deletions models/alibaba/qwen2.5-7b-instruct.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,25 @@
# https://huggingface.co/Qwen/Qwen2.5-7B-Instruct
name = "Qwen2.5 7B Instruct"
description = "Open Qwen instruction model for multilingual chat, coding, and structured responses"
family = "qwen"
release_date = "2024-09-19"
last_updated = "2024-09-19"
attachment = false
reasoning = false
temperature = true
tool_call = true
structured_output = true
open_weights = true
license = "Apache-2.0"

[limit]
context = 131_072
output = 8_192

[modalities]
input = ["text"]
output = ["text"]

[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen2.5-7B-Instruct"
24 changes: 24 additions & 0 deletions models/alibaba/qwen3-235b-a22b-thinking-2507.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,24 @@
# https://huggingface.co/Qwen/Qwen3-235B-A22B-Thinking-2507
name = "Qwen3 235B A22B Thinking 2507"
description = "Qwen reasoning model for deliberate problem solving, math, coding, and agentic workflows"
family = "qwen"
release_date = "2025-07-25"
last_updated = "2025-07-25"
attachment = false
reasoning = true
temperature = true
tool_call = true
knowledge = "2025-04"
open_weights = true

[limit]
context = 131_072
output = 32_768

[modalities]
input = ["text"]
output = ["text"]

[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3-235B-A22B-Thinking-2507"
20 changes: 20 additions & 0 deletions models/alibaba/qwen3-omni-30b-a3b-instruct.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,20 @@
name = "Qwen3 Omni 30B A3B Instruct"
description = "Qwen3 omni instruction model for text, vision, audio, and multimodal tasks"
family = "qwen"
release_date = "2025-09-24"
last_updated = "2025-09-24"
attachment = true
reasoning = false
temperature = true
knowledge = "2024-04"
tool_call = true
structured_output = true
open_weights = true

[limit]
context = 65_536
output = 16_384

[modalities]
input = ["text", "video", "audio", "image"]
output = ["text", "audio"]
19 changes: 19 additions & 0 deletions models/alibaba/qwen3-omni-30b-a3b-thinking.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,19 @@
name = "Qwen3 Omni 30B A3B Thinking"
description = "Qwen3 omni reasoning model for multimodal text, vision, and audio tasks"
family = "qwen"
release_date = "2025-09-24"
last_updated = "2025-09-24"
attachment = true
reasoning = true
temperature = true
tool_call = true
structured_output = true
open_weights = true

[limit]
context = 65_536
output = 16_384

[modalities]
input = ["text", "audio", "video", "image"]
output = ["text"]
24 changes: 24 additions & 0 deletions models/alibaba/qwen3-vl-30b-a3b-instruct.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,24 @@
# https://huggingface.co/Qwen/Qwen3-VL-30B-A3B-Instruct
name = "Qwen3 VL 30B A3B Instruct"
description = "Qwen vision-language instruction model for documents, visual understanding, and agent tasks"
family = "qwen"
release_date = "2025-10-11"
last_updated = "2025-10-11"
attachment = true
reasoning = false
temperature = true
tool_call = true
structured_output = true
open_weights = true

[limit]
context = 131_072
output = 32_768

[modalities]
input = ["text", "image"]
output = ["text"]

[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/Qwen/Qwen3-VL-30B-A3B-Instruct"
20 changes: 20 additions & 0 deletions models/baichuan/baichuan-m2-32b.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,20 @@
name = "Baichuan M2 32B"
description = "Open-weight Baichuan instruction model for chat and analysis"
family = "baichuan"
release_date = "2025-08-13"
last_updated = "2025-08-13"
attachment = false
reasoning = false
temperature = true
tool_call = false
structured_output = false
knowledge = "2024-12"
open_weights = true

[limit]
context = 131_072
output = 131_072

[modalities]
input = ["text"]
output = ["text"]
20 changes: 20 additions & 0 deletions models/baidu/ernie-4.5-21b-a3b.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,20 @@
name = "ERNIE 4.5 21B A3B"
description = "Baidu ERNIE 4.5 open-weight mixture-of-experts instruction model"
family = "ernie"
release_date = "2025-06-30"
last_updated = "2025-06-30"
attachment = false
reasoning = false
temperature = true
tool_call = true
structured_output = false
knowledge = "2025-03"
open_weights = true

[limit]
context = 120_000
output = 8_000

[modalities]
input = ["text"]
output = ["text"]
Original file line number Diff line number Diff line change
@@ -1,20 +1,18 @@
name = "L31 70B Euryale V2.2"
name = "ERNIE 4.5 300B A47B"
description = "Open-weight instruction model for adaptable chat and self-hosted production workloads"
release_date = "2024-09-19"
last_updated = "2024-09-19"
family = "ernie"
release_date = "2025-06-30"
last_updated = "2025-06-30"
attachment = false
reasoning = false
temperature = true
tool_call = true
structured_output = true
open_weights = true

[cost]
input = 1.48
output = 1.48

[limit]
context = 8_192
output = 8_192
context = 123_000
output = 12_000

[modalities]
input = ["text"]
Expand Down
19 changes: 19 additions & 0 deletions models/baidu/ernie-4.5-vl-424b-a47b.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,19 @@
name = "ERNIE 4.5 VL 424B A47B"
description = "Baidu ERNIE 4.5 vision-language model for visual analysis, planning, and tool use"
family = "ernie"
release_date = "2025-06-30"
last_updated = "2025-06-30"
attachment = true
reasoning = true
temperature = true
tool_call = false
structured_output = false
open_weights = true

[limit]
context = 123_000
output = 16_000

[modalities]
input = ["text", "image"]
output = ["text"]
25 changes: 25 additions & 0 deletions models/deepseek/deepseek-ocr.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,25 @@
# https://huggingface.co/deepseek-ai/DeepSeek-OCR
name = "DeepSeek-OCR"
description = "OCR model for extracting structured text from documents and screenshots"
family = "deepseek"
release_date = "2025-10-20"
last_updated = "2025-10-20"
attachment = true
reasoning = false
temperature = true
tool_call = false
structured_output = true
open_weights = true
license = "MIT"

[limit]
context = 8_192
output = 8_192

[modalities]
input = ["text", "image"]
output = ["text"]

[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/deepseek-ai/DeepSeek-OCR"
24 changes: 24 additions & 0 deletions models/deepseek/deepseek-r1-0528-qwen3-8b.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,24 @@
# https://huggingface.co/deepseek-ai/DeepSeek-R1-0528-Qwen3-8B
name = "DeepSeek R1 0528 Qwen3 8B"
description = "DeepSeek R1 0528 distilled into Qwen3 8B for compact math, coding, and reasoning"
family = "deepseek-thinking"
release_date = "2025-05-28"
last_updated = "2025-05-28"
attachment = false
reasoning = true
temperature = true
tool_call = false
structured_output = false
open_weights = true

[limit]
context = 128_000
output = 32_000

[modalities]
input = ["text"]
output = ["text"]

[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/deepseek-ai/DeepSeek-R1-0528-Qwen3-8B"
20 changes: 20 additions & 0 deletions models/deepseek/deepseek-r1-0528.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,20 @@
name = "DeepSeek R1 0528"
description = "DeepSeek R1 reasoning model updated in May 2025 for math and coding"
family = "deepseek-thinking"
release_date = "2025-05-28"
last_updated = "2025-05-28"
attachment = false
reasoning = true
temperature = true
knowledge = "2024-07"
tool_call = true
structured_output = true
open_weights = true

[limit]
context = 163_840
output = 32_768

[modalities]
input = ["text"]
output = ["text"]
24 changes: 24 additions & 0 deletions models/deepseek/deepseek-r1-distill-llama-70b.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,24 @@
# https://huggingface.co/deepseek-ai/DeepSeek-R1-Distill-Llama-70B
name = "DeepSeek R1 Distill Llama 70B"
description = "DeepSeek R1 reasoning distilled into Llama 3.3 70B for math, coding, and analysis"
family = "deepseek-thinking"
release_date = "2025-01-20"
last_updated = "2025-01-20"
attachment = false
reasoning = true
temperature = true
tool_call = false
structured_output = false
open_weights = true

[limit]
context = 8_192
output = 8_192

[modalities]
input = ["text"]
output = ["text"]

[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/deepseek-ai/DeepSeek-R1-Distill-Llama-70B"
26 changes: 26 additions & 0 deletions models/deepseek/deepseek-v3.1-terminus.toml
Original file line number Diff line number Diff line change
@@ -0,0 +1,26 @@
# https://huggingface.co/deepseek-ai/DeepSeek-V3.1-Terminus
name = "DeepSeek V3.1 Terminus"
description = "DeepSeek V3.1 Terminus hybrid-reasoning model for coding, analysis, and agent workflows"
family = "deepseek"
release_date = "2025-09-22"
last_updated = "2025-09-22"
attachment = false
reasoning = true
temperature = true
tool_call = true
structured_output = true
knowledge = "2025-03-31"
open_weights = true
license = "MIT License"

[limit]
context = 163_840
output = 32_768

[modalities]
input = ["text"]
output = ["text"]

[[weights]]
label = "Hugging Face"
url = "https://huggingface.co/deepseek-ai/DeepSeek-V3.1-Terminus"
Loading
Loading