Add dedicated Hermes compression model
This commit is contained in:
@@ -100,9 +100,15 @@ auxiliary:
|
||||
title_generation:
|
||||
enabled: false
|
||||
compression:
|
||||
provider: "main"
|
||||
model: ""
|
||||
provider: "openai-api"
|
||||
model: "qwen-compression"
|
||||
base_url: "http://192.168.1.212:8099/v1"
|
||||
api_key: "local"
|
||||
reasoning_effort: "none"
|
||||
extra_body:
|
||||
chat_template_kwargs:
|
||||
enable_thinking: false
|
||||
timeout: 600
|
||||
|
||||
# The full hermes-cli preset injects several large, overlapping schemas on
|
||||
# every turn. Athena already exposes browsing, orchestration and host services
|
||||
|
||||
Reference in New Issue
Block a user