Integrate Qwen image prompt enhancers
This commit is contained in:
+146
@@ -572,6 +572,8 @@ services:
|
||||
ALLOWED_PROFILES: fast,medium,large,ultra,uncensored
|
||||
IMAGE_WORKER: image
|
||||
RESTORE_WORKER: restore
|
||||
IMAGE_PROMPT_I2I_WORKER: image-prompt-i2i
|
||||
IMAGE_PROMPT_T2I_WORKER: image-prompt-t2i
|
||||
FLUX_STANDBY_WORKER: flux-standby
|
||||
TTS_WORKER: qwen3
|
||||
MUSIC_WORKER: acestep
|
||||
@@ -602,6 +604,8 @@ services:
|
||||
volumes:
|
||||
- ./router/router_profiles.json:/etc/mike-ai/router-profiles.json:ro
|
||||
- ./config/global-system-policy.txt:/etc/mike-ai/global-system-policy.txt:ro
|
||||
- "${QWEN_IMAGE_PE_I2I_MODEL_DIR:-/data/models/qwen-image-2.1-pe-i2i-q5}/system_prompt.txt:/etc/mike-ai/qwen-image-pe-i2i-system-prompt.txt:ro"
|
||||
- "${QWEN_IMAGE_PE_T2I_MODEL_DIR:-/data/models/qwen-image-2.1-pe-t2i-q5}/system_prompt.txt:/etc/mike-ai/qwen-image-pe-t2i-system-prompt.txt:ro"
|
||||
- router-state:/var/lib/mike-ai-profile-router
|
||||
- router-images:/data/images
|
||||
environment:
|
||||
@@ -631,6 +635,11 @@ services:
|
||||
IMAGE_DIR: /data/images
|
||||
IMAGE_WORKER_URL: http://image-worker:8086
|
||||
IMAGE_WORKER_TOKEN: "${CONTROLLER_TOKEN:?CONTROLLER_TOKEN is required}"
|
||||
IMAGE_PROMPT_I2I_URL: http://image-prompt-enhancer-i2i:8080
|
||||
IMAGE_PROMPT_T2I_URL: http://image-prompt-enhancer-t2i:8080
|
||||
IMAGE_PROMPT_I2I_SYSTEM_FILE: /etc/mike-ai/qwen-image-pe-i2i-system-prompt.txt
|
||||
IMAGE_PROMPT_T2I_SYSTEM_FILE: /etc/mike-ai/qwen-image-pe-t2i-system-prompt.txt
|
||||
IMAGE_PROMPT_ENHANCER_TIMEOUT: "180"
|
||||
IMAGE_MODEL_NAME: Qwen-Image-2.1-int8
|
||||
IMAGE_INFERENCE_STEPS: "25"
|
||||
CHAT_IMAGE_ALLOW_REMOTE_URLS: "false"
|
||||
@@ -716,6 +725,143 @@ services:
|
||||
retries: 90
|
||||
start_period: 10s
|
||||
|
||||
image-prompt-enhancer-i2i:
|
||||
image: mike-ai/llama.cpp:b10930
|
||||
container_name: mike-ai-image-prompt-enhancer-i2i
|
||||
restart: "no"
|
||||
profiles: [image]
|
||||
labels:
|
||||
com.mike-ai.image-worker: image-prompt-i2i
|
||||
deploy:
|
||||
resources:
|
||||
reservations:
|
||||
devices:
|
||||
- driver: nvidia
|
||||
device_ids: ["${QWEN_IMAGE_PE_GPU:-0}"]
|
||||
capabilities: [gpu]
|
||||
read_only: true
|
||||
tmpfs: ["/tmp:size=512m,mode=1777"]
|
||||
volumes:
|
||||
- "${QWEN_IMAGE_PE_I2I_MODEL_DIR:-/data/models/qwen-image-2.1-pe-i2i-q5}:/models:ro"
|
||||
command:
|
||||
- --model
|
||||
- /models/Qwen-Image-2.1-PE-I2I.Q5_K_M.gguf
|
||||
- --mmproj
|
||||
- /models/Qwen-Image-2.1-PE-I2I.mmproj-bf16.gguf
|
||||
- --mmproj-offload
|
||||
- --mmproj-device
|
||||
- CUDA0
|
||||
- --alias
|
||||
- qwen-image-pe-i2i
|
||||
- --ctx-size
|
||||
- "16384"
|
||||
- --flash-attn
|
||||
- "on"
|
||||
- --cache-type-k
|
||||
- q8_0
|
||||
- --cache-type-v
|
||||
- q8_0
|
||||
- --threads
|
||||
- "6"
|
||||
- --threads-batch
|
||||
- "6"
|
||||
- --batch-size
|
||||
- "1024"
|
||||
- --ubatch-size
|
||||
- "128"
|
||||
- --parallel
|
||||
- "1"
|
||||
- --jinja
|
||||
- --reasoning
|
||||
- auto
|
||||
- --host
|
||||
- 0.0.0.0
|
||||
- --port
|
||||
- "8080"
|
||||
- --metrics
|
||||
- --n-gpu-layers
|
||||
- all
|
||||
- --device
|
||||
- CUDA0
|
||||
- --split-mode
|
||||
- none
|
||||
- --no-ui
|
||||
networks: [inference]
|
||||
security_opt: ["no-new-privileges:true"]
|
||||
cap_drop: [ALL]
|
||||
healthcheck:
|
||||
test: [CMD, curl, -fsS, "http://127.0.0.1:8080/health"]
|
||||
interval: 5s
|
||||
timeout: 3s
|
||||
retries: 40
|
||||
start_period: 10s
|
||||
|
||||
image-prompt-enhancer-t2i:
|
||||
image: mike-ai/llama.cpp:b10930
|
||||
container_name: mike-ai-image-prompt-enhancer-t2i
|
||||
restart: "no"
|
||||
profiles: [image]
|
||||
labels:
|
||||
com.mike-ai.image-worker: image-prompt-t2i
|
||||
deploy:
|
||||
resources:
|
||||
reservations:
|
||||
devices:
|
||||
- driver: nvidia
|
||||
device_ids: ["${QWEN_IMAGE_PE_GPU:-0}"]
|
||||
capabilities: [gpu]
|
||||
read_only: true
|
||||
tmpfs: ["/tmp:size=512m,mode=1777"]
|
||||
volumes:
|
||||
- "${QWEN_IMAGE_PE_T2I_MODEL_DIR:-/data/models/qwen-image-2.1-pe-t2i-q5}:/models:ro"
|
||||
command:
|
||||
- --model
|
||||
- /models/Qwen-Image-2.1-PE-T2I.Q5_K_M.gguf
|
||||
- --alias
|
||||
- qwen-image-pe-t2i
|
||||
- --ctx-size
|
||||
- "16384"
|
||||
- --flash-attn
|
||||
- "on"
|
||||
- --cache-type-k
|
||||
- q8_0
|
||||
- --cache-type-v
|
||||
- q8_0
|
||||
- --threads
|
||||
- "6"
|
||||
- --threads-batch
|
||||
- "6"
|
||||
- --batch-size
|
||||
- "1024"
|
||||
- --ubatch-size
|
||||
- "128"
|
||||
- --parallel
|
||||
- "1"
|
||||
- --jinja
|
||||
- --reasoning
|
||||
- auto
|
||||
- --host
|
||||
- 0.0.0.0
|
||||
- --port
|
||||
- "8080"
|
||||
- --metrics
|
||||
- --n-gpu-layers
|
||||
- all
|
||||
- --device
|
||||
- CUDA0
|
||||
- --split-mode
|
||||
- none
|
||||
- --no-ui
|
||||
networks: [inference]
|
||||
security_opt: ["no-new-privileges:true"]
|
||||
cap_drop: [ALL]
|
||||
healthcheck:
|
||||
test: [CMD, curl, -fsS, "http://127.0.0.1:8080/health"]
|
||||
interval: 5s
|
||||
timeout: 3s
|
||||
retries: 40
|
||||
start_period: 10s
|
||||
|
||||
# Previous production image model, retained as a stopped rollback target.
|
||||
# It is outside the normal router path and can only be started through the
|
||||
# controller's allowlisted flux-standby endpoint.
|
||||
|
||||
Reference in New Issue
Block a user