Stage 1 step 0: train/.venv (py3.11, mlx-lm 0.32), MLX 4-bit Qwen3.8-27B, serve script, baseline runner, README
Co-Authored-By: Claude Opus 5.5 <noreply@anthropic.com> Claude-Session: https://claude.ai/code/session_014aUaQeLnwbb1zTpN7kHeat
This commit is contained in:
11
train/serve.sh
Executable file
11
train/serve.sh
Executable file
@@ -0,0 +1,11 @@
|
||||
#!/bin/sh
|
||||
# Serve the base model (no adapter) with mlx_lm.server. Settings: train/README.md.
|
||||
# Usage: train/serve.sh [--adapter-path PATH]
|
||||
cd "$(dirname "$0")/.."
|
||||
exec train/.venv/bin/mlx_lm.server \
|
||||
--model "$HOME/models/Qwen3.8-27B-4bit" \
|
||||
--host 127.0.0.1 --port 8080 \
|
||||
--temp 0.2 --top-p 0.95 --top-k 20 --min-p 0 \
|
||||
--max-tokens 32768 \
|
||||
--chat-template-args '{"enable_thinking": true, "reasoning_effort": "medium"}' \
|
||||
"$@"
|
||||
Reference in New Issue
Block a user