chore: scaffold buddhagpt, verify Qwen2.5-7B-4bit runs on M5 (PAI-87)

This commit is contained in:
marcuspaico
2026-08-14 17:01:33 -07:00
parent 4a5f2026c4
commit 76d1858ffe
7 changed files with 1638 additions and 0 deletions

View File

@@ -0,0 +1,8 @@
from mlx_lm import load, generate
model, tokenizer = load("mlx-community/Qwen2.5-7B-Instruct-4bit")
prompt = tokenizer.apply_chat_template(
[{"role": "user", "content": "What is the Second Noble Truth?"}],
tokenize=False, add_generation_prompt=True,
)
print(generate(model, tokenizer, prompt=prompt, max_tokens=200))