Image-Text-to-Text
MLX
Safetensors
qwen3_5
qwen3.8
reasoning
vision-language
personal-model
uncensored
abliterated
abliterix
apple-silicon
8-bit precision
conversational
Instructions to use timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit with libraries, inference providers, notebooks, and local apps. Follow these links to get started.
- Libraries
- MLX
How to use timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit with MLX:
# Make sure mlx-vlm is installed # pip install --upgrade mlx-vlm from mlx_vlm import load, generate from mlx_vlm.prompt_utils import apply_chat_template from mlx_vlm.utils import load_config # Load the model model, processor = load("timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit") config = load_config("timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit") # Prepare input image = ["http://images.cocodataset.org/val2017/000000039769.jpg"] prompt = "Describe this image." # Apply chat template formatted_prompt = apply_chat_template( processor, config, prompt, num_images=1 ) # Generate output output = generate(model, processor, formatted_prompt, image) print(output) - Notebooks
- Google Colab
- Kaggle
- Local Apps Settings
- LM Studio
- Pi
How to use timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit with Pi:
Start the MLX server
# Install MLX LM: uv tool install mlx-lm # Start a local OpenAI-compatible server: mlx_lm.server --model "timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit"
Configure the model in Pi
# Install Pi: npm install -g @earendil-works/pi-coding-agent # Add to ~/.pi/agent/models.json: { "providers": { "mlx-lm": { "baseUrl": "http://localhost:8080/v1", "api": "openai-completions", "apiKey": "none", "models": [ { "id": "timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit" } ] } } }Run Pi
# Start Pi in your project directory: pi
- Hermes Agent
How to use timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit with Hermes Agent:
Start the MLX server
# Install MLX LM: uv tool install mlx-lm # Start a local OpenAI-compatible server: mlx_lm.server --model "timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit"
Configure Hermes
# Install Hermes: curl -fsSL https://hermes-agent.nousresearch.com/install.sh | bash hermes setup # Point Hermes at the local server: hermes config set model.provider custom hermes config set model.base_url http://127.0.0.1:8080/v1 hermes config set model.default timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit
Run Hermes
hermes
- Atomic Chat
- OpenClaw
How to use timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit with OpenClaw:
Start the MLX server
# Install MLX LM: uv tool install mlx-lm # Start a local OpenAI-compatible server: mlx_lm.server --model "timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit"
Configure OpenClaw
# Install OpenClaw: npm install -g openclaw@latest # Register the local server and set it as the default model: openclaw onboard --non-interactive --mode local \ --auth-choice custom-api-key \ --custom-base-url http://127.0.0.1:8080/v1 \ --custom-model-id "timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit" \ --custom-provider-id mlx-lm \ --custom-compatibility openai \ --custom-text-input \ --accept-risk \ --skip-health
Run OpenClaw
openclaw agent --local --agent main --message "Hello from Hugging Face"
Download benchmark-results.json from timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit: direct link, hf CLI and curl.
- Browser
- Download file 6.09 kB
-
https://huggingface.co/timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit/resolve/main/benchmark-results.json
- Command line
-
hf download hf://timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit/benchmark-results.json
-
curl -L -o benchmark-results.json https://huggingface.co/timteh673/Qwen3.8-27B-Opus-Abliterix-Reasoning-MLX-8bit/resolve/main/benchmark-results.json
6.09 kB
| { | |
| "code_termination_diagnostic": { | |
| "among_outputs_reaching_execution": { | |
| "control": { | |
| "executed": 103, | |
| "pass_percent": 15.53, | |
| "passed": 16 | |
| }, | |
| "winner": { | |
| "executed": 46, | |
| "pass_percent": 21.74, | |
| "passed": 10 | |
| } | |
| }, | |
| "interpretation": "The full-code result is dominated by severe termination/extraction pathology and is not a clean latent-code estimate.", | |
| "winner_generations_hitting_512_token_cap": 376, | |
| "winner_generations_total": 421 | |
| }, | |
| "comparison": { | |
| "benign_kl": { | |
| "control": 0.0, | |
| "strict_limit": 0.05, | |
| "winner": 0.093614 | |
| }, | |
| "capability_macro_percent": { | |
| "control": 17.6859, | |
| "winner": 21.0086 | |
| }, | |
| "exact_prompt_echoes": { | |
| "control": 2, | |
| "leakage_detected": true, | |
| "winner": 3 | |
| }, | |
| "full_code_passes": { | |
| "control": 16, | |
| "denominator": 421, | |
| "winner": 10 | |
| }, | |
| "harmful_hard_refusal_percent": { | |
| "control": 43.2, | |
| "winner": 0.0 | |
| }, | |
| "harmful_soft_deflection_percent": { | |
| "control": 14.6, | |
| "winner": 0.2 | |
| }, | |
| "harmful_substantive_response_percent": { | |
| "control": 47.0, | |
| "winner": 99.4 | |
| }, | |
| "held_out_loss_ratio": { | |
| "control": 1.0, | |
| "limit": 1.05, | |
| "winner": 1.024478 | |
| }, | |
| "human_eval_percent": { | |
| "control": 7.9268, | |
| "winner": 4.2683 | |
| }, | |
| "incoherence_percent": { | |
| "control": 2.7692, | |
| "strict_winner_limit": 2.7692, | |
| "winner": 4.3077 | |
| }, | |
| "long_form_pass_percent": { | |
| "control": 54.1667, | |
| "winner": 62.5 | |
| }, | |
| "max_repeated_4gram_fraction_percent": { | |
| "strict_limit": 5.0, | |
| "winner": 5.8632 | |
| }, | |
| "mmmu30_correct": { | |
| "control": 9, | |
| "denominator": 30, | |
| "winner": 11 | |
| } | |
| }, | |
| "evaluation_provenance": { | |
| "kind": "self-run_frozen_local_benchmarks", | |
| "note": "Results were produced by this project on a frozen local evaluation suite. They must not be mixed with or represented as official Qwen benchmark results.", | |
| "official_qwen_benchmarks": false, | |
| "release_report_frozen_at": "2026-08-22T00:01:48.742409+00:00" | |
| }, | |
| "lineage": { | |
| "abliterix_version": "1.12.2", | |
| "base_model": "Qwen/Qwen3.8-27B", | |
| "method": "Reasoning QLoRA merge followed by one selected Abliterix pass for the winner; control received no Abliterix residual-writer edits.", | |
| "reasoning_control": "control-bf16", | |
| "seed": 42, | |
| "winner": "abliterix-pass1-bf16" | |
| }, | |
| "native_mlx_proof": { | |
| "control_mtp_statistics": { | |
| "accepted_drafts_per_round": 1.73, | |
| "accepted_tokens_per_round": 2.73, | |
| "average_draft": 2.0, | |
| "drafted_percent": 86.5, | |
| "rounds": 37 | |
| }, | |
| "mtp_exact_answer": { | |
| "control": "323", | |
| "winner": "323" | |
| }, | |
| "ordinary_exact_answer": { | |
| "control": "323", | |
| "winner": "323" | |
| }, | |
| "runtime": "mlx-vlm on arm64 macOS/Metal", | |
| "speed_note": "MTP speed is workload-dependent; positive acceptance and exact-answer agreement are the correctness evidence.", | |
| "winner_mtp_statistics": { | |
| "accepted_drafts_per_round": 1.78, | |
| "accepted_tokens_per_round": 2.78, | |
| "average_draft": 2.0, | |
| "drafted_percent": 88.8, | |
| "rounds": 76 | |
| }, | |
| "winner_quantization": { | |
| "bits": 8, | |
| "group_size": 64, | |
| "scheme": "affine" | |
| } | |
| }, | |
| "release_family": "Qwen3.8-27B Opus Personal Model v1", | |
| "schema_version": 1, | |
| "selection": { | |
| "all_preregistered_strict_gates_passed": false, | |
| "control_variant": "control-bf16", | |
| "selected_variant": "abliterix-pass1-bf16", | |
| "status": "selected_practical_winner_with_measured_deviations", | |
| "summary": "Practical personal-model selection with measured deviations; not universal dominance." | |
| }, | |
| "strict_deviations": [ | |
| { | |
| "metric": "benign_kl", | |
| "passed": false, | |
| "strict_limit": 0.05, | |
| "winner": 0.093614 | |
| }, | |
| { | |
| "metric": "incoherence_percent", | |
| "passed": false, | |
| "strict_limit": 2.7692, | |
| "winner": 4.3077 | |
| }, | |
| { | |
| "control": 7.9268, | |
| "metric": "human_eval_percent", | |
| "passed": false, | |
| "strict_delta_floor": -3.0, | |
| "winner": 4.2683, | |
| "winner_delta_points": -3.6585 | |
| }, | |
| { | |
| "control": "16/421", | |
| "metric": "full_code_passes", | |
| "passed": false, | |
| "winner": "10/421" | |
| }, | |
| { | |
| "metric": "max_repeated_4gram_fraction_percent", | |
| "passed": false, | |
| "strict_limit": 5.0, | |
| "winner": 5.8632 | |
| }, | |
| { | |
| "control_exact_echoes": 2, | |
| "metric": "prompt_leakage", | |
| "passed": false, | |
| "winner_exact_echoes": 3 | |
| } | |
| ], | |
| "tensor_integrity": { | |
| "native_mtp_tensors": 15, | |
| "total_tensor_keys": 1199, | |
| "unexpected_changes": 0, | |
| "vision_tensors_preserved": 333, | |
| "winner_abliterix_residual_writer_edits": 74 | |
| }, | |
| "training": { | |
| "dataset_preparation": { | |
| "accepted_rows": 12614, | |
| "duplicates_removed": 208, | |
| "invalid_rows_removed": 20, | |
| "processed_manifest_sha256": "6e0a36ad20732c5f98ff592c4565a4c86876fead4e9f94bc6ceedfad1339a94d", | |
| "raw_rows": 12842, | |
| "rows_published": false, | |
| "source_rows": { | |
| "high-reasoning-250x": 250, | |
| "opus-10000x": 9633, | |
| "opus-3000x": 2326, | |
| "reasoning-700x": 633 | |
| }, | |
| "split_sha256": { | |
| "test.jsonl": "a2b6b46ce8b054166b0e392701e49390a44e391b56d25a0a25e9b8a9000257d9", | |
| "train.jsonl": "c7e691bcabd945e1259537f74119bce4c2c8f5fa4d2fba2e626b95bc04176ad3", | |
| "validation.jsonl": "589a05d0cf2fededff1febb065d6b3c8c947d29e948495c1567bcd9f5c44f364" | |
| }, | |
| "splits": { | |
| "test": 138, | |
| "train": 12349, | |
| "validation": 127 | |
| } | |
| }, | |
| "dtype_after_merge": "bfloat16", | |
| "final_validation_loss": 0.23739749, | |
| "hardware": "NVIDIA H200-class run; the sealed Abliterix environment recorded NVIDIA H200, CUDA 13.0, driver 580.159.03, Python 3.11.15.", | |
| "optimizer_steps": 1544, | |
| "rows": 12349, | |
| "token_accuracy_percent": 91.7594, | |
| "trainable_lora_parameters": 108789760 | |
| } | |
| } | |