Nanthasit commited on
Commit
96c9e37
·
verified ·
1 Parent(s): 7cfc347

Upload .eval_results/inference-check-20260730.yaml with huggingface_hub

Browse files
.eval_results/inference-check-20260730.yaml ADDED
@@ -0,0 +1,26 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ ---
2
+ eval_type: "inference-check"
3
+ model: "Nanthasit/sakthai-coder-1.5b"
4
+ timestamp: "2026-07-30T23:18:45+00:00"
5
+ prompt: "Write a Python function to check if a number is prime."
6
+ max_new_tokens: 100
7
+ status: "BLOCKED"
8
+ blockers:
9
+ - "DNS failure: api-inference.huggingface.co does not resolve in this environment"
10
+ - "Model is GGUF-only (no SafeTensors) — not loadable by standard HF Inference API"
11
+ - "HF Inference Providers API: model_not_supported — not supported by any enabled provider"
12
+ - "No llama.cpp available locally to run GGUF inference"
13
+ recommendation: "Convert model to SafeTensors format and push to same repo, or deploy as HF Inference Endpoint, or install llama.cpp locally for GGUF evaluation"
14
+ attempted_endpoints:
15
+ - "POST api-inference.huggingface.co/models/Nanthasit/sakthai-coder-1.5b → DNS failure (exit code 6)"
16
+ - "POST router.huggingface.co/hf/v1/chat/completions → 404 Not Found"
17
+ - "InferenceClient.chat_completion(model='Nanthasit/sakthai-coder-1.5b') → BadRequestError: model_not_supported"
18
+ - "InferenceClient.text_generation(model='Nanthasit/sakthai-coder-1.5b') → StopIteration (model not loadable)"
19
+ base_model_check:
20
+ model: "Qwen/Qwen2.5-Coder-1.5B-Instruct"
21
+ result: "model_not_supported by any enabled provider"
22
+ model_metadata:
23
+ pipeline_tag: "text-generation"
24
+ library: "transformers"
25
+ format: "GGUF"
26
+ base_model: "Qwen/Qwen2.5-Coder-1.5B-Instruct"