Add logprobs workaround for harmony channel tokens
#3
by kndtran - opened
answerability/gpt-oss-20b/lora/io.yaml
CHANGED
|
@@ -20,8 +20,11 @@ transformations:
|
|
| 20 |
instruction: ~
|
| 21 |
parameters:
|
| 22 |
# "unanswerable" can be 6 tokens at high temperatures
|
| 23 |
-
|
|
|
|
| 24 |
# No sentence boundary detection
|
| 25 |
sentence_boundaries: ~
|
| 26 |
# RAG documents go in first message
|
| 27 |
docs_as_message: string
|
|
|
|
|
|
|
|
|
| 20 |
instruction: ~
|
| 21 |
parameters:
|
| 22 |
# "unanswerable" can be 6 tokens at high temperatures
|
| 23 |
+
# channel tokens add overhead on gpt-oss models
|
| 24 |
+
max_completion_tokens: 25
|
| 25 |
# No sentence boundary detection
|
| 26 |
sentence_boundaries: ~
|
| 27 |
# RAG documents go in first message
|
| 28 |
docs_as_message: string
|
| 29 |
+
# Use logprob tokens as ground truth for model output to handle harmony channel tokens
|
| 30 |
+
logprobs_workaround: true
|
citations/gpt-oss-20b/lora/io.yaml
CHANGED
|
@@ -95,3 +95,5 @@ sentence_boundaries:
|
|
| 95 |
documents: "c"
|
| 96 |
# gpt-oss base models have no "documents" argument
|
| 97 |
docs_as_message: json
|
|
|
|
|
|
|
|
|
| 95 |
documents: "c"
|
| 96 |
# gpt-oss base models have no "documents" argument
|
| 97 |
docs_as_message: json
|
| 98 |
+
# Use logprob tokens as ground truth for model output to handle harmony channel tokens
|
| 99 |
+
logprobs_workaround: true
|
hallucination_detection/gpt-oss-20b/lora/io.yaml
CHANGED
|
@@ -78,4 +78,6 @@ sentence_boundaries:
|
|
| 78 |
last_message: "i"
|
| 79 |
|
| 80 |
# gpt-oss base model has no "documents" argument
|
| 81 |
-
docs_as_message: json
|
|
|
|
|
|
|
|
|
| 78 |
last_message: "i"
|
| 79 |
|
| 80 |
# gpt-oss base model has no "documents" argument
|
| 81 |
+
docs_as_message: json
|
| 82 |
+
# Use logprob tokens as ground truth for model output to handle harmony channel tokens
|
| 83 |
+
logprobs_workaround: true
|
query_rewrite/gpt-oss-20b/lora/io.yaml
CHANGED
|
@@ -20,3 +20,5 @@ instruction: ~
|
|
| 20 |
parameters:
|
| 21 |
max_completion_tokens: 1024
|
| 22 |
sentence_boundaries: false
|
|
|
|
|
|
|
|
|
| 20 |
parameters:
|
| 21 |
max_completion_tokens: 1024
|
| 22 |
sentence_boundaries: false
|
| 23 |
+
# Use logprob tokens as ground truth for model output to handle harmony channel tokens
|
| 24 |
+
logprobs_workaround: true
|