Instructions to use ibm-granite/granitelib-rag-r1.0 with libraries, inference providers, notebooks, and local apps. Follow these links to get started.
- Libraries
- Granite Library
How to use ibm-granite/granitelib-rag-r1.0 with Granite Library:
# No code snippets available yet for this library. # To use this model, check the repository files and the library's documentation. # Want to help? PRs adding snippets are welcome at: # https://github.com/huggingface/huggingface.js
- Notebooks
- Google Colab
- Kaggle
feat: Add granite4:micro LoRA files.
Browse filesSigned-off-by: Khoi-Nguyen Tran <kndtran@ibm.com>
- .gitattributes +1 -0
- answer_relevance_classifier/granite4_micro/lora/Lora-q8_0.gguf +3 -0
- answer_relevance_classifier/granite4_micro/lora/Modelfile +2 -0
- answer_relevance_classifier/granite4_micro/lora/io.yaml +46 -0
- answer_relevance_rewriter/granite4_micro/lora/Lora-q8_0.gguf +3 -0
- answer_relevance_rewriter/granite4_micro/lora/Modelfile +2 -0
- answer_relevance_rewriter/granite4_micro/lora/io.yaml +28 -0
- answerability/granite4_micro/alora/io.yaml +21 -0
- answerability/granite4_micro/lora/Lora-q8_0.gguf +3 -0
- answerability/granite4_micro/lora/Modelfile +2 -0
- answerability/granite4_micro/lora/io.yaml +21 -0
- citations/granite4_micro/lora/Lora-q8_0.gguf +3 -0
- citations/granite4_micro/lora/Modelfile +2 -0
- citations/granite4_micro/lora/io.yaml +86 -0
- context_relevance/granite4_micro/lora/Lora-q8_0.gguf +3 -0
- context_relevance/granite4_micro/lora/Modelfile +2 -0
- context_relevance/granite4_micro/lora/io.yaml +29 -0
- hallucination_detection/granite4_micro/lora/Lora-q8_0.gguf +3 -0
- hallucination_detection/granite4_micro/lora/Modelfile +2 -0
- hallucination_detection/granite4_micro/lora/io.yaml +69 -0
- query_rewrite/granite4_micro/lora/Lora-q8_0.gguf +3 -0
- query_rewrite/granite4_micro/lora/Modelfile +2 -0
- query_rewrite/granite4_micro/lora/io.yaml +17 -0
.gitattributes
CHANGED
|
@@ -33,3 +33,4 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
|
|
| 33 |
*.zip filter=lfs diff=lfs merge=lfs -text
|
| 34 |
*.zst filter=lfs diff=lfs merge=lfs -text
|
| 35 |
*tfevents* filter=lfs diff=lfs merge=lfs -text
|
|
|
|
|
|
| 33 |
*.zip filter=lfs diff=lfs merge=lfs -text
|
| 34 |
*.zst filter=lfs diff=lfs merge=lfs -text
|
| 35 |
*tfevents* filter=lfs diff=lfs merge=lfs -text
|
| 36 |
+
*.gguf filter=lfs diff=lfs merge=lfs -text
|
answer_relevance_classifier/granite4_micro/lora/Lora-q8_0.gguf
ADDED
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:05ab02c60a86272867f2912a7e5be3938d66a2e1c21dd82dfb8101d9a3672a32
|
| 3 |
+
size 15335328
|
answer_relevance_classifier/granite4_micro/lora/Modelfile
ADDED
|
@@ -0,0 +1,2 @@
|
|
|
|
|
|
|
|
|
|
| 1 |
+
FROM granite4:micro
|
| 2 |
+
ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/answer_relevance_classifier/granite4_micro/lora/Lora-q8_0.gguf
|
answer_relevance_classifier/granite4_micro/lora/io.yaml
ADDED
|
@@ -0,0 +1,46 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
docs_as_message: roles
|
| 2 |
+
instruction: answer_relevance
|
| 3 |
+
model: null
|
| 4 |
+
parameters:
|
| 5 |
+
max_completion_tokens: 1024
|
| 6 |
+
response_format:
|
| 7 |
+
properties:
|
| 8 |
+
answer_relevance_analysis:
|
| 9 |
+
title: Answer Relevance Analysis
|
| 10 |
+
type: string
|
| 11 |
+
answer_relevance_category:
|
| 12 |
+
enum:
|
| 13 |
+
- Pertinent
|
| 14 |
+
- Pertinent with relevant extra
|
| 15 |
+
- Excessive unnecessary information
|
| 16 |
+
- Unduly restrictive
|
| 17 |
+
- Too vague or generic
|
| 18 |
+
- Contextual misalignment
|
| 19 |
+
- Misinterpreted inquiry
|
| 20 |
+
- No attempt
|
| 21 |
+
title: Answer Relevance Category
|
| 22 |
+
type: string
|
| 23 |
+
answer_relevance_judgment:
|
| 24 |
+
title: Answer Relevance Judgment
|
| 25 |
+
type: boolean
|
| 26 |
+
required:
|
| 27 |
+
- answer_relevance_analysis
|
| 28 |
+
- answer_relevance_category
|
| 29 |
+
- answer_relevance_judgment
|
| 30 |
+
title: AnswerRelevanceRawOutput
|
| 31 |
+
type: object
|
| 32 |
+
response_format: null
|
| 33 |
+
sentence_boundaries: null
|
| 34 |
+
transformations:
|
| 35 |
+
- categories_to_values:
|
| 36 |
+
false: 0.0
|
| 37 |
+
true: 1.0
|
| 38 |
+
input_path:
|
| 39 |
+
- answer_relevance_judgment
|
| 40 |
+
type: likelihood
|
| 41 |
+
- input_path: []
|
| 42 |
+
retained_fields:
|
| 43 |
+
answer_relevance_analysis: answer_relevance_analysis
|
| 44 |
+
answer_relevance_category: answer_relevance_category
|
| 45 |
+
answer_relevance_judgment: answer_relevance_likelihood
|
| 46 |
+
type: project
|
answer_relevance_rewriter/granite4_micro/lora/Lora-q8_0.gguf
ADDED
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:e42563c6d0af2650fbe62373614493e855bce43dc82e4862ce1d5a27f413818e
|
| 3 |
+
size 15335328
|
answer_relevance_rewriter/granite4_micro/lora/Modelfile
ADDED
|
@@ -0,0 +1,2 @@
|
|
|
|
|
|
|
|
|
|
| 1 |
+
FROM granite4:micro
|
| 2 |
+
ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/answer_relevance_rewriter/granite4_micro/lora/Lora-q8_0.gguf
|
answer_relevance_rewriter/granite4_micro/lora/io.yaml
ADDED
|
@@ -0,0 +1,28 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
docs_as_message: roles
|
| 2 |
+
instruction: "Rewrite the response for relevance.\nThe last assistant response is\
|
| 3 |
+
\ considered not fully relevant to the last user inquiry due to {answer_relevance_category}:\
|
| 4 |
+
\ {answer_relevance_analysis}\nDecide if you agree with this assessment, then act\
|
| 5 |
+
\ according to the following instructions: \nIf you disagree with the assessment,\
|
| 6 |
+
\ provide a verbatim copy of the original response. DO NOT attempt to correct any\
|
| 7 |
+
\ other perceived defects in the response. \nIf you agree with the assessment, provide\
|
| 8 |
+
\ an updated response that no longer fit the label {answer_relevance_category},\
|
| 9 |
+
\ by {correction_method}. Your response should be entirely based on the provided\
|
| 10 |
+
\ documents and should not rely on other prior knowledge. Your response should be\
|
| 11 |
+
\ suitable to be directly provided to the user. It should NOT contain meta information\
|
| 12 |
+
\ regarding the original response, its assessment, or this instruction. The user\
|
| 13 |
+
\ does not see any of these. Your response is the only response they will see to\
|
| 14 |
+
\ their inquiry, in place of the original response.\n"
|
| 15 |
+
model: null
|
| 16 |
+
parameters:
|
| 17 |
+
max_completion_tokens: 1024
|
| 18 |
+
response_format:
|
| 19 |
+
properties:
|
| 20 |
+
answer_relevance_rewrite:
|
| 21 |
+
title: Rewritten answer
|
| 22 |
+
type: string
|
| 23 |
+
required:
|
| 24 |
+
- answer_relevance_rewrite
|
| 25 |
+
type: object
|
| 26 |
+
response_format: null
|
| 27 |
+
sentence_boundaries: null
|
| 28 |
+
transformations: null
|
answerability/granite4_micro/alora/io.yaml
ADDED
|
@@ -0,0 +1,21 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
docs_as_message: roles
|
| 2 |
+
instruction: null
|
| 3 |
+
model: null
|
| 4 |
+
parameters:
|
| 5 |
+
max_completion_tokens: 6
|
| 6 |
+
response_format:
|
| 7 |
+
enum:
|
| 8 |
+
- answerable
|
| 9 |
+
- unanswerable
|
| 10 |
+
type: string
|
| 11 |
+
response_format: null
|
| 12 |
+
sentence_boundaries: null
|
| 13 |
+
transformations:
|
| 14 |
+
- categories_to_values:
|
| 15 |
+
answerable: 1.0
|
| 16 |
+
unanswerable: 0.0
|
| 17 |
+
input_path: []
|
| 18 |
+
type: likelihood
|
| 19 |
+
- field_name: answerability_likelihood
|
| 20 |
+
input_path: []
|
| 21 |
+
type: nest
|
answerability/granite4_micro/lora/Lora-q8_0.gguf
ADDED
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:88a80717c357a5165b249ba0350b00e09605597db74686430d3c924ca36c9184
|
| 3 |
+
size 15335776
|
answerability/granite4_micro/lora/Modelfile
ADDED
|
@@ -0,0 +1,2 @@
|
|
|
|
|
|
|
|
|
|
| 1 |
+
FROM granite4:micro
|
| 2 |
+
ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/answerability/granite4_micro/lora/Lora-q8_0.gguf
|
answerability/granite4_micro/lora/io.yaml
ADDED
|
@@ -0,0 +1,21 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
docs_as_message: roles
|
| 2 |
+
instruction: null
|
| 3 |
+
model: null
|
| 4 |
+
parameters:
|
| 5 |
+
max_completion_tokens: 6
|
| 6 |
+
response_format:
|
| 7 |
+
enum:
|
| 8 |
+
- answerable
|
| 9 |
+
- unanswerable
|
| 10 |
+
type: string
|
| 11 |
+
response_format: null
|
| 12 |
+
sentence_boundaries: null
|
| 13 |
+
transformations:
|
| 14 |
+
- categories_to_values:
|
| 15 |
+
answerable: 1.0
|
| 16 |
+
unanswerable: 0.0
|
| 17 |
+
input_path: []
|
| 18 |
+
type: likelihood
|
| 19 |
+
- field_name: answerability_likelihood
|
| 20 |
+
input_path: []
|
| 21 |
+
type: nest
|
citations/granite4_micro/lora/Lora-q8_0.gguf
ADDED
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:a1b8a549b23b204e4689e29cf957bd2cdc483e034d5494dfd94083e16324e3c4
|
| 3 |
+
size 14849568
|
citations/granite4_micro/lora/Modelfile
ADDED
|
@@ -0,0 +1,2 @@
|
|
|
|
|
|
|
|
|
|
| 1 |
+
FROM granite4:micro
|
| 2 |
+
ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/citations/granite4_micro/lora/Lora-q8_0.gguf
|
citations/granite4_micro/lora/io.yaml
ADDED
|
@@ -0,0 +1,86 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
docs_as_message: roles
|
| 2 |
+
instruction: 'Split the last assistant response into individual sentences. For each
|
| 3 |
+
sentence in the response, identify the statement IDs from the below documents that
|
| 4 |
+
it references. Ensure that your output includes all response sentence IDs, and
|
| 5 |
+
for each response sentence ID, provide the list of corresponding referring document
|
| 6 |
+
sentence IDs. The output must be a json structure.
|
| 7 |
+
|
| 8 |
+
'
|
| 9 |
+
model: null
|
| 10 |
+
parameters:
|
| 11 |
+
max_completion_tokens: 4096
|
| 12 |
+
response_format:
|
| 13 |
+
$defs:
|
| 14 |
+
_MODEL_OUTPUT_ENTRY:
|
| 15 |
+
properties:
|
| 16 |
+
c:
|
| 17 |
+
items:
|
| 18 |
+
minimum: 0
|
| 19 |
+
type: integer
|
| 20 |
+
title: C
|
| 21 |
+
type: array
|
| 22 |
+
r:
|
| 23 |
+
minimum: 0
|
| 24 |
+
title: R
|
| 25 |
+
type: integer
|
| 26 |
+
required:
|
| 27 |
+
- r
|
| 28 |
+
- c
|
| 29 |
+
title: _MODEL_OUTPUT_ENTRY
|
| 30 |
+
type: object
|
| 31 |
+
items:
|
| 32 |
+
$ref: '#/$defs/_MODEL_OUTPUT_ENTRY'
|
| 33 |
+
title: _MODEL_OUTPUT
|
| 34 |
+
type: array
|
| 35 |
+
response_format: null
|
| 36 |
+
sentence_boundaries:
|
| 37 |
+
documents: c
|
| 38 |
+
last_message: r
|
| 39 |
+
transformations:
|
| 40 |
+
- input_path: []
|
| 41 |
+
target_field: c
|
| 42 |
+
type: explode
|
| 43 |
+
- input_path: []
|
| 44 |
+
target_fields:
|
| 45 |
+
- r
|
| 46 |
+
- c
|
| 47 |
+
type: drop_duplicates
|
| 48 |
+
- input_path:
|
| 49 |
+
- null
|
| 50 |
+
- r
|
| 51 |
+
output_names:
|
| 52 |
+
begin: response_begin
|
| 53 |
+
end: response_end
|
| 54 |
+
text: response_text
|
| 55 |
+
source: last_message
|
| 56 |
+
type: decode_sentences
|
| 57 |
+
- input_path:
|
| 58 |
+
- null
|
| 59 |
+
- c
|
| 60 |
+
output_names:
|
| 61 |
+
begin: citation_begin
|
| 62 |
+
document_id: citation_doc_id
|
| 63 |
+
end: citation_end
|
| 64 |
+
text: citation_text
|
| 65 |
+
source: documents
|
| 66 |
+
type: decode_sentences
|
| 67 |
+
- input_path: []
|
| 68 |
+
retained_fields:
|
| 69 |
+
- response_begin
|
| 70 |
+
- response_end
|
| 71 |
+
- response_text
|
| 72 |
+
- citation_doc_id
|
| 73 |
+
- citation_begin
|
| 74 |
+
- citation_end
|
| 75 |
+
- citation_text
|
| 76 |
+
type: project
|
| 77 |
+
- begin_field: citation_begin
|
| 78 |
+
end_field: citation_end
|
| 79 |
+
group_fields:
|
| 80 |
+
- response_begin
|
| 81 |
+
- response_end
|
| 82 |
+
- response_text
|
| 83 |
+
- citation_doc_id
|
| 84 |
+
input_path: []
|
| 85 |
+
text_field: citation_text
|
| 86 |
+
type: merge_spans
|
context_relevance/granite4_micro/lora/Lora-q8_0.gguf
ADDED
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:ebafe44ff0e904b8dee4f909636f44631b65f912b17528983ef2e6cf1043bf5e
|
| 3 |
+
size 15335776
|
context_relevance/granite4_micro/lora/Modelfile
ADDED
|
@@ -0,0 +1,2 @@
|
|
|
|
|
|
|
|
|
|
| 1 |
+
FROM granite4:micro
|
| 2 |
+
ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/context_relevance/granite4_micro/lora/Lora-q8_0.gguf
|
context_relevance/granite4_micro/lora/io.yaml
ADDED
|
@@ -0,0 +1,29 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
docs_as_message: roles
|
| 2 |
+
instruction: 'DOCUMENT: {document_content}
|
| 3 |
+
|
| 4 |
+
'
|
| 5 |
+
model: null
|
| 6 |
+
parameters:
|
| 7 |
+
response_format:
|
| 8 |
+
properties:
|
| 9 |
+
context_relevance:
|
| 10 |
+
description: Context relevancy judgment.
|
| 11 |
+
enum:
|
| 12 |
+
- relevant
|
| 13 |
+
- irrelevant
|
| 14 |
+
- partially relevant
|
| 15 |
+
type: string
|
| 16 |
+
required:
|
| 17 |
+
- context_relevance
|
| 18 |
+
title: ContextRelevanceOutput
|
| 19 |
+
type: object
|
| 20 |
+
response_format: null
|
| 21 |
+
sentence_boundaries: null
|
| 22 |
+
transformations:
|
| 23 |
+
- categories_to_values:
|
| 24 |
+
irrelevant: 0.0
|
| 25 |
+
partially relevant: 0.5
|
| 26 |
+
relevant: 1.0
|
| 27 |
+
input_path:
|
| 28 |
+
- context_relevance
|
| 29 |
+
type: likelihood
|
hallucination_detection/granite4_micro/lora/Lora-q8_0.gguf
ADDED
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:04540ac6aff27b31fe3ab01e1e546c83c4605b35d9d83258a3984ca31d3cd1db
|
| 3 |
+
size 14849568
|
hallucination_detection/granite4_micro/lora/Modelfile
ADDED
|
@@ -0,0 +1,2 @@
|
|
|
|
|
|
|
|
|
|
| 1 |
+
FROM granite4:micro
|
| 2 |
+
ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/hallucination_detection/granite4_micro/lora/Lora-q8_0.gguf
|
hallucination_detection/granite4_micro/lora/io.yaml
ADDED
|
@@ -0,0 +1,69 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
docs_as_message: roles
|
| 2 |
+
instruction: 'Split the last assistant response into individual sentences. For each
|
| 3 |
+
sentence in the last assistant response, identify the faithfulness by comparing
|
| 4 |
+
with the provided documents and generate the faithfulness reasoning and faithfulness
|
| 5 |
+
decision. Ensure that your output includes all response sentence IDs, and for each
|
| 6 |
+
response sentence ID, provide the corresponding faithfulness reasoning and faithfulness
|
| 7 |
+
decision. The output must be a json structure.
|
| 8 |
+
|
| 9 |
+
'
|
| 10 |
+
model: null
|
| 11 |
+
parameters:
|
| 12 |
+
max_completion_tokens: 4096
|
| 13 |
+
response_format:
|
| 14 |
+
$defs:
|
| 15 |
+
HallucinationOutputEntry:
|
| 16 |
+
properties:
|
| 17 |
+
e:
|
| 18 |
+
title: Reasoning
|
| 19 |
+
type: string
|
| 20 |
+
f:
|
| 21 |
+
enum:
|
| 22 |
+
- faithful
|
| 23 |
+
- partial
|
| 24 |
+
- unfaithful
|
| 25 |
+
title: Is Faithful
|
| 26 |
+
type: string
|
| 27 |
+
r:
|
| 28 |
+
minimum: 0
|
| 29 |
+
title: Sentence Num
|
| 30 |
+
type: integer
|
| 31 |
+
required:
|
| 32 |
+
- r
|
| 33 |
+
- e
|
| 34 |
+
- f
|
| 35 |
+
title: HallucinationOutputEntry
|
| 36 |
+
type: object
|
| 37 |
+
items:
|
| 38 |
+
$ref: '#/$defs/HallucinationOutputEntry'
|
| 39 |
+
title: HallucinationOutput
|
| 40 |
+
type: array
|
| 41 |
+
response_format: null
|
| 42 |
+
sentence_boundaries:
|
| 43 |
+
last_message: i
|
| 44 |
+
transformations:
|
| 45 |
+
- categories_to_values:
|
| 46 |
+
faithful: 1.0
|
| 47 |
+
partial: 0.5
|
| 48 |
+
unfaithful: 0.0
|
| 49 |
+
input_path:
|
| 50 |
+
- null
|
| 51 |
+
- f
|
| 52 |
+
type: likelihood
|
| 53 |
+
- input_path:
|
| 54 |
+
- null
|
| 55 |
+
- r
|
| 56 |
+
output_names:
|
| 57 |
+
begin: response_begin
|
| 58 |
+
end: response_end
|
| 59 |
+
text: response_text
|
| 60 |
+
source: last_message
|
| 61 |
+
type: decode_sentences
|
| 62 |
+
- input_path: []
|
| 63 |
+
retained_fields:
|
| 64 |
+
e: explanation
|
| 65 |
+
f: faithfulness_likelihood
|
| 66 |
+
response_begin: response_begin
|
| 67 |
+
response_end: response_end
|
| 68 |
+
response_text: response_text
|
| 69 |
+
type: project
|
query_rewrite/granite4_micro/lora/Lora-q8_0.gguf
ADDED
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:d34e33d7849c07a323d94bc75aff2b0aa9dd333ddb4826647d5b6f46fd046b8c
|
| 3 |
+
size 15335776
|
query_rewrite/granite4_micro/lora/Modelfile
ADDED
|
@@ -0,0 +1,2 @@
|
|
|
|
|
|
|
|
|
|
| 1 |
+
FROM granite4:micro
|
| 2 |
+
ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/query_rewrite/granite4_micro/lora/Lora-q8_0.gguf
|
query_rewrite/granite4_micro/lora/io.yaml
ADDED
|
@@ -0,0 +1,17 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
docs_as_message: roles
|
| 2 |
+
instruction: null
|
| 3 |
+
model: null
|
| 4 |
+
parameters:
|
| 5 |
+
max_completion_tokens: 1024
|
| 6 |
+
response_format:
|
| 7 |
+
properties:
|
| 8 |
+
rewritten_question:
|
| 9 |
+
title: Rewritten Question
|
| 10 |
+
type: string
|
| 11 |
+
required:
|
| 12 |
+
- rewritten_question
|
| 13 |
+
title: QueryRewriteOutput
|
| 14 |
+
type: object
|
| 15 |
+
response_format: null
|
| 16 |
+
sentence_boundaries: false
|
| 17 |
+
transformations: null
|