Granite Library
Safetensors
GGUF
English
kndtran commited on
Commit
1fc300c
·
1 Parent(s): 3459def

feat: Add granite4:micro LoRA files.

Browse files

Signed-off-by: Khoi-Nguyen Tran <kndtran@ibm.com>

.gitattributes CHANGED
@@ -33,3 +33,4 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
 
 
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
36
+ *.gguf filter=lfs diff=lfs merge=lfs -text
answer_relevance_classifier/granite4_micro/lora/Lora-q8_0.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:05ab02c60a86272867f2912a7e5be3938d66a2e1c21dd82dfb8101d9a3672a32
3
+ size 15335328
answer_relevance_classifier/granite4_micro/lora/Modelfile ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ FROM granite4:micro
2
+ ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/answer_relevance_classifier/granite4_micro/lora/Lora-q8_0.gguf
answer_relevance_classifier/granite4_micro/lora/io.yaml ADDED
@@ -0,0 +1,46 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ docs_as_message: roles
2
+ instruction: answer_relevance
3
+ model: null
4
+ parameters:
5
+ max_completion_tokens: 1024
6
+ response_format:
7
+ properties:
8
+ answer_relevance_analysis:
9
+ title: Answer Relevance Analysis
10
+ type: string
11
+ answer_relevance_category:
12
+ enum:
13
+ - Pertinent
14
+ - Pertinent with relevant extra
15
+ - Excessive unnecessary information
16
+ - Unduly restrictive
17
+ - Too vague or generic
18
+ - Contextual misalignment
19
+ - Misinterpreted inquiry
20
+ - No attempt
21
+ title: Answer Relevance Category
22
+ type: string
23
+ answer_relevance_judgment:
24
+ title: Answer Relevance Judgment
25
+ type: boolean
26
+ required:
27
+ - answer_relevance_analysis
28
+ - answer_relevance_category
29
+ - answer_relevance_judgment
30
+ title: AnswerRelevanceRawOutput
31
+ type: object
32
+ response_format: null
33
+ sentence_boundaries: null
34
+ transformations:
35
+ - categories_to_values:
36
+ false: 0.0
37
+ true: 1.0
38
+ input_path:
39
+ - answer_relevance_judgment
40
+ type: likelihood
41
+ - input_path: []
42
+ retained_fields:
43
+ answer_relevance_analysis: answer_relevance_analysis
44
+ answer_relevance_category: answer_relevance_category
45
+ answer_relevance_judgment: answer_relevance_likelihood
46
+ type: project
answer_relevance_rewriter/granite4_micro/lora/Lora-q8_0.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e42563c6d0af2650fbe62373614493e855bce43dc82e4862ce1d5a27f413818e
3
+ size 15335328
answer_relevance_rewriter/granite4_micro/lora/Modelfile ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ FROM granite4:micro
2
+ ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/answer_relevance_rewriter/granite4_micro/lora/Lora-q8_0.gguf
answer_relevance_rewriter/granite4_micro/lora/io.yaml ADDED
@@ -0,0 +1,28 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ docs_as_message: roles
2
+ instruction: "Rewrite the response for relevance.\nThe last assistant response is\
3
+ \ considered not fully relevant to the last user inquiry due to {answer_relevance_category}:\
4
+ \ {answer_relevance_analysis}\nDecide if you agree with this assessment, then act\
5
+ \ according to the following instructions: \nIf you disagree with the assessment,\
6
+ \ provide a verbatim copy of the original response. DO NOT attempt to correct any\
7
+ \ other perceived defects in the response. \nIf you agree with the assessment, provide\
8
+ \ an updated response that no longer fit the label {answer_relevance_category},\
9
+ \ by {correction_method}. Your response should be entirely based on the provided\
10
+ \ documents and should not rely on other prior knowledge. Your response should be\
11
+ \ suitable to be directly provided to the user. It should NOT contain meta information\
12
+ \ regarding the original response, its assessment, or this instruction. The user\
13
+ \ does not see any of these. Your response is the only response they will see to\
14
+ \ their inquiry, in place of the original response.\n"
15
+ model: null
16
+ parameters:
17
+ max_completion_tokens: 1024
18
+ response_format:
19
+ properties:
20
+ answer_relevance_rewrite:
21
+ title: Rewritten answer
22
+ type: string
23
+ required:
24
+ - answer_relevance_rewrite
25
+ type: object
26
+ response_format: null
27
+ sentence_boundaries: null
28
+ transformations: null
answerability/granite4_micro/alora/io.yaml ADDED
@@ -0,0 +1,21 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ docs_as_message: roles
2
+ instruction: null
3
+ model: null
4
+ parameters:
5
+ max_completion_tokens: 6
6
+ response_format:
7
+ enum:
8
+ - answerable
9
+ - unanswerable
10
+ type: string
11
+ response_format: null
12
+ sentence_boundaries: null
13
+ transformations:
14
+ - categories_to_values:
15
+ answerable: 1.0
16
+ unanswerable: 0.0
17
+ input_path: []
18
+ type: likelihood
19
+ - field_name: answerability_likelihood
20
+ input_path: []
21
+ type: nest
answerability/granite4_micro/lora/Lora-q8_0.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:88a80717c357a5165b249ba0350b00e09605597db74686430d3c924ca36c9184
3
+ size 15335776
answerability/granite4_micro/lora/Modelfile ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ FROM granite4:micro
2
+ ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/answerability/granite4_micro/lora/Lora-q8_0.gguf
answerability/granite4_micro/lora/io.yaml ADDED
@@ -0,0 +1,21 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ docs_as_message: roles
2
+ instruction: null
3
+ model: null
4
+ parameters:
5
+ max_completion_tokens: 6
6
+ response_format:
7
+ enum:
8
+ - answerable
9
+ - unanswerable
10
+ type: string
11
+ response_format: null
12
+ sentence_boundaries: null
13
+ transformations:
14
+ - categories_to_values:
15
+ answerable: 1.0
16
+ unanswerable: 0.0
17
+ input_path: []
18
+ type: likelihood
19
+ - field_name: answerability_likelihood
20
+ input_path: []
21
+ type: nest
citations/granite4_micro/lora/Lora-q8_0.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a1b8a549b23b204e4689e29cf957bd2cdc483e034d5494dfd94083e16324e3c4
3
+ size 14849568
citations/granite4_micro/lora/Modelfile ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ FROM granite4:micro
2
+ ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/citations/granite4_micro/lora/Lora-q8_0.gguf
citations/granite4_micro/lora/io.yaml ADDED
@@ -0,0 +1,86 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ docs_as_message: roles
2
+ instruction: 'Split the last assistant response into individual sentences. For each
3
+ sentence in the response, identify the statement IDs from the below documents that
4
+ it references. Ensure that your output includes all response sentence IDs, and
5
+ for each response sentence ID, provide the list of corresponding referring document
6
+ sentence IDs. The output must be a json structure.
7
+
8
+ '
9
+ model: null
10
+ parameters:
11
+ max_completion_tokens: 4096
12
+ response_format:
13
+ $defs:
14
+ _MODEL_OUTPUT_ENTRY:
15
+ properties:
16
+ c:
17
+ items:
18
+ minimum: 0
19
+ type: integer
20
+ title: C
21
+ type: array
22
+ r:
23
+ minimum: 0
24
+ title: R
25
+ type: integer
26
+ required:
27
+ - r
28
+ - c
29
+ title: _MODEL_OUTPUT_ENTRY
30
+ type: object
31
+ items:
32
+ $ref: '#/$defs/_MODEL_OUTPUT_ENTRY'
33
+ title: _MODEL_OUTPUT
34
+ type: array
35
+ response_format: null
36
+ sentence_boundaries:
37
+ documents: c
38
+ last_message: r
39
+ transformations:
40
+ - input_path: []
41
+ target_field: c
42
+ type: explode
43
+ - input_path: []
44
+ target_fields:
45
+ - r
46
+ - c
47
+ type: drop_duplicates
48
+ - input_path:
49
+ - null
50
+ - r
51
+ output_names:
52
+ begin: response_begin
53
+ end: response_end
54
+ text: response_text
55
+ source: last_message
56
+ type: decode_sentences
57
+ - input_path:
58
+ - null
59
+ - c
60
+ output_names:
61
+ begin: citation_begin
62
+ document_id: citation_doc_id
63
+ end: citation_end
64
+ text: citation_text
65
+ source: documents
66
+ type: decode_sentences
67
+ - input_path: []
68
+ retained_fields:
69
+ - response_begin
70
+ - response_end
71
+ - response_text
72
+ - citation_doc_id
73
+ - citation_begin
74
+ - citation_end
75
+ - citation_text
76
+ type: project
77
+ - begin_field: citation_begin
78
+ end_field: citation_end
79
+ group_fields:
80
+ - response_begin
81
+ - response_end
82
+ - response_text
83
+ - citation_doc_id
84
+ input_path: []
85
+ text_field: citation_text
86
+ type: merge_spans
context_relevance/granite4_micro/lora/Lora-q8_0.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:ebafe44ff0e904b8dee4f909636f44631b65f912b17528983ef2e6cf1043bf5e
3
+ size 15335776
context_relevance/granite4_micro/lora/Modelfile ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ FROM granite4:micro
2
+ ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/context_relevance/granite4_micro/lora/Lora-q8_0.gguf
context_relevance/granite4_micro/lora/io.yaml ADDED
@@ -0,0 +1,29 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ docs_as_message: roles
2
+ instruction: 'DOCUMENT: {document_content}
3
+
4
+ '
5
+ model: null
6
+ parameters:
7
+ response_format:
8
+ properties:
9
+ context_relevance:
10
+ description: Context relevancy judgment.
11
+ enum:
12
+ - relevant
13
+ - irrelevant
14
+ - partially relevant
15
+ type: string
16
+ required:
17
+ - context_relevance
18
+ title: ContextRelevanceOutput
19
+ type: object
20
+ response_format: null
21
+ sentence_boundaries: null
22
+ transformations:
23
+ - categories_to_values:
24
+ irrelevant: 0.0
25
+ partially relevant: 0.5
26
+ relevant: 1.0
27
+ input_path:
28
+ - context_relevance
29
+ type: likelihood
hallucination_detection/granite4_micro/lora/Lora-q8_0.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:04540ac6aff27b31fe3ab01e1e546c83c4605b35d9d83258a3984ca31d3cd1db
3
+ size 14849568
hallucination_detection/granite4_micro/lora/Modelfile ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ FROM granite4:micro
2
+ ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/hallucination_detection/granite4_micro/lora/Lora-q8_0.gguf
hallucination_detection/granite4_micro/lora/io.yaml ADDED
@@ -0,0 +1,69 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ docs_as_message: roles
2
+ instruction: 'Split the last assistant response into individual sentences. For each
3
+ sentence in the last assistant response, identify the faithfulness by comparing
4
+ with the provided documents and generate the faithfulness reasoning and faithfulness
5
+ decision. Ensure that your output includes all response sentence IDs, and for each
6
+ response sentence ID, provide the corresponding faithfulness reasoning and faithfulness
7
+ decision. The output must be a json structure.
8
+
9
+ '
10
+ model: null
11
+ parameters:
12
+ max_completion_tokens: 4096
13
+ response_format:
14
+ $defs:
15
+ HallucinationOutputEntry:
16
+ properties:
17
+ e:
18
+ title: Reasoning
19
+ type: string
20
+ f:
21
+ enum:
22
+ - faithful
23
+ - partial
24
+ - unfaithful
25
+ title: Is Faithful
26
+ type: string
27
+ r:
28
+ minimum: 0
29
+ title: Sentence Num
30
+ type: integer
31
+ required:
32
+ - r
33
+ - e
34
+ - f
35
+ title: HallucinationOutputEntry
36
+ type: object
37
+ items:
38
+ $ref: '#/$defs/HallucinationOutputEntry'
39
+ title: HallucinationOutput
40
+ type: array
41
+ response_format: null
42
+ sentence_boundaries:
43
+ last_message: i
44
+ transformations:
45
+ - categories_to_values:
46
+ faithful: 1.0
47
+ partial: 0.5
48
+ unfaithful: 0.0
49
+ input_path:
50
+ - null
51
+ - f
52
+ type: likelihood
53
+ - input_path:
54
+ - null
55
+ - r
56
+ output_names:
57
+ begin: response_begin
58
+ end: response_end
59
+ text: response_text
60
+ source: last_message
61
+ type: decode_sentences
62
+ - input_path: []
63
+ retained_fields:
64
+ e: explanation
65
+ f: faithfulness_likelihood
66
+ response_begin: response_begin
67
+ response_end: response_end
68
+ response_text: response_text
69
+ type: project
query_rewrite/granite4_micro/lora/Lora-q8_0.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d34e33d7849c07a323d94bc75aff2b0aa9dd333ddb4826647d5b6f46fd046b8c
3
+ size 15335776
query_rewrite/granite4_micro/lora/Modelfile ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ FROM granite4:micro
2
+ ADAPTER /proj/dmfexp/8cc/kndtran/huggingface/ibm-granite/granite-lib-rag-r1.0/query_rewrite/granite4_micro/lora/Lora-q8_0.gguf
query_rewrite/granite4_micro/lora/io.yaml ADDED
@@ -0,0 +1,17 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ docs_as_message: roles
2
+ instruction: null
3
+ model: null
4
+ parameters:
5
+ max_completion_tokens: 1024
6
+ response_format:
7
+ properties:
8
+ rewritten_question:
9
+ title: Rewritten Question
10
+ type: string
11
+ required:
12
+ - rewritten_question
13
+ title: QueryRewriteOutput
14
+ type: object
15
+ response_format: null
16
+ sentence_boundaries: false
17
+ transformations: null