-
Notifications
You must be signed in to change notification settings - Fork 97
Expand file tree
/
Copy pathvalues.yaml
More file actions
144 lines (120 loc) · 3.37 KB
/
Copy pathvalues.yaml
File metadata and controls
144 lines (120 loc) · 3.37 KB
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
76
77
78
79
80
81
82
83
84
85
86
87
88
89
90
91
92
93
94
95
96
97
98
99
100
101
102
103
104
105
106
107
108
109
110
111
112
113
114
115
116
117
118
119
120
121
122
123
124
125
126
127
128
129
130
131
132
133
134
135
136
137
138
139
140
141
142
143
144
# Copyright (C) 2024 Intel Corporation
# SPDX-License-Identifier: Apache-2.0
# Default values for chatqna.
# This is a YAML-formatted file.
# Declare variables to be passed into your templates.
replicaCount: 1
image:
repository: opea/chatqna
# Uncomment the following line to set desired image pull policy if needed, as one of Always, IfNotPresent, Never.
# pullPolicy: ""
# Overrides the image tag whose default is the chart appVersion.
tag: "latest"
imagePullSecrets: []
nameOverride: ""
fullnameOverride: ""
serviceAccount:
# Specifies whether a service account should be created
create: true
# Automatically mount a ServiceAccount's API credentials?
automount: true
# Annotations to add to the service account
annotations: {}
# The name of the service account to use.
# If not set and create is true, a name is generated using the fullname template
name: ""
podAnnotations: {}
podSecurityContext: {}
# fsGroup: 2000
securityContext:
readOnlyRootFilesystem: true
allowPrivilegeEscalation: false
runAsNonRoot: true
runAsUser: 1000
capabilities:
drop:
- ALL
seccompProfile:
type: RuntimeDefault
port: 8888
service:
type: ClusterIP
port: 8888
nodeSelector: {}
tolerations: []
affinity: {}
# This is just to avoid Helm errors when HPA is NOT used
# (use hpa-values.yaml files to actually enable HPA).
autoscaling:
enabled: false
# Optional subcharts enablement and subcharts settings overwritten
# LLM choice, tgi by default.
tgi:
enabled: false
LLM_MODEL_ID: meta-llama/Meta-Llama-3-8B-Instruct
vllm:
enabled: true
LLM_MODEL_ID: meta-llama/Meta-Llama-3-8B-Instruct
shmSize: 128Gi
VLLM_TORCH_PROFILER_DIR: "/tmp/vllm_profile"
data-prep:
DATAPREP_BACKEND: "REDIS"
INDEX_NAME: "rag-redis"
retriever-usvc:
RETRIEVER_BACKEND: "REDIS"
INDEX_NAME: "rag-redis"
# disable guardrails by default
# See guardrails-values.yaml for guardrail related options
guardrails-usvc:
enabled: false
tgi-guardrails:
enabled: false
# reranking
teirerank:
enabled: true
# vector db choice, redis by default.
redis-vector-db:
enabled: true
opensearch:
enabled: false
# Microservice layer, disabled by default
llm-uservice:
enabled: false
TEXTGEN_BACKEND: "vLLM"
LLM_MODEL_ID: meta-llama/Meta-Llama-3-8B-Instruct
embedding-usvc:
enabled: false
EMBEDDING_BACKEND: "TEI"
reranking-usvc:
enabled: false
RERANK_BACKEND: "TEI"
nginx:
service:
type: NodePort
# Uncomment the following lines
chatqna-ui:
image:
repository: opea/chatqna-ui
tag: "latest"
containerPort: "5173"
dashboard:
prefix: "OPEA ChatQnA"
global:
http_proxy: ""
https_proxy: ""
no_proxy: ""
HUGGINGFACEHUB_API_TOKEN: "insert-your-huggingface-token-here"
# service account name to be shared with all parent/child charts.
# If set, it will overwrite serviceAccount.name.
# If set, and serviceAccount.create is false, it will assume this service account is already created by others.
sharedSAName: "chatqna"
# set modelUseHostPath or modelUsePVC to use model cache.
modelUseHostPath: ""
# modelUseHostPath: /mnt/opea-models
# modelUsePVC: model-volume
# Prometheus monitoring + Grafana dashboard(s) for service components?
monitoring: false
# Prometheus/Grafana namespace for Dashboard installation
prometheusNamespace: monitoring
# Prometheus Helm install release name needed for serviceMonitors
prometheusRelease: prometheus-stack