Update example K8s manifests (#40)

This commit is contained in:
Tyler Gillson 2023-04-20 09:31:11 -07:00 committed by GitHub
parent 1254951fab
commit c905512bb0
No known key found for this signature in database
GPG key ID: 4AEE18F83AFDEB23
2 changed files with 56 additions and 11 deletions

View file

@ -1,38 +1,55 @@
apiVersion: v1
kind: Namespace
metadata:
name: llama
name: local-ai
---
apiVersion: apps/v1
kind: Deployment
metadata:
name: llama
namespace: llama
name: local-ai
namespace: local-ai
labels:
app: llama
app: local-ai
spec:
selector:
matchLabels:
app: llama
app: local-ai
replicas: 1
template:
metadata:
labels:
app: llama
name: llama
app: local-ai
name: local-ai
spec:
containers:
- name: llama
- name: local-ai
image: quay.io/go-skynet/local-ai:latest
env:
- name: THREADS
value: "14"
- name: CONTEXT_SIZE
value: "512"
- name: MODELS_PATH
value: /models
volumeMounts:
- mountPath: /models
name: models
volumes:
- name: models
persistentVolumeClaim:
claimName: models
---
apiVersion: v1
kind: Service
metadata:
name: llama
namespace: llama
name: local-ai
namespace: local-ai
# If using AWS, you'll need to override the default 60s load balancer idle timeout
# annotations:
# service.beta.kubernetes.io/aws-load-balancer-connection-idle-timeout: "1200"
spec:
selector:
app: llama
app: local-ai
type: LoadBalancer
ports:
- protocol: TCP