Merge pull request #67 from sirius0xdev/feat/optimize-rtx6000-vllm-keda-scaling
fix(keda): simplify HTTPScaledObject to valid schema fields only
This commit is contained in:
commit
b9e8448956
2 changed files with 2 additions and 6 deletions
|
|
@ -5,15 +5,14 @@ metadata:
|
||||||
namespace: customer1
|
namespace: customer1
|
||||||
spec:
|
spec:
|
||||||
hosts:
|
hosts:
|
||||||
- a100-vllm.internal.cluster # Replace with actual internal routing host if needed
|
- a100-vllm.internal.cluster
|
||||||
scaleTargetRef:
|
scaleTargetRef:
|
||||||
name: openclaw-brain-vllm
|
name: openclaw-brain-vllm
|
||||||
kind: Deployment
|
kind: Deployment
|
||||||
apiVersion: apps/v1
|
apiVersion: apps/v1
|
||||||
service: openclaw-brain-service
|
service: openclaw-brain-service
|
||||||
port: 8000
|
port: 8000
|
||||||
|
|
||||||
replicas:
|
replicas:
|
||||||
min: 0
|
min: 0
|
||||||
max: 1
|
max: 1
|
||||||
scaledownPeriod: 900 # 15 minutes of idle time before scaling to zero
|
pollingInterval: 10
|
||||||
|
|
|
||||||
|
|
@ -12,10 +12,7 @@ spec:
|
||||||
apiVersion: apps/v1
|
apiVersion: apps/v1
|
||||||
service: rtx6000-brain-service
|
service: rtx6000-brain-service
|
||||||
port: 8000
|
port: 8000
|
||||||
|
|
||||||
replicas:
|
replicas:
|
||||||
min: 0
|
min: 0
|
||||||
max: 1
|
max: 1
|
||||||
pollingInterval: 10
|
pollingInterval: 10
|
||||||
cooldownPeriod: 30
|
|
||||||
scaledownPeriod: 900
|
|
||||||
|
|
|
||||||
Loading…
Add table
Reference in a new issue