(feat) working semantic cache on proxy

This commit is contained in:
ishaan-jaff
2024-02-06 13:30:59 -08:00
committed by Krrish Dholakia
parent ef32a5da1b
commit e9d00a92ed
+6 -4
View File
@@ -73,10 +73,12 @@ litellm_settings:
max_budget: 1.5000
models: ["azure-gpt-3.5"]
duration: None
upperbound_key_generate_params:
max_budget: 100
duration: "30d"
# cache: True
cache: True # set cache responses to True
cache_params:
type: "redis-semantic"
similarity_threshold: 0.8
redis_semantic_cache_use_async: True
# cache: True
# setting callback class
# callbacks: custom_callbacks.proxy_handler_instance # sets litellm.callbacks = [proxy_handler_instance]