mirror of
https://github.com/tiennm99/litellm.git
synced 2026-08-27 18:31:15 +00:00
(feat) working semantic cache on proxy
This commit is contained in:
committed by
Krrish Dholakia
parent
ef32a5da1b
commit
e9d00a92ed
@@ -73,10 +73,12 @@ litellm_settings:
|
||||
max_budget: 1.5000
|
||||
models: ["azure-gpt-3.5"]
|
||||
duration: None
|
||||
upperbound_key_generate_params:
|
||||
max_budget: 100
|
||||
duration: "30d"
|
||||
# cache: True
|
||||
cache: True # set cache responses to True
|
||||
cache_params:
|
||||
type: "redis-semantic"
|
||||
similarity_threshold: 0.8
|
||||
redis_semantic_cache_use_async: True
|
||||
# cache: True
|
||||
# setting callback class
|
||||
# callbacks: custom_callbacks.proxy_handler_instance # sets litellm.callbacks = [proxy_handler_instance]
|
||||
|
||||
|
||||
Reference in New Issue
Block a user