plano/demos/function_calling/arch_config.yaml

90 lines
2.2 KiB
YAML
Raw Normal View History

2024-09-30 17:49:05 -07:00
version: "0.1-beta"
2024-09-30 17:49:05 -07:00
listener:
address: 0.0.0.0
port: 10000
message_format: huggingface
connect_timeout: 0.005s
2024-09-30 17:49:05 -07:00
endpoints:
api_server:
endpoint: host.docker.internal:18083
2024-09-30 17:49:05 -07:00
connect_timeout: 0.005s
overrides:
# confidence threshold for prompt target intent matching
prompt_target_intent_matching_threshold: 0.6
2024-09-30 17:49:05 -07:00
llm_providers:
2024-10-09 15:47:32 -07:00
- name: gpt-4o
access_key: OPENAI_API_KEY
provider: openai
model: gpt-4o
default: true
2024-10-09 15:47:32 -07:00
- name: mistral-large-latest
access_key: MISTRAL_API_KEY
provider: mistral
model: mistral-large-latest
system_prompt: |
You are a helpful assistant.
2024-09-30 17:49:05 -07:00
prompt_targets:
2024-09-30 17:49:05 -07:00
- name: weather_forecast
2024-10-07 16:01:12 -07:00
description: Check weather information for a given city.
parameters:
- name: city
2024-10-07 16:01:12 -07:00
description: the name of the city
required: true
2024-10-07 16:01:12 -07:00
type: str
- name: days
2024-10-07 16:01:12 -07:00
description: the number of days
type: int
required: true
- name: units
2024-10-07 16:01:12 -07:00
description: the temperature unit, e.g., Celsius and Fahrenheit
type: str
default: Fahrenheit
endpoint:
2024-09-30 17:49:05 -07:00
name: api_server
path: /weather
2024-09-30 17:49:05 -07:00
- name: insurance_claim_details
2024-10-07 16:01:12 -07:00
description: Get the details of the insurance claim for a given policy number
parameters:
- name: policy_number
2024-10-07 16:01:12 -07:00
type: str
description: the policy number for the insurance claim
required: true
- name: include_expired
2024-10-07 16:01:12 -07:00
description: indicate whether to include expired insurance claims
2024-09-27 13:34:10 -07:00
type: bool
required: true
endpoint:
2024-09-30 17:49:05 -07:00
name: api_server
path: /insurance_claim_details
- name: default_target
default: true
description: This is the default target for all unmatched prompts.
endpoint:
name: api_server
path: /default_target
system_prompt: |
You are a helpful assistant. Use the information that is provided to you.
# if it is set to false arch will send response that it received from this prompt target to the user
# if true arch will forward the response to the default LLM
auto_llm_dispatch_on_response: true
2024-09-30 17:49:05 -07:00
ratelimits:
- model: gpt-4
2024-09-30 17:49:05 -07:00
selector:
key: selector-key
value: selector-value
limit:
tokens: 1
unit: minute
2024-10-08 16:24:08 -07:00
tracing:
random_sampling: 100