diff --git a/.buildinfo b/.buildinfo
index d2102ed1..a21e115d 100755
--- a/.buildinfo
+++ b/.buildinfo
@@ -1,4 +1,4 @@
 # Sphinx build info version 1
 # This file records the configuration used when building these files. When it is not found, a full rebuild will be done.
-config: 9db6364d8186d5eaae6b4148a1288a59
+config: d54d7379a33d5f6b3c1acac04a881e38
 tags: 645f666f9bcd5a90fca523b33c5a78b7
diff --git a/CNAME b/CNAME
index 5a6f76d1..f7a19b2f 100644
--- a/CNAME
+++ b/CNAME
@@ -1 +1 @@
-docs.planoai.dev
\ No newline at end of file
+docs.planoai.dev
diff --git a/_downloads/ca9d3b7116524473d8adbde7cf15d167/arch_config_full_reference.yaml b/_downloads/ca9d3b7116524473d8adbde7cf15d167/arch_config_full_reference.yaml
index c9d5e4ff..aa186c26 100755
--- a/_downloads/ca9d3b7116524473d8adbde7cf15d167/arch_config_full_reference.yaml
+++ b/_downloads/ca9d3b7116524473d8adbde7cf15d167/arch_config_full_reference.yaml
@@ -1,100 +1,110 @@
-version: v0.1
 
+# Arch Gateway configuration version
+version: v0.3.0
+
+
+# External HTTP agents - API type is controlled by request path (/v1/responses, /v1/messages, /v1/chat/completions)
+agents:
+  - id: weather_agent  # Example agent for weather
+    url: http://host.docker.internal:10510
+
+  - id: flight_agent   # Example agent for flights
+    url: http://host.docker.internal:10520
+
+
+# MCP filters applied to requests/responses (e.g., input validation, query rewriting)
+filters:
+  - id: input_guards  # Example filter for input validation
+    url: http://host.docker.internal:10500
+    # type: mcp (default)
+    # transport: streamable-http (default)
+    # tool: input_guards (default - same as filter id)
+
+
+# LLM provider configurations with API keys and model routing
+model_providers:
+  - model: openai/gpt-4o
+    access_key: $OPENAI_API_KEY
+    default: true
+
+  - model: openai/gpt-4o-mini
+    access_key: $OPENAI_API_KEY
+
+  - model: anthropic/claude-sonnet-4-0
+    access_key: $ANTHROPIC_API_KEY
+
+  - model: mistral/ministral-3b-latest
+    access_key: $MISTRAL_API_KEY
+
+
+# Model aliases - use friendly names instead of full provider model names
+model_aliases:
+  fast-llm:
+    target: gpt-4o-mini
+
+  smart-llm:
+    target: gpt-4o
+
+
+# HTTP listeners - entry points for agent routing, prompt targets, and direct LLM access
 listeners:
-  ingress_traffic:
+  # Agent listener for routing requests to multiple agents
+  - type: agent
+    name: travel_booking_service
+    port: 8001
+    router: plano_orchestrator_v1
     address: 0.0.0.0
-    port: 10000
-    message_format: openai
-    timeout: 5s
-  egress_traffic:
+    agents:
+      - id: rag_agent
+        description: virtual assistant for retrieval augmented generation tasks
+        filter_chain:
+          - input_guards
+
+  # Model listener for direct LLM access
+  - type: model
+    name: model_1
     address: 0.0.0.0
     port: 12000
-    message_format: openai
-    timeout: 5s
 
-# Arch creates a round-robin load balancing between different endpoints, managed via the cluster subsystem.
+  # Prompt listener for function calling (for prompt_targets)
+  - type: prompt
+    name: prompt_function_listener
+    address: 0.0.0.0
+    port: 10000
+    # This listener is used for prompt_targets and function calling
+
+
+# Reusable service endpoints
 endpoints:
   app_server:
-    # value could be ip address or a hostname with port
-    # this could also be a list of endpoints for load balancing
-    # for example endpoint: [ ip1:port, ip2:port ]
     endpoint: 127.0.0.1:80
-    # max time to wait for a connection to be established
     connect_timeout: 0.005s
 
   mistral_local:
     endpoint: 127.0.0.1:8001
 
-  error_target:
-    endpoint: error_target_1
-
-# Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way
-llm_providers:
-  - name: openai/gpt-4o
-    access_key: $OPENAI_API_KEY
-    model: openai/gpt-4o
-    default: true
-
-  - access_key: $MISTRAL_API_KEY
-    model: mistral/mistral-8x7b
-
-  - model: mistral/mistral-7b-instruct
-    base_url: http://mistral_local
-
-# Model aliases - friendly names that map to actual provider names
-model_aliases:
-  # Alias for summarization tasks -> fast/cheap model
-  arch.summarize.v1:
-    target: gpt-4o
-
-  # Alias for general purpose tasks -> latest model
-  arch.v1:
-    target: mistral-8x7b
-
-# provides a way to override default settings for the arch system
-overrides:
-  # By default Arch uses an NLI + embedding approach to match an incoming prompt to a prompt target.
-  # The intent matching threshold is kept at 0.80, you can override this behavior if you would like
-  prompt_target_intent_matching_threshold: 0.60
-
-# default system prompt used by all prompt targets
-system_prompt: You are a network assistant that just offers facts; not advice on manufacturers or purchasing decisions.
-
-prompt_guards:
-  input_guards:
-    jailbreak:
-      on_exception:
-        message: Looks like you're curious about my abilities, but I can only provide assistance within my programmed parameters.
 
+# Prompt targets for function calling and API orchestration
 prompt_targets:
-  - name: information_extraction
-    default: true
-    description: handel all scenarios that are question and answer in nature. Like summarization, information extraction, etc.
-    endpoint:
-      name: app_server
-      path: /agent/summary
-      http_method: POST
-    # Arch uses the default LLM and treats the response from the endpoint as the prompt to send to the LLM
-    auto_llm_dispatch_on_response: true
-    # override system prompt for this prompt target
-    system_prompt: You are a helpful information extraction assistant. Use the information that is provided to you.
-
-  - name: reboot_network_device
-    description: Reboot a specific network device
-    endpoint:
-      name: app_server
-      path: /agent/action
+  - name: get_current_weather
+    description: Get current weather at a location.
     parameters:
-      - name: device_id
-        type: str
-        description: Identifier of the network device to reboot.
+      - name: location
+        description: The location to get the weather for
         required: true
-      - name: confirmation
-        type: bool
-        description: Confirmation flag to proceed with reboot.
-        default: false
-        enum: [true, false]
+        type: string
+        format: City, State
+      - name: days
+        description: the number of days for the request
+        required: true
+        type: int
+    endpoint:
+      name: app_server
+      path: /weather
+      http_method: POST
 
+
+# OpenTelemetry tracing configuration
 tracing:
-  # sampling rate. Note by default Arch works on OpenTelemetry compatible tracing.
-  sampling_rate: 0.1
+  # Random sampling percentage (1-100)
+  random_sampling: 100
diff --git a/_images/PlanoTagline.svg b/_images/PlanoTagline.svg
new file mode 100755
index 00000000..c0c10548
--- /dev/null
+++ b/_images/PlanoTagline.svg
@@ -0,0 +1,56 @@
+<svg width="1769" height="693" viewBox="0 0 1769 693" fill="none" xmlns="http://www.w3.org/2000/svg">
+<rect width="1769" height="693" fill="#2A2E4F" style="fill:#2A2E4F;fill:color(display-p3 0.1647 0.1804 0.3098);fill-opacity:1;"/>
+<path d="M331.711 418.187V271.429H356.238V292.538H361.867L356.238 298.368C356.238 289.388 358.919 282.352 364.28 277.259C369.641 272.032 376.878 269.419 385.992 269.419C397.116 269.419 406.028 273.238 412.73 280.878C419.431 288.517 422.782 298.77 422.782 311.637V341.591C422.782 350.169 421.24 357.674 418.158 364.108C415.209 370.407 410.987 375.299 405.492 378.783C399.997 382.268 393.497 384.01 385.992 384.01C376.878 384.01 369.641 381.464 364.28 376.371C358.919 371.144 356.238 364.041 356.238 355.061L361.867 360.891H356.037L356.841 387.227V418.187H331.711ZM377.146 362.298C383.579 362.298 388.605 360.422 392.224 356.669C395.843 352.916 397.652 347.555 397.652 340.586V312.843C397.652 305.874 395.843 300.513 392.224 296.76C388.605 293.007 383.579 291.131 377.146 291.131C370.847 291.131 365.888 293.074 362.269 296.961C358.65 300.714 356.841 306.008 356.841 312.843V340.586C356.841 347.421 358.65 352.782 362.269 356.669C365.888 360.422 370.847 362.298 377.146 362.298ZM497.891 382C490.654 382 484.288 380.526 478.793 377.577C473.432 374.629 469.21 370.474 466.127 365.113C463.045 359.618 461.503 353.319 461.503 346.215V257.959H426.121V235.242H486.633V346.215C486.633 350.236 487.772 353.453 490.051 355.865C492.463 358.143 495.68 359.283 499.701 359.283H533.073V382H497.891ZM573.202 384.01C561.81 384.01 552.83 380.995 546.263 374.964C539.696 368.933 536.412 360.824 536.412 350.638C536.412 339.782 540.031 331.405 547.268 325.508C554.506 319.611 564.759 316.663 578.027 316.663H605.569V307.214C605.569 301.853 603.827 297.698 600.342 294.749C596.858 291.667 592.1 290.126 586.069 290.126C580.574 290.126 576.017 291.332 572.398 293.744C568.779 296.157 566.635 299.44 565.965 303.595H541.438C542.644 293.141 547.335 284.832 555.511 278.666C563.686 272.501 574.14 269.419 586.873 269.419C600.409 269.419 611.064 272.836 618.838 279.672C626.745 286.373 630.699 295.487 630.699 307.013V382H606.373V362.7H602.353L606.373 357.272C606.373 365.448 603.358 371.948 597.327 376.773C591.296 381.598 583.254 384.01 573.202 384.01ZM581.445 365.113C588.548 365.113 594.311 363.303 598.734 359.685C603.291 356.066 605.569 351.375 605.569 345.612V332.143H578.429C573.336 332.143 569.248 333.617 566.166 336.565C563.083 339.514 561.542 343.401 561.542 348.226C561.542 353.453 563.284 357.607 566.769 360.69C570.388 363.639 575.28 365.113 581.445 365.113ZM645.095 382V271.429H669.622V292.538H676.457L669.622 298.368C669.622 289.254 672.235 282.151 677.462 277.058C682.823 271.965 690.128 269.419 699.376 269.419C710.232 269.419 718.876 273.037 725.31 280.275C731.877 287.512 735.16 297.229 735.16 309.425V382H710.031V312.039C710.031 305.337 708.288 300.177 704.804 296.559C701.319 292.94 696.427 291.131 690.128 291.131C683.963 291.131 679.071 293.007 675.452 296.76C671.967 300.513 670.225 305.874 670.225 312.843V382H645.095ZM794.388 383.809C785.006 383.809 776.831 382.067 769.861 378.582C763.026 374.964 757.665 369.938 753.778 363.504C750.026 356.937 748.149 349.231 748.149 340.385V313.044C748.149 304.198 750.026 296.559 753.778 290.126C757.665 283.558 763.026 278.532 769.861 275.048C776.831 271.429 785.006 269.62 794.388 269.62C803.904 269.62 812.079 271.429 818.915 275.048C825.75 278.532 831.044 283.558 834.797 290.126C838.684 296.559 840.627 304.131 840.627 312.843V340.385C840.627 349.231 838.684 356.937 834.797 363.504C831.044 369.938 825.75 374.964 818.915 378.582C812.079 382.067 803.904 383.809 794.388 383.809ZM794.388 361.896C801.089 361.896 806.249 360.087 809.868 356.468C813.621 352.715 815.497 347.354 815.497 340.385V313.044C815.497 305.941 813.621 300.58 809.868 296.961C806.249 293.342 801.089 291.533 794.388 291.533C787.821 291.533 782.661 293.342 778.908 296.961C775.155 300.58 773.279 305.941 773.279 313.044V340.385C773.279 347.354 775.155 352.715 778.908 356.468C782.661 360.087 787.821 361.896 794.388 361.896Z" fill="#7780D9" style="fill:#7780D9;fill:color(display-p3 0.4667 0.5020 0.8510);fill-opacity:1;"/>
+<rect x="140" y="235.402" width="147.807" height="147.807" fill="#7780D9" style="fill:#7780D9;fill:color(display-p3 0.4667 0.5020 0.8510);fill-opacity:1;"/>
+<rect x="151.37" y="246.771" width="125.067" height="125.067" fill="#B9BFFF" style="fill:#B9BFFF;fill:color(display-p3 0.7246 0.7499 1.0000);fill-opacity:1;"/>
+<path d="M157.055 252.456H168.424V263.825H157.055V252.456Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M157.055 269.51H168.424V280.88H157.055V269.51Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M157.055 286.565H168.424V297.935H157.055V286.565Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M157.055 303.62H168.424V314.989H157.055V303.62Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M157.055 320.675H168.424V332.044H157.055V320.675Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M157.055 337.729H168.424V349.099H157.055V337.729Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M157.055 354.784H168.424V366.153H157.055V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M208.219 252.456H219.589V263.825H208.219V252.456Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M208.219 269.51H219.589V280.88H208.219V269.51Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M208.219 286.565H219.589V297.935H208.219V286.565Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M208.219 303.62H219.589V314.989H208.219V303.62Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M208.219 320.675H219.589V332.044H208.219V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M208.219 337.729H219.589V349.099H208.219V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M208.219 354.784H219.589V366.153H208.219V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M174.11 252.456H185.48V263.825H174.11V252.456Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M174.11 269.51H185.48V280.88H174.11V269.51Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M174.11 286.565H185.48V297.935H174.11V286.565Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M174.11 303.62H185.48V314.989H174.11V303.62Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M174.11 320.675H185.48V332.044H174.11V320.675Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M174.11 337.729H185.48V349.099H174.11V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M174.11 354.784H185.48V366.153H174.11V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 252.456H236.643V263.825H225.273V252.456Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M225.273 269.51H236.643V280.88H225.273V269.51Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M225.273 286.565H236.643V297.935H225.273V286.565Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 303.62H236.643V314.989H225.273V303.62Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 320.675H236.643V332.044H225.273V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 337.729H236.643V349.099H225.273V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 354.784H236.643V366.153H225.273V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M191.164 252.456H202.534V263.825H191.164V252.456Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M191.164 269.51H202.534V280.88H191.164V269.51Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M191.164 286.565H202.534V297.935H191.164V286.565Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M191.164 303.62H202.534V314.989H191.164V303.62Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M191.164 320.675H202.534V332.044H191.164V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M191.164 337.729H202.534V349.099H191.164V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M191.164 354.784H202.534V366.153H191.164V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 252.456H253.697V263.825H242.328V252.456Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M242.328 269.51H253.697V280.88H242.328V269.51Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 286.565H253.697V297.935H242.328V286.565Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 303.62H253.697V314.989H242.328V303.62Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 320.675H253.697V332.044H242.328V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 337.729H253.697V349.099H242.328V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 354.784H253.697V366.153H242.328V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 252.456H270.753V263.825H259.383V252.456Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 269.51H270.753V280.88H259.383V269.51Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 286.565H270.753V297.935H259.383V286.565Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 303.62H270.753V314.989H259.383V303.62Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 320.675H270.753V332.044H259.383V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 337.729H270.753V349.099H259.383V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 354.784H270.753V366.153H259.383V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M139.988 460.516H157.91C161.274 460.516 163.903 461.463 165.798 463.358C167.731 465.253 168.698 467.785 168.698 470.956C168.698 472.464 168.485 473.759 168.06 474.842C167.635 475.886 167.093 476.756 166.436 477.452C165.779 478.109 165.025 478.612 164.174 478.96C163.323 479.269 162.492 479.463 161.68 479.54V479.888C162.492 479.927 163.381 480.12 164.348 480.468C165.353 480.816 166.281 481.377 167.132 482.15C167.983 482.885 168.698 483.851 169.278 485.05C169.858 486.21 170.148 487.641 170.148 489.342C170.148 490.966 169.877 492.493 169.336 493.924C168.833 495.355 168.118 496.592 167.19 497.636C166.262 498.68 165.16 499.511 163.884 500.13C162.608 500.71 161.216 501 159.708 501H139.988V460.516ZM146.542 495.374H157.794C159.495 495.374 160.829 494.929 161.796 494.04C162.763 493.151 163.246 491.875 163.246 490.212V488.24C163.246 486.577 162.763 485.301 161.796 484.412C160.829 483.523 159.495 483.078 157.794 483.078H146.542V495.374ZM146.542 477.626H156.692C158.316 477.626 159.573 477.22 160.462 476.408C161.351 475.557 161.796 474.359 161.796 472.812V470.956C161.796 469.409 161.351 468.23 160.462 467.418C159.573 466.567 158.316 466.142 156.692 466.142H146.542V477.626ZM191.966 496.012H191.734C191.464 496.747 191.096 497.462 190.632 498.158C190.207 498.854 189.646 499.473 188.95 500.014C188.293 500.517 187.481 500.923 186.514 501.232C185.586 501.541 184.504 501.696 183.266 501.696C180.134 501.696 177.698 500.691 175.958 498.68C174.257 496.669 173.406 493.789 173.406 490.038V470.84H179.728V489.226C179.728 493.905 181.662 496.244 185.528 496.244C186.34 496.244 187.133 496.147 187.906 495.954C188.68 495.722 189.356 495.393 189.936 494.968C190.555 494.543 191.038 494.001 191.386 493.344C191.773 492.687 191.966 491.913 191.966 491.024V470.84H198.288V501H191.966V496.012ZM206.471 465.446C205.156 465.446 204.19 465.137 203.571 464.518C202.991 463.899 202.701 463.107 202.701 462.14V461.154C202.701 460.187 202.991 459.395 203.571 458.776C204.19 458.157 205.156 457.848 206.471 457.848C207.786 457.848 208.733 458.157 209.313 458.776C209.893 459.395 210.183 460.187 210.183 461.154V462.14C210.183 463.107 209.893 463.899 209.313 464.518C208.733 465.137 207.786 465.446 206.471 465.446ZM203.281 470.84H209.603V501H203.281V470.84ZM221.009 501C218.843 501 217.219 500.459 216.137 499.376C215.093 498.255 214.571 496.708 214.571 494.736V458.08H220.893V495.838H225.069V501H221.009ZM246.016 496.012H245.726C245.068 497.791 243.986 499.183 242.478 500.188C241.008 501.193 239.268 501.696 237.258 501.696C233.43 501.696 230.472 500.323 228.384 497.578C226.296 494.794 225.252 490.908 225.252 485.92C225.252 480.932 226.296 477.065 228.384 474.32C230.472 471.536 233.43 470.144 237.258 470.144C239.268 470.144 241.008 470.647 242.478 471.652C243.986 472.619 245.068 474.011 245.726 475.828H246.016V458.08H252.338V501H246.016V496.012ZM239.172 496.244C241.105 496.244 242.729 495.78 244.044 494.852C245.358 493.885 246.016 492.629 246.016 491.082V480.758C246.016 479.211 245.358 477.974 244.044 477.046C242.729 476.079 241.105 475.596 239.172 475.596C236.968 475.596 235.208 476.311 233.894 477.742C232.579 479.134 231.922 480.99 231.922 483.31V488.53C231.922 490.85 232.579 492.725 233.894 494.156C235.208 495.548 236.968 496.244 239.172 496.244ZM289.059 501C287.396 501 286.12 500.536 285.231 499.608C284.342 498.641 283.8 497.423 283.607 495.954H283.317C282.737 497.849 281.674 499.279 280.127 500.246C278.58 501.213 276.705 501.696 274.501 501.696C271.369 501.696 268.952 500.884 267.251 499.26C265.588 497.636 264.757 495.451 264.757 492.706C264.757 489.69 265.84 487.428 268.005 485.92C270.209 484.412 273.418 483.658 277.633 483.658H283.085V481.106C283.085 479.25 282.582 477.819 281.577 476.814C280.572 475.809 279.006 475.306 276.879 475.306C275.1 475.306 273.65 475.693 272.529 476.466C271.408 477.239 270.46 478.225 269.687 479.424L265.917 476.002C266.922 474.301 268.334 472.909 270.151 471.826C271.968 470.705 274.346 470.144 277.285 470.144C281.19 470.144 284.187 471.053 286.275 472.87C288.363 474.687 289.407 477.297 289.407 480.7V495.838H292.597V501H289.059ZM276.299 496.882C278.271 496.882 279.895 496.457 281.171 495.606C282.447 494.717 283.085 493.537 283.085 492.068V487.718H277.749C273.38 487.718 271.195 489.071 271.195 491.778V492.822C271.195 494.175 271.64 495.2 272.529 495.896C273.457 496.553 274.714 496.882 276.299 496.882ZM321.181 503.842C321.181 506.974 319.982 509.333 317.585 510.918C315.188 512.503 311.282 513.296 305.869 513.296C303.394 513.296 301.287 513.122 299.547 512.774C297.846 512.465 296.434 512.001 295.313 511.382C294.23 510.763 293.438 510.009 292.935 509.12C292.432 508.231 292.181 507.206 292.181 506.046C292.181 504.383 292.626 503.088 293.515 502.16C294.443 501.232 295.719 500.594 297.343 500.246V499.608C295.062 498.873 293.921 497.327 293.921 494.968C293.921 493.421 294.443 492.242 295.487 491.43C296.531 490.579 297.788 489.98 299.257 489.632V489.4C297.478 488.549 296.106 487.37 295.139 485.862C294.211 484.315 293.747 482.517 293.747 480.468C293.747 477.375 294.772 474.881 296.821 472.986C298.909 471.091 301.886 470.144 305.753 470.144C307.88 470.144 309.697 470.453 311.205 471.072V470.26C311.205 468.907 311.514 467.863 312.133 467.128C312.79 466.355 313.796 465.968 315.149 465.968H319.789V471.072H313.641V472.29C314.994 473.179 316.038 474.339 316.773 475.77C317.508 477.162 317.875 478.728 317.875 480.468C317.875 483.523 316.831 485.997 314.743 487.892C312.655 489.748 309.678 490.676 305.811 490.676C304.342 490.676 303.027 490.521 301.867 490.212C301.094 490.483 300.417 490.869 299.837 491.372C299.257 491.836 298.967 492.455 298.967 493.228C298.967 494.04 299.334 494.62 300.069 494.968C300.804 495.316 301.848 495.49 303.201 495.49H310.625C314.337 495.49 317.024 496.244 318.687 497.752C320.35 499.221 321.181 501.251 321.181 503.842ZM315.265 504.538C315.265 503.494 314.859 502.663 314.047 502.044C313.274 501.425 311.843 501.116 309.755 501.116H299.547C298.155 501.928 297.459 503.088 297.459 504.596C297.459 505.833 297.942 506.819 298.909 507.554C299.914 508.327 301.596 508.714 303.955 508.714H307.899C312.81 508.714 315.265 507.322 315.265 504.538ZM305.811 486.094C307.667 486.094 309.098 485.688 310.103 484.876C311.147 484.025 311.669 482.73 311.669 480.99V479.83C311.669 478.09 311.147 476.814 310.103 476.002C309.098 475.151 307.667 474.726 305.811 474.726C303.955 474.726 302.505 475.151 301.461 476.002C300.456 476.814 299.953 478.09 299.953 479.83V480.99C299.953 482.73 300.456 484.025 301.461 484.876C302.505 485.688 303.955 486.094 305.811 486.094ZM333.317 501.696C331.152 501.696 329.219 501.329 327.517 500.594C325.816 499.859 324.366 498.815 323.167 497.462C321.969 496.07 321.041 494.407 320.383 492.474C319.765 490.502 319.455 488.317 319.455 485.92C319.455 483.523 319.765 481.357 320.383 479.424C321.041 477.452 321.969 475.789 323.167 474.436C324.366 473.044 325.816 471.981 327.517 471.246C329.219 470.511 331.152 470.144 333.317 470.144C335.521 470.144 337.455 470.531 339.117 471.304C340.819 472.077 342.23 473.16 343.351 474.552C344.473 475.905 345.304 477.491 345.845 479.308C346.425 481.125 346.715 483.078 346.715 485.166V487.544H326.009V488.53C326.009 490.85 326.686 492.764 328.039 494.272C329.431 495.741 331.403 496.476 333.955 496.476C335.811 496.476 337.377 496.07 338.653 495.258C339.929 494.446 341.012 493.344 341.901 491.952L345.613 495.606C344.492 497.462 342.868 498.951 340.741 500.072C338.615 501.155 336.14 501.696 333.317 501.696ZM333.317 475.074C332.235 475.074 331.229 475.267 330.301 475.654C329.412 476.041 328.639 476.582 327.981 477.278C327.363 477.974 326.879 478.805 326.531 479.772C326.183 480.739 326.009 481.802 326.009 482.962V483.368H340.045V482.788C340.045 480.468 339.446 478.612 338.247 477.22C337.049 475.789 335.405 475.074 333.317 475.074ZM349.654 501V470.84H355.976V475.828H356.266C356.923 474.204 357.909 472.851 359.224 471.768C360.577 470.685 362.414 470.144 364.734 470.144C367.827 470.144 370.225 471.169 371.926 473.218C373.666 475.229 374.536 478.109 374.536 481.86V501H368.214V482.672C368.214 477.955 366.319 475.596 362.53 475.596C361.718 475.596 360.906 475.712 360.094 475.944C359.321 476.137 358.625 476.447 358.006 476.872C357.387 477.297 356.885 477.839 356.498 478.496C356.15 479.153 355.976 479.927 355.976 480.816V501H349.654ZM387.59 501C385.386 501 383.724 500.439 382.602 499.318C381.481 498.158 380.92 496.534 380.92 494.446V476.002H376.222V470.84H378.774C379.818 470.84 380.534 470.608 380.92 470.144C381.346 469.68 381.558 468.926 381.558 467.882V462.604H387.242V470.84H393.564V476.002H387.242V495.838H393.1V501H387.59ZM405.932 501.696C403.071 501.696 400.673 501.193 398.74 500.188C396.807 499.183 395.125 497.791 393.694 496.012L397.754 492.3C398.875 493.653 400.113 494.717 401.466 495.49C402.858 496.225 404.463 496.592 406.28 496.592C408.136 496.592 409.509 496.244 410.398 495.548C411.326 494.813 411.79 493.808 411.79 492.532C411.79 491.565 411.461 490.753 410.804 490.096C410.185 489.4 409.083 488.955 407.498 488.762L404.714 488.414C401.621 488.027 399.185 487.138 397.406 485.746C395.666 484.315 394.796 482.208 394.796 479.424C394.796 477.955 395.067 476.659 395.608 475.538C396.149 474.378 396.903 473.411 397.87 472.638C398.875 471.826 400.055 471.207 401.408 470.782C402.8 470.357 404.327 470.144 405.99 470.144C408.697 470.144 410.901 470.569 412.602 471.42C414.342 472.271 415.889 473.45 417.242 474.958L413.356 478.67C412.583 477.742 411.558 476.949 410.282 476.292C409.045 475.596 407.614 475.248 405.99 475.248C404.25 475.248 402.955 475.596 402.104 476.292C401.292 476.988 400.886 477.897 400.886 479.018C400.886 480.178 401.253 481.048 401.988 481.628C402.723 482.208 403.902 482.633 405.526 482.904L408.31 483.252C411.635 483.755 414.052 484.741 415.56 486.21C417.107 487.641 417.88 489.574 417.88 492.01C417.88 493.479 417.59 494.813 417.01 496.012C416.469 497.172 415.676 498.177 414.632 499.028C413.588 499.879 412.331 500.536 410.862 501C409.393 501.464 407.749 501.696 405.932 501.696ZM432.175 476.002H427.593V470.84H432.175V465.852C432.175 463.416 432.813 461.521 434.089 460.168C435.365 458.776 437.279 458.08 439.831 458.08H444.819V463.242H438.497V470.84H444.819V476.002H438.497V501H432.175V476.002ZM468.358 501C466.696 501 465.42 500.536 464.53 499.608C463.641 498.641 463.1 497.423 462.906 495.954H462.616C462.036 497.849 460.973 499.279 459.426 500.246C457.88 501.213 456.004 501.696 453.8 501.696C450.668 501.696 448.252 500.884 446.55 499.26C444.888 497.636 444.056 495.451 444.056 492.706C444.056 489.69 445.139 487.428 447.304 485.92C449.508 484.412 452.718 483.658 456.932 483.658H462.384V481.106C462.384 479.25 461.882 477.819 460.876 476.814C459.871 475.809 458.305 475.306 456.178 475.306C454.4 475.306 452.95 475.693 451.828 476.466C450.707 477.239 449.76 478.225 448.986 479.424L445.216 476.002C446.222 474.301 447.633 472.909 449.45 471.826C451.268 470.705 453.646 470.144 456.584 470.144C460.49 470.144 463.486 471.053 465.574 472.87C467.662 474.687 468.706 477.297 468.706 480.7V495.838H471.896V501H468.358ZM455.598 496.882C457.57 496.882 459.194 496.457 460.47 495.606C461.746 494.717 462.384 493.537 462.384 492.068V487.718H457.048C452.679 487.718 450.494 489.071 450.494 491.778V492.822C450.494 494.175 450.939 495.2 451.828 495.896C452.756 496.553 454.013 496.882 455.598 496.882ZM483.602 501.696C480.741 501.696 478.344 501.193 476.41 500.188C474.477 499.183 472.795 497.791 471.364 496.012L475.424 492.3C476.546 493.653 477.783 494.717 479.136 495.49C480.528 496.225 482.133 496.592 483.95 496.592C485.806 496.592 487.179 496.244 488.068 495.548C488.996 494.813 489.46 493.808 489.46 492.532C489.46 491.565 489.132 490.753 488.474 490.096C487.856 489.4 486.754 488.955 485.168 488.762L482.384 488.414C479.291 488.027 476.855 487.138 475.076 485.746C473.336 484.315 472.466 482.208 472.466 479.424C472.466 477.955 472.737 476.659 473.278 475.538C473.82 474.378 474.574 473.411 475.54 472.638C476.546 471.826 477.725 471.207 479.078 470.782C480.47 470.357 481.998 470.144 483.66 470.144C486.367 470.144 488.571 470.569 490.272 471.42C492.012 472.271 493.559 473.45 494.912 474.958L491.026 478.67C490.253 477.742 489.228 476.949 487.952 476.292C486.715 475.596 485.284 475.248 483.66 475.248C481.92 475.248 480.625 475.596 479.774 476.292C478.962 476.988 478.556 477.897 478.556 479.018C478.556 480.178 478.924 481.048 479.658 481.628C480.393 482.208 481.572 482.633 483.196 482.904L485.98 483.252C489.306 483.755 491.722 484.741 493.23 486.21C494.777 487.641 495.55 489.574 495.55 492.01C495.55 493.479 495.26 494.813 494.68 496.012C494.139 497.172 493.346 498.177 492.302 499.028C491.258 499.879 490.002 500.536 488.532 501C487.063 501.464 485.42 501.696 483.602 501.696ZM506.588 501C504.384 501 502.721 500.439 501.6 499.318C500.479 498.158 499.918 496.534 499.918 494.446V476.002H495.22V470.84H497.772C498.816 470.84 499.531 470.608 499.918 470.144C500.343 469.68 500.556 468.926 500.556 467.882V462.604H506.24V470.84H512.562V476.002H506.24V495.838H512.098V501H506.588ZM526.965 501.696C524.8 501.696 522.866 501.329 521.165 500.594C519.464 499.859 518.014 498.815 516.815 497.462C515.616 496.07 514.688 494.407 514.031 492.474C513.412 490.502 513.103 488.317 513.103 485.92C513.103 483.523 513.412 481.357 514.031 479.424C514.688 477.452 515.616 475.789 516.815 474.436C518.014 473.044 519.464 471.981 521.165 471.246C522.866 470.511 524.8 470.144 526.965 470.144C529.169 470.144 531.102 470.531 532.765 471.304C534.466 472.077 535.878 473.16 536.999 474.552C538.12 475.905 538.952 477.491 539.493 479.308C540.073 481.125 540.363 483.078 540.363 485.166V487.544H519.657V488.53C519.657 490.85 520.334 492.764 521.687 494.272C523.079 495.741 525.051 496.476 527.603 496.476C529.459 496.476 531.025 496.07 532.301 495.258C533.577 494.446 534.66 493.344 535.549 491.952L539.261 495.606C538.14 497.462 536.516 498.951 534.389 500.072C532.262 501.155 529.788 501.696 526.965 501.696ZM526.965 475.074C525.882 475.074 524.877 475.267 523.949 475.654C523.06 476.041 522.286 476.582 521.629 477.278C521.01 477.974 520.527 478.805 520.179 479.772C519.831 480.739 519.657 481.802 519.657 482.962V483.368H533.693V482.788C533.693 480.468 533.094 478.612 531.895 477.22C530.696 475.789 529.053 475.074 526.965 475.074ZM543.301 501V470.84H549.623V476.64H549.913C550.339 475.093 551.228 473.74 552.581 472.58C553.935 471.42 555.81 470.84 558.207 470.84H559.889V476.93H557.395C554.882 476.93 552.949 477.336 551.595 478.148C550.281 478.96 549.623 480.159 549.623 481.744V501H543.301ZM560.686 492.938C562.117 492.938 563.18 493.305 563.876 494.04C564.572 494.775 564.92 495.722 564.92 496.882V497.81C564.92 499.705 564.398 501.677 563.354 503.726C562.349 505.814 560.976 507.573 559.236 509.004H554.538C555.775 507.767 556.8 506.568 557.612 505.408C558.424 504.248 559.043 502.953 559.468 501.522C558.463 501.329 557.709 500.903 557.206 500.246C556.703 499.55 556.452 498.738 556.452 497.81V496.882C556.452 495.722 556.8 494.775 557.496 494.04C558.192 493.305 559.255 492.938 560.686 492.938ZM601.307 501C599.645 501 598.369 500.536 597.479 499.608C596.59 498.641 596.049 497.423 595.855 495.954H595.565C594.985 497.849 593.922 499.279 592.375 500.246C590.829 501.213 588.953 501.696 586.749 501.696C583.617 501.696 581.201 500.884 579.499 499.26C577.837 497.636 577.005 495.451 577.005 492.706C577.005 489.69 578.088 487.428 580.253 485.92C582.457 484.412 585.667 483.658 589.881 483.658H595.333V481.106C595.333 479.25 594.831 477.819 593.825 476.814C592.82 475.809 591.254 475.306 589.127 475.306C587.349 475.306 585.899 475.693 584.777 476.466C583.656 477.239 582.709 478.225 581.935 479.424L578.165 476.002C579.171 474.301 580.582 472.909 582.399 471.826C584.217 470.705 586.595 470.144 589.533 470.144C593.439 470.144 596.435 471.053 598.523 472.87C600.611 474.687 601.655 477.297 601.655 480.7V495.838H604.845V501H601.307ZM588.547 496.882C590.519 496.882 592.143 496.457 593.419 495.606C594.695 494.717 595.333 493.537 595.333 492.068V487.718H589.997C585.628 487.718 583.443 489.071 583.443 491.778V492.822C583.443 494.175 583.888 495.2 584.777 495.896C585.705 496.553 586.962 496.882 588.547 496.882ZM606.981 501V470.84H613.303V475.828H613.593C614.251 474.204 615.237 472.851 616.551 471.768C617.905 470.685 619.741 470.144 622.061 470.144C625.155 470.144 627.552 471.169 629.253 473.218C630.993 475.229 631.863 478.109 631.863 481.86V501H625.541V482.672C625.541 477.955 623.647 475.596 619.857 475.596C619.045 475.596 618.233 475.712 617.421 475.944C616.648 476.137 615.952 476.447 615.333 476.872C614.715 477.297 614.212 477.839 613.825 478.496C613.477 479.153 613.303 479.927 613.303 480.816V501H606.981ZM655.532 496.012H655.242C654.585 497.791 653.502 499.183 651.994 500.188C650.525 501.193 648.785 501.696 646.774 501.696C642.946 501.696 639.988 500.323 637.9 497.578C635.812 494.794 634.768 490.908 634.768 485.92C634.768 480.932 635.812 477.065 637.9 474.32C639.988 471.536 642.946 470.144 646.774 470.144C648.785 470.144 650.525 470.647 651.994 471.652C653.502 472.619 654.585 474.011 655.242 475.828H655.532V458.08H661.854V501H655.532V496.012ZM648.688 496.244C650.621 496.244 652.245 495.78 653.56 494.852C654.875 493.885 655.532 492.629 655.532 491.082V480.758C655.532 479.211 654.875 477.974 653.56 477.046C652.245 476.079 650.621 475.596 648.688 475.596C646.484 475.596 644.725 476.311 643.41 477.742C642.095 479.134 641.438 480.99 641.438 483.31V488.53C641.438 490.85 642.095 492.725 643.41 494.156C644.725 495.548 646.484 496.244 648.688 496.244ZM695.443 496.012H695.153C694.496 497.791 693.413 499.183 691.905 500.188C690.436 501.193 688.696 501.696 686.685 501.696C682.857 501.696 679.899 500.323 677.811 497.578C675.723 494.794 674.679 490.908 674.679 485.92C674.679 480.932 675.723 477.065 677.811 474.32C679.899 471.536 682.857 470.144 686.685 470.144C688.696 470.144 690.436 470.647 691.905 471.652C693.413 472.619 694.496 474.011 695.153 475.828H695.443V458.08H701.765V501H695.443V496.012ZM688.599 496.244C690.532 496.244 692.156 495.78 693.471 494.852C694.786 493.885 695.443 492.629 695.443 491.082V480.758C695.443 479.211 694.786 477.974 693.471 477.046C692.156 476.079 690.532 475.596 688.599 475.596C686.395 475.596 684.636 476.311 683.321 477.742C682.006 479.134 681.349 480.99 681.349 483.31V488.53C681.349 490.85 682.006 492.725 683.321 494.156C684.636 495.548 686.395 496.244 688.599 496.244ZM718.573 501.696C716.408 501.696 714.475 501.329 712.773 500.594C711.072 499.859 709.622 498.815 708.423 497.462C707.225 496.07 706.297 494.407 705.639 492.474C705.021 490.502 704.711 488.317 704.711 485.92C704.711 483.523 705.021 481.357 705.639 479.424C706.297 477.452 707.225 475.789 708.423 474.436C709.622 473.044 711.072 471.981 712.773 471.246C714.475 470.511 716.408 470.144 718.573 470.144C720.777 470.144 722.711 470.531 724.373 471.304C726.075 472.077 727.486 473.16 728.607 474.552C729.729 475.905 730.56 477.491 731.101 479.308C731.681 481.125 731.971 483.078 731.971 485.166V487.544H711.265V488.53C711.265 490.85 711.942 492.764 713.295 494.272C714.687 495.741 716.659 496.476 719.211 496.476C721.067 496.476 722.633 496.07 723.909 495.258C725.185 494.446 726.268 493.344 727.157 491.952L730.869 495.606C729.748 497.462 728.124 498.951 725.997 500.072C723.871 501.155 721.396 501.696 718.573 501.696ZM718.573 475.074C717.491 475.074 716.485 475.267 715.557 475.654C714.668 476.041 713.895 476.582 713.237 477.278C712.619 477.974 712.135 478.805 711.787 479.772C711.439 480.739 711.265 481.802 711.265 482.962V483.368H725.301V482.788C725.301 480.468 724.702 478.612 723.503 477.22C722.305 475.789 720.661 475.074 718.573 475.074ZM741.348 501C739.183 501 737.559 500.459 736.476 499.376C735.432 498.255 734.91 496.708 734.91 494.736V458.08H741.232V495.838H745.408V501H741.348ZM750.579 465.446C749.264 465.446 748.298 465.137 747.679 464.518C747.099 463.899 746.809 463.107 746.809 462.14V461.154C746.809 460.187 747.099 459.395 747.679 458.776C748.298 458.157 749.264 457.848 750.579 457.848C751.894 457.848 752.841 458.157 753.421 458.776C754.001 459.395 754.291 460.187 754.291 461.154V462.14C754.291 463.107 754.001 463.899 753.421 464.518C752.841 465.137 751.894 465.446 750.579 465.446ZM747.389 470.84H753.711V501H747.389V470.84ZM765.233 501L754.967 470.84H761.231L765.813 484.586L768.887 495.142H769.235L772.309 484.586L777.007 470.84H783.039L772.657 501H765.233ZM795.677 501.696C793.512 501.696 791.578 501.329 789.877 500.594C788.176 499.859 786.726 498.815 785.527 497.462C784.328 496.07 783.4 494.407 782.743 492.474C782.124 490.502 781.815 488.317 781.815 485.92C781.815 483.523 782.124 481.357 782.743 479.424C783.4 477.452 784.328 475.789 785.527 474.436C786.726 473.044 788.176 471.981 789.877 471.246C791.578 470.511 793.512 470.144 795.677 470.144C797.881 470.144 799.814 470.531 801.477 471.304C803.178 472.077 804.59 473.16 805.711 474.552C806.832 475.905 807.664 477.491 808.205 479.308C808.785 481.125 809.075 483.078 809.075 485.166V487.544H788.369V488.53C788.369 490.85 789.046 492.764 790.399 494.272C791.791 495.741 793.763 496.476 796.315 496.476C798.171 496.476 799.737 496.07 801.013 495.258C802.289 494.446 803.372 493.344 804.261 491.952L807.973 495.606C806.852 497.462 805.228 498.951 803.101 500.072C800.974 501.155 798.5 501.696 795.677 501.696ZM795.677 475.074C794.594 475.074 793.589 475.267 792.661 475.654C791.772 476.041 790.998 476.582 790.341 477.278C789.722 477.974 789.239 478.805 788.891 479.772C788.543 480.739 788.369 481.802 788.369 482.962V483.368H802.405V482.788C802.405 480.468 801.806 478.612 800.607 477.22C799.408 475.789 797.765 475.074 795.677 475.074ZM812.014 501V470.84H818.336V476.64H818.626C819.051 475.093 819.94 473.74 821.294 472.58C822.647 471.42 824.522 470.84 826.92 470.84H828.602V476.93H826.108C823.594 476.93 821.661 477.336 820.308 478.148C818.993 478.96 818.336 480.159 818.336 481.744V501H812.014ZM848.156 501C845.952 501 844.289 500.439 843.168 499.318C842.046 498.158 841.486 496.534 841.486 494.446V476.002H836.788V470.84H839.34C840.384 470.84 841.099 470.608 841.486 470.144C841.911 469.68 842.124 468.926 842.124 467.882V462.604H847.808V470.84H854.13V476.002H847.808V495.838H853.666V501H848.156ZM856.927 458.08H863.249V475.828H863.539C864.197 474.204 865.183 472.851 866.497 471.768C867.851 470.685 869.687 470.144 872.007 470.144C875.101 470.144 877.498 471.169 879.199 473.218C880.939 475.229 881.809 478.109 881.809 481.86V501H875.487V482.614C875.487 477.935 873.593 475.596 869.803 475.596C868.991 475.596 868.179 475.712 867.367 475.944C866.594 476.137 865.898 476.447 865.279 476.872C864.661 477.297 864.158 477.839 863.771 478.496C863.423 479.153 863.249 479.907 863.249 480.758V501H856.927V458.08ZM898.344 501.696C896.179 501.696 894.245 501.329 892.544 500.594C890.843 499.859 889.393 498.815 888.194 497.462C886.995 496.07 886.067 494.407 885.41 492.474C884.791 490.502 884.482 488.317 884.482 485.92C884.482 483.523 884.791 481.357 885.41 479.424C886.067 477.452 886.995 475.789 888.194 474.436C889.393 473.044 890.843 471.981 892.544 471.246C894.245 470.511 896.179 470.144 898.344 470.144C900.548 470.144 902.481 470.531 904.144 471.304C905.845 472.077 907.257 473.16 908.378 474.552C909.499 475.905 910.331 477.491 910.872 479.308C911.452 481.125 911.742 483.078 911.742 485.166V487.544H891.036V488.53C891.036 490.85 891.713 492.764 893.066 494.272C894.458 495.741 896.43 496.476 898.982 496.476C900.838 496.476 902.404 496.07 903.68 495.258C904.956 494.446 906.039 493.344 906.928 491.952L910.64 495.606C909.519 497.462 907.895 498.951 905.768 500.072C903.641 501.155 901.167 501.696 898.344 501.696ZM898.344 475.074C897.261 475.074 896.256 475.267 895.328 475.654C894.439 476.041 893.665 476.582 893.008 477.278C892.389 477.974 891.906 478.805 891.558 479.772C891.21 480.739 891.036 481.802 891.036 482.962V483.368H905.072V482.788C905.072 480.468 904.473 478.612 903.274 477.22C902.075 475.789 900.432 475.074 898.344 475.074ZM914.68 501V470.84H921.002V475.828H921.292C921.602 475.055 921.969 474.32 922.394 473.624C922.858 472.928 923.4 472.329 924.018 471.826C924.676 471.285 925.43 470.879 926.28 470.608C927.17 470.299 928.194 470.144 929.354 470.144C931.404 470.144 933.221 470.647 934.806 471.652C936.392 472.657 937.552 474.204 938.286 476.292H938.46C939.002 474.591 940.046 473.141 941.592 471.942C943.139 470.743 945.13 470.144 947.566 470.144C950.582 470.144 952.922 471.169 954.584 473.218C956.247 475.229 957.078 478.109 957.078 481.86V501H950.756V482.614C950.756 480.294 950.312 478.554 949.422 477.394C948.533 476.195 947.122 475.596 945.188 475.596C944.376 475.596 943.603 475.712 942.868 475.944C942.134 476.137 941.476 476.447 940.896 476.872C940.355 477.297 939.91 477.839 939.562 478.496C939.214 479.153 939.04 479.907 939.04 480.758V501H932.718V482.614C932.718 477.935 930.882 475.596 927.208 475.596C926.435 475.596 925.662 475.712 924.888 475.944C924.154 476.137 923.496 476.447 922.916 476.872C922.336 477.297 921.872 477.839 921.524 478.496C921.176 479.153 921.002 479.907 921.002 480.758V501H914.68ZM971.414 501V470.84H977.736V476.64H978.026C978.451 475.093 979.341 473.74 980.694 472.58C982.047 471.42 983.923 470.84 986.32 470.84H988.002V476.93H985.508C982.995 476.93 981.061 477.336 979.708 478.148C978.393 478.96 977.736 480.159 977.736 481.744V501H971.414ZM1001.11 501.696C998.94 501.696 997.007 501.329 995.306 500.594C993.604 499.859 992.154 498.815 990.956 497.462C989.757 496.07 988.829 494.407 988.172 492.474C987.553 490.502 987.244 488.317 987.244 485.92C987.244 483.523 987.553 481.357 988.172 479.424C988.829 477.452 989.757 475.789 990.956 474.436C992.154 473.044 993.604 471.981 995.306 471.246C997.007 470.511 998.94 470.144 1001.11 470.144C1003.31 470.144 1005.24 470.531 1006.91 471.304C1008.61 472.077 1010.02 473.16 1011.14 474.552C1012.26 475.905 1013.09 477.491 1013.63 479.308C1014.21 481.125 1014.5 483.078 1014.5 485.166V487.544H993.798V488.53C993.798 490.85 994.474 492.764 995.828 494.272C997.22 495.741 999.192 496.476 1001.74 496.476C1003.6 496.476 1005.17 496.07 1006.44 495.258C1007.72 494.446 1008.8 493.344 1009.69 491.952L1013.4 495.606C1012.28 497.462 1010.66 498.951 1008.53 500.072C1006.4 501.155 1003.93 501.696 1001.11 501.696ZM1001.11 475.074C1000.02 475.074 999.018 475.267 998.09 475.654C997.2 476.041 996.427 476.582 995.77 477.278C995.151 477.974 994.668 478.805 994.32 479.772C993.972 480.739 993.798 481.802 993.798 482.962V483.368H1007.83V482.788C1007.83 480.468 1007.23 478.612 1006.04 477.22C1004.84 475.789 1003.19 475.074 1001.11 475.074ZM1023.88 501C1021.72 501 1020.09 500.459 1019.01 499.376C1017.96 498.255 1017.44 496.708 1017.44 494.736V458.08H1023.76V495.838H1027.94V501H1023.88ZM1033.11 465.446C1031.8 465.446 1030.83 465.137 1030.21 464.518C1029.63 463.899 1029.34 463.107 1029.34 462.14V461.154C1029.34 460.187 1029.63 459.395 1030.21 458.776C1030.83 458.157 1031.8 457.848 1033.11 457.848C1034.43 457.848 1035.37 458.157 1035.95 458.776C1036.53 459.395 1036.82 460.187 1036.82 461.154V462.14C1036.82 463.107 1036.53 463.899 1035.95 464.518C1035.37 465.137 1034.43 465.446 1033.11 465.446ZM1029.92 470.84H1036.24V501H1029.92V470.84ZM1063.31 501C1061.65 501 1060.37 500.536 1059.48 499.608C1058.59 498.641 1058.05 497.423 1057.86 495.954H1057.57C1056.99 497.849 1055.92 499.279 1054.38 500.246C1052.83 501.213 1050.96 501.696 1048.75 501.696C1045.62 501.696 1043.2 500.884 1041.5 499.26C1039.84 497.636 1039.01 495.451 1039.01 492.706C1039.01 489.69 1040.09 487.428 1042.26 485.92C1044.46 484.412 1047.67 483.658 1051.88 483.658H1057.34V481.106C1057.34 479.25 1056.83 477.819 1055.83 476.814C1054.82 475.809 1053.26 475.306 1051.13 475.306C1049.35 475.306 1047.9 475.693 1046.78 476.466C1045.66 477.239 1044.71 478.225 1043.94 479.424L1040.17 476.002C1041.17 474.301 1042.58 472.909 1044.4 471.826C1046.22 470.705 1048.6 470.144 1051.54 470.144C1055.44 470.144 1058.44 471.053 1060.53 472.87C1062.61 474.687 1063.66 477.297 1063.66 480.7V495.838H1066.85V501H1063.31ZM1050.55 496.882C1052.52 496.882 1054.15 496.457 1055.42 495.606C1056.7 494.717 1057.34 493.537 1057.34 492.068V487.718H1052C1047.63 487.718 1045.45 489.071 1045.45 491.778V492.822C1045.45 494.175 1045.89 495.2 1046.78 495.896C1047.71 496.553 1048.96 496.882 1050.55 496.882ZM1068.98 458.08H1075.31V475.828H1075.6C1076.25 474.011 1077.32 472.619 1078.79 471.652C1080.29 470.647 1082.05 470.144 1084.06 470.144C1087.89 470.144 1090.85 471.536 1092.94 474.32C1095.03 477.065 1096.07 480.932 1096.07 485.92C1096.07 490.908 1095.03 494.794 1092.94 497.578C1090.85 500.323 1087.89 501.696 1084.06 501.696C1082.05 501.696 1080.29 501.193 1078.79 500.188C1077.32 499.183 1076.25 497.791 1075.6 496.012H1075.31V501H1068.98V458.08ZM1082.15 496.244C1084.35 496.244 1086.11 495.548 1087.43 494.156C1088.74 492.725 1089.4 490.85 1089.4 488.53V483.31C1089.4 480.99 1088.74 479.134 1087.43 477.742C1086.11 476.311 1084.35 475.596 1082.15 475.596C1080.22 475.596 1078.59 476.079 1077.28 477.046C1075.96 477.974 1075.31 479.211 1075.31 480.758V491.082C1075.31 492.629 1075.96 493.885 1077.28 494.852C1078.59 495.78 1080.22 496.244 1082.15 496.244ZM1105.69 501C1103.52 501 1101.9 500.459 1100.81 499.376C1099.77 498.255 1099.25 496.708 1099.25 494.736V458.08H1105.57V495.838H1109.75V501H1105.69ZM1129.78 470.84H1135.87L1123.16 506.974C1122.82 507.979 1122.41 508.83 1121.95 509.526C1121.52 510.261 1121 510.841 1120.38 511.266C1119.8 511.73 1119.08 512.059 1118.23 512.252C1117.38 512.484 1116.38 512.6 1115.22 512.6H1111.56V507.438H1116.67L1118.41 502.334L1107.45 470.84H1113.77L1119.8 488.588L1121.54 495.142H1121.83L1123.74 488.588L1129.78 470.84ZM1155.01 501C1152.8 501 1151.14 500.439 1150.02 499.318C1148.9 498.158 1148.34 496.534 1148.34 494.446V476.002H1143.64V470.84H1146.19C1147.23 470.84 1147.95 470.608 1148.34 470.144C1148.76 469.68 1148.97 468.926 1148.97 467.882V462.604H1154.66V470.84H1160.98V476.002H1154.66V495.838H1160.52V501H1155.01ZM1175.38 501.696C1173.29 501.696 1171.38 501.329 1169.64 500.594C1167.94 499.859 1166.49 498.815 1165.29 497.462C1164.09 496.07 1163.16 494.407 1162.51 492.474C1161.85 490.502 1161.52 488.317 1161.52 485.92C1161.52 483.523 1161.85 481.357 1162.51 479.424C1163.16 477.452 1164.09 475.789 1165.29 474.436C1166.49 473.044 1167.94 471.981 1169.64 471.246C1171.38 470.511 1173.29 470.144 1175.38 470.144C1177.47 470.144 1179.36 470.511 1181.07 471.246C1182.81 471.981 1184.28 473.044 1185.47 474.436C1186.67 475.789 1187.6 477.452 1188.26 479.424C1188.92 481.357 1189.24 483.523 1189.24 485.92C1189.24 488.317 1188.92 490.502 1188.26 492.474C1187.6 494.407 1186.67 496.07 1185.47 497.462C1184.28 498.815 1182.81 499.859 1181.07 500.594C1179.36 501.329 1177.47 501.696 1175.38 501.696ZM1175.38 496.476C1177.55 496.476 1179.29 495.819 1180.6 494.504C1181.92 493.151 1182.57 491.14 1182.57 488.472V483.368C1182.57 480.7 1181.92 478.709 1180.6 477.394C1179.29 476.041 1177.55 475.364 1175.38 475.364C1173.22 475.364 1171.48 476.041 1170.16 477.394C1168.85 478.709 1168.19 480.7 1168.19 483.368V488.472C1168.19 491.14 1168.85 493.151 1170.16 494.504C1171.48 495.819 1173.22 496.476 1175.38 496.476ZM1201.88 470.84H1208.2V475.828H1208.49C1209.14 474.011 1210.21 472.619 1211.68 471.652C1213.19 470.647 1214.94 470.144 1216.96 470.144C1220.78 470.144 1223.74 471.536 1225.83 474.32C1227.92 477.065 1228.96 480.932 1228.96 485.92C1228.96 490.908 1227.92 494.794 1225.83 497.578C1223.74 500.323 1220.78 501.696 1216.96 501.696C1214.94 501.696 1213.19 501.193 1211.68 500.188C1210.21 499.183 1209.14 497.791 1208.49 496.012H1208.2V512.6H1201.88V470.84ZM1215.04 496.244C1217.25 496.244 1219 495.548 1220.32 494.156C1221.63 492.725 1222.29 490.85 1222.29 488.53V483.31C1222.29 480.99 1221.63 479.134 1220.32 477.742C1219 476.311 1217.25 475.596 1215.04 475.596C1213.11 475.596 1211.48 476.079 1210.17 477.046C1208.85 477.974 1208.2 479.211 1208.2 480.758V491.082C1208.2 492.629 1208.85 493.885 1210.17 494.852C1211.48 495.78 1213.11 496.244 1215.04 496.244ZM1232.14 501V470.84H1238.46V476.64H1238.75C1239.18 475.093 1240.07 473.74 1241.42 472.58C1242.77 471.42 1244.65 470.84 1247.05 470.84H1248.73V476.93H1246.23C1243.72 476.93 1241.79 477.336 1240.43 478.148C1239.12 478.96 1238.46 480.159 1238.46 481.744V501H1232.14ZM1261.83 501.696C1259.74 501.696 1257.83 501.329 1256.09 500.594C1254.39 499.859 1252.94 498.815 1251.74 497.462C1250.54 496.07 1249.61 494.407 1248.96 492.474C1248.3 490.502 1247.97 488.317 1247.97 485.92C1247.97 483.523 1248.3 481.357 1248.96 479.424C1249.61 477.452 1250.54 475.789 1251.74 474.436C1252.94 473.044 1254.39 471.981 1256.09 471.246C1257.83 470.511 1259.74 470.144 1261.83 470.144C1263.92 470.144 1265.81 470.511 1267.52 471.246C1269.26 471.981 1270.73 473.044 1271.92 474.436C1273.12 475.789 1274.05 477.452 1274.71 479.424C1275.37 481.357 1275.69 483.523 1275.69 485.92C1275.69 488.317 1275.37 490.502 1274.71 492.474C1274.05 494.407 1273.12 496.07 1271.92 497.462C1270.73 498.815 1269.26 499.859 1267.52 500.594C1265.81 501.329 1263.92 501.696 1261.83 501.696ZM1261.83 496.476C1264 496.476 1265.74 495.819 1267.05 494.504C1268.37 493.151 1269.02 491.14 1269.02 488.472V483.368C1269.02 480.7 1268.37 478.709 1267.05 477.394C1265.74 476.041 1264 475.364 1261.83 475.364C1259.67 475.364 1257.93 476.041 1256.61 477.394C1255.3 478.709 1254.64 480.7 1254.64 483.368V488.472C1254.64 491.14 1255.3 493.151 1256.61 494.504C1257.93 495.819 1259.67 496.476 1261.83 496.476ZM1297.64 496.012H1297.35C1296.7 497.791 1295.61 499.183 1294.11 500.188C1292.64 501.193 1290.9 501.696 1288.89 501.696C1285.06 501.696 1282.1 500.323 1280.01 497.578C1277.92 494.794 1276.88 490.908 1276.88 485.92C1276.88 480.932 1277.92 477.065 1280.01 474.32C1282.1 471.536 1285.06 470.144 1288.89 470.144C1290.9 470.144 1292.64 470.647 1294.11 471.652C1295.61 472.619 1296.7 474.011 1297.35 475.828H1297.64V458.08H1303.97V501H1297.64V496.012ZM1290.8 496.244C1292.73 496.244 1294.36 495.78 1295.67 494.852C1296.99 493.885 1297.64 492.629 1297.64 491.082V480.758C1297.64 479.211 1296.99 477.974 1295.67 477.046C1294.36 476.079 1292.73 475.596 1290.8 475.596C1288.6 475.596 1286.84 476.311 1285.52 477.742C1284.21 479.134 1283.55 480.99 1283.55 483.31V488.53C1283.55 490.85 1284.21 492.725 1285.52 494.156C1286.84 495.548 1288.6 496.244 1290.8 496.244ZM1327.21 496.012H1326.98C1326.71 496.747 1326.34 497.462 1325.88 498.158C1325.45 498.854 1324.89 499.473 1324.2 500.014C1323.54 500.517 1322.73 500.923 1321.76 501.232C1320.83 501.541 1319.75 501.696 1318.51 501.696C1315.38 501.696 1312.94 500.691 1311.2 498.68C1309.5 496.669 1308.65 493.789 1308.65 490.038V470.84H1314.97V489.226C1314.97 493.905 1316.91 496.244 1320.77 496.244C1321.59 496.244 1322.38 496.147 1323.15 495.954C1323.93 495.722 1324.6 495.393 1325.18 494.968C1325.8 494.543 1326.28 494.001 1326.63 493.344C1327.02 492.687 1327.21 491.913 1327.21 491.024V470.84H1333.53V501H1327.21V496.012ZM1350.18 501.696C1348.02 501.696 1346.09 501.329 1344.38 500.594C1342.68 499.859 1341.25 498.815 1340.09 497.462C1338.93 496.07 1338.04 494.407 1337.42 492.474C1336.81 490.502 1336.5 488.317 1336.5 485.92C1336.5 483.523 1336.81 481.357 1337.42 479.424C1338.04 477.452 1338.93 475.789 1340.09 474.436C1341.25 473.044 1342.68 471.981 1344.38 471.246C1346.09 470.511 1348.02 470.144 1350.18 470.144C1353.2 470.144 1355.68 470.821 1357.61 472.174C1359.54 473.527 1360.95 475.325 1361.84 477.568L1356.62 480.004C1356.2 478.612 1355.44 477.51 1354.36 476.698C1353.32 475.847 1351.92 475.422 1350.18 475.422C1347.86 475.422 1346.11 476.157 1344.91 477.626C1343.75 479.057 1343.17 480.932 1343.17 483.252V488.646C1343.17 490.966 1343.75 492.861 1344.91 494.33C1346.11 495.761 1347.86 496.476 1350.18 496.476C1352.04 496.476 1353.51 496.031 1354.59 495.142C1355.71 494.214 1356.6 492.996 1357.26 491.488L1362.07 494.04C1361.07 496.515 1359.56 498.409 1357.55 499.724C1355.54 501.039 1353.08 501.696 1350.18 501.696ZM1372.39 501C1370.18 501 1368.52 500.439 1367.4 499.318C1366.28 498.158 1365.72 496.534 1365.72 494.446V476.002H1361.02V470.84H1363.57C1364.61 470.84 1365.33 470.608 1365.72 470.144C1366.14 469.68 1366.35 468.926 1366.35 467.882V462.604H1372.04V470.84H1378.36V476.002H1372.04V495.838H1377.9V501H1372.39ZM1384.35 465.446C1383.03 465.446 1382.07 465.137 1381.45 464.518C1380.87 463.899 1380.58 463.107 1380.58 462.14V461.154C1380.58 460.187 1380.87 459.395 1381.45 458.776C1382.07 458.157 1383.03 457.848 1384.35 457.848C1385.66 457.848 1386.61 458.157 1387.19 458.776C1387.77 459.395 1388.06 460.187 1388.06 461.154V462.14C1388.06 463.107 1387.77 463.899 1387.19 464.518C1386.61 465.137 1385.66 465.446 1384.35 465.446ZM1381.16 470.84H1387.48V501H1381.16V470.84ZM1404.28 501.696C1402.19 501.696 1400.28 501.329 1398.54 500.594C1396.84 499.859 1395.39 498.815 1394.19 497.462C1392.99 496.07 1392.06 494.407 1391.4 492.474C1390.75 490.502 1390.42 488.317 1390.42 485.92C1390.42 483.523 1390.75 481.357 1391.4 479.424C1392.06 477.452 1392.99 475.789 1394.19 474.436C1395.39 473.044 1396.84 471.981 1398.54 471.246C1400.28 470.511 1402.19 470.144 1404.28 470.144C1406.37 470.144 1408.26 470.511 1409.96 471.246C1411.7 471.981 1413.17 473.044 1414.37 474.436C1415.57 475.789 1416.5 477.452 1417.15 479.424C1417.81 481.357 1418.14 483.523 1418.14 485.92C1418.14 488.317 1417.81 490.502 1417.15 492.474C1416.5 494.407 1415.57 496.07 1414.37 497.462C1413.17 498.815 1411.7 499.859 1409.96 500.594C1408.26 501.329 1406.37 501.696 1404.28 501.696ZM1404.28 496.476C1406.44 496.476 1408.18 495.819 1409.5 494.504C1410.81 493.151 1411.47 491.14 1411.47 488.472V483.368C1411.47 480.7 1410.81 478.709 1409.5 477.394C1408.18 476.041 1406.44 475.364 1404.28 475.364C1402.11 475.364 1400.37 476.041 1399.06 477.394C1397.74 478.709 1397.09 480.7 1397.09 483.368V488.472C1397.09 491.14 1397.74 493.151 1399.06 494.504C1400.37 495.819 1402.11 496.476 1404.28 496.476ZM1421.12 501V470.84H1427.45V475.828H1427.74C1428.39 474.204 1429.38 472.851 1430.69 471.768C1432.05 470.685 1433.88 470.144 1436.2 470.144C1439.3 470.144 1441.7 471.169 1443.4 473.218C1445.14 475.229 1446.01 478.109 1446.01 481.86V501H1439.68V482.672C1439.68 477.955 1437.79 475.596 1434 475.596C1433.19 475.596 1432.38 475.712 1431.56 475.944C1430.79 476.137 1430.1 476.447 1429.48 476.872C1428.86 477.297 1428.36 477.839 1427.97 478.496C1427.62 479.153 1427.45 479.927 1427.45 480.816V501H1421.12Z" fill="white" style="fill:white;fill-opacity:1;"/>
+</svg>
diff --git a/_images/arch-logo.png b/_images/arch-logo.png
deleted file mode 100755
index bbffb318..00000000
Binary files a/_images/arch-logo.png and /dev/null differ
diff --git a/_images/arch-system-architecture.jpg b/_images/arch-system-architecture.jpg
deleted file mode 100755
index 3c8839a7..00000000
Binary files a/_images/arch-system-architecture.jpg and /dev/null differ
diff --git a/_images/arch_network_diagram_high_level.png b/_images/arch_network_diagram_high_level.png
deleted file mode 100755
index e83e7165..00000000
Binary files a/_images/arch_network_diagram_high_level.png and /dev/null differ
diff --git a/_images/function-calling-flow.jpg b/_images/function-calling-flow.jpg
deleted file mode 100755
index 9f0f4a59..00000000
Binary files a/_images/function-calling-flow.jpg and /dev/null differ
diff --git a/_images/network-topology-agent.jpg b/_images/network-topology-agent.jpg
deleted file mode 100755
index 50ba9a64..00000000
Binary files a/_images/network-topology-agent.jpg and /dev/null differ
diff --git a/_images/network-topology-ingress-egress.jpg b/_images/network-topology-ingress-egress.jpg
deleted file mode 100755
index 03e36e77..00000000
Binary files a/_images/network-topology-ingress-egress.jpg and /dev/null differ
diff --git a/_images/network-topology-ingress-egress.png b/_images/network-topology-ingress-egress.png
new file mode 100755
index 00000000..c2e55584
Binary files /dev/null and b/_images/network-topology-ingress-egress.png differ
diff --git a/_images/plano-system-architecture.png b/_images/plano-system-architecture.png
new file mode 100755
index 00000000..792477d5
Binary files /dev/null and b/_images/plano-system-architecture.png differ
diff --git a/_images/plano_network_diagram_high_level.png b/_images/plano_network_diagram_high_level.png
new file mode 100755
index 00000000..da1c5b92
Binary files /dev/null and b/_images/plano_network_diagram_high_level.png differ
diff --git a/_images/tracing.png b/_images/tracing.png
index 91d6a82b..bb34db91 100755
Binary files a/_images/tracing.png and b/_images/tracing.png differ
diff --git a/_static/css/custom.css b/_static/css/custom.css
new file mode 100755
index 00000000..b7ccb7aa
--- /dev/null
+++ b/_static/css/custom.css
@@ -0,0 +1,6 @@
+/* Prevent sphinxawesome-theme's Tailwind utility `dark:invert` from inverting the header logo. */
+.dark header img[alt="Logo"],
+.dark #left-sidebar img[alt="Logo"] {
+  --tw-invert: invert(0%) !important;
+  filter: none !important;
+}
diff --git a/_static/documentation_options.js b/_static/documentation_options.js
index 0ba5ec80..caaa7577 100755
--- a/_static/documentation_options.js
+++ b/_static/documentation_options.js
@@ -1,5 +1,5 @@
 const DOCUMENTATION_OPTIONS = {
-    VERSION: ' v0.3.22',
+    VERSION: ' v0.4',
     LANGUAGE: 'en',
     COLLAPSE_INDEX: false,
     BUILDER: 'html',
diff --git a/_static/favicon.ico b/_static/favicon.ico
index 29dc5902..a1b75eb4 100755
Binary files a/_static/favicon.ico and b/_static/favicon.ico differ
diff --git a/_static/img/PlanoTagline.svg b/_static/img/PlanoTagline.svg
new file mode 100755
index 00000000..c0c10548
--- /dev/null
+++ b/_static/img/PlanoTagline.svg
@@ -0,0 +1,56 @@
+<svg width="1769" height="693" viewBox="0 0 1769 693" fill="none" xmlns="http://www.w3.org/2000/svg">
+<rect width="1769" height="693" fill="#2A2E4F" style="fill:#2A2E4F;fill:color(display-p3 0.1647 0.1804 0.3098);fill-opacity:1;"/>
+<path d="M331.711 418.187V271.429H356.238V292.538H361.867L356.238 298.368C356.238 289.388 358.919 282.352 364.28 277.259C369.641 272.032 376.878 269.419 385.992 269.419C397.116 269.419 406.028 273.238 412.73 280.878C419.431 288.517 422.782 298.77 422.782 311.637V341.591C422.782 350.169 421.24 357.674 418.158 364.108C415.209 370.407 410.987 375.299 405.492 378.783C399.997 382.268 393.497 384.01 385.992 384.01C376.878 384.01 369.641 381.464 364.28 376.371C358.919 371.144 356.238 364.041 356.238 355.061L361.867 360.891H356.037L356.841 387.227V418.187H331.711ZM377.146 362.298C383.579 362.298 388.605 360.422 392.224 356.669C395.843 352.916 397.652 347.555 397.652 340.586V312.843C397.652 305.874 395.843 300.513 392.224 296.76C388.605 293.007 383.579 291.131 377.146 291.131C370.847 291.131 365.888 293.074 362.269 296.961C358.65 300.714 356.841 306.008 356.841 312.843V340.586C356.841 347.421 358.65 352.782 362.269 356.669C365.888 360.422 370.847 362.298 377.146 362.298ZM497.891 382C490.654 382 484.288 380.526 478.793 377.577C473.432 374.629 469.21 370.474 466.127 365.113C463.045 359.618 461.503 353.319 461.503 346.215V257.959H426.121V235.242H486.633V346.215C486.633 350.236 487.772 353.453 490.051 355.865C492.463 358.143 495.68 359.283 499.701 359.283H533.073V382H497.891ZM573.202 384.01C561.81 384.01 552.83 380.995 546.263 374.964C539.696 368.933 536.412 360.824 536.412 350.638C536.412 339.782 540.031 331.405 547.268 325.508C554.506 319.611 564.759 316.663 578.027 316.663H605.569V307.214C605.569 301.853 603.827 297.698 600.342 294.749C596.858 291.667 592.1 290.126 586.069 290.126C580.574 290.126 576.017 291.332 572.398 293.744C568.779 296.157 566.635 299.44 565.965 303.595H541.438C542.644 293.141 547.335 284.832 555.511 278.666C563.686 272.501 574.14 269.419 586.873 269.419C600.409 269.419 611.064 272.836 618.838 279.672C626.745 286.373 630.699 295.487 630.699 307.013V382H606.373V362.7H602.353L606.373 357.272C606.373 365.448 603.358 371.948 597.327 376.773C591.296 381.598 583.254 384.01 573.202 384.01ZM581.445 365.113C588.548 365.113 594.311 363.303 598.734 359.685C603.291 356.066 605.569 351.375 605.569 345.612V332.143H578.429C573.336 332.143 569.248 333.617 566.166 336.565C563.083 339.514 561.542 343.401 561.542 348.226C561.542 353.453 563.284 357.607 566.769 360.69C570.388 363.639 575.28 365.113 581.445 365.113ZM645.095 382V271.429H669.622V292.538H676.457L669.622 298.368C669.622 289.254 672.235 282.151 677.462 277.058C682.823 271.965 690.128 269.419 699.376 269.419C710.232 269.419 718.876 273.037 725.31 280.275C731.877 287.512 735.16 297.229 735.16 309.425V382H710.031V312.039C710.031 305.337 708.288 300.177 704.804 296.559C701.319 292.94 696.427 291.131 690.128 291.131C683.963 291.131 679.071 293.007 675.452 296.76C671.967 300.513 670.225 305.874 670.225 312.843V382H645.095ZM794.388 383.809C785.006 383.809 776.831 382.067 769.861 378.582C763.026 374.964 757.665 369.938 753.778 363.504C750.026 356.937 748.149 349.231 748.149 340.385V313.044C748.149 304.198 750.026 296.559 753.778 290.126C757.665 283.558 763.026 278.532 769.861 275.048C776.831 271.429 785.006 269.62 794.388 269.62C803.904 269.62 812.079 271.429 818.915 275.048C825.75 278.532 831.044 283.558 834.797 290.126C838.684 296.559 840.627 304.131 840.627 312.843V340.385C840.627 349.231 838.684 356.937 834.797 363.504C831.044 369.938 825.75 374.964 818.915 378.582C812.079 382.067 803.904 383.809 794.388 383.809ZM794.388 361.896C801.089 361.896 806.249 360.087 809.868 356.468C813.621 352.715 815.497 347.354 815.497 340.385V313.044C815.497 305.941 813.621 300.58 809.868 296.961C806.249 293.342 801.089 291.533 794.388 291.533C787.821 291.533 782.661 293.342 778.908 296.961C775.155 300.58 773.279 305.941 773.279 313.044V340.385C773.279 347.354 775.155 352.715 778.908 356.468C782.661 360.087 787.821 361.896 794.388 361.896Z" fill="#7780D9" style="fill:#7780D9;fill:color(display-p3 0.4667 0.5020 0.8510);fill-opacity:1;"/>
+<rect x="140" y="235.402" width="147.807" height="147.807" fill="#7780D9" style="fill:#7780D9;fill:color(display-p3 0.4667 0.5020 0.8510);fill-opacity:1;"/>
+<rect x="151.37" y="246.771" width="125.067" height="125.067" fill="#B9BFFF" style="fill:#B9BFFF;fill:color(display-p3 0.7246 0.7499 1.0000);fill-opacity:1;"/>
+<path d="M157.055 252.456H168.424V263.825H157.055V252.456Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M157.055 269.51H168.424V280.88H157.055V269.51Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M157.055 286.565H168.424V297.935H157.055V286.565Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M157.055 303.62H168.424V314.989H157.055V303.62Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M157.055 320.675H168.424V332.044H157.055V320.675Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M157.055 337.729H168.424V349.099H157.055V337.729Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M157.055 354.784H168.424V366.153H157.055V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M208.219 252.456H219.589V263.825H208.219V252.456Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M208.219 269.51H219.589V280.88H208.219V269.51Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M208.219 286.565H219.589V297.935H208.219V286.565Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M208.219 303.62H219.589V314.989H208.219V303.62Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M208.219 320.675H219.589V332.044H208.219V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M208.219 337.729H219.589V349.099H208.219V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M208.219 354.784H219.589V366.153H208.219V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M174.11 252.456H185.48V263.825H174.11V252.456Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M174.11 269.51H185.48V280.88H174.11V269.51Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M174.11 286.565H185.48V297.935H174.11V286.565Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M174.11 303.62H185.48V314.989H174.11V303.62Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M174.11 320.675H185.48V332.044H174.11V320.675Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M174.11 337.729H185.48V349.099H174.11V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M174.11 354.784H185.48V366.153H174.11V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 252.456H236.643V263.825H225.273V252.456Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M225.273 269.51H236.643V280.88H225.273V269.51Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M225.273 286.565H236.643V297.935H225.273V286.565Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 303.62H236.643V314.989H225.273V303.62Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 320.675H236.643V332.044H225.273V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 337.729H236.643V349.099H225.273V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M225.273 354.784H236.643V366.153H225.273V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M191.164 252.456H202.534V263.825H191.164V252.456Z" fill="#B0B7FF" style="fill:#B0B7FF;fill:color(display-p3 0.6897 0.7182 1.0000);fill-opacity:1;"/>
+<path d="M191.164 269.51H202.534V280.88H191.164V269.51Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M191.164 286.565H202.534V297.935H191.164V286.565Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M191.164 303.62H202.534V314.989H191.164V303.62Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M191.164 320.675H202.534V332.044H191.164V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M191.164 337.729H202.534V349.099H191.164V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M191.164 354.784H202.534V366.153H191.164V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 252.456H253.697V263.825H242.328V252.456Z" fill="#ABB2FA" style="fill:#ABB2FA;fill:color(display-p3 0.6706 0.6990 0.9799);fill-opacity:1;"/>
+<path d="M242.328 269.51H253.697V280.88H242.328V269.51Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 286.565H253.697V297.935H242.328V286.565Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 303.62H253.697V314.989H242.328V303.62Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 320.675H253.697V332.044H242.328V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 337.729H253.697V349.099H242.328V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M242.328 354.784H253.697V366.153H242.328V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 252.456H270.753V263.825H259.383V252.456Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 269.51H270.753V280.88H259.383V269.51Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 286.565H270.753V297.935H259.383V286.565Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 303.62H270.753V314.989H259.383V303.62Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 320.675H270.753V332.044H259.383V320.675Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 337.729H270.753V349.099H259.383V337.729Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M259.383 354.784H270.753V366.153H259.383V354.784Z" fill="#969FF4" style="fill:#969FF4;fill:color(display-p3 0.5882 0.6220 0.9566);fill-opacity:1;"/>
+<path d="M139.988 460.516H157.91C161.274 460.516 163.903 461.463 165.798 463.358C167.731 465.253 168.698 467.785 168.698 470.956C168.698 472.464 168.485 473.759 168.06 474.842C167.635 475.886 167.093 476.756 166.436 477.452C165.779 478.109 165.025 478.612 164.174 478.96C163.323 479.269 162.492 479.463 161.68 479.54V479.888C162.492 479.927 163.381 480.12 164.348 480.468C165.353 480.816 166.281 481.377 167.132 482.15C167.983 482.885 168.698 483.851 169.278 485.05C169.858 486.21 170.148 487.641 170.148 489.342C170.148 490.966 169.877 492.493 169.336 493.924C168.833 495.355 168.118 496.592 167.19 497.636C166.262 498.68 165.16 499.511 163.884 500.13C162.608 500.71 161.216 501 159.708 501H139.988V460.516ZM146.542 495.374H157.794C159.495 495.374 160.829 494.929 161.796 494.04C162.763 493.151 163.246 491.875 163.246 490.212V488.24C163.246 486.577 162.763 485.301 161.796 484.412C160.829 483.523 159.495 483.078 157.794 483.078H146.542V495.374ZM146.542 477.626H156.692C158.316 477.626 159.573 477.22 160.462 476.408C161.351 475.557 161.796 474.359 161.796 472.812V470.956C161.796 469.409 161.351 468.23 160.462 467.418C159.573 466.567 158.316 466.142 156.692 466.142H146.542V477.626ZM191.966 496.012H191.734C191.464 496.747 191.096 497.462 190.632 498.158C190.207 498.854 189.646 499.473 188.95 500.014C188.293 500.517 187.481 500.923 186.514 501.232C185.586 501.541 184.504 501.696 183.266 501.696C180.134 501.696 177.698 500.691 175.958 498.68C174.257 496.669 173.406 493.789 173.406 490.038V470.84H179.728V489.226C179.728 493.905 181.662 496.244 185.528 496.244C186.34 496.244 187.133 496.147 187.906 495.954C188.68 495.722 189.356 495.393 189.936 494.968C190.555 494.543 191.038 494.001 191.386 493.344C191.773 492.687 191.966 491.913 191.966 491.024V470.84H198.288V501H191.966V496.012ZM206.471 465.446C205.156 465.446 204.19 465.137 203.571 464.518C202.991 463.899 202.701 463.107 202.701 462.14V461.154C202.701 460.187 202.991 459.395 203.571 458.776C204.19 458.157 205.156 457.848 206.471 457.848C207.786 457.848 208.733 458.157 209.313 458.776C209.893 459.395 210.183 460.187 210.183 461.154V462.14C210.183 463.107 209.893 463.899 209.313 464.518C208.733 465.137 207.786 465.446 206.471 465.446ZM203.281 470.84H209.603V501H203.281V470.84ZM221.009 501C218.843 501 217.219 500.459 216.137 499.376C215.093 498.255 214.571 496.708 214.571 494.736V458.08H220.893V495.838H225.069V501H221.009ZM246.016 496.012H245.726C245.068 497.791 243.986 499.183 242.478 500.188C241.008 501.193 239.268 501.696 237.258 501.696C233.43 501.696 230.472 500.323 228.384 497.578C226.296 494.794 225.252 490.908 225.252 485.92C225.252 480.932 226.296 477.065 228.384 474.32C230.472 471.536 233.43 470.144 237.258 470.144C239.268 470.144 241.008 470.647 242.478 471.652C243.986 472.619 245.068 474.011 245.726 475.828H246.016V458.08H252.338V501H246.016V496.012ZM239.172 496.244C241.105 496.244 242.729 495.78 244.044 494.852C245.358 493.885 246.016 492.629 246.016 491.082V480.758C246.016 479.211 245.358 477.974 244.044 477.046C242.729 476.079 241.105 475.596 239.172 475.596C236.968 475.596 235.208 476.311 233.894 477.742C232.579 479.134 231.922 480.99 231.922 483.31V488.53C231.922 490.85 232.579 492.725 233.894 494.156C235.208 495.548 236.968 496.244 239.172 496.244ZM289.059 501C287.396 501 286.12 500.536 285.231 499.608C284.342 498.641 283.8 497.423 283.607 495.954H283.317C282.737 497.849 281.674 499.279 280.127 500.246C278.58 501.213 276.705 501.696 274.501 501.696C271.369 501.696 268.952 500.884 267.251 499.26C265.588 497.636 264.757 495.451 264.757 492.706C264.757 489.69 265.84 487.428 268.005 485.92C270.209 484.412 273.418 483.658 277.633 483.658H283.085V481.106C283.085 479.25 282.582 477.819 281.577 476.814C280.572 475.809 279.006 475.306 276.879 475.306C275.1 475.306 273.65 475.693 272.529 476.466C271.408 477.239 270.46 478.225 269.687 479.424L265.917 476.002C266.922 474.301 268.334 472.909 270.151 471.826C271.968 470.705 274.346 470.144 277.285 470.144C281.19 470.144 284.187 471.053 286.275 472.87C288.363 474.687 289.407 477.297 289.407 480.7V495.838H292.597V501H289.059ZM276.299 496.882C278.271 496.882 279.895 496.457 281.171 495.606C282.447 494.717 283.085 493.537 283.085 492.068V487.718H277.749C273.38 487.718 271.195 489.071 271.195 491.778V492.822C271.195 494.175 271.64 495.2 272.529 495.896C273.457 496.553 274.714 496.882 276.299 496.882ZM321.181 503.842C321.181 506.974 319.982 509.333 317.585 510.918C315.188 512.503 311.282 513.296 305.869 513.296C303.394 513.296 301.287 513.122 299.547 512.774C297.846 512.465 296.434 512.001 295.313 511.382C294.23 510.763 293.438 510.009 292.935 509.12C292.432 508.231 292.181 507.206 292.181 506.046C292.181 504.383 292.626 503.088 293.515 502.16C294.443 501.232 295.719 500.594 297.343 500.246V499.608C295.062 498.873 293.921 497.327 293.921 494.968C293.921 493.421 294.443 492.242 295.487 491.43C296.531 490.579 297.788 489.98 299.257 489.632V489.4C297.478 488.549 296.106 487.37 295.139 485.862C294.211 484.315 293.747 482.517 293.747 480.468C293.747 477.375 294.772 474.881 296.821 472.986C298.909 471.091 301.886 470.144 305.753 470.144C307.88 470.144 309.697 470.453 311.205 471.072V470.26C311.205 468.907 311.514 467.863 312.133 467.128C312.79 466.355 313.796 465.968 315.149 465.968H319.789V471.072H313.641V472.29C314.994 473.179 316.038 474.339 316.773 475.77C317.508 477.162 317.875 478.728 317.875 480.468C317.875 483.523 316.831 485.997 314.743 487.892C312.655 489.748 309.678 490.676 305.811 490.676C304.342 490.676 303.027 490.521 301.867 490.212C301.094 490.483 300.417 490.869 299.837 491.372C299.257 491.836 298.967 492.455 298.967 493.228C298.967 494.04 299.334 494.62 300.069 494.968C300.804 495.316 301.848 495.49 303.201 495.49H310.625C314.337 495.49 317.024 496.244 318.687 497.752C320.35 499.221 321.181 501.251 321.181 503.842ZM315.265 504.538C315.265 503.494 314.859 502.663 314.047 502.044C313.274 501.425 311.843 501.116 309.755 501.116H299.547C298.155 501.928 297.459 503.088 297.459 504.596C297.459 505.833 297.942 506.819 298.909 507.554C299.914 508.327 301.596 508.714 303.955 508.714H307.899C312.81 508.714 315.265 507.322 315.265 504.538ZM305.811 486.094C307.667 486.094 309.098 485.688 310.103 484.876C311.147 484.025 311.669 482.73 311.669 480.99V479.83C311.669 478.09 311.147 476.814 310.103 476.002C309.098 475.151 307.667 474.726 305.811 474.726C303.955 474.726 302.505 475.151 301.461 476.002C300.456 476.814 299.953 478.09 299.953 479.83V480.99C299.953 482.73 300.456 484.025 301.461 484.876C302.505 485.688 303.955 486.094 305.811 486.094ZM333.317 501.696C331.152 501.696 329.219 501.329 327.517 500.594C325.816 499.859 324.366 498.815 323.167 497.462C321.969 496.07 321.041 494.407 320.383 492.474C319.765 490.502 319.455 488.317 319.455 485.92C319.455 483.523 319.765 481.357 320.383 479.424C321.041 477.452 321.969 475.789 323.167 474.436C324.366 473.044 325.816 471.981 327.517 471.246C329.219 470.511 331.152 470.144 333.317 470.144C335.521 470.144 337.455 470.531 339.117 471.304C340.819 472.077 342.23 473.16 343.351 474.552C344.473 475.905 345.304 477.491 345.845 479.308C346.425 481.125 346.715 483.078 346.715 485.166V487.544H326.009V488.53C326.009 490.85 326.686 492.764 328.039 494.272C329.431 495.741 331.403 496.476 333.955 496.476C335.811 496.476 337.377 496.07 338.653 495.258C339.929 494.446 341.012 493.344 341.901 491.952L345.613 495.606C344.492 497.462 342.868 498.951 340.741 500.072C338.615 501.155 336.14 501.696 333.317 501.696ZM333.317 475.074C332.235 475.074 331.229 475.267 330.301 475.654C329.412 476.041 328.639 476.582 327.981 477.278C327.363 477.974 326.879 478.805 326.531 479.772C326.183 480.739 326.009 481.802 326.009 482.962V483.368H340.045V482.788C340.045 480.468 339.446 478.612 338.247 477.22C337.049 475.789 335.405 475.074 333.317 475.074ZM349.654 501V470.84H355.976V475.828H356.266C356.923 474.204 357.909 472.851 359.224 471.768C360.577 470.685 362.414 470.144 364.734 470.144C367.827 470.144 370.225 471.169 371.926 473.218C373.666 475.229 374.536 478.109 374.536 481.86V501H368.214V482.672C368.214 477.955 366.319 475.596 362.53 475.596C361.718 475.596 360.906 475.712 360.094 475.944C359.321 476.137 358.625 476.447 358.006 476.872C357.387 477.297 356.885 477.839 356.498 478.496C356.15 479.153 355.976 479.927 355.976 480.816V501H349.654ZM387.59 501C385.386 501 383.724 500.439 382.602 499.318C381.481 498.158 380.92 496.534 380.92 494.446V476.002H376.222V470.84H378.774C379.818 470.84 380.534 470.608 380.92 470.144C381.346 469.68 381.558 468.926 381.558 467.882V462.604H387.242V470.84H393.564V476.002H387.242V495.838H393.1V501H387.59ZM405.932 501.696C403.071 501.696 400.673 501.193 398.74 500.188C396.807 499.183 395.125 497.791 393.694 496.012L397.754 492.3C398.875 493.653 400.113 494.717 401.466 495.49C402.858 496.225 404.463 496.592 406.28 496.592C408.136 496.592 409.509 496.244 410.398 495.548C411.326 494.813 411.79 493.808 411.79 492.532C411.79 491.565 411.461 490.753 410.804 490.096C410.185 489.4 409.083 488.955 407.498 488.762L404.714 488.414C401.621 488.027 399.185 487.138 397.406 485.746C395.666 484.315 394.796 482.208 394.796 479.424C394.796 477.955 395.067 476.659 395.608 475.538C396.149 474.378 396.903 473.411 397.87 472.638C398.875 471.826 400.055 471.207 401.408 470.782C402.8 470.357 404.327 470.144 405.99 470.144C408.697 470.144 410.901 470.569 412.602 471.42C414.342 472.271 415.889 473.45 417.242 474.958L413.356 478.67C412.583 477.742 411.558 476.949 410.282 476.292C409.045 475.596 407.614 475.248 405.99 475.248C404.25 475.248 402.955 475.596 402.104 476.292C401.292 476.988 400.886 477.897 400.886 479.018C400.886 480.178 401.253 481.048 401.988 481.628C402.723 482.208 403.902 482.633 405.526 482.904L408.31 483.252C411.635 483.755 414.052 484.741 415.56 486.21C417.107 487.641 417.88 489.574 417.88 492.01C417.88 493.479 417.59 494.813 417.01 496.012C416.469 497.172 415.676 498.177 414.632 499.028C413.588 499.879 412.331 500.536 410.862 501C409.393 501.464 407.749 501.696 405.932 501.696ZM432.175 476.002H427.593V470.84H432.175V465.852C432.175 463.416 432.813 461.521 434.089 460.168C435.365 458.776 437.279 458.08 439.831 458.08H444.819V463.242H438.497V470.84H444.819V476.002H438.497V501H432.175V476.002ZM468.358 501C466.696 501 465.42 500.536 464.53 499.608C463.641 498.641 463.1 497.423 462.906 495.954H462.616C462.036 497.849 460.973 499.279 459.426 500.246C457.88 501.213 456.004 501.696 453.8 501.696C450.668 501.696 448.252 500.884 446.55 499.26C444.888 497.636 444.056 495.451 444.056 492.706C444.056 489.69 445.139 487.428 447.304 485.92C449.508 484.412 452.718 483.658 456.932 483.658H462.384V481.106C462.384 479.25 461.882 477.819 460.876 476.814C459.871 475.809 458.305 475.306 456.178 475.306C454.4 475.306 452.95 475.693 451.828 476.466C450.707 477.239 449.76 478.225 448.986 479.424L445.216 476.002C446.222 474.301 447.633 472.909 449.45 471.826C451.268 470.705 453.646 470.144 456.584 470.144C460.49 470.144 463.486 471.053 465.574 472.87C467.662 474.687 468.706 477.297 468.706 480.7V495.838H471.896V501H468.358ZM455.598 496.882C457.57 496.882 459.194 496.457 460.47 495.606C461.746 494.717 462.384 493.537 462.384 492.068V487.718H457.048C452.679 487.718 450.494 489.071 450.494 491.778V492.822C450.494 494.175 450.939 495.2 451.828 495.896C452.756 496.553 454.013 496.882 455.598 496.882ZM483.602 501.696C480.741 501.696 478.344 501.193 476.41 500.188C474.477 499.183 472.795 497.791 471.364 496.012L475.424 492.3C476.546 493.653 477.783 494.717 479.136 495.49C480.528 496.225 482.133 496.592 483.95 496.592C485.806 496.592 487.179 496.244 488.068 495.548C488.996 494.813 489.46 493.808 489.46 492.532C489.46 491.565 489.132 490.753 488.474 490.096C487.856 489.4 486.754 488.955 485.168 488.762L482.384 488.414C479.291 488.027 476.855 487.138 475.076 485.746C473.336 484.315 472.466 482.208 472.466 479.424C472.466 477.955 472.737 476.659 473.278 475.538C473.82 474.378 474.574 473.411 475.54 472.638C476.546 471.826 477.725 471.207 479.078 470.782C480.47 470.357 481.998 470.144 483.66 470.144C486.367 470.144 488.571 470.569 490.272 471.42C492.012 472.271 493.559 473.45 494.912 474.958L491.026 478.67C490.253 477.742 489.228 476.949 487.952 476.292C486.715 475.596 485.284 475.248 483.66 475.248C481.92 475.248 480.625 475.596 479.774 476.292C478.962 476.988 478.556 477.897 478.556 479.018C478.556 480.178 478.924 481.048 479.658 481.628C480.393 482.208 481.572 482.633 483.196 482.904L485.98 483.252C489.306 483.755 491.722 484.741 493.23 486.21C494.777 487.641 495.55 489.574 495.55 492.01C495.55 493.479 495.26 494.813 494.68 496.012C494.139 497.172 493.346 498.177 492.302 499.028C491.258 499.879 490.002 500.536 488.532 501C487.063 501.464 485.42 501.696 483.602 501.696ZM506.588 501C504.384 501 502.721 500.439 501.6 499.318C500.479 498.158 499.918 496.534 499.918 494.446V476.002H495.22V470.84H497.772C498.816 470.84 499.531 470.608 499.918 470.144C500.343 469.68 500.556 468.926 500.556 467.882V462.604H506.24V470.84H512.562V476.002H506.24V495.838H512.098V501H506.588ZM526.965 501.696C524.8 501.696 522.866 501.329 521.165 500.594C519.464 499.859 518.014 498.815 516.815 497.462C515.616 496.07 514.688 494.407 514.031 492.474C513.412 490.502 513.103 488.317 513.103 485.92C513.103 483.523 513.412 481.357 514.031 479.424C514.688 477.452 515.616 475.789 516.815 474.436C518.014 473.044 519.464 471.981 521.165 471.246C522.866 470.511 524.8 470.144 526.965 470.144C529.169 470.144 531.102 470.531 532.765 471.304C534.466 472.077 535.878 473.16 536.999 474.552C538.12 475.905 538.952 477.491 539.493 479.308C540.073 481.125 540.363 483.078 540.363 485.166V487.544H519.657V488.53C519.657 490.85 520.334 492.764 521.687 494.272C523.079 495.741 525.051 496.476 527.603 496.476C529.459 496.476 531.025 496.07 532.301 495.258C533.577 494.446 534.66 493.344 535.549 491.952L539.261 495.606C538.14 497.462 536.516 498.951 534.389 500.072C532.262 501.155 529.788 501.696 526.965 501.696ZM526.965 475.074C525.882 475.074 524.877 475.267 523.949 475.654C523.06 476.041 522.286 476.582 521.629 477.278C521.01 477.974 520.527 478.805 520.179 479.772C519.831 480.739 519.657 481.802 519.657 482.962V483.368H533.693V482.788C533.693 480.468 533.094 478.612 531.895 477.22C530.696 475.789 529.053 475.074 526.965 475.074ZM543.301 501V470.84H549.623V476.64H549.913C550.339 475.093 551.228 473.74 552.581 472.58C553.935 471.42 555.81 470.84 558.207 470.84H559.889V476.93H557.395C554.882 476.93 552.949 477.336 551.595 478.148C550.281 478.96 549.623 480.159 549.623 481.744V501H543.301ZM560.686 492.938C562.117 492.938 563.18 493.305 563.876 494.04C564.572 494.775 564.92 495.722 564.92 496.882V497.81C564.92 499.705 564.398 501.677 563.354 503.726C562.349 505.814 560.976 507.573 559.236 509.004H554.538C555.775 507.767 556.8 506.568 557.612 505.408C558.424 504.248 559.043 502.953 559.468 501.522C558.463 501.329 557.709 500.903 557.206 500.246C556.703 499.55 556.452 498.738 556.452 497.81V496.882C556.452 495.722 556.8 494.775 557.496 494.04C558.192 493.305 559.255 492.938 560.686 492.938ZM601.307 501C599.645 501 598.369 500.536 597.479 499.608C596.59 498.641 596.049 497.423 595.855 495.954H595.565C594.985 497.849 593.922 499.279 592.375 500.246C590.829 501.213 588.953 501.696 586.749 501.696C583.617 501.696 581.201 500.884 579.499 499.26C577.837 497.636 577.005 495.451 577.005 492.706C577.005 489.69 578.088 487.428 580.253 485.92C582.457 484.412 585.667 483.658 589.881 483.658H595.333V481.106C595.333 479.25 594.831 477.819 593.825 476.814C592.82 475.809 591.254 475.306 589.127 475.306C587.349 475.306 585.899 475.693 584.777 476.466C583.656 477.239 582.709 478.225 581.935 479.424L578.165 476.002C579.171 474.301 580.582 472.909 582.399 471.826C584.217 470.705 586.595 470.144 589.533 470.144C593.439 470.144 596.435 471.053 598.523 472.87C600.611 474.687 601.655 477.297 601.655 480.7V495.838H604.845V501H601.307ZM588.547 496.882C590.519 496.882 592.143 496.457 593.419 495.606C594.695 494.717 595.333 493.537 595.333 492.068V487.718H589.997C585.628 487.718 583.443 489.071 583.443 491.778V492.822C583.443 494.175 583.888 495.2 584.777 495.896C585.705 496.553 586.962 496.882 588.547 496.882ZM606.981 501V470.84H613.303V475.828H613.593C614.251 474.204 615.237 472.851 616.551 471.768C617.905 470.685 619.741 470.144 622.061 470.144C625.155 470.144 627.552 471.169 629.253 473.218C630.993 475.229 631.863 478.109 631.863 481.86V501H625.541V482.672C625.541 477.955 623.647 475.596 619.857 475.596C619.045 475.596 618.233 475.712 617.421 475.944C616.648 476.137 615.952 476.447 615.333 476.872C614.715 477.297 614.212 477.839 613.825 478.496C613.477 479.153 613.303 479.927 613.303 480.816V501H606.981ZM655.532 496.012H655.242C654.585 497.791 653.502 499.183 651.994 500.188C650.525 501.193 648.785 501.696 646.774 501.696C642.946 501.696 639.988 500.323 637.9 497.578C635.812 494.794 634.768 490.908 634.768 485.92C634.768 480.932 635.812 477.065 637.9 474.32C639.988 471.536 642.946 470.144 646.774 470.144C648.785 470.144 650.525 470.647 651.994 471.652C653.502 472.619 654.585 474.011 655.242 475.828H655.532V458.08H661.854V501H655.532V496.012ZM648.688 496.244C650.621 496.244 652.245 495.78 653.56 494.852C654.875 493.885 655.532 492.629 655.532 491.082V480.758C655.532 479.211 654.875 477.974 653.56 477.046C652.245 476.079 650.621 475.596 648.688 475.596C646.484 475.596 644.725 476.311 643.41 477.742C642.095 479.134 641.438 480.99 641.438 483.31V488.53C641.438 490.85 642.095 492.725 643.41 494.156C644.725 495.548 646.484 496.244 648.688 496.244ZM695.443 496.012H695.153C694.496 497.791 693.413 499.183 691.905 500.188C690.436 501.193 688.696 501.696 686.685 501.696C682.857 501.696 679.899 500.323 677.811 497.578C675.723 494.794 674.679 490.908 674.679 485.92C674.679 480.932 675.723 477.065 677.811 474.32C679.899 471.536 682.857 470.144 686.685 470.144C688.696 470.144 690.436 470.647 691.905 471.652C693.413 472.619 694.496 474.011 695.153 475.828H695.443V458.08H701.765V501H695.443V496.012ZM688.599 496.244C690.532 496.244 692.156 495.78 693.471 494.852C694.786 493.885 695.443 492.629 695.443 491.082V480.758C695.443 479.211 694.786 477.974 693.471 477.046C692.156 476.079 690.532 475.596 688.599 475.596C686.395 475.596 684.636 476.311 683.321 477.742C682.006 479.134 681.349 480.99 681.349 483.31V488.53C681.349 490.85 682.006 492.725 683.321 494.156C684.636 495.548 686.395 496.244 688.599 496.244ZM718.573 501.696C716.408 501.696 714.475 501.329 712.773 500.594C711.072 499.859 709.622 498.815 708.423 497.462C707.225 496.07 706.297 494.407 705.639 492.474C705.021 490.502 704.711 488.317 704.711 485.92C704.711 483.523 705.021 481.357 705.639 479.424C706.297 477.452 707.225 475.789 708.423 474.436C709.622 473.044 711.072 471.981 712.773 471.246C714.475 470.511 716.408 470.144 718.573 470.144C720.777 470.144 722.711 470.531 724.373 471.304C726.075 472.077 727.486 473.16 728.607 474.552C729.729 475.905 730.56 477.491 731.101 479.308C731.681 481.125 731.971 483.078 731.971 485.166V487.544H711.265V488.53C711.265 490.85 711.942 492.764 713.295 494.272C714.687 495.741 716.659 496.476 719.211 496.476C721.067 496.476 722.633 496.07 723.909 495.258C725.185 494.446 726.268 493.344 727.157 491.952L730.869 495.606C729.748 497.462 728.124 498.951 725.997 500.072C723.871 501.155 721.396 501.696 718.573 501.696ZM718.573 475.074C717.491 475.074 716.485 475.267 715.557 475.654C714.668 476.041 713.895 476.582 713.237 477.278C712.619 477.974 712.135 478.805 711.787 479.772C711.439 480.739 711.265 481.802 711.265 482.962V483.368H725.301V482.788C725.301 480.468 724.702 478.612 723.503 477.22C722.305 475.789 720.661 475.074 718.573 475.074ZM741.348 501C739.183 501 737.559 500.459 736.476 499.376C735.432 498.255 734.91 496.708 734.91 494.736V458.08H741.232V495.838H745.408V501H741.348ZM750.579 465.446C749.264 465.446 748.298 465.137 747.679 464.518C747.099 463.899 746.809 463.107 746.809 462.14V461.154C746.809 460.187 747.099 459.395 747.679 458.776C748.298 458.157 749.264 457.848 750.579 457.848C751.894 457.848 752.841 458.157 753.421 458.776C754.001 459.395 754.291 460.187 754.291 461.154V462.14C754.291 463.107 754.001 463.899 753.421 464.518C752.841 465.137 751.894 465.446 750.579 465.446ZM747.389 470.84H753.711V501H747.389V470.84ZM765.233 501L754.967 470.84H761.231L765.813 484.586L768.887 495.142H769.235L772.309 484.586L777.007 470.84H783.039L772.657 501H765.233ZM795.677 501.696C793.512 501.696 791.578 501.329 789.877 500.594C788.176 499.859 786.726 498.815 785.527 497.462C784.328 496.07 783.4 494.407 782.743 492.474C782.124 490.502 781.815 488.317 781.815 485.92C781.815 483.523 782.124 481.357 782.743 479.424C783.4 477.452 784.328 475.789 785.527 474.436C786.726 473.044 788.176 471.981 789.877 471.246C791.578 470.511 793.512 470.144 795.677 470.144C797.881 470.144 799.814 470.531 801.477 471.304C803.178 472.077 804.59 473.16 805.711 474.552C806.832 475.905 807.664 477.491 808.205 479.308C808.785 481.125 809.075 483.078 809.075 485.166V487.544H788.369V488.53C788.369 490.85 789.046 492.764 790.399 494.272C791.791 495.741 793.763 496.476 796.315 496.476C798.171 496.476 799.737 496.07 801.013 495.258C802.289 494.446 803.372 493.344 804.261 491.952L807.973 495.606C806.852 497.462 805.228 498.951 803.101 500.072C800.974 501.155 798.5 501.696 795.677 501.696ZM795.677 475.074C794.594 475.074 793.589 475.267 792.661 475.654C791.772 476.041 790.998 476.582 790.341 477.278C789.722 477.974 789.239 478.805 788.891 479.772C788.543 480.739 788.369 481.802 788.369 482.962V483.368H802.405V482.788C802.405 480.468 801.806 478.612 800.607 477.22C799.408 475.789 797.765 475.074 795.677 475.074ZM812.014 501V470.84H818.336V476.64H818.626C819.051 475.093 819.94 473.74 821.294 472.58C822.647 471.42 824.522 470.84 826.92 470.84H828.602V476.93H826.108C823.594 476.93 821.661 477.336 820.308 478.148C818.993 478.96 818.336 480.159 818.336 481.744V501H812.014ZM848.156 501C845.952 501 844.289 500.439 843.168 499.318C842.046 498.158 841.486 496.534 841.486 494.446V476.002H836.788V470.84H839.34C840.384 470.84 841.099 470.608 841.486 470.144C841.911 469.68 842.124 468.926 842.124 467.882V462.604H847.808V470.84H854.13V476.002H847.808V495.838H853.666V501H848.156ZM856.927 458.08H863.249V475.828H863.539C864.197 474.204 865.183 472.851 866.497 471.768C867.851 470.685 869.687 470.144 872.007 470.144C875.101 470.144 877.498 471.169 879.199 473.218C880.939 475.229 881.809 478.109 881.809 481.86V501H875.487V482.614C875.487 477.935 873.593 475.596 869.803 475.596C868.991 475.596 868.179 475.712 867.367 475.944C866.594 476.137 865.898 476.447 865.279 476.872C864.661 477.297 864.158 477.839 863.771 478.496C863.423 479.153 863.249 479.907 863.249 480.758V501H856.927V458.08ZM898.344 501.696C896.179 501.696 894.245 501.329 892.544 500.594C890.843 499.859 889.393 498.815 888.194 497.462C886.995 496.07 886.067 494.407 885.41 492.474C884.791 490.502 884.482 488.317 884.482 485.92C884.482 483.523 884.791 481.357 885.41 479.424C886.067 477.452 886.995 475.789 888.194 474.436C889.393 473.044 890.843 471.981 892.544 471.246C894.245 470.511 896.179 470.144 898.344 470.144C900.548 470.144 902.481 470.531 904.144 471.304C905.845 472.077 907.257 473.16 908.378 474.552C909.499 475.905 910.331 477.491 910.872 479.308C911.452 481.125 911.742 483.078 911.742 485.166V487.544H891.036V488.53C891.036 490.85 891.713 492.764 893.066 494.272C894.458 495.741 896.43 496.476 898.982 496.476C900.838 496.476 902.404 496.07 903.68 495.258C904.956 494.446 906.039 493.344 906.928 491.952L910.64 495.606C909.519 497.462 907.895 498.951 905.768 500.072C903.641 501.155 901.167 501.696 898.344 501.696ZM898.344 475.074C897.261 475.074 896.256 475.267 895.328 475.654C894.439 476.041 893.665 476.582 893.008 477.278C892.389 477.974 891.906 478.805 891.558 479.772C891.21 480.739 891.036 481.802 891.036 482.962V483.368H905.072V482.788C905.072 480.468 904.473 478.612 903.274 477.22C902.075 475.789 900.432 475.074 898.344 475.074ZM914.68 501V470.84H921.002V475.828H921.292C921.602 475.055 921.969 474.32 922.394 473.624C922.858 472.928 923.4 472.329 924.018 471.826C924.676 471.285 925.43 470.879 926.28 470.608C927.17 470.299 928.194 470.144 929.354 470.144C931.404 470.144 933.221 470.647 934.806 471.652C936.392 472.657 937.552 474.204 938.286 476.292H938.46C939.002 474.591 940.046 473.141 941.592 471.942C943.139 470.743 945.13 470.144 947.566 470.144C950.582 470.144 952.922 471.169 954.584 473.218C956.247 475.229 957.078 478.109 957.078 481.86V501H950.756V482.614C950.756 480.294 950.312 478.554 949.422 477.394C948.533 476.195 947.122 475.596 945.188 475.596C944.376 475.596 943.603 475.712 942.868 475.944C942.134 476.137 941.476 476.447 940.896 476.872C940.355 477.297 939.91 477.839 939.562 478.496C939.214 479.153 939.04 479.907 939.04 480.758V501H932.718V482.614C932.718 477.935 930.882 475.596 927.208 475.596C926.435 475.596 925.662 475.712 924.888 475.944C924.154 476.137 923.496 476.447 922.916 476.872C922.336 477.297 921.872 477.839 921.524 478.496C921.176 479.153 921.002 479.907 921.002 480.758V501H914.68ZM971.414 501V470.84H977.736V476.64H978.026C978.451 475.093 979.341 473.74 980.694 472.58C982.047 471.42 983.923 470.84 986.32 470.84H988.002V476.93H985.508C982.995 476.93 981.061 477.336 979.708 478.148C978.393 478.96 977.736 480.159 977.736 481.744V501H971.414ZM1001.11 501.696C998.94 501.696 997.007 501.329 995.306 500.594C993.604 499.859 992.154 498.815 990.956 497.462C989.757 496.07 988.829 494.407 988.172 492.474C987.553 490.502 987.244 488.317 987.244 485.92C987.244 483.523 987.553 481.357 988.172 479.424C988.829 477.452 989.757 475.789 990.956 474.436C992.154 473.044 993.604 471.981 995.306 471.246C997.007 470.511 998.94 470.144 1001.11 470.144C1003.31 470.144 1005.24 470.531 1006.91 471.304C1008.61 472.077 1010.02 473.16 1011.14 474.552C1012.26 475.905 1013.09 477.491 1013.63 479.308C1014.21 481.125 1014.5 483.078 1014.5 485.166V487.544H993.798V488.53C993.798 490.85 994.474 492.764 995.828 494.272C997.22 495.741 999.192 496.476 1001.74 496.476C1003.6 496.476 1005.17 496.07 1006.44 495.258C1007.72 494.446 1008.8 493.344 1009.69 491.952L1013.4 495.606C1012.28 497.462 1010.66 498.951 1008.53 500.072C1006.4 501.155 1003.93 501.696 1001.11 501.696ZM1001.11 475.074C1000.02 475.074 999.018 475.267 998.09 475.654C997.2 476.041 996.427 476.582 995.77 477.278C995.151 477.974 994.668 478.805 994.32 479.772C993.972 480.739 993.798 481.802 993.798 482.962V483.368H1007.83V482.788C1007.83 480.468 1007.23 478.612 1006.04 477.22C1004.84 475.789 1003.19 475.074 1001.11 475.074ZM1023.88 501C1021.72 501 1020.09 500.459 1019.01 499.376C1017.96 498.255 1017.44 496.708 1017.44 494.736V458.08H1023.76V495.838H1027.94V501H1023.88ZM1033.11 465.446C1031.8 465.446 1030.83 465.137 1030.21 464.518C1029.63 463.899 1029.34 463.107 1029.34 462.14V461.154C1029.34 460.187 1029.63 459.395 1030.21 458.776C1030.83 458.157 1031.8 457.848 1033.11 457.848C1034.43 457.848 1035.37 458.157 1035.95 458.776C1036.53 459.395 1036.82 460.187 1036.82 461.154V462.14C1036.82 463.107 1036.53 463.899 1035.95 464.518C1035.37 465.137 1034.43 465.446 1033.11 465.446ZM1029.92 470.84H1036.24V501H1029.92V470.84ZM1063.31 501C1061.65 501 1060.37 500.536 1059.48 499.608C1058.59 498.641 1058.05 497.423 1057.86 495.954H1057.57C1056.99 497.849 1055.92 499.279 1054.38 500.246C1052.83 501.213 1050.96 501.696 1048.75 501.696C1045.62 501.696 1043.2 500.884 1041.5 499.26C1039.84 497.636 1039.01 495.451 1039.01 492.706C1039.01 489.69 1040.09 487.428 1042.26 485.92C1044.46 484.412 1047.67 483.658 1051.88 483.658H1057.34V481.106C1057.34 479.25 1056.83 477.819 1055.83 476.814C1054.82 475.809 1053.26 475.306 1051.13 475.306C1049.35 475.306 1047.9 475.693 1046.78 476.466C1045.66 477.239 1044.71 478.225 1043.94 479.424L1040.17 476.002C1041.17 474.301 1042.58 472.909 1044.4 471.826C1046.22 470.705 1048.6 470.144 1051.54 470.144C1055.44 470.144 1058.44 471.053 1060.53 472.87C1062.61 474.687 1063.66 477.297 1063.66 480.7V495.838H1066.85V501H1063.31ZM1050.55 496.882C1052.52 496.882 1054.15 496.457 1055.42 495.606C1056.7 494.717 1057.34 493.537 1057.34 492.068V487.718H1052C1047.63 487.718 1045.45 489.071 1045.45 491.778V492.822C1045.45 494.175 1045.89 495.2 1046.78 495.896C1047.71 496.553 1048.96 496.882 1050.55 496.882ZM1068.98 458.08H1075.31V475.828H1075.6C1076.25 474.011 1077.32 472.619 1078.79 471.652C1080.29 470.647 1082.05 470.144 1084.06 470.144C1087.89 470.144 1090.85 471.536 1092.94 474.32C1095.03 477.065 1096.07 480.932 1096.07 485.92C1096.07 490.908 1095.03 494.794 1092.94 497.578C1090.85 500.323 1087.89 501.696 1084.06 501.696C1082.05 501.696 1080.29 501.193 1078.79 500.188C1077.32 499.183 1076.25 497.791 1075.6 496.012H1075.31V501H1068.98V458.08ZM1082.15 496.244C1084.35 496.244 1086.11 495.548 1087.43 494.156C1088.74 492.725 1089.4 490.85 1089.4 488.53V483.31C1089.4 480.99 1088.74 479.134 1087.43 477.742C1086.11 476.311 1084.35 475.596 1082.15 475.596C1080.22 475.596 1078.59 476.079 1077.28 477.046C1075.96 477.974 1075.31 479.211 1075.31 480.758V491.082C1075.31 492.629 1075.96 493.885 1077.28 494.852C1078.59 495.78 1080.22 496.244 1082.15 496.244ZM1105.69 501C1103.52 501 1101.9 500.459 1100.81 499.376C1099.77 498.255 1099.25 496.708 1099.25 494.736V458.08H1105.57V495.838H1109.75V501H1105.69ZM1129.78 470.84H1135.87L1123.16 506.974C1122.82 507.979 1122.41 508.83 1121.95 509.526C1121.52 510.261 1121 510.841 1120.38 511.266C1119.8 511.73 1119.08 512.059 1118.23 512.252C1117.38 512.484 1116.38 512.6 1115.22 512.6H1111.56V507.438H1116.67L1118.41 502.334L1107.45 470.84H1113.77L1119.8 488.588L1121.54 495.142H1121.83L1123.74 488.588L1129.78 470.84ZM1155.01 501C1152.8 501 1151.14 500.439 1150.02 499.318C1148.9 498.158 1148.34 496.534 1148.34 494.446V476.002H1143.64V470.84H1146.19C1147.23 470.84 1147.95 470.608 1148.34 470.144C1148.76 469.68 1148.97 468.926 1148.97 467.882V462.604H1154.66V470.84H1160.98V476.002H1154.66V495.838H1160.52V501H1155.01ZM1175.38 501.696C1173.29 501.696 1171.38 501.329 1169.64 500.594C1167.94 499.859 1166.49 498.815 1165.29 497.462C1164.09 496.07 1163.16 494.407 1162.51 492.474C1161.85 490.502 1161.52 488.317 1161.52 485.92C1161.52 483.523 1161.85 481.357 1162.51 479.424C1163.16 477.452 1164.09 475.789 1165.29 474.436C1166.49 473.044 1167.94 471.981 1169.64 471.246C1171.38 470.511 1173.29 470.144 1175.38 470.144C1177.47 470.144 1179.36 470.511 1181.07 471.246C1182.81 471.981 1184.28 473.044 1185.47 474.436C1186.67 475.789 1187.6 477.452 1188.26 479.424C1188.92 481.357 1189.24 483.523 1189.24 485.92C1189.24 488.317 1188.92 490.502 1188.26 492.474C1187.6 494.407 1186.67 496.07 1185.47 497.462C1184.28 498.815 1182.81 499.859 1181.07 500.594C1179.36 501.329 1177.47 501.696 1175.38 501.696ZM1175.38 496.476C1177.55 496.476 1179.29 495.819 1180.6 494.504C1181.92 493.151 1182.57 491.14 1182.57 488.472V483.368C1182.57 480.7 1181.92 478.709 1180.6 477.394C1179.29 476.041 1177.55 475.364 1175.38 475.364C1173.22 475.364 1171.48 476.041 1170.16 477.394C1168.85 478.709 1168.19 480.7 1168.19 483.368V488.472C1168.19 491.14 1168.85 493.151 1170.16 494.504C1171.48 495.819 1173.22 496.476 1175.38 496.476ZM1201.88 470.84H1208.2V475.828H1208.49C1209.14 474.011 1210.21 472.619 1211.68 471.652C1213.19 470.647 1214.94 470.144 1216.96 470.144C1220.78 470.144 1223.74 471.536 1225.83 474.32C1227.92 477.065 1228.96 480.932 1228.96 485.92C1228.96 490.908 1227.92 494.794 1225.83 497.578C1223.74 500.323 1220.78 501.696 1216.96 501.696C1214.94 501.696 1213.19 501.193 1211.68 500.188C1210.21 499.183 1209.14 497.791 1208.49 496.012H1208.2V512.6H1201.88V470.84ZM1215.04 496.244C1217.25 496.244 1219 495.548 1220.32 494.156C1221.63 492.725 1222.29 490.85 1222.29 488.53V483.31C1222.29 480.99 1221.63 479.134 1220.32 477.742C1219 476.311 1217.25 475.596 1215.04 475.596C1213.11 475.596 1211.48 476.079 1210.17 477.046C1208.85 477.974 1208.2 479.211 1208.2 480.758V491.082C1208.2 492.629 1208.85 493.885 1210.17 494.852C1211.48 495.78 1213.11 496.244 1215.04 496.244ZM1232.14 501V470.84H1238.46V476.64H1238.75C1239.18 475.093 1240.07 473.74 1241.42 472.58C1242.77 471.42 1244.65 470.84 1247.05 470.84H1248.73V476.93H1246.23C1243.72 476.93 1241.79 477.336 1240.43 478.148C1239.12 478.96 1238.46 480.159 1238.46 481.744V501H1232.14ZM1261.83 501.696C1259.74 501.696 1257.83 501.329 1256.09 500.594C1254.39 499.859 1252.94 498.815 1251.74 497.462C1250.54 496.07 1249.61 494.407 1248.96 492.474C1248.3 490.502 1247.97 488.317 1247.97 485.92C1247.97 483.523 1248.3 481.357 1248.96 479.424C1249.61 477.452 1250.54 475.789 1251.74 474.436C1252.94 473.044 1254.39 471.981 1256.09 471.246C1257.83 470.511 1259.74 470.144 1261.83 470.144C1263.92 470.144 1265.81 470.511 1267.52 471.246C1269.26 471.981 1270.73 473.044 1271.92 474.436C1273.12 475.789 1274.05 477.452 1274.71 479.424C1275.37 481.357 1275.69 483.523 1275.69 485.92C1275.69 488.317 1275.37 490.502 1274.71 492.474C1274.05 494.407 1273.12 496.07 1271.92 497.462C1270.73 498.815 1269.26 499.859 1267.52 500.594C1265.81 501.329 1263.92 501.696 1261.83 501.696ZM1261.83 496.476C1264 496.476 1265.74 495.819 1267.05 494.504C1268.37 493.151 1269.02 491.14 1269.02 488.472V483.368C1269.02 480.7 1268.37 478.709 1267.05 477.394C1265.74 476.041 1264 475.364 1261.83 475.364C1259.67 475.364 1257.93 476.041 1256.61 477.394C1255.3 478.709 1254.64 480.7 1254.64 483.368V488.472C1254.64 491.14 1255.3 493.151 1256.61 494.504C1257.93 495.819 1259.67 496.476 1261.83 496.476ZM1297.64 496.012H1297.35C1296.7 497.791 1295.61 499.183 1294.11 500.188C1292.64 501.193 1290.9 501.696 1288.89 501.696C1285.06 501.696 1282.1 500.323 1280.01 497.578C1277.92 494.794 1276.88 490.908 1276.88 485.92C1276.88 480.932 1277.92 477.065 1280.01 474.32C1282.1 471.536 1285.06 470.144 1288.89 470.144C1290.9 470.144 1292.64 470.647 1294.11 471.652C1295.61 472.619 1296.7 474.011 1297.35 475.828H1297.64V458.08H1303.97V501H1297.64V496.012ZM1290.8 496.244C1292.73 496.244 1294.36 495.78 1295.67 494.852C1296.99 493.885 1297.64 492.629 1297.64 491.082V480.758C1297.64 479.211 1296.99 477.974 1295.67 477.046C1294.36 476.079 1292.73 475.596 1290.8 475.596C1288.6 475.596 1286.84 476.311 1285.52 477.742C1284.21 479.134 1283.55 480.99 1283.55 483.31V488.53C1283.55 490.85 1284.21 492.725 1285.52 494.156C1286.84 495.548 1288.6 496.244 1290.8 496.244ZM1327.21 496.012H1326.98C1326.71 496.747 1326.34 497.462 1325.88 498.158C1325.45 498.854 1324.89 499.473 1324.2 500.014C1323.54 500.517 1322.73 500.923 1321.76 501.232C1320.83 501.541 1319.75 501.696 1318.51 501.696C1315.38 501.696 1312.94 500.691 1311.2 498.68C1309.5 496.669 1308.65 493.789 1308.65 490.038V470.84H1314.97V489.226C1314.97 493.905 1316.91 496.244 1320.77 496.244C1321.59 496.244 1322.38 496.147 1323.15 495.954C1323.93 495.722 1324.6 495.393 1325.18 494.968C1325.8 494.543 1326.28 494.001 1326.63 493.344C1327.02 492.687 1327.21 491.913 1327.21 491.024V470.84H1333.53V501H1327.21V496.012ZM1350.18 501.696C1348.02 501.696 1346.09 501.329 1344.38 500.594C1342.68 499.859 1341.25 498.815 1340.09 497.462C1338.93 496.07 1338.04 494.407 1337.42 492.474C1336.81 490.502 1336.5 488.317 1336.5 485.92C1336.5 483.523 1336.81 481.357 1337.42 479.424C1338.04 477.452 1338.93 475.789 1340.09 474.436C1341.25 473.044 1342.68 471.981 1344.38 471.246C1346.09 470.511 1348.02 470.144 1350.18 470.144C1353.2 470.144 1355.68 470.821 1357.61 472.174C1359.54 473.527 1360.95 475.325 1361.84 477.568L1356.62 480.004C1356.2 478.612 1355.44 477.51 1354.36 476.698C1353.32 475.847 1351.92 475.422 1350.18 475.422C1347.86 475.422 1346.11 476.157 1344.91 477.626C1343.75 479.057 1343.17 480.932 1343.17 483.252V488.646C1343.17 490.966 1343.75 492.861 1344.91 494.33C1346.11 495.761 1347.86 496.476 1350.18 496.476C1352.04 496.476 1353.51 496.031 1354.59 495.142C1355.71 494.214 1356.6 492.996 1357.26 491.488L1362.07 494.04C1361.07 496.515 1359.56 498.409 1357.55 499.724C1355.54 501.039 1353.08 501.696 1350.18 501.696ZM1372.39 501C1370.18 501 1368.52 500.439 1367.4 499.318C1366.28 498.158 1365.72 496.534 1365.72 494.446V476.002H1361.02V470.84H1363.57C1364.61 470.84 1365.33 470.608 1365.72 470.144C1366.14 469.68 1366.35 468.926 1366.35 467.882V462.604H1372.04V470.84H1378.36V476.002H1372.04V495.838H1377.9V501H1372.39ZM1384.35 465.446C1383.03 465.446 1382.07 465.137 1381.45 464.518C1380.87 463.899 1380.58 463.107 1380.58 462.14V461.154C1380.58 460.187 1380.87 459.395 1381.45 458.776C1382.07 458.157 1383.03 457.848 1384.35 457.848C1385.66 457.848 1386.61 458.157 1387.19 458.776C1387.77 459.395 1388.06 460.187 1388.06 461.154V462.14C1388.06 463.107 1387.77 463.899 1387.19 464.518C1386.61 465.137 1385.66 465.446 1384.35 465.446ZM1381.16 470.84H1387.48V501H1381.16V470.84ZM1404.28 501.696C1402.19 501.696 1400.28 501.329 1398.54 500.594C1396.84 499.859 1395.39 498.815 1394.19 497.462C1392.99 496.07 1392.06 494.407 1391.4 492.474C1390.75 490.502 1390.42 488.317 1390.42 485.92C1390.42 483.523 1390.75 481.357 1391.4 479.424C1392.06 477.452 1392.99 475.789 1394.19 474.436C1395.39 473.044 1396.84 471.981 1398.54 471.246C1400.28 470.511 1402.19 470.144 1404.28 470.144C1406.37 470.144 1408.26 470.511 1409.96 471.246C1411.7 471.981 1413.17 473.044 1414.37 474.436C1415.57 475.789 1416.5 477.452 1417.15 479.424C1417.81 481.357 1418.14 483.523 1418.14 485.92C1418.14 488.317 1417.81 490.502 1417.15 492.474C1416.5 494.407 1415.57 496.07 1414.37 497.462C1413.17 498.815 1411.7 499.859 1409.96 500.594C1408.26 501.329 1406.37 501.696 1404.28 501.696ZM1404.28 496.476C1406.44 496.476 1408.18 495.819 1409.5 494.504C1410.81 493.151 1411.47 491.14 1411.47 488.472V483.368C1411.47 480.7 1410.81 478.709 1409.5 477.394C1408.18 476.041 1406.44 475.364 1404.28 475.364C1402.11 475.364 1400.37 476.041 1399.06 477.394C1397.74 478.709 1397.09 480.7 1397.09 483.368V488.472C1397.09 491.14 1397.74 493.151 1399.06 494.504C1400.37 495.819 1402.11 496.476 1404.28 496.476ZM1421.12 501V470.84H1427.45V475.828H1427.74C1428.39 474.204 1429.38 472.851 1430.69 471.768C1432.05 470.685 1433.88 470.144 1436.2 470.144C1439.3 470.144 1441.7 471.169 1443.4 473.218C1445.14 475.229 1446.01 478.109 1446.01 481.86V501H1439.68V482.672C1439.68 477.955 1437.79 475.596 1434 475.596C1433.19 475.596 1432.38 475.712 1431.56 475.944C1430.79 476.137 1430.1 476.447 1429.48 476.872C1428.86 477.297 1428.36 477.839 1427.97 478.496C1427.62 479.153 1427.45 479.927 1427.45 480.816V501H1421.12Z" fill="white" style="fill:white;fill-opacity:1;"/>
+</svg>
diff --git a/_static/img/arch-logo.png b/_static/img/arch-logo.png
deleted file mode 100755
index bbffb318..00000000
Binary files a/_static/img/arch-logo.png and /dev/null differ
diff --git a/_static/img/arch-nav-logo.png b/_static/img/arch-nav-logo.png
deleted file mode 100755
index 5a1a7776..00000000
Binary files a/_static/img/arch-nav-logo.png and /dev/null differ
diff --git a/_static/img/arch-system-architecture.jpg b/_static/img/arch-system-architecture.jpg
deleted file mode 100755
index 3c8839a7..00000000
Binary files a/_static/img/arch-system-architecture.jpg and /dev/null differ
diff --git a/_static/img/arch_network_diagram_high_level.png b/_static/img/arch_network_diagram_high_level.png
deleted file mode 100755
index e83e7165..00000000
Binary files a/_static/img/arch_network_diagram_high_level.png and /dev/null differ
diff --git a/_static/img/network-topology-agent.jpg b/_static/img/network-topology-agent.jpg
deleted file mode 100755
index 50ba9a64..00000000
Binary files a/_static/img/network-topology-agent.jpg and /dev/null differ
diff --git a/_static/img/network-topology-ingress-egress.jpg b/_static/img/network-topology-ingress-egress.jpg
deleted file mode 100755
index 03e36e77..00000000
Binary files a/_static/img/network-topology-ingress-egress.jpg and /dev/null differ
diff --git a/_static/img/network-topology-ingress-egress.png b/_static/img/network-topology-ingress-egress.png
new file mode 100755
index 00000000..c2e55584
Binary files /dev/null and b/_static/img/network-topology-ingress-egress.png differ
diff --git a/_static/img/plano-system-architecture.png b/_static/img/plano-system-architecture.png
new file mode 100755
index 00000000..792477d5
Binary files /dev/null and b/_static/img/plano-system-architecture.png differ
diff --git a/_static/img/plano_network_diagram_high_level.png b/_static/img/plano_network_diagram_high_level.png
new file mode 100755
index 00000000..da1c5b92
Binary files /dev/null and b/_static/img/plano_network_diagram_high_level.png differ
diff --git a/_static/img/tracing.png b/_static/img/tracing.png
index 91d6a82b..bb34db91 100755
Binary files a/_static/img/tracing.png and b/_static/img/tracing.png differ
diff --git a/build_with_arch/agent.html b/build_with_arch/agent.html
deleted file mode 100755
index 58e724de..00000000
--- a/build_with_arch/agent.html
+++ /dev/null
@@ -1,432 +0,0 @@
-<!DOCTYPE html>
-
-<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
-<head>
-<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
-<meta charset="utf-8"/>
-<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
-<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
-<meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Agentic Apps | Arch Docs v0.3.22</title>
-<meta content="Agentic Apps | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Agentic Apps | Arch Docs v0.3.22" name="twitter:title"/>
-<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
-<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
-<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
-<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/build_with_arch/agent.html" rel="canonical"/>
-<link href="../_static/favicon.ico" rel="icon"/>
-<link href="../search.html" rel="search" title="Search"/>
-<link href="rag.html" rel="next" title="RAG Apps"/>
-<link href="../guides/observability/access_logging.html" rel="prev" title="Access Logging"/>
-<script>
-    <!-- Prevent Flash of wrong theme -->
-      const userPreference = localStorage.getItem('darkMode');
-      let mode;
-      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
-        mode = 'dark';
-        document.documentElement.classList.add('dark');
-      } else {
-        mode = 'light';
-      }
-      if (!userPreference) {localStorage.setItem('darkMode', mode)}
-    </script>
-</head>
-<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
-<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
-      Skip to content
-    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
-<div class="hidden mr-4 md:flex">
-<a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
-<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
-</svg>
-<span class="sr-only">Toggle navigation menu</span>
-</button>
-<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
-<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
-<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
-<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
-<span class="text-xs">⌘</span>
-    K
-  </kbd>
-</form>
-</div>
-<nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
-<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
-<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
-</div>
-</a>
-<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
-<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
-</svg>
-<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
-</svg>
-</button>
-</nav>
-</div>
-</div>
-</header>
-<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
-<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a>
-<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
-<div class="overflow-y-auto h-full w-full relative pr-6">
-
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
-<script>
-  window.dataLayer = window.dataLayer || [];
-  function gtag(){dataLayer.push(arguments);}
-  gtag('js', new Date());
-
-  gtag('config', 'G-K2LXXSX6HB');
-</script>
-<nav class="table w-full min-w-full my-6 lg:my-8">
-<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../concepts/prompt_target.html">Prompt Target</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
-</ul>
-</li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul class="current">
-<li class="toctree-l1 current"><a class="current reference internal" href="#">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="multi_turn.html">Multi-Turn</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
-</ul>
-</nav>
-</div>
-</div>
-<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
-<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
-</svg>
-</button>
-</aside>
-<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
-<div class="w-full min-w-0 mx-auto">
-<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
-<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
-<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
-<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
-</svg>
-</a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Agentic Apps</span>
-</nav>
-<div id="content" role="main">
-<section id="agentic-apps">
-<span id="arch-agent-guide"></span><h1>Agentic Apps<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#agentic-apps"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Arch helps you build personalized agentic applications by calling application-specific (API) functions via user prompts.
-This involves any predefined functions or APIs you want to expose to users to perform tasks, gather information,
-or manipulate data. This capability is generally referred to as <a class="reference internal" href="../guides/function_calling.html#function-calling"><span class="std std-ref">function calling</span></a>, where
-you can support “agentic” apps tailored to specific use cases - from updating insurance claims to creating ad campaigns - via prompts.</p>
-<p>Arch analyzes prompts, extracts critical information from prompts, engages in lightweight conversation with the user to
-gather any missing parameters and makes API calls so that you can focus on writing business logic. Arch does this via its
-purpose-built <a class="reference external" href="https://huggingface.co/collections/katanemo/arch-function-66f209a693ea8df14317ad68" rel="nofollow noopener">Arch-Function<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> -
-the fastest (200ms p50 - 12x faser than GPT-4o) and cheapest (44x than GPT-4o) function calling LLM that matches or outperforms
-frontier LLMs.</p>
-<a class="reference internal image-reference" href="../_images/function-calling-flow.jpg"><img alt="../_images/function-calling-flow.jpg" class="align-center" src="../_images/function-calling-flow.jpg" style="width: 100%;"/>
-</a>
-<section id="single-function-call">
-<h2>Single Function Call<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#single-function-call" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#single-function-call'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>In the most common scenario, users will request a single action via prompts, and Arch efficiently processes the
-request by extracting relevant parameters, validating the input, and calling the designated function or API. Here
-is how you would go about enabling this scenario with Arch:</p>
-<section id="step-1-define-prompt-targets">
-<h3>Step 1: Define Prompt Targets<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-1-define-prompt-targets" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-1-define-prompt-targets'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<div class="literal-block-wrapper docutils container" id="id2">
-<div class="code-block-caption"><span class="caption-text">Prompt Target Example Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1</span>
-</span><span id="line-2"><span class="linenos"> 2</span><span class="nt">listener</span><span class="p">:</span>
-</span><span id="line-3"><span class="linenos"> 3</span><span class="w">  </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1</span>
-</span><span id="line-4"><span class="linenos"> 4</span><span class="w">  </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">8080</span><span class="w"> </span><span class="c1">#If you configure port 443, you'll need to update the listener with tls_certificates</span>
-</span><span id="line-5"><span class="linenos"> 5</span><span class="w">  </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">huggingface</span>
-</span><span id="line-6"><span class="linenos"> 6</span>
-</span><span id="line-7"><span class="linenos"> 7</span><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
-</span><span id="line-8"><span class="linenos"> 8</span><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-9"><span class="linenos"> 9</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">OpenAI</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="w">    </span><span class="nt">provider</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-11"><span class="linenos">11</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-12"><span class="linenos">12</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-3.5-turbo</span>
-</span><span id="line-13"><span class="linenos">13</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-14"><span class="linenos">14</span>
-</span><span id="line-15"><span class="linenos">15</span><span class="c1"># default system prompt used by all prompt targets</span>
-</span><span id="line-16"><span class="linenos">16</span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
-</span><span id="line-17"><span class="linenos">17</span><span class="w">  </span><span class="no">You are a network assistant that just offers facts; not advice on manufacturers or purchasing decisions.</span>
-</span><span id="line-18"><span class="linenos">18</span>
-</span><span id="line-19"><mark><span class="linenos">19</span><span class="nt">prompt_targets</span><span class="p">:</span>
-</mark></span><span id="line-20"><mark><span class="linenos">20</span><span class="w">    </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">network_qa</span>
-</mark></span><span id="line-21"><mark><span class="linenos">21</span><span class="w">      </span><span class="nt">endpoint</span><span class="p">:</span>
-</mark></span><span id="line-22"><mark><span class="linenos">22</span><span class="w">        </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</mark></span><span id="line-23"><mark><span class="linenos">23</span><span class="w">        </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/network_summary</span>
-</mark></span><span id="line-24"><mark><span class="linenos">24</span><span class="w">      </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Handle general Q/A related to networking.</span>
-</mark></span><span id="line-25"><mark><span class="linenos">25</span><span class="w">      </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</mark></span><span id="line-26"><mark><span class="linenos">26</span><span class="w">    </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">reboot_devices</span>
-</mark></span><span id="line-27"><mark><span class="linenos">27</span><span class="w">      </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Reboot specific devices or device groups</span>
-</mark></span><span id="line-28"><mark><span class="linenos">28</span><span class="w">      </span><span class="nt">endpoint</span><span class="p">:</span>
-</mark></span><span id="line-29"><mark><span class="linenos">29</span><span class="w">        </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</mark></span><span id="line-30"><mark><span class="linenos">30</span><span class="w">        </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/device_reboot</span>
-</mark></span><span id="line-31"><mark><span class="linenos">31</span><span class="w">      </span><span class="nt">parameters</span><span class="p">:</span>
-</mark></span><span id="line-32"><mark><span class="linenos">32</span><span class="w">        </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_ids</span>
-</mark></span><span id="line-33"><mark><span class="linenos">33</span><span class="w">          </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">list</span>
-</mark></span><span id="line-34"><mark><span class="linenos">34</span><span class="w">          </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">A list of device identifiers (IDs) to reboot.</span>
-</mark></span><span id="line-35"><mark><span class="linenos">35</span><span class="w">          </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</mark></span><span id="line-36"><mark><span class="linenos">36</span><span class="w">    </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_summary</span>
-</mark></span><span id="line-37"><mark><span class="linenos">37</span><span class="w">      </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Retrieve statistics for specific devices within a time range</span>
-</mark></span><span id="line-38"><mark><span class="linenos">38</span><span class="w">      </span><span class="nt">endpoint</span><span class="p">:</span>
-</mark></span><span id="line-39"><mark><span class="linenos">39</span><span class="w">        </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</mark></span><span id="line-40"><mark><span class="linenos">40</span><span class="w">        </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/device_summary</span>
-</mark></span><span id="line-41"><mark><span class="linenos">41</span><span class="w">      </span><span class="nt">parameters</span><span class="p">:</span>
-</mark></span><span id="line-42"><mark><span class="linenos">42</span><span class="w">        </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_ids</span>
-</mark></span><span id="line-43"><mark><span class="linenos">43</span><span class="w">          </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">list</span>
-</mark></span><span id="line-44"><mark><span class="linenos">44</span><span class="w">          </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">A list of device identifiers (IDs) to retrieve statistics for.</span>
-</mark></span><span id="line-45"><mark><span class="linenos">45</span><span class="w">          </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span><span class="w">  </span><span class="c1"># device_ids are required to get device statistics</span>
-</mark></span><span id="line-46"><mark><span class="linenos">46</span><span class="w">        </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">time_range</span>
-</mark></span><span id="line-47"><mark><span class="linenos">47</span><span class="w">          </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">int</span>
-</mark></span><span id="line-48"><mark><span class="linenos">48</span><span class="w">          </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Time range in days for which to gather device statistics. Defaults to 7.</span>
-</mark></span><span id="line-49"><mark><span class="linenos">49</span><span class="w">          </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">7</span>
-</mark></span><span id="line-50"><span class="linenos">50</span>
-</span><span id="line-51"><span class="linenos">51</span><span class="c1"># Arch creates a round-robin load balancing between different endpoints, managed via the cluster subsystem.</span>
-</span><span id="line-52"><span class="linenos">52</span><span class="nt">endpoints</span><span class="p">:</span>
-</span><span id="line-53"><span class="linenos">53</span><span class="w">  </span><span class="nt">app_server</span><span class="p">:</span>
-</span><span id="line-54"><span class="linenos">54</span><span class="w">    </span><span class="c1"># value could be ip address or a hostname with port</span>
-</span><span id="line-55"><span class="linenos">55</span><span class="w">    </span><span class="c1"># this could also be a list of endpoints for load balancing</span>
-</span><span id="line-56"><span class="linenos">56</span><span class="w">    </span><span class="c1"># for example endpoint: [ ip1:port, ip2:port ]</span>
-</span><span id="line-57"><span class="linenos">57</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">host.docker.internal:18083</span>
-</span><span id="line-58"><span class="linenos">58</span><span class="w">    </span><span class="c1"># max time to wait for a connection to be established</span>
-</span><span id="line-59"><span class="linenos">59</span><span class="w">    </span><span class="nt">connect_timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.005s</span>
-</span></code></pre></div>
-</div>
-</div>
-</section>
-<section id="step-2-process-request-parameters">
-<h3>Step 2: Process Request Parameters<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-process-request-parameters" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-2-process-request-parameters'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Once the prompt targets are configured as above, handling those parameters is</p>
-<div class="literal-block-wrapper docutils container" id="id3">
-<div class="code-block-caption"><span class="caption-text">Parameter handling with Flask</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id3"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="kn">from</span><span class="w"> </span><span class="nn">flask</span><span class="w"> </span><span class="kn">import</span> <span class="n">Flask</span><span class="p">,</span> <span class="n">request</span><span class="p">,</span> <span class="n">jsonify</span>
-</span><span id="line-2"><span class="linenos"> 2</span>
-</span><span id="line-3"><span class="linenos"> 3</span><span class="n">app</span> <span class="o">=</span> <span class="n">Flask</span><span class="p">(</span><span class="vm">__name__</span><span class="p">)</span>
-</span><span id="line-4"><span class="linenos"> 4</span>
-</span><span id="line-5"><span class="linenos"> 5</span>
-</span><span id="line-6"><span class="linenos"> 6</span><span class="nd">@app</span><span class="o">.</span><span class="n">route</span><span class="p">(</span><span class="s2">"/agent/device_summary"</span><span class="p">,</span> <span class="n">methods</span><span class="o">=</span><span class="p">[</span><span class="s2">"POST"</span><span class="p">])</span>
-</span><span id="line-7"><span class="linenos"> 7</span><span class="k">def</span><span class="w"> </span><span class="nf">get_device_summary</span><span class="p">():</span>
-</span><span id="line-8"><span class="linenos"> 8</span><span class="w">    </span><span class="sd">"""</span>
-</span><span id="line-9"><span class="linenos"> 9</span><span class="sd">    Endpoint to retrieve device statistics based on device IDs and an optional time range.</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="sd">    """</span>
-</span><span id="line-11"><span class="linenos">11</span>    <span class="n">data</span> <span class="o">=</span> <span class="n">request</span><span class="o">.</span><span class="n">get_json</span><span class="p">()</span>
-</span><span id="line-12"><span class="linenos">12</span>
-</span><span id="line-13"><span class="linenos">13</span>    <span class="c1"># Validate 'device_ids' parameter</span>
-</span><span id="line-14"><span class="linenos">14</span>    <span class="n">device_ids</span> <span class="o">=</span> <span class="n">data</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"device_ids"</span><span class="p">)</span>
-</span><span id="line-15"><span class="linenos">15</span>    <span class="k">if</span> <span class="ow">not</span> <span class="n">device_ids</span> <span class="ow">or</span> <span class="ow">not</span> <span class="nb">isinstance</span><span class="p">(</span><span class="n">device_ids</span><span class="p">,</span> <span class="nb">list</span><span class="p">):</span>
-</span><span id="line-16"><span class="linenos">16</span>        <span class="k">return</span> <span class="p">(</span>
-</span><span id="line-17"><span class="linenos">17</span>            <span class="n">jsonify</span><span class="p">({</span><span class="s2">"error"</span><span class="p">:</span> <span class="s2">"'device_ids' parameter is required and must be a list"</span><span class="p">}),</span>
-</span><span id="line-18"><span class="linenos">18</span>            <span class="mi">400</span><span class="p">,</span>
-</span><span id="line-19"><span class="linenos">19</span>        <span class="p">)</span>
-</span><span id="line-20"><span class="linenos">20</span>
-</span><span id="line-21"><span class="linenos">21</span>    <span class="c1"># Validate 'time_range' parameter (optional, defaults to 7)</span>
-</span><span id="line-22"><span class="linenos">22</span>    <span class="n">time_range</span> <span class="o">=</span> <span class="n">data</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"time_range"</span><span class="p">,</span> <span class="mi">7</span><span class="p">)</span>
-</span><span id="line-23"><span class="linenos">23</span>    <span class="k">if</span> <span class="ow">not</span> <span class="nb">isinstance</span><span class="p">(</span><span class="n">time_range</span><span class="p">,</span> <span class="nb">int</span><span class="p">):</span>
-</span><span id="line-24"><span class="linenos">24</span>        <span class="k">return</span> <span class="n">jsonify</span><span class="p">({</span><span class="s2">"error"</span><span class="p">:</span> <span class="s2">"'time_range' must be an integer"</span><span class="p">}),</span> <span class="mi">400</span>
-</span><span id="line-25"><span class="linenos">25</span>
-</span><span id="line-26"><span class="linenos">26</span>    <span class="c1"># Simulate retrieving statistics for the given device IDs and time range</span>
-</span><span id="line-27"><span class="linenos">27</span>    <span class="c1"># In a real application, you would query your database or external service here</span>
-</span><span id="line-28"><span class="linenos">28</span>    <span class="n">statistics</span> <span class="o">=</span> <span class="p">[]</span>
-</span><span id="line-29"><span class="linenos">29</span>    <span class="k">for</span> <span class="n">device_id</span> <span class="ow">in</span> <span class="n">device_ids</span><span class="p">:</span>
-</span><span id="line-30"><span class="linenos">30</span>        <span class="c1"># Placeholder for actual data retrieval</span>
-</span><span id="line-31"><span class="linenos">31</span>        <span class="n">stats</span> <span class="o">=</span> <span class="p">{</span>
-</span><span id="line-32"><span class="linenos">32</span>            <span class="s2">"device_id"</span><span class="p">:</span> <span class="n">device_id</span><span class="p">,</span>
-</span><span id="line-33"><span class="linenos">33</span>            <span class="s2">"time_range"</span><span class="p">:</span> <span class="sa">f</span><span class="s2">"Last </span><span class="si">{</span><span class="n">time_range</span><span class="si">}</span><span class="s2"> days"</span><span class="p">,</span>
-</span><span id="line-34"><span class="linenos">34</span>            <span class="s2">"data"</span><span class="p">:</span> <span class="sa">f</span><span class="s2">"Statistics data for device </span><span class="si">{</span><span class="n">device_id</span><span class="si">}</span><span class="s2"> over the last </span><span class="si">{</span><span class="n">time_range</span><span class="si">}</span><span class="s2"> days."</span><span class="p">,</span>
-</span><span id="line-35"><span class="linenos">35</span>        <span class="p">}</span>
-</span><span id="line-36"><span class="linenos">36</span>        <span class="n">statistics</span><span class="o">.</span><span class="n">append</span><span class="p">(</span><span class="n">stats</span><span class="p">)</span>
-</span><span id="line-37"><span class="linenos">37</span>
-</span><span id="line-38"><span class="linenos">38</span>    <span class="n">response</span> <span class="o">=</span> <span class="p">{</span><span class="s2">"statistics"</span><span class="p">:</span> <span class="n">statistics</span><span class="p">}</span>
-</span><span id="line-39"><span class="linenos">39</span>
-</span><span id="line-40"><span class="linenos">40</span>    <span class="k">return</span> <span class="n">jsonify</span><span class="p">(</span><span class="n">response</span><span class="p">),</span> <span class="mi">200</span>
-</span><span id="line-41"><span class="linenos">41</span>
-</span><span id="line-42"><span class="linenos">42</span>
-</span><span id="line-43"><span class="linenos">43</span><span class="k">if</span> <span class="vm">__name__</span> <span class="o">==</span> <span class="s2">"__main__"</span><span class="p">:</span>
-</span><span id="line-44"><span class="linenos">44</span>    <span class="n">app</span><span class="o">.</span><span class="n">run</span><span class="p">(</span><span class="n">debug</span><span class="o">=</span><span class="kc">True</span><span class="p">)</span>
-</span></code></pre></div>
-</div>
-</div>
-</section>
-</section>
-<section id="parallel-multiple-function-calling">
-<h2>Parallel &amp; Multiple Function Calling<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#parallel-multiple-function-calling" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#parallel-multiple-function-calling'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>In more complex use cases, users may request multiple actions or need multiple APIs/functions to be called
-simultaneously or sequentially. With Arch, you can handle these scenarios efficiently using parallel or multiple
-function calling. This allows your application to engage in a broader range of interactions, such as updating
-different datasets, triggering events across systems, or collecting results from multiple services in one prompt.</p>
-<p>Arch-FC1B is built to manage these parallel tasks efficiently, ensuring low latency and high throughput, even
-when multiple functions are invoked. It provides two mechanisms to handle these cases:</p>
-<section id="id1">
-<h3>Step 1: Define Prompt Targets<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id1" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#id1'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>When enabling multiple function calling, define the prompt targets in a way that supports multiple functions or
-API calls based on the user’s prompt. These targets can be triggered in parallel or sequentially, depending on
-the user’s intent.</p>
-<p>Example of Multiple Prompt Targets in YAML:</p>
-<div class="literal-block-wrapper docutils container" id="id4">
-<div class="code-block-caption"><span class="caption-text">Prompt Target Example Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id4"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1</span>
-</span><span id="line-2"><span class="linenos"> 2</span><span class="nt">listener</span><span class="p">:</span>
-</span><span id="line-3"><span class="linenos"> 3</span><span class="w">  </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1</span>
-</span><span id="line-4"><span class="linenos"> 4</span><span class="w">  </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">8080</span><span class="w"> </span><span class="c1">#If you configure port 443, you'll need to update the listener with tls_certificates</span>
-</span><span id="line-5"><span class="linenos"> 5</span><span class="w">  </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">huggingface</span>
-</span><span id="line-6"><span class="linenos"> 6</span>
-</span><span id="line-7"><span class="linenos"> 7</span><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
-</span><span id="line-8"><span class="linenos"> 8</span><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-9"><span class="linenos"> 9</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">OpenAI</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="w">    </span><span class="nt">provider</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-11"><span class="linenos">11</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-12"><span class="linenos">12</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-3.5-turbo</span>
-</span><span id="line-13"><span class="linenos">13</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-14"><span class="linenos">14</span>
-</span><span id="line-15"><span class="linenos">15</span><span class="c1"># default system prompt used by all prompt targets</span>
-</span><span id="line-16"><span class="linenos">16</span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
-</span><span id="line-17"><span class="linenos">17</span><span class="w">  </span><span class="no">You are a network assistant that just offers facts; not advice on manufacturers or purchasing decisions.</span>
-</span><span id="line-18"><span class="linenos">18</span>
-</span><span id="line-19"><mark><span class="linenos">19</span><span class="nt">prompt_targets</span><span class="p">:</span>
-</mark></span><span id="line-20"><mark><span class="linenos">20</span><span class="w">    </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">network_qa</span>
-</mark></span><span id="line-21"><mark><span class="linenos">21</span><span class="w">      </span><span class="nt">endpoint</span><span class="p">:</span>
-</mark></span><span id="line-22"><mark><span class="linenos">22</span><span class="w">        </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</mark></span><span id="line-23"><mark><span class="linenos">23</span><span class="w">        </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/network_summary</span>
-</mark></span><span id="line-24"><mark><span class="linenos">24</span><span class="w">      </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Handle general Q/A related to networking.</span>
-</mark></span><span id="line-25"><mark><span class="linenos">25</span><span class="w">      </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</mark></span><span id="line-26"><mark><span class="linenos">26</span><span class="w">    </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">reboot_devices</span>
-</mark></span><span id="line-27"><mark><span class="linenos">27</span><span class="w">      </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Reboot specific devices or device groups</span>
-</mark></span><span id="line-28"><mark><span class="linenos">28</span><span class="w">      </span><span class="nt">endpoint</span><span class="p">:</span>
-</mark></span><span id="line-29"><mark><span class="linenos">29</span><span class="w">        </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</mark></span><span id="line-30"><mark><span class="linenos">30</span><span class="w">        </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/device_reboot</span>
-</mark></span><span id="line-31"><mark><span class="linenos">31</span><span class="w">      </span><span class="nt">parameters</span><span class="p">:</span>
-</mark></span><span id="line-32"><mark><span class="linenos">32</span><span class="w">        </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_ids</span>
-</mark></span><span id="line-33"><mark><span class="linenos">33</span><span class="w">          </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">list</span>
-</mark></span><span id="line-34"><mark><span class="linenos">34</span><span class="w">          </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">A list of device identifiers (IDs) to reboot.</span>
-</mark></span><span id="line-35"><mark><span class="linenos">35</span><span class="w">          </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</mark></span><span id="line-36"><mark><span class="linenos">36</span><span class="w">    </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_summary</span>
-</mark></span><span id="line-37"><mark><span class="linenos">37</span><span class="w">      </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Retrieve statistics for specific devices within a time range</span>
-</mark></span><span id="line-38"><mark><span class="linenos">38</span><span class="w">      </span><span class="nt">endpoint</span><span class="p">:</span>
-</mark></span><span id="line-39"><mark><span class="linenos">39</span><span class="w">        </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</mark></span><span id="line-40"><mark><span class="linenos">40</span><span class="w">        </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/device_summary</span>
-</mark></span><span id="line-41"><mark><span class="linenos">41</span><span class="w">      </span><span class="nt">parameters</span><span class="p">:</span>
-</mark></span><span id="line-42"><mark><span class="linenos">42</span><span class="w">        </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_ids</span>
-</mark></span><span id="line-43"><mark><span class="linenos">43</span><span class="w">          </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">list</span>
-</mark></span><span id="line-44"><mark><span class="linenos">44</span><span class="w">          </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">A list of device identifiers (IDs) to retrieve statistics for.</span>
-</mark></span><span id="line-45"><mark><span class="linenos">45</span><span class="w">          </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span><span class="w">  </span><span class="c1"># device_ids are required to get device statistics</span>
-</mark></span><span id="line-46"><mark><span class="linenos">46</span><span class="w">        </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">time_range</span>
-</mark></span><span id="line-47"><mark><span class="linenos">47</span><span class="w">          </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">int</span>
-</mark></span><span id="line-48"><mark><span class="linenos">48</span><span class="w">          </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Time range in days for which to gather device statistics. Defaults to 7.</span>
-</mark></span><span id="line-49"><mark><span class="linenos">49</span><span class="w">          </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">7</span>
-</mark></span><span id="line-50"><span class="linenos">50</span>
-</span><span id="line-51"><span class="linenos">51</span><span class="c1"># Arch creates a round-robin load balancing between different endpoints, managed via the cluster subsystem.</span>
-</span><span id="line-52"><span class="linenos">52</span><span class="nt">endpoints</span><span class="p">:</span>
-</span><span id="line-53"><span class="linenos">53</span><span class="w">  </span><span class="nt">app_server</span><span class="p">:</span>
-</span><span id="line-54"><span class="linenos">54</span><span class="w">    </span><span class="c1"># value could be ip address or a hostname with port</span>
-</span><span id="line-55"><span class="linenos">55</span><span class="w">    </span><span class="c1"># this could also be a list of endpoints for load balancing</span>
-</span><span id="line-56"><span class="linenos">56</span><span class="w">    </span><span class="c1"># for example endpoint: [ ip1:port, ip2:port ]</span>
-</span><span id="line-57"><span class="linenos">57</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">host.docker.internal:18083</span>
-</span><span id="line-58"><span class="linenos">58</span><span class="w">    </span><span class="c1"># max time to wait for a connection to be established</span>
-</span><span id="line-59"><span class="linenos">59</span><span class="w">    </span><span class="nt">connect_timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.005s</span>
-</span></code></pre></div>
-</div>
-</div>
-</section>
-</section>
-</section>
-</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
-<div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../guides/observability/access_logging.html">
-<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="15 18 9 12 15 6"></polyline>
-</svg>
-        Access Logging
-      </a>
-</div>
-<div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="rag.html">
-        RAG Apps
-        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="9 18 15 12 9 6"></polyline>
-</svg>
-</a>
-</div>
-</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
-<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
-<ul>
-<li><a :data-current="activeSection === '#single-function-call'" class="reference internal" href="#single-function-call">Single Function Call</a><ul>
-<li><a :data-current="activeSection === '#step-1-define-prompt-targets'" class="reference internal" href="#step-1-define-prompt-targets">Step 1: Define Prompt Targets</a></li>
-<li><a :data-current="activeSection === '#step-2-process-request-parameters'" class="reference internal" href="#step-2-process-request-parameters">Step 2: Process Request Parameters</a></li>
-</ul>
-</li>
-<li><a :data-current="activeSection === '#parallel-multiple-function-calling'" class="reference internal" href="#parallel-multiple-function-calling">Parallel &amp; Multiple Function Calling</a><ul>
-<li><a :data-current="activeSection === '#id1'" class="reference internal" href="#id1">Step 1: Define Prompt Targets</a></li>
-</ul>
-</li>
-</ul>
-</div>
-</aside>
-</main>
-</div>
-</div><footer class="py-6 border-t border-border md:py-0">
-<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
-<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
-</div>
-</div>
-</footer>
-</div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
-<script src="../_static/doctools.js?v=9bcbadda"></script>
-<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
-<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
-<script src="../_static/design-tabs.js?v=f930bc37"></script>
-</body>
-</html>
\ No newline at end of file
diff --git a/build_with_arch/multi_turn.html b/build_with_arch/multi_turn.html
deleted file mode 100755
index 93c52eba..00000000
--- a/build_with_arch/multi_turn.html
+++ /dev/null
@@ -1,368 +0,0 @@
-<!DOCTYPE html>
-
-<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
-<head>
-<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
-<meta charset="utf-8"/>
-<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
-<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
-<meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Multi-Turn | Arch Docs v0.3.22</title>
-<meta content="Multi-Turn | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Multi-Turn | Arch Docs v0.3.22" name="twitter:title"/>
-<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
-<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
-<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
-<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/build_with_arch/multi_turn.html" rel="canonical"/>
-<link href="../_static/favicon.ico" rel="icon"/>
-<link href="../search.html" rel="search" title="Search"/>
-<link href="../resources/deployment.html" rel="next" title="Deployment"/>
-<link href="rag.html" rel="prev" title="RAG Apps"/>
-<script>
-    <!-- Prevent Flash of wrong theme -->
-      const userPreference = localStorage.getItem('darkMode');
-      let mode;
-      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
-        mode = 'dark';
-        document.documentElement.classList.add('dark');
-      } else {
-        mode = 'light';
-      }
-      if (!userPreference) {localStorage.setItem('darkMode', mode)}
-    </script>
-</head>
-<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
-<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
-      Skip to content
-    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
-<div class="hidden mr-4 md:flex">
-<a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
-<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
-</svg>
-<span class="sr-only">Toggle navigation menu</span>
-</button>
-<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
-<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
-<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
-<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
-<span class="text-xs">⌘</span>
-    K
-  </kbd>
-</form>
-</div>
-<nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
-<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
-<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
-</div>
-</a>
-<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
-<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
-</svg>
-<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
-</svg>
-</button>
-</nav>
-</div>
-</div>
-</header>
-<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
-<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a>
-<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
-<div class="overflow-y-auto h-full w-full relative pr-6">
-
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
-<script>
-  window.dataLayer = window.dataLayer || [];
-  function gtag(){dataLayer.push(arguments);}
-  gtag('js', new Date());
-
-  gtag('config', 'G-K2LXXSX6HB');
-</script>
-<nav class="table w-full min-w-full my-6 lg:my-8">
-<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../concepts/prompt_target.html">Prompt Target</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
-</ul>
-</li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul class="current">
-<li class="toctree-l1"><a class="reference internal" href="agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="rag.html">RAG Apps</a></li>
-<li class="toctree-l1 current"><a class="current reference internal" href="#">Multi-Turn</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
-</ul>
-</nav>
-</div>
-</div>
-<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
-<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
-</svg>
-</button>
-</aside>
-<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
-<div class="w-full min-w-0 mx-auto">
-<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
-<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
-<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
-<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
-</svg>
-</a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Multi-Turn</span>
-</nav>
-<div id="content" role="main">
-<section id="multi-turn">
-<span id="arch-multi-turn-guide"></span><h1>Multi-Turn<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#multi-turn"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Developers often <a class="reference external" href="https://www.reddit.com/r/LocalLLaMA/comments/18mqwg6/best_practice_for_rag_with_followup_chat/" rel="nofollow noopener">struggle<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> to efficiently handle
-<code class="docutils literal notranslate"><span class="pre">follow-up</span></code> or <code class="docutils literal notranslate"><span class="pre">clarification</span></code> questions. Specifically, when users ask for changes or additions to previous responses, it requires developers to
-re-write prompts using LLMs with precise prompt engineering techniques. This process is slow, manual, error prone and adds latency and token cost for
-common scenarios that can be managed more efficiently.</p>
-<p>Arch is highly capable of accurately detecting and processing prompts in multi-turn scenarios so that you can buil fast and accurate agents in minutes.
-Below are some cnversational examples that you can build via Arch. Each example is enriched with annotations (via ** [Arch] ** ) that illustrates how Arch
-processess conversational messages on your behalf.</p>
-<div class="admonition note">
-<p class="admonition-title">Note</p>
-<p>The following section assumes that you have some knowledge about the core concepts of Arch, such as <a class="reference internal" href="../concepts/tech_overview/prompt.html#arch-overview-prompt-handling"><span class="std std-ref">prompt_targets</span></a>.
-If you haven’t familizaried yourself with Arch’s concepts, we recommend you first read the <a class="reference internal" href="../concepts/tech_overview/tech_overview.html#tech-overview"><span class="std std-ref">tech overview</span></a> section firtst.
-Additionally, the conversation examples below assume the usage of the following <a class="reference internal" href="#multi-turn-subsection-prompt-target"><span class="std std-ref">arch_config.yaml</span></a> file.</p>
-</div>
-<section id="example-1-adjusting-retrieval">
-<h2>Example 1: Adjusting Retrieval<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-1-adjusting-retrieval" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-1-adjusting-retrieval'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<div class="highlight-text notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">User: What are the benefits of renewable energy?
-</span><span id="line-2">**[Arch]**: Check if there is an available &lt;prompt_target&gt; that can handle this user query.
-</span><span id="line-3">**[Arch]**: Found "get_info_for_energy_source" prompt_target in arch_config.yaml. Forward prompt to the endpoint configured in "get_info_for_energy_source"
-</span><span id="line-4">...
-</span><span id="line-5">Assistant: Renewable energy reduces greenhouse gas emissions, lowers air pollution, and provides sustainable power sources like solar and wind.
-</span><span id="line-6">
-</span><span id="line-7">User: Include cost considerations in the response.
-</span><span id="line-8">**[Arch]**: Follow-up detected. Forward prompt history to the "get_info_for_energy_source" prompt_target and post the following parameters consideration="cost"
-</span><span id="line-9">...
-</span><span id="line-10">Assistant: Renewable energy reduces greenhouse gas emissions, lowers air pollution, and provides sustainable power sources like solar and wind. While the initial setup costs can be high, long-term savings from reduced fuel expenses and government incentives make it cost-effective.
-</span></code></pre></div>
-</div>
-</section>
-<section id="example-2-switching-intent">
-<h2>Example 2: Switching Intent<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-2-switching-intent" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-2-switching-intent'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<div class="highlight-text notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">User: What are the symptoms of diabetes?
-</span><span id="line-2">**[Arch]**: Check if there is an available &lt;prompt_target&gt; that can handle this user query.
-</span><span id="line-3">**[Arch]**: Found "diseases_symptoms" prompt_target in arch_config.yaml. Forward disease=diabeteres to "diseases_symptoms" prompt target
-</span><span id="line-4">...
-</span><span id="line-5">Assistant: Common symptoms include frequent urination, excessive thirst, fatigue, and blurry vision.
-</span><span id="line-6">
-</span><span id="line-7">User: How is it diagnosed?
-</span><span id="line-8">**[Arch]**: New intent detected.
-</span><span id="line-9">**[Arch]**: Found "disease_diagnoses" prompt_target in arch_config.yaml. Forward disease=diabeteres to "disease_diagnoses" prompt target
-</span><span id="line-10">...
-</span><span id="line-11">Assistant: Diabetes is diagnosed through blood tests like fasting blood sugar, A1C, or an oral glucose tolerance test.
-</span></code></pre></div>
-</div>
-</section>
-<section id="build-multi-turn-rag-apps">
-<h2>Build Multi-Turn RAG Apps<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#build-multi-turn-rag-apps" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#build-multi-turn-rag-apps'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>The following section describes how you can easilly add support for multi-turn scenarios via Arch. You process and manage multi-turn prompts
-just like you manage single-turn ones. Arch handles the conpleixity of detecting the correct intent based on the last user prompt and
-the covnersational history, extracts relevant parameters needed by downstream APIs, and dipatches calls to any upstream LLMs to summarize the
-response from your APIs.</p>
-<section id="step-1-define-arch-config">
-<span id="multi-turn-subsection-prompt-target"></span><h3>Step 1: Define Arch Config<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-1-define-arch-config" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-1-define-arch-config'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<div class="literal-block-wrapper docutils container" id="id2">
-<div class="code-block-caption"><span class="caption-text">Arch Config</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1</span>
-</span><span id="line-2"><span class="linenos"> 2</span><span class="nt">listener</span><span class="p">:</span>
-</span><span id="line-3"><span class="linenos"> 3</span><span class="w">  </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1</span>
-</span><span id="line-4"><span class="linenos"> 4</span><span class="w">  </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">8080</span><span class="w"> </span><span class="c1">#If you configure port 443, you'll need to update the listener with tls_certificates</span>
-</span><span id="line-5"><span class="linenos"> 5</span><span class="w">  </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">huggingface</span>
-</span><span id="line-6"><span class="linenos"> 6</span>
-</span><span id="line-7"><span class="linenos"> 7</span><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
-</span><span id="line-8"><span class="linenos"> 8</span><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-9"><span class="linenos"> 9</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">OpenAI</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="w">    </span><span class="nt">provider</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-11"><span class="linenos">11</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-12"><span class="linenos">12</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-3.5-turbo</span>
-</span><span id="line-13"><span class="linenos">13</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-14"><span class="linenos">14</span>
-</span><span id="line-15"><span class="linenos">15</span><span class="c1"># default system prompt used by all prompt targets</span>
-</span><span id="line-16"><span class="linenos">16</span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
-</span><span id="line-17"><span class="linenos">17</span><span class="w">   </span><span class="no">You are a helpful assistant and can offer information about energy sources. You will get a JSON object with energy_source and consideration fields. Focus on answering using those fields</span>
-</span><span id="line-18"><span class="linenos">18</span>
-</span><span id="line-19"><span class="linenos">19</span><span class="nt">prompt_targets</span><span class="p">:</span>
-</span><span id="line-20"><span class="linenos">20</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">get_info_for_energy_source</span>
-</span><span id="line-21"><span class="linenos">21</span><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">get information about an energy source</span>
-</span><span id="line-22"><span class="linenos">22</span><span class="w">    </span><span class="nt">parameters</span><span class="p">:</span>
-</span><span id="line-23"><span class="linenos">23</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">energy_source</span>
-</span><span id="line-24"><span class="linenos">24</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
-</span><span id="line-25"><span class="linenos">25</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">a source of energy</span>
-</span><span id="line-26"><span class="linenos">26</span><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-27"><span class="linenos">27</span><span class="w">        </span><span class="nt">enum</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">[</span><span class="nv">renewable</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">fossil</span><span class="p p-Indicator">]</span>
-</span><span id="line-28"><span class="linenos">28</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">consideration</span>
-</span><span id="line-29"><span class="linenos">29</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
-</span><span id="line-30"><span class="linenos">30</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">a specific type of consideration for an energy source</span>
-</span><span id="line-31"><span class="linenos">31</span><span class="w">        </span><span class="nt">enum</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">[</span><span class="nv">cost</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">economic</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">technology</span><span class="p p-Indicator">]</span>
-</span><span id="line-32"><span class="linenos">32</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
-</span><span id="line-33"><span class="linenos">33</span><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">rag_energy_source_agent</span>
-</span><span id="line-34"><span class="linenos">34</span><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/energy_source_info</span>
-</span><span id="line-35"><span class="linenos">35</span><span class="w">      </span><span class="nt">http_method</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">POST</span>
-</span></code></pre></div>
-</div>
-</div>
-</section>
-<section id="step-2-process-request-in-flask">
-<h3>Step 2: Process Request in Flask<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-process-request-in-flask" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-2-process-request-in-flask'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Once the prompt targets are configured as above, handle parameters across multi-turn as if its a single-turn request</p>
-<div class="literal-block-wrapper docutils container" id="id3">
-<div class="code-block-caption"><span class="caption-text">Parameter handling with Flask</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id3"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="kn">import</span><span class="w"> </span><span class="nn">os</span>
-</span><span id="line-2"><span class="linenos"> 2</span><span class="kn">import</span><span class="w"> </span><span class="nn">gradio</span><span class="w"> </span><span class="k">as</span><span class="w"> </span><span class="nn">gr</span>
-</span><span id="line-3"><span class="linenos"> 3</span>
-</span><span id="line-4"><span class="linenos"> 4</span><span class="kn">from</span><span class="w"> </span><span class="nn">fastapi</span><span class="w"> </span><span class="kn">import</span> <span class="n">FastAPI</span><span class="p">,</span> <span class="n">HTTPException</span>
-</span><span id="line-5"><span class="linenos"> 5</span><span class="kn">from</span><span class="w"> </span><span class="nn">pydantic</span><span class="w"> </span><span class="kn">import</span> <span class="n">BaseModel</span>
-</span><span id="line-6"><span class="linenos"> 6</span><span class="kn">from</span><span class="w"> </span><span class="nn">typing</span><span class="w"> </span><span class="kn">import</span> <span class="n">Optional</span>
-</span><span id="line-7"><span class="linenos"> 7</span><span class="kn">from</span><span class="w"> </span><span class="nn">openai</span><span class="w"> </span><span class="kn">import</span> <span class="n">OpenAI</span>
-</span><span id="line-8"><span class="linenos"> 8</span><span class="kn">from</span><span class="w"> </span><span class="nn">common</span><span class="w"> </span><span class="kn">import</span> <span class="n">create_gradio_app</span>
-</span><span id="line-9"><span class="linenos"> 9</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="n">app</span> <span class="o">=</span> <span class="n">FastAPI</span><span class="p">()</span>
-</span><span id="line-11"><span class="linenos">11</span>
-</span><span id="line-12"><span class="linenos">12</span>
-</span><span id="line-13"><span class="linenos">13</span><span class="c1"># Define the request model</span>
-</span><span id="line-14"><span class="linenos">14</span><span class="k">class</span><span class="w"> </span><span class="nc">EnergySourceRequest</span><span class="p">(</span><span class="n">BaseModel</span><span class="p">):</span>
-</span><span id="line-15"><span class="linenos">15</span>    <span class="n">energy_source</span><span class="p">:</span> <span class="nb">str</span>
-</span><span id="line-16"><span class="linenos">16</span>    <span class="n">consideration</span><span class="p">:</span> <span class="n">Optional</span><span class="p">[</span><span class="nb">str</span><span class="p">]</span> <span class="o">=</span> <span class="kc">None</span>
-</span><span id="line-17"><span class="linenos">17</span>
-</span><span id="line-18"><span class="linenos">18</span>
-</span><span id="line-19"><span class="linenos">19</span><span class="k">class</span><span class="w"> </span><span class="nc">EnergySourceResponse</span><span class="p">(</span><span class="n">BaseModel</span><span class="p">):</span>
-</span><span id="line-20"><span class="linenos">20</span>    <span class="n">energy_source</span><span class="p">:</span> <span class="nb">str</span>
-</span><span id="line-21"><span class="linenos">21</span>    <span class="n">consideration</span><span class="p">:</span> <span class="n">Optional</span><span class="p">[</span><span class="nb">str</span><span class="p">]</span> <span class="o">=</span> <span class="kc">None</span>
-</span><span id="line-22"><span class="linenos">22</span>
-</span><span id="line-23"><span class="linenos">23</span>
-</span><span id="line-24"><span class="linenos">24</span><span class="c1"># Post method for device summary</span>
-</span><span id="line-25"><span class="linenos">25</span><span class="nd">@app</span><span class="o">.</span><span class="n">post</span><span class="p">(</span><span class="s2">"/agent/energy_source_info"</span><span class="p">)</span>
-</span><span id="line-26"><span class="linenos">26</span><span class="k">def</span><span class="w"> </span><span class="nf">get_workforce</span><span class="p">(</span><span class="n">request</span><span class="p">:</span> <span class="n">EnergySourceRequest</span><span class="p">):</span>
-</span><span id="line-27"><span class="linenos">27</span><span class="w">    </span><span class="sd">"""</span>
-</span><span id="line-28"><span class="linenos">28</span><span class="sd">    Endpoint to get details about energy source</span>
-</span><span id="line-29"><span class="linenos">29</span><span class="sd">    """</span>
-</span><span id="line-30"><span class="linenos">30</span>    <span class="n">considertion</span> <span class="o">=</span> <span class="s2">"You don't have any specific consideration. Feel free to talk in a more open ended fashion"</span>
-</span><span id="line-31"><span class="linenos">31</span>
-</span><span id="line-32"><span class="linenos">32</span>    <span class="k">if</span> <span class="n">request</span><span class="o">.</span><span class="n">consideration</span> <span class="ow">is</span> <span class="ow">not</span> <span class="kc">None</span><span class="p">:</span>
-</span><span id="line-33"><span class="linenos">33</span>        <span class="n">considertion</span> <span class="o">=</span> <span class="sa">f</span><span class="s2">"Add specific focus on the following consideration when you summarize the content for the energy source: </span><span class="si">{</span><span class="n">request</span><span class="o">.</span><span class="n">consideration</span><span class="si">}</span><span class="s2">"</span>
-</span><span id="line-34"><span class="linenos">34</span>
-</span><span id="line-35"><span class="linenos">35</span>    <span class="n">response</span> <span class="o">=</span> <span class="p">{</span>
-</span><span id="line-36"><span class="linenos">36</span>        <span class="s2">"energy_source"</span><span class="p">:</span> <span class="n">request</span><span class="o">.</span><span class="n">energy_source</span><span class="p">,</span>
-</span><span id="line-37"><span class="linenos">37</span>        <span class="s2">"consideration"</span><span class="p">:</span> <span class="n">considertion</span><span class="p">,</span>
-</span><span id="line-38"><span class="linenos">38</span>    <span class="p">}</span>
-</span><span id="line-39"><span class="linenos">39</span>    <span class="k">return</span> <span class="n">response</span>
-</span></code></pre></div>
-</div>
-</div>
-</section>
-<section id="demo-app">
-<h3>Demo App<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#demo-app" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#demo-app'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>For your convenience, we’ve built a <a class="reference external" href="https://github.com/katanemo/archgw/tree/main/demos/samples_python/multi_turn_rag_agent" rel="nofollow noopener">demo app<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>
-that you can test and modify locally for multi-turn RAG scenarios.</p>
-<figure class="align-center" id="id4">
-<a class="reference internal image-reference" href="../_images/mutli-turn-example.png"><img alt="../_images/mutli-turn-example.png" src="../_images/mutli-turn-example.png" style="width: 100%;"/>
-</a>
-<figcaption>
-<p><span class="caption-text">Example multi-turn user conversation showing adjusting retrieval</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id4"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></p>
-</figcaption>
-</figure>
-</section>
-</section>
-</section>
-</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
-<div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="rag.html">
-<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="15 18 9 12 15 6"></polyline>
-</svg>
-        RAG Apps
-      </a>
-</div>
-<div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../resources/deployment.html">
-        Deployment
-        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="9 18 15 12 9 6"></polyline>
-</svg>
-</a>
-</div>
-</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
-<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
-<ul>
-<li><a :data-current="activeSection === '#example-1-adjusting-retrieval'" class="reference internal" href="#example-1-adjusting-retrieval">Example 1: Adjusting Retrieval</a></li>
-<li><a :data-current="activeSection === '#example-2-switching-intent'" class="reference internal" href="#example-2-switching-intent">Example 2: Switching Intent</a></li>
-<li><a :data-current="activeSection === '#build-multi-turn-rag-apps'" class="reference internal" href="#build-multi-turn-rag-apps">Build Multi-Turn RAG Apps</a><ul>
-<li><a :data-current="activeSection === '#step-1-define-arch-config'" class="reference internal" href="#step-1-define-arch-config">Step 1: Define Arch Config</a></li>
-<li><a :data-current="activeSection === '#step-2-process-request-in-flask'" class="reference internal" href="#step-2-process-request-in-flask">Step 2: Process Request in Flask</a></li>
-<li><a :data-current="activeSection === '#demo-app'" class="reference internal" href="#demo-app">Demo App</a></li>
-</ul>
-</li>
-</ul>
-</div>
-</aside>
-</main>
-</div>
-</div><footer class="py-6 border-t border-border md:py-0">
-<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
-<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
-</div>
-</div>
-</footer>
-</div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
-<script src="../_static/doctools.js?v=9bcbadda"></script>
-<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
-<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
-<script src="../_static/design-tabs.js?v=f930bc37"></script>
-</body>
-</html>
\ No newline at end of file
diff --git a/build_with_arch/rag.html b/build_with_arch/rag.html
deleted file mode 100755
index fa72216f..00000000
--- a/build_with_arch/rag.html
+++ /dev/null
@@ -1,316 +0,0 @@
-<!DOCTYPE html>
-
-<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
-<head>
-<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
-<meta charset="utf-8"/>
-<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
-<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
-<meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>RAG Apps | Arch Docs v0.3.22</title>
-<meta content="RAG Apps | Arch Docs v0.3.22" property="og:title"/>
-<meta content="RAG Apps | Arch Docs v0.3.22" name="twitter:title"/>
-<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
-<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
-<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
-<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/build_with_arch/rag.html" rel="canonical"/>
-<link href="../_static/favicon.ico" rel="icon"/>
-<link href="../search.html" rel="search" title="Search"/>
-<link href="multi_turn.html" rel="next" title="Multi-Turn"/>
-<link href="agent.html" rel="prev" title="Agentic Apps"/>
-<script>
-    <!-- Prevent Flash of wrong theme -->
-      const userPreference = localStorage.getItem('darkMode');
-      let mode;
-      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
-        mode = 'dark';
-        document.documentElement.classList.add('dark');
-      } else {
-        mode = 'light';
-      }
-      if (!userPreference) {localStorage.setItem('darkMode', mode)}
-    </script>
-</head>
-<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
-<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
-      Skip to content
-    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
-<div class="hidden mr-4 md:flex">
-<a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
-<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
-</svg>
-<span class="sr-only">Toggle navigation menu</span>
-</button>
-<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
-<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
-<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
-<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
-<span class="text-xs">⌘</span>
-    K
-  </kbd>
-</form>
-</div>
-<nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
-<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
-<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
-</div>
-</a>
-<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
-<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
-</svg>
-<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
-</svg>
-</button>
-</nav>
-</div>
-</div>
-</header>
-<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
-<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a>
-<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
-<div class="overflow-y-auto h-full w-full relative pr-6">
-
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
-<script>
-  window.dataLayer = window.dataLayer || [];
-  function gtag(){dataLayer.push(arguments);}
-  gtag('js', new Date());
-
-  gtag('config', 'G-K2LXXSX6HB');
-</script>
-<nav class="table w-full min-w-full my-6 lg:my-8">
-<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../concepts/prompt_target.html">Prompt Target</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
-</ul>
-</li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul class="current">
-<li class="toctree-l1"><a class="reference internal" href="agent.html">Agentic Apps</a></li>
-<li class="toctree-l1 current"><a class="current reference internal" href="#">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="multi_turn.html">Multi-Turn</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
-</ul>
-</nav>
-</div>
-</div>
-<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
-<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
-</svg>
-</button>
-</aside>
-<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
-<div class="w-full min-w-0 mx-auto">
-<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
-<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
-<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
-<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
-</svg>
-</a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">RAG Apps</span>
-</nav>
-<div id="content" role="main">
-<section id="rag-apps">
-<span id="arch-rag-guide"></span><h1>RAG Apps<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#rag-apps"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>The following section describes how Arch can help you build faster, smarter and more accurate
-Retrieval-Augmented Generation (RAG) applications, including fast and accurate RAG in multi-turn
-converational scenarios.</p>
-<section id="what-is-retrieval-augmented-generation-rag">
-<h2>What is Retrieval-Augmented Generation (RAG)?<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#what-is-retrieval-augmented-generation-rag" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#what-is-retrieval-augmented-generation-rag'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>RAG applications combine retrieval-based methods with generative AI models to provide more accurate,
-contextually relevant, and reliable outputs. These applications leverage external data sources to augment
-the capabilities of Large Language Models (LLMs), enabling them to retrieve and integrate specific information
-rather than relying solely on the LLM’s internal knowledge.</p>
-</section>
-<section id="parameter-extraction-for-rag">
-<h2>Parameter Extraction for RAG<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#parameter-extraction-for-rag" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#parameter-extraction-for-rag'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>To build RAG (Retrieval Augmented Generation) applications, you can configure prompt targets with parameters,
-enabling Arch to retrieve critical information in a structured way for processing. This approach improves the
-retrieval quality and speed of your application. By extracting parameters from the conversation, you can pull
-the appropriate chunks from a vector database or SQL-like data store to enhance accuracy. With Arch, you can
-streamline data retrieval and processing to build more efficient and precise RAG applications.</p>
-<section id="step-1-define-prompt-targets">
-<h3>Step 1: Define Prompt Targets<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-1-define-prompt-targets" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-1-define-prompt-targets'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<div class="literal-block-wrapper docutils container" id="id1">
-<div class="code-block-caption"><span class="caption-text">Prompt Targets</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id1"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">prompt_targets</span><span class="p">:</span>
-</span><span id="line-2"><span class="linenos"> 2</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">get_device_statistics</span>
-</span><span id="line-3"><span class="linenos"> 3</span><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Retrieve and present the relevant data based on the specified devices and time range</span>
-</span><span id="line-4"><span class="linenos"> 4</span>
-</span><span id="line-5"><span class="linenos"> 5</span><span class="w">    </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/device_summary</span>
-</span><span id="line-6"><span class="linenos"> 6</span><span class="w">    </span><span class="nt">parameters</span><span class="p">:</span>
-</span><span id="line-7"><span class="linenos"> 7</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_ids</span>
-</span><span id="line-8"><span class="linenos"> 8</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">list</span>
-</span><span id="line-9"><span class="linenos"> 9</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">A list of device identifiers (IDs) to reboot.</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-11"><span class="linenos">11</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">time_range</span>
-</span><span id="line-12"><span class="linenos">12</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">int</span>
-</span><span id="line-13"><span class="linenos">13</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">The number of days in the past over which to retrieve device statistics</span>
-</span><span id="line-14"><span class="linenos">14</span><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">false</span>
-</span><span id="line-15"><span class="linenos">15</span><span class="w">        </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">7</span>
-</span></code></pre></div>
-</div>
-</div>
-</section>
-<section id="step-2-process-request-parameters-in-flask">
-<h3>Step 2: Process Request Parameters in Flask<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-process-request-parameters-in-flask" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-2-process-request-parameters-in-flask'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Once the prompt targets are configured as above, handling those parameters is</p>
-<div class="literal-block-wrapper docutils container" id="id2">
-<div class="code-block-caption"><span class="caption-text">Parameter handling with Flask</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="kn">from</span><span class="w"> </span><span class="nn">flask</span><span class="w"> </span><span class="kn">import</span> <span class="n">Flask</span><span class="p">,</span> <span class="n">request</span><span class="p">,</span> <span class="n">jsonify</span>
-</span><span id="line-2"><span class="linenos"> 2</span>
-</span><span id="line-3"><span class="linenos"> 3</span><span class="n">app</span> <span class="o">=</span> <span class="n">Flask</span><span class="p">(</span><span class="vm">__name__</span><span class="p">)</span>
-</span><span id="line-4"><span class="linenos"> 4</span>
-</span><span id="line-5"><span class="linenos"> 5</span>
-</span><span id="line-6"><span class="linenos"> 6</span><span class="nd">@app</span><span class="o">.</span><span class="n">route</span><span class="p">(</span><span class="s2">"/agent/device_summary"</span><span class="p">,</span> <span class="n">methods</span><span class="o">=</span><span class="p">[</span><span class="s2">"POST"</span><span class="p">])</span>
-</span><span id="line-7"><span class="linenos"> 7</span><span class="k">def</span><span class="w"> </span><span class="nf">get_device_summary</span><span class="p">():</span>
-</span><span id="line-8"><span class="linenos"> 8</span><span class="w">    </span><span class="sd">"""</span>
-</span><span id="line-9"><span class="linenos"> 9</span><span class="sd">    Endpoint to retrieve device statistics based on device IDs and an optional time range.</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="sd">    """</span>
-</span><span id="line-11"><span class="linenos">11</span>    <span class="n">data</span> <span class="o">=</span> <span class="n">request</span><span class="o">.</span><span class="n">get_json</span><span class="p">()</span>
-</span><span id="line-12"><span class="linenos">12</span>
-</span><span id="line-13"><span class="linenos">13</span>    <span class="c1"># Validate 'device_ids' parameter</span>
-</span><span id="line-14"><span class="linenos">14</span>    <span class="n">device_ids</span> <span class="o">=</span> <span class="n">data</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"device_ids"</span><span class="p">)</span>
-</span><span id="line-15"><span class="linenos">15</span>    <span class="k">if</span> <span class="ow">not</span> <span class="n">device_ids</span> <span class="ow">or</span> <span class="ow">not</span> <span class="nb">isinstance</span><span class="p">(</span><span class="n">device_ids</span><span class="p">,</span> <span class="nb">list</span><span class="p">):</span>
-</span><span id="line-16"><span class="linenos">16</span>        <span class="k">return</span> <span class="p">(</span>
-</span><span id="line-17"><span class="linenos">17</span>            <span class="n">jsonify</span><span class="p">({</span><span class="s2">"error"</span><span class="p">:</span> <span class="s2">"'device_ids' parameter is required and must be a list"</span><span class="p">}),</span>
-</span><span id="line-18"><span class="linenos">18</span>            <span class="mi">400</span><span class="p">,</span>
-</span><span id="line-19"><span class="linenos">19</span>        <span class="p">)</span>
-</span><span id="line-20"><span class="linenos">20</span>
-</span><span id="line-21"><span class="linenos">21</span>    <span class="c1"># Validate 'time_range' parameter (optional, defaults to 7)</span>
-</span><span id="line-22"><span class="linenos">22</span>    <span class="n">time_range</span> <span class="o">=</span> <span class="n">data</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"time_range"</span><span class="p">,</span> <span class="mi">7</span><span class="p">)</span>
-</span><span id="line-23"><span class="linenos">23</span>    <span class="k">if</span> <span class="ow">not</span> <span class="nb">isinstance</span><span class="p">(</span><span class="n">time_range</span><span class="p">,</span> <span class="nb">int</span><span class="p">):</span>
-</span><span id="line-24"><span class="linenos">24</span>        <span class="k">return</span> <span class="n">jsonify</span><span class="p">({</span><span class="s2">"error"</span><span class="p">:</span> <span class="s2">"'time_range' must be an integer"</span><span class="p">}),</span> <span class="mi">400</span>
-</span><span id="line-25"><span class="linenos">25</span>
-</span><span id="line-26"><span class="linenos">26</span>    <span class="c1"># Simulate retrieving statistics for the given device IDs and time range</span>
-</span><span id="line-27"><span class="linenos">27</span>    <span class="c1"># In a real application, you would query your database or external service here</span>
-</span><span id="line-28"><span class="linenos">28</span>    <span class="n">statistics</span> <span class="o">=</span> <span class="p">[]</span>
-</span><span id="line-29"><span class="linenos">29</span>    <span class="k">for</span> <span class="n">device_id</span> <span class="ow">in</span> <span class="n">device_ids</span><span class="p">:</span>
-</span><span id="line-30"><span class="linenos">30</span>        <span class="c1"># Placeholder for actual data retrieval</span>
-</span><span id="line-31"><span class="linenos">31</span>        <span class="n">stats</span> <span class="o">=</span> <span class="p">{</span>
-</span><span id="line-32"><span class="linenos">32</span>            <span class="s2">"device_id"</span><span class="p">:</span> <span class="n">device_id</span><span class="p">,</span>
-</span><span id="line-33"><span class="linenos">33</span>            <span class="s2">"time_range"</span><span class="p">:</span> <span class="sa">f</span><span class="s2">"Last </span><span class="si">{</span><span class="n">time_range</span><span class="si">}</span><span class="s2"> days"</span><span class="p">,</span>
-</span><span id="line-34"><span class="linenos">34</span>            <span class="s2">"data"</span><span class="p">:</span> <span class="sa">f</span><span class="s2">"Statistics data for device </span><span class="si">{</span><span class="n">device_id</span><span class="si">}</span><span class="s2"> over the last </span><span class="si">{</span><span class="n">time_range</span><span class="si">}</span><span class="s2"> days."</span><span class="p">,</span>
-</span><span id="line-35"><span class="linenos">35</span>        <span class="p">}</span>
-</span><span id="line-36"><span class="linenos">36</span>        <span class="n">statistics</span><span class="o">.</span><span class="n">append</span><span class="p">(</span><span class="n">stats</span><span class="p">)</span>
-</span><span id="line-37"><span class="linenos">37</span>
-</span><span id="line-38"><span class="linenos">38</span>    <span class="n">response</span> <span class="o">=</span> <span class="p">{</span><span class="s2">"statistics"</span><span class="p">:</span> <span class="n">statistics</span><span class="p">}</span>
-</span><span id="line-39"><span class="linenos">39</span>
-</span><span id="line-40"><span class="linenos">40</span>    <span class="k">return</span> <span class="n">jsonify</span><span class="p">(</span><span class="n">response</span><span class="p">),</span> <span class="mi">200</span>
-</span><span id="line-41"><span class="linenos">41</span>
-</span><span id="line-42"><span class="linenos">42</span>
-</span><span id="line-43"><span class="linenos">43</span><span class="k">if</span> <span class="vm">__name__</span> <span class="o">==</span> <span class="s2">"__main__"</span><span class="p">:</span>
-</span><span id="line-44"><span class="linenos">44</span>    <span class="n">app</span><span class="o">.</span><span class="n">run</span><span class="p">(</span><span class="n">debug</span><span class="o">=</span><span class="kc">True</span><span class="p">)</span>
-</span></code></pre></div>
-</div>
-</div>
-</section>
-</section>
-<section id="multi-turn-rag-follow-up-questions">
-<h2>Multi-Turn RAG (Follow-up Questions)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#multi-turn-rag-follow-up-questions" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#multi-turn-rag-follow-up-questions'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Developers often <a class="reference external" href="https://www.reddit.com/r/LocalLLaMA/comments/18mqwg6/best_practice_for_rag_with_followup_chat/" rel="nofollow noopener">struggle<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> to efficiently handle
-<code class="docutils literal notranslate"><span class="pre">follow-up</span></code> or <code class="docutils literal notranslate"><span class="pre">clarification</span></code> questions. Specifically, when users ask for changes or additions to previous responses, it requires developers to
-re-write prompts using LLMs with precise prompt engineering techniques. This process is slow, manual, error prone and adds signifcant latency to the
-user experience.</p>
-<p>Arch is highly capable of accurately detecting and processing prompts in a multi-turn scenarios so that you can buil fast and accurate RAG apps in
-minutes. For additional details on how to build multi-turn RAG applications please refer to our <a class="reference internal" href="multi_turn.html#arch-multi-turn-guide"><span class="std std-ref">multi-turn</span></a> docs.</p>
-</section>
-</section>
-</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
-<div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="agent.html">
-<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="15 18 9 12 15 6"></polyline>
-</svg>
-        Agentic Apps
-      </a>
-</div>
-<div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="multi_turn.html">
-        Multi-Turn
-        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="9 18 15 12 9 6"></polyline>
-</svg>
-</a>
-</div>
-</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
-<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
-<ul>
-<li><a :data-current="activeSection === '#what-is-retrieval-augmented-generation-rag'" class="reference internal" href="#what-is-retrieval-augmented-generation-rag">What is Retrieval-Augmented Generation (RAG)?</a></li>
-<li><a :data-current="activeSection === '#parameter-extraction-for-rag'" class="reference internal" href="#parameter-extraction-for-rag">Parameter Extraction for RAG</a><ul>
-<li><a :data-current="activeSection === '#step-1-define-prompt-targets'" class="reference internal" href="#step-1-define-prompt-targets">Step 1: Define Prompt Targets</a></li>
-<li><a :data-current="activeSection === '#step-2-process-request-parameters-in-flask'" class="reference internal" href="#step-2-process-request-parameters-in-flask">Step 2: Process Request Parameters in Flask</a></li>
-</ul>
-</li>
-<li><a :data-current="activeSection === '#multi-turn-rag-follow-up-questions'" class="reference internal" href="#multi-turn-rag-follow-up-questions">Multi-Turn RAG (Follow-up Questions)</a></li>
-</ul>
-</div>
-</aside>
-</main>
-</div>
-</div><footer class="py-6 border-t border-border md:py-0">
-<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
-<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
-</div>
-</div>
-</footer>
-</div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
-<script src="../_static/doctools.js?v=9bcbadda"></script>
-<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
-<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
-<script src="../_static/design-tabs.js?v=f930bc37"></script>
-</body>
-</html>
\ No newline at end of file
diff --git a/concepts/agents.html b/concepts/agents.html
new file mode 100755
index 00000000..2eb301fb
--- /dev/null
+++ b/concepts/agents.html
@@ -0,0 +1,277 @@
+<!DOCTYPE html>
+
+<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
+<head>
+<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
+<meta charset="utf-8"/>
+<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
+<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
+<meta content="width=device-width, initial-scale=1" name="viewport"/>
+<title>Agents | Plano Docs v0.4</title>
+<meta content="Agents | Plano Docs v0.4" property="og:title"/>
+<meta content="Agents | Plano Docs v0.4" name="twitter:title"/>
+<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
+<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
+<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
+<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
+<link href="./docs/concepts/agents.html" rel="canonical"/>
+<link href="../_static/favicon.ico" rel="icon"/>
+<link href="../search.html" rel="search" title="Search"/>
+<link href="filter_chain.html" rel="next" title="Filter Chains"/>
+<link href="listeners.html" rel="prev" title="Listeners"/>
+<script>
+    <!-- Prevent Flash of wrong theme -->
+      const userPreference = localStorage.getItem('darkMode');
+      let mode;
+      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
+        mode = 'dark';
+        document.documentElement.classList.add('dark');
+      } else {
+        mode = 'light';
+      }
+      if (!userPreference) {localStorage.setItem('darkMode', mode)}
+    </script>
+</head>
+<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
+<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
+      Skip to content
+    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
+<div class="hidden mr-4 md:flex">
+<a class="flex items-center mr-6" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
+<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
+</svg>
+<span class="sr-only">Toggle navigation menu</span>
+</button>
+<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
+<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
+<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
+<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
+<span class="text-xs">⌘</span>
+    K
+  </kbd>
+</form>
+</div>
+<nav class="flex items-center space-x-1">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
+<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
+<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
+</div>
+</a>
+<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
+<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
+</svg>
+<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
+</svg>
+</button>
+</nav>
+</div>
+</div>
+</header>
+<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
+<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a>
+<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
+<div class="overflow-y-auto h-full w-full relative pr-6">
+
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
+<script>
+  window.dataLayer = window.dataLayer || [];
+  function gtag(){dataLayer.push(arguments);}
+  gtag('js', new Date());
+
+  gtag('config', 'G-EH2VW19FXE');
+</script>
+<nav class="table w-full min-w-full my-6 lg:my-8">
+<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
+<ul class="current">
+<li class="toctree-l1"><a class="reference internal" href="listeners.html">Listeners</a></li>
+<li class="toctree-l1 current"><a class="current reference internal" href="#">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="llm_providers/model_aliases.html">Model Aliases</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="prompt_target.html">Prompt Target</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
+<ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
+</ul>
+</nav>
+</div>
+</div>
+<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
+<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
+</svg>
+</button>
+</aside>
+<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
+<div class="w-full min-w-0 mx-auto">
+<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
+<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
+<span class="hidden md:inline">Plano Docs v0.4</span>
+<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
+<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
+</svg>
+</a>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Agents</span>
+</nav>
+<div id="content" role="main">
+<section id="agents">
+<span id="id1"></span><h1>Agents<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#agents"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p>Agents are autonomous systems that handle wide-ranging, open-ended tasks by calling models in a loop until the work is complete. Unlike deterministic <a class="reference internal" href="prompt_target.html#prompt-target"><span class="std std-ref">prompt targets</span></a>, agents have access to tools, reason about which actions to take, and adapt their behavior based on intermediate results—making them ideal for complex workflows that require multi-step reasoning, external API calls, and dynamic decision-making.</p>
+<p>Plano helps developers build and scale multi-agent systems by managing the orchestration layer—deciding which agent(s) or LLM(s) should handle each request, and in what sequence—while developers focus on implementing agent logic in any language or framework they choose.</p>
+<section id="agent-orchestration">
+<h2>Agent Orchestration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#agent-orchestration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#agent-orchestration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p><strong>Plano-Orchestrator</strong> is a family of state-of-the-art routing and orchestration models that decide which agent(s) should handle each request, and in what sequence. Built for real-world multi-agent deployments, it analyzes user intent and conversation context to make precise routing and orchestration decisions while remaining efficient enough for low-latency production use across general chat, coding, and long-context multi-turn conversations.</p>
+<p>This allows development teams to:</p>
+<ul class="simple">
+<li><p><strong>Scale multi-agent systems</strong>: Route requests across multiple specialized agents without hardcoding routing logic in application code.</p></li>
+<li><p><strong>Improve performance</strong>: Direct requests to the most appropriate agent based on intent, reducing unnecessary handoffs and improving response quality.</p></li>
+<li><p><strong>Enhance debuggability</strong>: Centralized routing decisions are observable through Plano’s tracing and logging, making it easier to understand why a particular agent was selected.</p></li>
+</ul>
+</section>
+<section id="inner-loop-vs-outer-loop">
+<h2>Inner Loop vs. Outer Loop<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#inner-loop-vs-outer-loop" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#inner-loop-vs-outer-loop'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Plano distinguishes between the <strong>inner loop</strong> (agent implementation logic) and the <strong>outer loop</strong> (orchestration and routing):</p>
+<section id="inner-loop-agent-logic">
+<h3>Inner Loop (Agent Logic)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#inner-loop-agent-logic" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#inner-loop-agent-logic'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>The inner loop is where your agent lives—the business logic that decides which tools to call, how to interpret results, and when the task is complete. You implement this in any language or framework:</p>
+<ul class="simple">
+<li><p><strong>Python agents</strong>: Using frameworks like LangChain, LlamaIndex, CrewAI, or custom Python code.</p></li>
+<li><p><strong>JavaScript/TypeScript agents</strong>: Using frameworks like LangChain.js or custom Node.js implementations.</p></li>
+<li><p><strong>Any other AI famreowkr</strong>: Agents are just HTTP services that Plano can route to.</p></li>
+</ul>
+<p>Your agent controls:</p>
+<ul class="simple">
+<li><p>Which tools or APIs to call in response to a prompt.</p></li>
+<li><p>How to interpret tool results and decide next steps.</p></li>
+<li><p>When to call the LLM for reasoning or summarization.</p></li>
+<li><p>When the task is complete and what response to return.</p></li>
+</ul>
+<div class="admonition note">
+<p class="admonition-title">Note</p>
+<p><strong>Making LLM Calls from Agents</strong></p>
+<p>When your agent needs to call an LLM for reasoning, summarization, or completion, you should route those calls through Plano’s Model Proxy rather than calling LLM providers directly. This gives you:</p>
+<ul class="simple">
+<li><p><strong>Consistent responses</strong>: Normalized response formats across all <a class="reference internal" href="llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">LLM providers</span></a>, whether you’re using OpenAI, Anthropic, Azure OpenAI, or any OpenAI-compatible provider.</p></li>
+<li><p><strong>Rich agentic signals</strong>: Automatic capture of function calls, tool usage, reasoning steps, and model behavior—surfaced through traces and metrics without instrumenting your agent code.</p></li>
+<li><p><strong>Smart model routing</strong>: Leverage <a class="reference internal" href="llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">model-based, alias-based, or preference-aligned routing</span></a> to dynamically select the best model for each task based on cost, performance, or custom policies.</p></li>
+</ul>
+<p>By routing LLM calls through the Model Proxy, your agents remain decoupled from specific providers and can benefit from centralized policy enforcement, observability, and intelligent routing—all managed in the outer loop. For a step-by-step guide, see <a class="reference internal" href="../guides/llm_router.html#llm-router"><span class="std std-ref">LLM Routing</span></a> in the LLM Router guide.</p>
+</div>
+</section>
+<section id="outer-loop-orchestration">
+<h3>Outer Loop (Orchestration)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#outer-loop-orchestration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#outer-loop-orchestration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>The outer loop is Plano’s orchestration layer—it manages the lifecycle of requests across agents and LLMs:</p>
+<ul class="simple">
+<li><p><strong>Intent analysis</strong>: Plano-Orchestrator analyzes incoming prompts to determine user intent and conversation context.</p></li>
+<li><p><strong>Routing decisions</strong>: Routes requests to the appropriate agent(s) or LLM(s) based on capabilities, context, and availability.</p></li>
+<li><p><strong>Sequencing</strong>: Determines whether multiple agents need to collaborate and in what order.</p></li>
+<li><p><strong>Lifecycle management</strong>: Handles retries, failover, circuit breaking, and load balancing across agent instances.</p></li>
+</ul>
+<p>By managing the outer loop, Plano allows you to:</p>
+<ul class="simple">
+<li><p>Add new agents without changing routing logic in existing agents.</p></li>
+<li><p>Run multiple versions or variants of agents for A/B testing or canary deployments.</p></li>
+<li><p>Apply consistent <a class="reference internal" href="filter_chain.html#filter-chain"><span class="std std-ref">filter chains</span></a> (guardrails, context enrichment) before requests reach agents.</p></li>
+<li><p>Monitor and debug multi-agent workflows through centralized observability.</p></li>
+</ul>
+</section>
+</section>
+<section id="key-benefits">
+<h2>Key Benefits<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#key-benefits" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#key-benefits'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<ul class="simple">
+<li><p><strong>Language and framework agnostic</strong>: Write agents in any language; Plano orchestrates them via HTTP.</p></li>
+<li><p><strong>Reduced complexity</strong>: Agents focus on task logic; Plano handles routing, retries, and cross-cutting concerns.</p></li>
+<li><p><strong>Better observability</strong>: Centralized tracing shows which agents were called, in what sequence, and why.</p></li>
+<li><p><strong>Easier scaling</strong>: Add more agent instances or new agent types without refactoring existing code.</p></li>
+</ul>
+</section>
+</section>
+</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
+<div class="mr-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="listeners.html">
+<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="15 18 9 12 15 6"></polyline>
+</svg>
+        Listeners
+      </a>
+</div>
+<div class="ml-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="filter_chain.html">
+        Filter Chains
+        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="9 18 15 12 9 6"></polyline>
+</svg>
+</a>
+</div>
+</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
+<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
+<ul>
+<li><a :data-current="activeSection === '#agent-orchestration'" class="reference internal" href="#agent-orchestration">Agent Orchestration</a></li>
+<li><a :data-current="activeSection === '#inner-loop-vs-outer-loop'" class="reference internal" href="#inner-loop-vs-outer-loop">Inner Loop vs. Outer Loop</a><ul>
+<li><a :data-current="activeSection === '#inner-loop-agent-logic'" class="reference internal" href="#inner-loop-agent-logic">Inner Loop (Agent Logic)</a></li>
+<li><a :data-current="activeSection === '#outer-loop-orchestration'" class="reference internal" href="#outer-loop-orchestration">Outer Loop (Orchestration)</a></li>
+</ul>
+</li>
+<li><a :data-current="activeSection === '#key-benefits'" class="reference internal" href="#key-benefits">Key Benefits</a></li>
+</ul>
+</div>
+</aside>
+</main>
+</div>
+</div><footer class="py-6 border-t border-border md:py-0">
+<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
+<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
+</div>
+</div>
+</footer>
+</div>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
+<script src="../_static/doctools.js?v=9bcbadda"></script>
+<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
+<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
+<script src="../_static/design-tabs.js?v=f930bc37"></script>
+</body>
+</html>
\ No newline at end of file
diff --git a/concepts/filter_chain.html b/concepts/filter_chain.html
new file mode 100755
index 00000000..39f1b218
--- /dev/null
+++ b/concepts/filter_chain.html
@@ -0,0 +1,304 @@
+<!DOCTYPE html>
+
+<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
+<head>
+<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
+<meta charset="utf-8"/>
+<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
+<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
+<meta content="width=device-width, initial-scale=1" name="viewport"/>
+<title>Filter Chains | Plano Docs v0.4</title>
+<meta content="Filter Chains | Plano Docs v0.4" property="og:title"/>
+<meta content="Filter Chains | Plano Docs v0.4" name="twitter:title"/>
+<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
+<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
+<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
+<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
+<link href="./docs/concepts/filter_chain.html" rel="canonical"/>
+<link href="../_static/favicon.ico" rel="icon"/>
+<link href="../search.html" rel="search" title="Search"/>
+<link href="llm_providers/llm_providers.html" rel="next" title="Model (LLM) Providers"/>
+<link href="agents.html" rel="prev" title="Agents"/>
+<script>
+    <!-- Prevent Flash of wrong theme -->
+      const userPreference = localStorage.getItem('darkMode');
+      let mode;
+      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
+        mode = 'dark';
+        document.documentElement.classList.add('dark');
+      } else {
+        mode = 'light';
+      }
+      if (!userPreference) {localStorage.setItem('darkMode', mode)}
+    </script>
+</head>
+<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
+<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
+      Skip to content
+    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
+<div class="hidden mr-4 md:flex">
+<a class="flex items-center mr-6" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
+<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
+</svg>
+<span class="sr-only">Toggle navigation menu</span>
+</button>
+<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
+<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
+<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
+<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
+<span class="text-xs">⌘</span>
+    K
+  </kbd>
+</form>
+</div>
+<nav class="flex items-center space-x-1">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
+<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
+<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
+</div>
+</a>
+<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
+<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
+</svg>
+<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
+</svg>
+</button>
+</nav>
+</div>
+</div>
+</header>
+<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
+<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a>
+<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
+<div class="overflow-y-auto h-full w-full relative pr-6">
+
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
+<script>
+  window.dataLayer = window.dataLayer || [];
+  function gtag(){dataLayer.push(arguments);}
+  gtag('js', new Date());
+
+  gtag('config', 'G-EH2VW19FXE');
+</script>
+<nav class="table w-full min-w-full my-6 lg:my-8">
+<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
+<ul class="current">
+<li class="toctree-l1"><a class="reference internal" href="listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="agents.html">Agents</a></li>
+<li class="toctree-l1 current"><a class="current reference internal" href="#">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="llm_providers/model_aliases.html">Model Aliases</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="prompt_target.html">Prompt Target</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
+<ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
+</ul>
+</nav>
+</div>
+</div>
+<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
+<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
+</svg>
+</button>
+</aside>
+<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
+<div class="w-full min-w-0 mx-auto">
+<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
+<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
+<span class="hidden md:inline">Plano Docs v0.4</span>
+<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
+<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
+</svg>
+</a>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Filter Chains</span>
+</nav>
+<div id="content" role="main">
+<section id="filter-chains">
+<span id="filter-chain"></span><h1>Filter Chains<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#filter-chains"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p>Filter chains are Plano’s way of capturing <strong>reusable workflow steps</strong> in the dataplane, without duplication and coupling logic into application code. A filter chain is an ordered list of <strong>mutations</strong> that a request flows through before reaching its final destination —such as an agent, an LLM, or a tool backend. Each filter is a network-addressable service/path that can:</p>
+<ol class="arabic simple">
+<li><p>Inspect the incoming prompt, metadata, and conversation state.</p></li>
+<li><p>Mutate or enrich the request (for example, rewrite queries or build context).</p></li>
+<li><p>Short-circuit the flow and return a response early (for example, block a request on a compliance failure).</p></li>
+<li><p>Emit structured logs and traces so you can debug and continuously improve your agents.</p></li>
+</ol>
+<p>In other words, filter chains provide a lightweight programming model over HTTP for building reusable steps
+in your agent architectures.</p>
+<section id="typical-use-cases">
+<h2>Typical Use Cases<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#typical-use-cases" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#typical-use-cases'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Without a dataplane programming model, teams tend to spread logic like query rewriting, compliance checks,
+context building, and routing decisions across many agents and frameworks. This quickly becomes hard to reason
+about and even harder to evolve.</p>
+<p>Filter chains show up most often in patterns like:</p>
+<ul class="simple">
+<li><p><strong>Guardrails and Compliance</strong>: Enforcing content policies, stripping or masking sensitive data, and blocking obviously unsafe or off-topic requests before they reach an agent.</p></li>
+<li><p><strong>Query rewriting, RAG, and Memory</strong>: Rewriting user queries for retrieval, normalizing entities, and assembling RAG context envelopes while pulling in relevant memory (for example, conversation history, user profiles, or prior tool results) before calling a model or tool.</p></li>
+<li><p><strong>Cross-cutting Observability</strong>: Injecting correlation IDs, sampling traces, or logging enriched request metadata at consistent points in the request path.</p></li>
+</ul>
+<p>Because these behaviors live in the dataplane rather than inside individual agents, you define them once, attach them to many agents and prompt targets, and can add, remove, or reorder them without changing application code.</p>
+</section>
+<section id="configuration-example">
+<h2>Configuration example<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuration-example" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configuration-example'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>The example below shows a configuration where an agent uses a filter chain with two filters: a query rewriter,
+and a context builder that prepares retrieval context before the agent runs.</p>
+<div class="literal-block-wrapper docutils container" id="id1">
+<div class="code-block-caption"><span class="caption-text">Example Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id1"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.3.0</span>
+</span><span id="line-2"><span class="linenos"> 2</span>
+</span><span id="line-3"><span class="linenos"> 3</span><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-4"><span class="linenos"> 4</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">rag_agent</span>
+</span><span id="line-5"><span class="linenos"> 5</span><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10505</span>
+</span><span id="line-6"><span class="linenos"> 6</span>
+</span><span id="line-7"><mark><span class="linenos"> 7</span><span class="nt">filters</span><span class="p">:</span>
+</mark></span><span id="line-8"><mark><span class="linenos"> 8</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">query_rewriter</span>
+</mark></span><span id="line-9"><mark><span class="linenos"> 9</span><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10501</span>
+</mark></span><span id="line-10"><mark><span class="linenos">10</span><span class="w">    </span><span class="c1"># type: mcp # default is mcp</span>
+</mark></span><span id="line-11"><mark><span class="linenos">11</span><span class="w">    </span><span class="c1"># transport: streamable-http # default is streamable-http</span>
+</mark></span><span id="line-12"><mark><span class="linenos">12</span><span class="w">    </span><span class="c1"># tool: query_rewriter # default name is the filter id</span>
+</mark></span><span id="line-13"><mark><span class="linenos">13</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">context_builder</span>
+</mark></span><span id="line-14"><mark><span class="linenos">14</span><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10502</span>
+</mark></span><span id="line-15"><span class="linenos">15</span>
+</span><span id="line-16"><span class="linenos">16</span><span class="nt">model_providers</span><span class="p">:</span>
+</span><span id="line-17"><span class="linenos">17</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o-mini</span>
+</span><span id="line-18"><span class="linenos">18</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-19"><span class="linenos">19</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-20"><span class="linenos">20</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
+</span><span id="line-21"><span class="linenos">21</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-22"><span class="linenos">22</span>
+</span><span id="line-23"><span class="linenos">23</span><span class="nt">model_aliases</span><span class="p">:</span>
+</span><span id="line-24"><span class="linenos">24</span><span class="w">  </span><span class="nt">fast-llm</span><span class="p">:</span>
+</span><span id="line-25"><span class="linenos">25</span><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-4o-mini</span>
+</span><span id="line-26"><span class="linenos">26</span><span class="w">  </span><span class="nt">smart-llm</span><span class="p">:</span>
+</span><span id="line-27"><span class="linenos">27</span><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-4o</span>
+</span><span id="line-28"><span class="linenos">28</span>
+</span><span id="line-29"><span class="linenos">29</span><span class="nt">listeners</span><span class="p">:</span>
+</span><span id="line-30"><span class="linenos">30</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">agent</span>
+</span><span id="line-31"><span class="linenos">31</span><span class="w">    </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">agent_1</span>
+</span><span id="line-32"><span class="linenos">32</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">8001</span>
+</span><span id="line-33"><span class="linenos">33</span><span class="w">    </span><span class="nt">router</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">arch_agent_router</span>
+</span><span id="line-34"><span class="linenos">34</span><span class="w">    </span><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-35"><span class="linenos">35</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">rag_agent</span>
+</span><span id="line-36"><span class="linenos">36</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">virtual assistant for retrieval augmented generation tasks</span>
+</span><span id="line-37"><mark><span class="linenos">37</span><span class="w">        </span><span class="nt">filter_chain</span><span class="p">:</span>
+</mark></span><span id="line-38"><mark><span class="linenos">38</span><span class="w">          </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">query_rewriter</span>
+</mark></span><span id="line-39"><mark><span class="linenos">39</span><span class="w">          </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">context_builder</span>
+</mark></span><span id="line-40"><span class="linenos">40</span><span class="nt">tracing</span><span class="p">:</span>
+</span><span id="line-41"><span class="linenos">41</span><span class="w">  </span><span class="nt">random_sampling</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">100</span>
+</span></code></pre></div>
+</div>
+</div>
+<p>In this setup:</p>
+<ul class="simple">
+<li><p>The <code class="docutils literal notranslate"><span class="pre">filters</span></code> section defines the reusable filters, each running as its own HTTP/MCP service.</p></li>
+<li><p>The <code class="docutils literal notranslate"><span class="pre">listeners</span></code> section wires the <code class="docutils literal notranslate"><span class="pre">rag_agent</span></code> behind an <code class="docutils literal notranslate"><span class="pre">agent</span></code> listener and attaches a <code class="docutils literal notranslate"><span class="pre">filter_chain</span></code> with <code class="docutils literal notranslate"><span class="pre">query_rewriter</span></code> followed by <code class="docutils literal notranslate"><span class="pre">context_builder</span></code>.</p></li>
+<li><p>When a request arrives at <code class="docutils literal notranslate"><span class="pre">agent_1</span></code>, Plano executes the filters in order before handing control to <code class="docutils literal notranslate"><span class="pre">rag_agent</span></code>.</p></li>
+</ul>
+</section>
+<section id="filter-chain-programming-model-http-and-mcp">
+<h2>Filter Chain Programming Model (HTTP and MCP)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#filter-chain-programming-model-http-and-mcp" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#filter-chain-programming-model-http-and-mcp'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Filters are implemented as simple RESTful endpoints reachable via HTTP. If you want to use the <a class="reference external" href="https://modelcontextprotocol.io/" rel="nofollow noopener">Model Context Protocol (MCP)<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, you can configure that as well, which makes it easy to write filters in any language. However, you can also write a filter as a plain HTTP service.</p>
+<p>When defining a filter in Plano configuration, the following fields are optional:</p>
+<ul class="simple">
+<li><p><code class="docutils literal notranslate"><span class="pre">type</span></code>: Controls the filter runtime. Use <code class="docutils literal notranslate"><span class="pre">mcp</span></code> for Model Context Protocol filters, or <code class="docutils literal notranslate"><span class="pre">http</span></code> for plain HTTP filters. Defaults to <code class="docutils literal notranslate"><span class="pre">mcp</span></code>.</p></li>
+<li><p><code class="docutils literal notranslate"><span class="pre">transport</span></code>: Controls how Plano talks to the filter (defaults to <code class="docutils literal notranslate"><span class="pre">streamable-http</span></code> for efficient streaming interactions over HTTP). You can omit this for standard HTTP transport.</p></li>
+<li><p><code class="docutils literal notranslate"><span class="pre">tool</span></code>: Names the MCP tool Plano will invoke (by default, the filter <code class="docutils literal notranslate"><span class="pre">id</span></code>). You can omit this if the tool name matches your filter id.</p></li>
+</ul>
+<p>In practice, you typically only need to specify <code class="docutils literal notranslate"><span class="pre">id</span></code> and <code class="docutils literal notranslate"><span class="pre">url</span></code> to get started. Plano’s sensible defaults mean a filter can be as simple as an HTTP endpoint. If you want to customize the runtime or protocol, those fields are there, but they’re optional.</p>
+<p>Filters communicate the outcome of their work via HTTP status codes:</p>
+<ul class="simple">
+<li><p><strong>HTTP 200 (Success)</strong>: The filter successfully processed the request. If the filter mutated the request (e.g., rewrote a query or enriched context), those mutations are passed downstream.</p></li>
+<li><p><strong>HTTP 4xx (User Error)</strong>: The request violates a filter’s rules or constraints—for example, content moderation policies or compliance checks. The request is terminated, and the error is returned to the caller. This is <em>not</em> a fatal error; it represents expected user-facing policy enforcement.</p></li>
+<li><p><strong>HTTP 5xx (Fatal Error)</strong>: An unexpected failure in the filter itself (for example, a crash or misconfiguration). Plano will surface the error back to the caller and record it in logs and traces.</p></li>
+</ul>
+<p>This semantics allows filters to enforce guardrails and policies (4xx) without blocking the entire system, while still surfacing critical failures (5xx) for investigation.</p>
+<p>If any filter fails or decides to terminate the request early (for example, after a policy violation), Plano will
+surface that outcome back to the caller and record it in logs and traces. This makes filter chains a safe and
+powerful abstraction for evolving your agent workflows over time.</p>
+</section>
+</section>
+</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
+<div class="mr-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="agents.html">
+<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="15 18 9 12 15 6"></polyline>
+</svg>
+        Agents
+      </a>
+</div>
+<div class="ml-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="llm_providers/llm_providers.html">
+        Model (LLM) Providers
+        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="9 18 15 12 9 6"></polyline>
+</svg>
+</a>
+</div>
+</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
+<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
+<ul>
+<li><a :data-current="activeSection === '#typical-use-cases'" class="reference internal" href="#typical-use-cases">Typical Use Cases</a></li>
+<li><a :data-current="activeSection === '#configuration-example'" class="reference internal" href="#configuration-example">Configuration example</a></li>
+<li><a :data-current="activeSection === '#filter-chain-programming-model-http-and-mcp'" class="reference internal" href="#filter-chain-programming-model-http-and-mcp">Filter Chain Programming Model (HTTP and MCP)</a></li>
+</ul>
+</div>
+</aside>
+</main>
+</div>
+</div><footer class="py-6 border-t border-border md:py-0">
+<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
+<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
+</div>
+</div>
+</footer>
+</div>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
+<script src="../_static/doctools.js?v=9bcbadda"></script>
+<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
+<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
+<script src="../_static/design-tabs.js?v=f930bc37"></script>
+</body>
+</html>
\ No newline at end of file
diff --git a/concepts/tech_overview/listener.html b/concepts/listeners.html
similarity index 50%
rename from concepts/tech_overview/listener.html
rename to concepts/listeners.html
index 74079e48..a979454f 100755
--- a/concepts/tech_overview/listener.html
+++ b/concepts/listeners.html
@@ -1,25 +1,25 @@
 <!DOCTYPE html>
 
-<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
+<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
 <head>
 <meta content="width=device-width, initial-scale=1.0" name="viewport"/>
 <meta charset="utf-8"/>
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Listener | Arch Docs v0.3.22</title>
-<meta content="Listener | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Listener | Arch Docs v0.3.22" name="twitter:title"/>
-<link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
-<link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
-<link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
-<link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/concepts/tech_overview/listener.html" rel="canonical"/>
-<link href="../../_static/favicon.ico" rel="icon"/>
-<link href="../../search.html" rel="search" title="Search"/>
-<link href="prompt.html" rel="next" title="Prompts"/>
-<link href="threading_model.html" rel="prev" title="Threading Model"/>
+<title>Listeners | Plano Docs v0.4</title>
+<meta content="Listeners | Plano Docs v0.4" property="og:title"/>
+<meta content="Listeners | Plano Docs v0.4" name="twitter:title"/>
+<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
+<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
+<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
+<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
+<link href="./docs/concepts/listeners.html" rel="canonical"/>
+<link href="../_static/favicon.ico" rel="icon"/>
+<link href="../search.html" rel="search" title="Search"/>
+<link href="agents.html" rel="next" title="Agents"/>
+<link href="../get_started/quickstart.html" rel="prev" title="Quickstart"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -38,8 +38,8 @@
       Skip to content
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
-<a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<a class="flex items-center mr-6" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -47,7 +47,7 @@
 <span class="sr-only">Toggle navigation menu</span>
 </button>
 <div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
-<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../../search.html" class="relative flex items-center group" id="searchbox" method="get">
+<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
 <input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
 <kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
 <span class="text-xs">⌘</span>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -74,71 +74,66 @@
 </div>
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
-<a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul class="current">
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
-<li class="toctree-l2 current"><a class="current reference internal" href="#">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html">Error Target</a></li>
+<li class="toctree-l1 current"><a class="current reference internal" href="#">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="llm_providers/model_aliases.html">Model Aliases</a></li>
 </ul>
 </li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../prompt_target.html">Prompt Target</a></li>
+<li class="toctree-l1"><a class="reference internal" href="prompt_target.html">Prompt Target</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -152,77 +147,102 @@
 <main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
-<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
 </a>
-<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="tech_overview.html">Tech Overview</a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Listener</span>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Listeners</span>
 </nav>
 <div id="content" role="main">
-<section id="listener">
-<span id="arch-overview-listeners"></span><h1>Listener<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#listener"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p><strong>Listener</strong> is a top level primitive in Arch, which simplifies the configuration required to bind incoming
-connections from downstream clients, and for egress connections to LLMs (hosted or API)</p>
-<p>Arch builds on Envoy’s Listener subsystem to streamline connection management for developers. Arch minimizes
-the complexity of Envoy’s listener setup by using best-practices and exposing only essential settings,
-making it easier for developers to bind connections without deep knowledge of Envoy’s configuration model. This
-simplification ensures that connections are secure, reliable, and optimized for performance.</p>
-<section id="downstream-ingress">
-<h2>Downstream (Ingress)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#downstream-ingress" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#downstream-ingress'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Developers can configure Arch to accept connections from downstream clients. A downstream listener acts as the
-primary entry point for incoming traffic, handling initial connection setup, including network filtering, guardrails,
-and additional network security checks. For more details on prompt security and safety,
-see <a class="reference internal" href="prompt.html#arch-overview-prompt-handling"><span class="std std-ref">here</span></a>.</p>
+<section id="listeners">
+<span id="plano-overview-listeners"></span><h1>Listeners<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#listeners"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p><strong>Listeners</strong> are a top-level primitive in Plano that bind network traffic to the dataplane. They simplify the
+configuration required to accept incoming connections from downstream clients (edge) and to expose a unified egress
+endpoint for calls from your applications to upstream LLMs.</p>
+<p>Plano builds on Envoy’s Listener subsystem to streamline connection management for developers. It hides most of
+Envoy’s complexity behind sensible defaults and a focused configuration surface, so you can bind listeners without
+deep knowledge of Envoy’s configuration model while still getting secure, reliable, and performant connections.</p>
+<p>Listeners are modular building blocks: you can configure only inbound listeners (for edge proxying and guardrails),
+only outbound/model-proxy listeners (for LLM routing from your services), or both together. This lets you fit Plano
+cleanly into existing architectures, whether you need it at the edge, behind the firewall, or across the full
+request path.</p>
+<section id="network-topology">
+<h2>Network Topology<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#network-topology" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#network-topology'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>The diagram below shows how inbound and outbound traffic flow through Plano and how listeners relate to agents,
+prompt targets, and upstream LLMs:</p>
+<a class="reference internal image-reference" href="../_images/network-topology-ingress-egress.png"><img alt="../_images/network-topology-ingress-egress.png" class="align-center" src="../_images/network-topology-ingress-egress.png" style="width: 100%;"/>
+</a>
 </section>
-<section id="upstream-egress">
-<h2>Upstream (Egress)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#upstream-egress" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#upstream-egress'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch automatically configures a listener to route requests from your application to upstream LLM API providers (or hosts).
-When you start Arch, it creates a listener for egress traffic based on the presence of the <code class="docutils literal notranslate"><span class="pre">listener</span></code> configuration
-section in the configuration file. Arch binds itself to a local address such as <code class="docutils literal notranslate"><span class="pre">127.0.0.1:12000/v1</span></code> or a DNS-based
-address like <code class="docutils literal notranslate"><span class="pre">arch.local:12000/v1</span></code> for outgoing traffic. For more details on LLM providers, read <a class="reference internal" href="../llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">here</span></a>.</p>
+<section id="inbound-agent-prompt-target">
+<h2>Inbound (Agent &amp; Prompt Target)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#inbound-agent-prompt-target" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#inbound-agent-prompt-target'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Developers configure <strong>inbound listeners</strong> to accept connections from clients such as web frontends, backend
+services, or other gateways. An inbound listener acts as the primary entry point for prompt traffic, handling
+initial connection setup, TLS termination, guardrails, and forwarding incoming traffic to the appropriate prompt
+targets or agents.</p>
+<p>There are two primary types of inbound connections exposed via listeners:</p>
+<ul class="simple">
+<li><p><strong>Agent Inbound (Edge)</strong>: Clients (web/mobile apps or other services) connect to Plano, send prompts, and receive
+responses. This is typically your public/edge listener where Plano applies guardrails, routing, and orchestration
+before returning results to the caller.</p></li>
+<li><p><strong>Prompt Target Inbound (Edge)</strong>: Your application server calls Plano’s internal listener targeting
+<a class="reference internal" href="prompt_target.html#prompt-target"><span class="std std-ref">prompt targets</span></a> that can invoke tools and LLMs directly on its behalf.</p></li>
+</ul>
+<p>Inbound listeners are where you attach <a class="reference internal" href="filter_chain.html#filter-chain"><span class="std std-ref">Filter Chains</span></a> so that safety and context-building happen
+consistently at the edge.</p>
 </section>
-<section id="configure-listener">
-<h2>Configure Listener<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configure-listener" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configure-listener'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>To configure a Downstream (Ingress) Listener, simply add the <code class="docutils literal notranslate"><span class="pre">listener</span></code> directive to your configuration file:</p>
+<section id="outbound-model-proxy-egress">
+<h2>Outbound (Model Proxy &amp; Egress)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#outbound-model-proxy-egress" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#outbound-model-proxy-egress'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Plano also exposes an <strong>egress listener</strong> that your applications call when sending requests to upstream LLM providers
+or self-hosted models. From your application’s perspective this looks like a single OpenAI-compatible HTTP endpoint
+(for example, <code class="docutils literal notranslate"><span class="pre">http://127.0.0.1:12000/v1</span></code>), while Plano handles provider selection, retries, and failover behind
+the scenes.</p>
+<p>Under the hood, Plano opens outbound HTTP(S) connections to upstream LLM providers using its unified API surface and
+smart model routing. For more details on how Plano talks to models and how providers are configured, see
+<a class="reference internal" href="llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">LLM providers</span></a>.</p>
+</section>
+<section id="configure-listeners">
+<h2>Configure Listeners<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configure-listeners" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configure-listeners'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Listeners are configured via the <code class="docutils literal notranslate"><span class="pre">listeners</span></code> block in your Plano configuration. You can define one or more inbound
+listeners (for example, <code class="docutils literal notranslate"><span class="pre">type:edge</span></code>) or one or more outbound/model listeners (for example, <code class="docutils literal notranslate"><span class="pre">type:model</span></code>), or both
+in the same deployment.</p>
+<p>To configure an inbound (edge) listener, add a <code class="docutils literal notranslate"><span class="pre">listeners</span></code> block to your configuration file and define at least one
+listener with address, port, and protocol details:</p>
 <div class="literal-block-wrapper docutils container" id="id1">
 <div class="code-block-caption"><span class="caption-text">Example Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id1"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1.0</span>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.2.0</span>
 </span><span id="line-2"><span class="linenos"> 2</span>
 </span><span id="line-3"><mark><span class="linenos"> 3</span><span class="nt">listeners</span><span class="p">:</span>
 </mark></span><span id="line-4"><mark><span class="linenos"> 4</span><span class="w">  </span><span class="nt">ingress_traffic</span><span class="p">:</span>
 </mark></span><span id="line-5"><mark><span class="linenos"> 5</span><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
 </mark></span><span id="line-6"><mark><span class="linenos"> 6</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">10000</span>
-</mark></span><span id="line-7"><mark><span class="linenos"> 7</span><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</mark></span><span id="line-8"><span class="linenos"> 8</span><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
-</span><span id="line-9"><span class="linenos"> 9</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
-</span><span id="line-11"><span class="linenos">11</span><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-12"><span class="linenos">12</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-13"><span class="linenos">13</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
-</span><span id="line-14"><span class="linenos">14</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-15"><span class="linenos">15</span>
-</span><span id="line-16"><span class="linenos">16</span><span class="c1"># default system prompt used by all prompt targets</span>
-</span><span id="line-17"><span class="linenos">17</span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">You are a network assistant that just offers facts; not advice on manufacturers or purchasing decisions.</span>
+</mark></span><span id="line-7"><mark><span class="linenos"> 7</span>
+</mark></span><span id="line-8"><span class="linenos"> 8</span><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
+</span><span id="line-9"><span class="linenos"> 9</span><span class="nt">model_providers</span><span class="p">:</span>
+</span><span id="line-10"><span class="linenos">10</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-11"><span class="linenos">11</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
+</span><span id="line-12"><span class="linenos">12</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
 </span></code></pre></div>
 </div>
 </div>
+<p>When you start Plano, you specify a listener address/port that you want to bind downstream. Plano also exposes a
+predefined internal listener (<code class="docutils literal notranslate"><span class="pre">127.0.0.1:12000</span></code>) that you can use to proxy egress calls originating from your
+application to LLMs (API-based or hosted) via prompt targets.</p>
 </section>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="threading_model.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../get_started/quickstart.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Threading Model
+        Quickstart
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="prompt.html">
-        Prompts
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="agents.html">
+        Agents
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -231,9 +251,10 @@ address like <code class="docutils literal notranslate"><span class="pre">arch.l
 </div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
 <div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
 <ul>
-<li><a :data-current="activeSection === '#downstream-ingress'" class="reference internal" href="#downstream-ingress">Downstream (Ingress)</a></li>
-<li><a :data-current="activeSection === '#upstream-egress'" class="reference internal" href="#upstream-egress">Upstream (Egress)</a></li>
-<li><a :data-current="activeSection === '#configure-listener'" class="reference internal" href="#configure-listener">Configure Listener</a></li>
+<li><a :data-current="activeSection === '#network-topology'" class="reference internal" href="#network-topology">Network Topology</a></li>
+<li><a :data-current="activeSection === '#inbound-agent-prompt-target'" class="reference internal" href="#inbound-agent-prompt-target">Inbound (Agent &amp; Prompt Target)</a></li>
+<li><a :data-current="activeSection === '#outbound-model-proxy-egress'" class="reference internal" href="#outbound-model-proxy-egress">Outbound (Model Proxy &amp; Egress)</a></li>
+<li><a :data-current="activeSection === '#configure-listeners'" class="reference internal" href="#configure-listeners">Configure Listeners</a></li>
 </ul>
 </div>
 </aside>
@@ -242,15 +263,15 @@ address like <code class="docutils literal notranslate"><span class="pre">arch.l
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
-<script src="../../_static/doctools.js?v=9bcbadda"></script>
-<script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
-<script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
-<script src="../../_static/design-tabs.js?v=f930bc37"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
+<script src="../_static/doctools.js?v=9bcbadda"></script>
+<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
+<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
+<script src="../_static/design-tabs.js?v=f930bc37"></script>
 </body>
 </html>
\ No newline at end of file
diff --git a/concepts/llm_providers/client_libraries.html b/concepts/llm_providers/client_libraries.html
index 2a7e6163..050880af 100755
--- a/concepts/llm_providers/client_libraries.html
+++ b/concepts/llm_providers/client_libraries.html
@@ -7,13 +7,13 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Client Libraries | Arch Docs v0.3.22</title>
-<meta content="Client Libraries | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Client Libraries | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Client Libraries | Plano Docs v0.4</title>
+<meta content="Client Libraries | Plano Docs v0.4" property="og:title"/>
+<meta content="Client Libraries | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/concepts/llm_providers/client_libraries.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul class="current">
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2 current"><a class="current reference internal" href="#">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,18 +148,18 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
 </a>
-<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="llm_providers.html">LLM Providers</a>
+<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="llm_providers.html">Model (LLM) Providers</a>
 <div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Client Libraries</span>
 </nav>
 <div id="content" role="main">
 <section id="client-libraries">
 <span id="id1"></span><h1>Client Libraries<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#client-libraries"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Arch provides a unified interface that works seamlessly with multiple client libraries and tools. You can use your preferred client library without changing your existing code - just point it to Arch’s gateway endpoints.</p>
+<p>Plano provides a unified interface that works seamlessly with multiple client libraries and tools. You can use your preferred client library without changing your existing code - just point it to Plano’s gateway endpoints.</p>
 <section id="supported-clients">
 <h2>Supported Clients<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#supported-clients" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#supported-clients'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <ul class="simple">
@@ -176,7 +171,7 @@
 </section>
 <section id="gateway-endpoints">
 <h2>Gateway Endpoints<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#gateway-endpoints" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#gateway-endpoints'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch exposes two main endpoints:</p>
+<p>Plano exposes three main endpoints:</p>
 <table class="docutils align-default">
 <colgroup>
 <col style="width: 40.0%"/>
@@ -191,7 +186,10 @@
 <tr class="row-even"><td><p><code class="docutils literal notranslate"><span class="pre">http://127.0.0.1:12000/v1/chat/completions</span></code></p></td>
 <td><p>OpenAI-compatible chat completions (LLM Gateway)</p></td>
 </tr>
-<tr class="row-odd"><td><p><code class="docutils literal notranslate"><span class="pre">http://127.0.0.1:12000/v1/messages</span></code></p></td>
+<tr class="row-odd"><td><p><code class="docutils literal notranslate"><span class="pre">http://127.0.0.1:12000/v1/responses</span></code></p></td>
+<td><p>OpenAI Responses API with <a class="reference internal" href="../../guides/state.html#managing-conversational-state"><span class="std std-ref">conversational state management</span></a> (LLM Gateway)</p></td>
+</tr>
+<tr class="row-even"><td><p><code class="docutils literal notranslate"><span class="pre">http://127.0.0.1:12000/v1/messages</span></code></p></td>
 <td><p>Anthropic-compatible messages (LLM Gateway)</p></td>
 </tr>
 </tbody>
@@ -199,7 +197,7 @@
 </section>
 <section id="openai-python-sdk">
 <h2>OpenAI (Python) SDK<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#openai-python-sdk" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#openai-python-sdk'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>The OpenAI SDK works with any provider through Arch’s OpenAI-compatible endpoint.</p>
+<p>The OpenAI SDK works with any provider through Plano’s OpenAI-compatible endpoint.</p>
 <p><strong>Installation:</strong></p>
 <div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">pip<span class="w"> </span>install<span class="w"> </span>openai
 </span></code></pre></div>
@@ -207,7 +205,7 @@
 <p><strong>Basic Usage:</strong></p>
 <div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="kn">from</span><span class="w"> </span><span class="nn">openai</span><span class="w"> </span><span class="kn">import</span> <span class="n">OpenAI</span>
 </span><span id="line-2">
-</span><span id="line-3"><span class="c1"># Point to Arch's LLM Gateway</span>
+</span><span id="line-3"><span class="c1"># Point to Plano's LLM Gateway</span>
 </span><span id="line-4"><span class="n">client</span> <span class="o">=</span> <span class="n">OpenAI</span><span class="p">(</span>
 </span><span id="line-5">    <span class="n">api_key</span><span class="o">=</span><span class="s2">"test-key"</span><span class="p">,</span>  <span class="c1"># Can be any value for local testing</span>
 </span><span id="line-6">    <span class="n">base_url</span><span class="o">=</span><span class="s2">"http://127.0.0.1:12000/v1"</span>
@@ -255,7 +253,7 @@
 </span></code></pre></div>
 </div>
 <p><strong>Using with Non-OpenAI Models:</strong></p>
-<p>The OpenAI SDK can be used with any provider configured in Arch:</p>
+<p>The OpenAI SDK can be used with any provider configured in Plano:</p>
 <div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Using Claude model through OpenAI SDK</span>
 </span><span id="line-2"><span class="n">completion</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
 </span><span id="line-3">    <span class="n">model</span><span class="o">=</span><span class="s2">"claude-3-5-sonnet-20241022"</span><span class="p">,</span>
@@ -282,9 +280,82 @@
 </span></code></pre></div>
 </div>
 </section>
+<section id="openai-responses-api-conversational-state">
+<h2>OpenAI Responses API (Conversational State)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#openai-responses-api-conversational-state" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#openai-responses-api-conversational-state'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>The OpenAI Responses API (<code class="docutils literal notranslate"><span class="pre">v1/responses</span></code>) enables multi-turn conversations with automatic state management. Plano handles conversation history for you, so you don’t need to manually include previous messages in each request.</p>
+<p>See <a class="reference internal" href="../../guides/state.html#managing-conversational-state"><span class="std std-ref">Conversational State</span></a> for detailed configuration and storage backend options.</p>
+<p><strong>Installation:</strong></p>
+<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">pip<span class="w"> </span>install<span class="w"> </span>openai
+</span></code></pre></div>
+</div>
+<p><strong>Basic Multi-Turn Conversation:</strong></p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="kn">from</span><span class="w"> </span><span class="nn">openai</span><span class="w"> </span><span class="kn">import</span> <span class="n">OpenAI</span>
+</span><span id="line-2">
+</span><span id="line-3"><span class="c1"># Point to Plano's LLM Gateway</span>
+</span><span id="line-4"><span class="n">client</span> <span class="o">=</span> <span class="n">OpenAI</span><span class="p">(</span>
+</span><span id="line-5">    <span class="n">api_key</span><span class="o">=</span><span class="s2">"test-key"</span><span class="p">,</span>
+</span><span id="line-6">    <span class="n">base_url</span><span class="o">=</span><span class="s2">"http://127.0.0.1:12000/v1"</span>
+</span><span id="line-7"><span class="p">)</span>
+</span><span id="line-8">
+</span><span id="line-9"><span class="c1"># First turn - creates a new conversation</span>
+</span><span id="line-10"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-11">    <span class="n">model</span><span class="o">=</span><span class="s2">"gpt-4o-mini"</span><span class="p">,</span>
+</span><span id="line-12">    <span class="n">messages</span><span class="o">=</span><span class="p">[</span>
+</span><span id="line-13">        <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"My name is Alice"</span><span class="p">}</span>
+</span><span id="line-14">    <span class="p">]</span>
+</span><span id="line-15"><span class="p">)</span>
+</span><span id="line-16">
+</span><span id="line-17"><span class="c1"># Extract response_id for conversation continuity</span>
+</span><span id="line-18"><span class="n">response_id</span> <span class="o">=</span> <span class="n">response</span><span class="o">.</span><span class="n">id</span>
+</span><span id="line-19"><span class="nb">print</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Assistant: </span><span class="si">{</span><span class="n">response</span><span class="o">.</span><span class="n">choices</span><span class="p">[</span><span class="mi">0</span><span class="p">]</span><span class="o">.</span><span class="n">message</span><span class="o">.</span><span class="n">content</span><span class="si">}</span><span class="s2">"</span><span class="p">)</span>
+</span><span id="line-20">
+</span><span id="line-21"><span class="c1"># Second turn - continues the conversation</span>
+</span><span id="line-22"><span class="c1"># Plano automatically retrieves and merges previous context</span>
+</span><span id="line-23"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-24">    <span class="n">model</span><span class="o">=</span><span class="s2">"gpt-4o-mini"</span><span class="p">,</span>
+</span><span id="line-25">    <span class="n">messages</span><span class="o">=</span><span class="p">[</span>
+</span><span id="line-26">        <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"What's my name?"</span><span class="p">}</span>
+</span><span id="line-27">    <span class="p">],</span>
+</span><span id="line-28">    <span class="n">metadata</span><span class="o">=</span><span class="p">{</span><span class="s2">"response_id"</span><span class="p">:</span> <span class="n">response_id</span><span class="p">}</span>  <span class="c1"># Reference previous conversation</span>
+</span><span id="line-29"><span class="p">)</span>
+</span><span id="line-30">
+</span><span id="line-31"><span class="nb">print</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Assistant: </span><span class="si">{</span><span class="n">response</span><span class="o">.</span><span class="n">choices</span><span class="p">[</span><span class="mi">0</span><span class="p">]</span><span class="o">.</span><span class="n">message</span><span class="o">.</span><span class="n">content</span><span class="si">}</span><span class="s2">"</span><span class="p">)</span>
+</span><span id="line-32"><span class="c1"># Output: "Your name is Alice"</span>
+</span></code></pre></div>
+</div>
+<p><strong>Using with Any Provider:</strong></p>
+<p>The Responses API works with any LLM provider configured in Plano:</p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Multi-turn conversation with Claude</span>
+</span><span id="line-2"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-3">    <span class="n">model</span><span class="o">=</span><span class="s2">"claude-3-5-sonnet-20241022"</span><span class="p">,</span>
+</span><span id="line-4">    <span class="n">messages</span><span class="o">=</span><span class="p">[</span>
+</span><span id="line-5">        <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Let's discuss quantum physics"</span><span class="p">}</span>
+</span><span id="line-6">    <span class="p">]</span>
+</span><span id="line-7"><span class="p">)</span>
+</span><span id="line-8">
+</span><span id="line-9"><span class="n">response_id</span> <span class="o">=</span> <span class="n">response</span><span class="o">.</span><span class="n">id</span>
+</span><span id="line-10">
+</span><span id="line-11"><span class="c1"># Continue conversation - Plano manages state regardless of provider</span>
+</span><span id="line-12"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-13">    <span class="n">model</span><span class="o">=</span><span class="s2">"claude-3-5-sonnet-20241022"</span><span class="p">,</span>
+</span><span id="line-14">    <span class="n">messages</span><span class="o">=</span><span class="p">[</span>
+</span><span id="line-15">        <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Tell me more about entanglement"</span><span class="p">}</span>
+</span><span id="line-16">    <span class="p">],</span>
+</span><span id="line-17">    <span class="n">metadata</span><span class="o">=</span><span class="p">{</span><span class="s2">"response_id"</span><span class="p">:</span> <span class="n">response_id</span><span class="p">}</span>
+</span><span id="line-18"><span class="p">)</span>
+</span></code></pre></div>
+</div>
+<p><strong>Key Benefits:</strong></p>
+<ul class="simple">
+<li><p><strong>Reduced payload size</strong>: No need to send full conversation history in each request</p></li>
+<li><p><strong>Provider flexibility</strong>: Use any configured LLM provider with state management</p></li>
+<li><p><strong>Automatic context merging</strong>: Plano handles conversation continuity behind the scenes</p></li>
+<li><p><strong>Production-ready storage</strong>: Configure <a class="reference internal" href="../../guides/state.html#managing-conversational-state"><span class="std std-ref">PostgreSQL or memory storage</span></a> based on your needs</p></li>
+</ul>
+</section>
 <section id="anthropic-python-sdk">
 <h2>Anthropic (Python) SDK<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#anthropic-python-sdk" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#anthropic-python-sdk'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>The Anthropic SDK works with any provider through Arch’s Anthropic-compatible endpoint.</p>
+<p>The Anthropic SDK works with any provider through Plano’s Anthropic-compatible endpoint.</p>
 <p><strong>Installation:</strong></p>
 <div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">pip<span class="w"> </span>install<span class="w"> </span>anthropic
 </span></code></pre></div>
@@ -292,7 +363,7 @@
 <p><strong>Basic Usage:</strong></p>
 <div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="kn">import</span><span class="w"> </span><span class="nn">anthropic</span>
 </span><span id="line-2">
-</span><span id="line-3"><span class="c1"># Point to Arch's LLM Gateway</span>
+</span><span id="line-3"><span class="c1"># Point to Plano's LLM Gateway</span>
 </span><span id="line-4"><span class="n">client</span> <span class="o">=</span> <span class="n">anthropic</span><span class="o">.</span><span class="n">Anthropic</span><span class="p">(</span>
 </span><span id="line-5">    <span class="n">api_key</span><span class="o">=</span><span class="s2">"test-key"</span><span class="p">,</span>  <span class="c1"># Can be any value for local testing</span>
 </span><span id="line-6">    <span class="n">base_url</span><span class="o">=</span><span class="s2">"http://127.0.0.1:12000"</span>
@@ -341,7 +412,7 @@
 </span></code></pre></div>
 </div>
 <p><strong>Using with Non-Anthropic Models:</strong></p>
-<p>The Anthropic SDK can be used with any provider configured in Arch:</p>
+<p>The Anthropic SDK can be used with any provider configured in Plano:</p>
 <div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Using OpenAI model through Anthropic SDK</span>
 </span><span id="line-2"><span class="n">message</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">messages</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
 </span><span id="line-3">    <span class="n">model</span><span class="o">=</span><span class="s2">"gpt-4o-mini"</span><span class="p">,</span>
@@ -426,7 +497,7 @@
 </section>
 <section id="cross-client-compatibility">
 <h2>Cross-Client Compatibility<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#cross-client-compatibility" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#cross-client-compatibility'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>One of Arch’s key features is cross-client compatibility. You can:</p>
+<p>One of Plano’s key features is cross-client compatibility. You can:</p>
 <p><strong>Use OpenAI SDK with Claude Models:</strong></p>
 <div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># OpenAI client calling Claude model</span>
 </span><span id="line-2"><span class="kn">from</span><span class="w"> </span><span class="nn">openai</span><span class="w"> </span><span class="kn">import</span> <span class="n">OpenAI</span>
@@ -570,6 +641,7 @@ Implement fallback logic for better reliability:</p>
 <li><a :data-current="activeSection === '#supported-clients'" class="reference internal" href="#supported-clients">Supported Clients</a></li>
 <li><a :data-current="activeSection === '#gateway-endpoints'" class="reference internal" href="#gateway-endpoints">Gateway Endpoints</a></li>
 <li><a :data-current="activeSection === '#openai-python-sdk'" class="reference internal" href="#openai-python-sdk">OpenAI (Python) SDK</a></li>
+<li><a :data-current="activeSection === '#openai-responses-api-conversational-state'" class="reference internal" href="#openai-responses-api-conversational-state">OpenAI Responses API (Conversational State)</a></li>
 <li><a :data-current="activeSection === '#anthropic-python-sdk'" class="reference internal" href="#anthropic-python-sdk">Anthropic (Python) SDK</a></li>
 <li><a :data-current="activeSection === '#curl-examples'" class="reference internal" href="#curl-examples">cURL Examples</a></li>
 <li><a :data-current="activeSection === '#cross-client-compatibility'" class="reference internal" href="#cross-client-compatibility">Cross-Client Compatibility</a></li>
@@ -584,12 +656,12 @@ Implement fallback logic for better reliability:</p>
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/concepts/llm_providers/llm_providers.html b/concepts/llm_providers/llm_providers.html
index 0822086b..962e202c 100755
--- a/concepts/llm_providers/llm_providers.html
+++ b/concepts/llm_providers/llm_providers.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>LLM Providers | Arch Docs v0.3.22</title>
-<meta content="LLM Providers | Arch Docs v0.3.22" property="og:title"/>
-<meta content="LLM Providers | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Model (LLM) Providers | Plano Docs v0.4</title>
+<meta content="Model (LLM) Providers | Plano Docs v0.4" property="og:title"/>
+<meta content="Model (LLM) Providers | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/concepts/llm_providers/llm_providers.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
 <link href="../../search.html" rel="search" title="Search"/>
 <link href="supported_providers.html" rel="next" title="Supported Providers &amp; Configuration"/>
-<link href="../tech_overview/error_target.html" rel="prev" title="Error Target"/>
+<link href="../filter_chain.html" rel="prev" title="Filter Chains"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul class="current">
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="current reference internal expandable" href="#">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="current reference internal expandable" href="#">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,53 +148,54 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
 </a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">LLM Providers</span>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Model (LLM) Providers</span>
 </nav>
 <div id="content" role="main">
-<section id="llm-providers">
-<span id="id1"></span><h1>LLM Providers<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#llm-providers"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p><strong>LLM Providers</strong> are a top-level primitive in Arch, helping developers centrally define, secure, observe,
-and manage the usage of their LLMs. Arch builds on Envoy’s reliable <a class="reference external" href="https://www.envoyproxy.io/docs/envoy/v1.31.2/intro/arch_overview/upstream/cluster_manager" rel="nofollow noopener">cluster subsystem<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>
-to manage egress traffic to LLMs, which includes intelligent routing, retry and fail-over mechanisms,
-ensuring high availability and fault tolerance. This abstraction also enables developers to seamlessly
-switch between LLM providers or upgrade LLM versions, simplifying the integration and scaling of LLMs
-across applications.</p>
-<p>Today, we are enabling you to connect to 11+ different AI providers through a unified interface with advanced routing and management capabilities.
-Whether you’re using OpenAI, Anthropic, Azure OpenAI, local Ollama models, or any OpenAI-compatible provider, Arch provides seamless integration with enterprise-grade features.</p>
+<section id="model-llm-providers">
+<span id="llm-providers"></span><h1>Model (LLM) Providers<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#model-llm-providers"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p><strong>Model Providers</strong> are a top-level primitive in Plano, helping developers centrally define, secure, observe,
+and manage the usage of their models. Plano builds on Envoy’s reliable <a class="reference external" href="https://www.envoyproxy.io/docs/envoy/v1.31.2/intro/arch_overview/upstream/cluster_manager" rel="nofollow noopener">cluster subsystem<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> to manage egress traffic to models, which includes intelligent routing, retry and fail-over mechanisms,
+ensuring high availability and fault tolerance. This abstraction also enables developers to seamlessly switch between model providers or upgrade model versions, simplifying the integration and scaling of models across applications.</p>
+<p>Today, we are enable you to connect to 15+ different AI providers through a unified interface with advanced routing and management capabilities.
+Whether you’re using OpenAI, Anthropic, Azure OpenAI, local Ollama models, or any OpenAI-compatible provider, Plano provides seamless integration with enterprise-grade features.</p>
+<div class="admonition note">
+<p class="admonition-title">Note</p>
+<p>Please refer to the quickstart guide <a class="reference internal" href="../../get_started/quickstart.html#llm-routing-quickstart"><span class="std std-ref">here</span></a> to configure and use LLM providers via common client libraries like OpenAI and Anthropic Python SDKs, or via direct HTTP/cURL requests.</p>
+</div>
 <section id="core-capabilities">
 <h2>Core Capabilities<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#core-capabilities" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#core-capabilities'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <p><strong>Multi-Provider Support</strong>
 Connect to any combination of providers simultaneously (see <a class="reference internal" href="supported_providers.html#supported-providers"><span class="std std-ref">Supported Providers &amp; Configuration</span></a> for full details):</p>
 <ul class="simple">
-<li><p><strong>First-Class Providers</strong>: Native integrations with OpenAI, Anthropic, DeepSeek, Mistral, Groq, Google Gemini, Together AI, xAI, Azure OpenAI, and Ollama</p></li>
-<li><p><strong>OpenAI-Compatible Providers</strong>: Any provider implementing the OpenAI Chat Completions API standard</p></li>
+<li><p>First-Class Providers: Native integrations with OpenAI, Anthropic, DeepSeek, Mistral, Groq, Google Gemini, Together AI, xAI, Azure OpenAI, and Ollama</p></li>
+<li><p>OpenAI-Compatible Providers: Any provider implementing the OpenAI Chat Completions API standard</p></li>
 </ul>
 <p><strong>Intelligent Routing</strong>
 Three powerful routing approaches to optimize model selection:</p>
 <ul class="simple">
-<li><p><strong>Model-based Routing</strong>: Direct routing to specific models using provider/model names (see <a class="reference internal" href="supported_providers.html#supported-providers"><span class="std std-ref">Supported Providers &amp; Configuration</span></a>)</p></li>
-<li><p><strong>Alias-based Routing</strong>: Semantic routing using custom aliases (see <a class="reference internal" href="model_aliases.html#model-aliases"><span class="std std-ref">Model Aliases</span></a>)</p></li>
-<li><p><strong>Preference-aligned Routing</strong>: Intelligent routing using the Arch-Router model (see <a class="reference internal" href="../../guides/llm_router.html#preference-aligned-routing"><span class="std std-ref">Preference-aligned Routing (Arch-Router)</span></a>)</p></li>
+<li><p>Model-based Routing: Direct routing to specific models using provider/model names (see <a class="reference internal" href="supported_providers.html#supported-providers"><span class="std std-ref">Supported Providers &amp; Configuration</span></a>)</p></li>
+<li><p>Alias-based Routing: Semantic routing using custom aliases (see <a class="reference internal" href="model_aliases.html#model-aliases"><span class="std std-ref">Model Aliases</span></a>)</p></li>
+<li><p>Preference-aligned Routing: Intelligent routing using the Plano-Router model (see <a class="reference internal" href="../../guides/llm_router.html#preference-aligned-routing"><span class="std std-ref">Preference-aligned routing (Arch-Router)</span></a>)</p></li>
 </ul>
 <p><strong>Unified Client Interface</strong>
 Use your preferred client library without changing existing code (see <a class="reference internal" href="client_libraries.html#client-libraries"><span class="std std-ref">Client Libraries</span></a> for details):</p>
 <ul class="simple">
-<li><p><strong>OpenAI Python SDK</strong>: Full compatibility with all providers</p></li>
-<li><p><strong>Anthropic Python SDK</strong>: Native support with cross-provider capabilities</p></li>
-<li><p><strong>cURL &amp; HTTP Clients</strong>: Direct REST API access for any programming language</p></li>
-<li><p><strong>Custom Integrations</strong>: Standard HTTP interfaces for seamless integration</p></li>
+<li><p>OpenAI Python SDK: Full compatibility with all providers</p></li>
+<li><p>Anthropic Python SDK: Native support with cross-provider capabilities</p></li>
+<li><p>cURL &amp; HTTP Clients: Direct REST API access for any programming language</p></li>
+<li><p>Custom Integrations: Standard HTTP interfaces for seamless integration</p></li>
 </ul>
 </section>
 <section id="key-benefits">
 <h2>Key Benefits<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#key-benefits" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#key-benefits'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <ul class="simple">
 <li><p><strong>Provider Flexibility</strong>: Switch between providers without changing client code</p></li>
-<li><p><strong>Three Routing Methods</strong>: Choose from model-based, alias-based, or preference-aligned routing (using <a class="reference external" href="https://huggingface.co/katanemo/Arch-Router-1.5B" rel="nofollow noopener">Arch-Router-1.5B<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>) strategies</p></li>
+<li><p><strong>Three Routing Methods</strong>: Choose from model-based, alias-based, or preference-aligned routing (using <a class="reference external" href="https://huggingface.co/katanemo/Plano-Router-1.5B" rel="nofollow noopener">Plano-Router-1.5B<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>) strategies</p></li>
 <li><p><strong>Cost Optimization</strong>: Route requests to cost-effective models based on complexity</p></li>
 <li><p><strong>Performance Optimization</strong>: Use fast models for simple tasks, powerful models for complex reasoning</p></li>
 <li><p><strong>Environment Management</strong>: Configure different models for different environments</p></li>
@@ -224,7 +220,7 @@ Use your preferred client library without changing existing code (see <a class="
 <section id="advanced-features">
 <h2>Advanced Features<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#advanced-features" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#advanced-features'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <ul class="simple">
-<li><p><a class="reference internal" href="../../guides/llm_router.html#preference-aligned-routing"><span class="std std-ref">Preference-aligned Routing (Arch-Router)</span></a> - Learn about preference-aligned dynamic routing and intelligent model selection</p></li>
+<li><p><a class="reference internal" href="../../guides/llm_router.html#preference-aligned-routing"><span class="std std-ref">Preference-aligned routing (Arch-Router)</span></a> - Learn about preference-aligned dynamic routing and intelligent model selection</p></li>
 </ul>
 </section>
 <section id="getting-started">
@@ -248,6 +244,7 @@ Use your preferred client library without changing existing code (see <a class="
 <li class="toctree-l2"><a class="reference internal" href="client_libraries.html#supported-clients">Supported Clients</a></li>
 <li class="toctree-l2"><a class="reference internal" href="client_libraries.html#gateway-endpoints">Gateway Endpoints</a></li>
 <li class="toctree-l2"><a class="reference internal" href="client_libraries.html#openai-python-sdk">OpenAI (Python) SDK</a></li>
+<li class="toctree-l2"><a class="reference internal" href="client_libraries.html#openai-responses-api-conversational-state">OpenAI Responses API (Conversational State)</a></li>
 <li class="toctree-l2"><a class="reference internal" href="client_libraries.html#anthropic-python-sdk">Anthropic (Python) SDK</a></li>
 <li class="toctree-l2"><a class="reference internal" href="client_libraries.html#curl-examples">cURL Examples</a></li>
 <li class="toctree-l2"><a class="reference internal" href="client_libraries.html#cross-client-compatibility">Cross-Client Compatibility</a></li>
@@ -271,11 +268,11 @@ Use your preferred client library without changing existing code (see <a class="
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../tech_overview/error_target.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../filter_chain.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Error Target
+        Filter Chains
       </a>
 </div>
 <div class="ml-auto">
@@ -302,12 +299,12 @@ Use your preferred client library without changing existing code (see <a class="
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/concepts/llm_providers/model_aliases.html b/concepts/llm_providers/model_aliases.html
index 3468aa42..6c212c6d 100755
--- a/concepts/llm_providers/model_aliases.html
+++ b/concepts/llm_providers/model_aliases.html
@@ -7,13 +7,13 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Model Aliases | Arch Docs v0.3.22</title>
-<meta content="Model Aliases | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Model Aliases | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Model Aliases | Plano Docs v0.4</title>
+<meta content="Model Aliases | Plano Docs v0.4" property="og:title"/>
+<meta content="Model Aliases | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/concepts/llm_providers/model_aliases.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul class="current">
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2 current"><a class="current reference internal" href="#">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,12 +148,12 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
 </a>
-<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="llm_providers.html">LLM Providers</a>
+<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="llm_providers.html">Model (LLM) Providers</a>
 <div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Model Aliases</span>
 </nav>
 <div id="content" role="main">
@@ -396,7 +391,7 @@
 <section id="see-also">
 <h2>See Also<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#see-also" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#see-also'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <ul class="simple">
-<li><p><a class="reference internal" href="llm_providers.html#llm-providers"><span class="std std-ref">LLM Providers</span></a> - Learn about configuring LLM providers</p></li>
+<li><p><a class="reference internal" href="llm_providers.html#llm-providers"><span class="std std-ref">Model (LLM) Providers</span></a> - Learn about configuring LLM providers</p></li>
 <li><p><a class="reference internal" href="../../guides/llm_router.html#llm-router"><span class="std std-ref">LLM Routing</span></a> - Understand how aliases work with intelligent routing</p></li>
 </ul>
 </section>
@@ -435,12 +430,12 @@
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/concepts/llm_providers/supported_providers.html b/concepts/llm_providers/supported_providers.html
index 7f9d383b..58cd8cf8 100755
--- a/concepts/llm_providers/supported_providers.html
+++ b/concepts/llm_providers/supported_providers.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Supported Providers &amp; Configuration | Arch Docs v0.3.22</title>
-<meta content="Supported Providers &amp; Configuration | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Supported Providers &amp; Configuration | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Supported Providers &amp; Configuration | Plano Docs v0.4</title>
+<meta content="Supported Providers &amp; Configuration | Plano Docs v0.4" property="og:title"/>
+<meta content="Supported Providers &amp; Configuration | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/concepts/llm_providers/supported_providers.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
 <link href="../../search.html" rel="search" title="Search"/>
 <link href="client_libraries.html" rel="next" title="Client Libraries"/>
-<link href="llm_providers.html" rel="prev" title="LLM Providers"/>
+<link href="llm_providers.html" rel="prev" title="Model (LLM) Providers"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul class="current">
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
 <li class="toctree-l2 current"><a class="current reference internal" href="#">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,39 +148,31 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
 </a>
-<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="llm_providers.html">LLM Providers</a>
+<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="llm_providers.html">Model (LLM) Providers</a>
 <div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Supported Providers &amp; Configuration</span>
 </nav>
 <div id="content" role="main">
 <section id="supported-providers-configuration">
 <span id="supported-providers"></span><h1>Supported Providers &amp; Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#supported-providers-configuration"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Arch provides first-class support for multiple LLM providers through native integrations and OpenAI-compatible interfaces. This comprehensive guide covers all supported providers, their available chat models, and detailed configuration instructions.</p>
+<p>Plano provides first-class support for multiple LLM providers through native integrations and OpenAI-compatible interfaces. This comprehensive guide covers all supported providers, their available chat models, and detailed configuration instructions.</p>
 <div class="admonition note">
 <p class="admonition-title">Note</p>
-<p><strong>Model Support:</strong> Arch supports all chat models from each provider, not just the examples shown in this guide. The configurations below demonstrate common models for reference, but you can use any chat model available from your chosen provider.</p>
+<p><strong>Model Support:</strong> Plano supports all chat models from each provider, not just the examples shown in this guide. The configurations below demonstrate common models for reference, but you can use any chat model available from your chosen provider.</p>
+<p>Please refer to the quuickstart guide <a class="reference internal" href="../../get_started/quickstart.html#llm-routing-quickstart"><span class="std std-ref">here</span></a> to configure and use LLM providers via common client libraries like OpenAI and Anthropic Python SDKs, or via direct HTTP/cURL requests.</p>
 </div>
 <section id="configuration-structure">
 <h2>Configuration Structure<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuration-structure" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configuration-structure'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>All providers are configured in the <code class="docutils literal notranslate"><span class="pre">llm_providers</span></code> section of your <code class="docutils literal notranslate"><span class="pre">arch_config.yaml</span></code> file:</p>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1</span>
-</span><span id="line-2">
-</span><span id="line-3"><span class="nt">listeners</span><span class="p">:</span>
-</span><span id="line-4"><span class="w">  </span><span class="nt">egress_traffic</span><span class="p">:</span>
-</span><span id="line-5"><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
-</span><span id="line-6"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">12000</span>
-</span><span id="line-7"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-8"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
-</span><span id="line-9">
-</span><span id="line-10"><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-11"><span class="w">  </span><span class="c1"># Provider configurations go here</span>
-</span><span id="line-12"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">provider/model-name</span>
-</span><span id="line-13"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$API_KEY</span>
-</span><span id="line-14"><span class="w">    </span><span class="c1"># Additional provider-specific options</span>
+<p>All providers are configured in the <code class="docutils literal notranslate"><span class="pre">llm_providers</span></code> section of your <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code> file:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">llm_providers</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="c1"># Provider configurations go here</span>
+</span><span id="line-3"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">provider/model-name</span>
+</span><span id="line-4"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$API_KEY</span>
+</span><span id="line-5"><span class="w">    </span><span class="c1"># Additional provider-specific options</span>
 </span></code></pre></div>
 </div>
 <p><strong>Common Configuration Fields:</strong></p>
@@ -206,7 +193,7 @@ Any provider that implements the OpenAI API interface can be configured using cu
 </section>
 <section id="supported-api-endpoints">
 <h2>Supported API Endpoints<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#supported-api-endpoints" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#supported-api-endpoints'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch supports the following standardized endpoints across providers:</p>
+<p>Plano supports the following standardized endpoints across providers:</p>
 <table class="docutils align-default">
 <colgroup>
 <col style="width: 30.0%"/>
@@ -228,6 +215,10 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <td><p>Anthropic-style messages</p></td>
 <td><p>Anthropic SDK, cURL, custom clients</p></td>
 </tr>
+<tr class="row-even"><td><p><code class="docutils literal notranslate"><span class="pre">/v1/responses</span></code></p></td>
+<td><p>Unified response endpoint for agentic apps</p></td>
+<td><p>All SDKs, cURL, custom clients</p></td>
+</tr>
 </tbody>
 </table>
 </section>
@@ -238,7 +229,7 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <p><strong>Provider Prefix:</strong> <code class="docutils literal notranslate"><span class="pre">openai/</span></code></p>
 <p><strong>API Endpoint:</strong> <code class="docutils literal notranslate"><span class="pre">/v1/chat/completions</span></code></p>
 <p><strong>Authentication:</strong> API Key - Get your OpenAI API key from <a class="reference external" href="https://platform.openai.com/api-keys" rel="nofollow noopener">OpenAI Platform<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.</p>
-<p><strong>Supported Chat Models:</strong> All OpenAI chat models including GPT-5, GPT-4o, GPT-4, GPT-3.5-turbo, and all future releases.</p>
+<p><strong>Supported Chat Models:</strong> All OpenAI chat models including GPT-5.2, GPT-5, GPT-4o, and all future releases.</p>
 <table class="docutils align-default">
 <colgroup>
 <col style="width: 30.0%"/>
@@ -252,31 +243,27 @@ Any provider that implements the OpenAI API interface can be configured using cu
 </tr>
 </thead>
 <tbody>
-<tr class="row-even"><td><p>GPT-5</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">openai/gpt-5</span></code></p></td>
+<tr class="row-even"><td><p>GPT-5.2</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">openai/gpt-5.2</span></code></p></td>
 <td><p>Next-generation model (use any model name from OpenAI’s API)</p></td>
 </tr>
-<tr class="row-odd"><td><p>GPT-4o</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">openai/gpt-4o</span></code></p></td>
+<tr class="row-odd"><td><p>GPT-5</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">openai/gpt-5</span></code></p></td>
 <td><p>Latest multimodal model</p></td>
 </tr>
 <tr class="row-even"><td><p>GPT-4o mini</p></td>
 <td><p><code class="docutils literal notranslate"><span class="pre">openai/gpt-4o-mini</span></code></p></td>
 <td><p>Fast, cost-effective model</p></td>
 </tr>
-<tr class="row-odd"><td><p>GPT-4</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">openai/gpt-4</span></code></p></td>
+<tr class="row-odd"><td><p>GPT-4o</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">openai/gpt-4o</span></code></p></td>
 <td><p>High-capability reasoning model</p></td>
 </tr>
-<tr class="row-even"><td><p>GPT-3.5 Turbo</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">openai/gpt-3.5-turbo</span></code></p></td>
-<td><p>Balanced performance and cost</p></td>
-</tr>
-<tr class="row-odd"><td><p>o3-mini</p></td>
+<tr class="row-even"><td><p>o3-mini</p></td>
 <td><p><code class="docutils literal notranslate"><span class="pre">openai/o3-mini</span></code></p></td>
 <td><p>Reasoning-focused model (preview)</p></td>
 </tr>
-<tr class="row-even"><td><p>o3</p></td>
+<tr class="row-odd"><td><p>o3</p></td>
 <td><p><code class="docutils literal notranslate"><span class="pre">openai/o3</span></code></p></td>
 <td><p>Advanced reasoning model (preview)</p></td>
 </tr>
@@ -285,15 +272,15 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <p><strong>Configuration Examples:</strong></p>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">llm_providers</span><span class="p">:</span>
 </span><span id="line-2"><span class="w">  </span><span class="c1"># Latest models (examples - use any OpenAI chat model)</span>
-</span><span id="line-3"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o-mini</span>
+</span><span id="line-3"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5.2</span>
 </span><span id="line-4"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
 </span><span id="line-5"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
 </span><span id="line-6">
-</span><span id="line-7"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
+</span><span id="line-7"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5</span>
 </span><span id="line-8"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
 </span><span id="line-9">
 </span><span id="line-10"><span class="w">  </span><span class="c1"># Use any model name from OpenAI's API</span>
-</span><span id="line-11"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5</span>
+</span><span id="line-11"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
 </span><span id="line-12"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
 </span></code></pre></div>
 </div>
@@ -303,7 +290,7 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <p><strong>Provider Prefix:</strong> <code class="docutils literal notranslate"><span class="pre">anthropic/</span></code></p>
 <p><strong>API Endpoint:</strong> <code class="docutils literal notranslate"><span class="pre">/v1/messages</span></code></p>
 <p><strong>Authentication:</strong> API Key - Get your Anthropic API key from <a class="reference external" href="https://console.anthropic.com/settings/keys" rel="nofollow noopener">Anthropic Console<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.</p>
-<p><strong>Supported Chat Models:</strong> All Anthropic Claude models including Claude Sonnet 4, Claude 3.5 Sonnet, Claude 3.5 Haiku, Claude 3 Opus, and all future releases.</p>
+<p><strong>Supported Chat Models:</strong> All Anthropic Claude models including Claude Sonnet 4.5, Claude Opus 4.5, Claude Haiku 4.5, and all future releases.</p>
 <table class="docutils align-default">
 <colgroup>
 <col style="width: 30.0%"/>
@@ -317,43 +304,35 @@ Any provider that implements the OpenAI API interface can be configured using cu
 </tr>
 </thead>
 <tbody>
-<tr class="row-even"><td><p>Claude Sonnet 4</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-sonnet-4</span></code></p></td>
-<td><p>Next-generation model (use any model name from Anthropic’s API)</p></td>
-</tr>
-<tr class="row-odd"><td><p>Claude 3.5 Sonnet</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-3-5-sonnet-20241022</span></code></p></td>
-<td><p>Latest high-performance model</p></td>
-</tr>
-<tr class="row-even"><td><p>Claude 3.5 Haiku</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-3-5-haiku-20241022</span></code></p></td>
-<td><p>Fast and efficient model</p></td>
-</tr>
-<tr class="row-odd"><td><p>Claude 3 Opus</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-3-opus-20240229</span></code></p></td>
+<tr class="row-even"><td><p>Claude Opus 4.5</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-opus-4-5</span></code></p></td>
 <td><p>Most capable model for complex tasks</p></td>
 </tr>
-<tr class="row-even"><td><p>Claude 3 Sonnet</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-3-sonnet-20240229</span></code></p></td>
+<tr class="row-odd"><td><p>Claude Sonnet 4.5</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-sonnet-4-5</span></code></p></td>
 <td><p>Balanced performance model</p></td>
 </tr>
-<tr class="row-odd"><td><p>Claude 3 Haiku</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-3-haiku-20240307</span></code></p></td>
-<td><p>Fastest model</p></td>
+<tr class="row-even"><td><p>Claude Haiku 4.5</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-haiku-4-5</span></code></p></td>
+<td><p>Fast and efficient model</p></td>
+</tr>
+<tr class="row-odd"><td><p>Claude Sonnet 3.5</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">anthropic/claude-sonnet-3-5</span></code></p></td>
+<td><p>Complex agents and coding</p></td>
 </tr>
 </tbody>
 </table>
 <p><strong>Configuration Examples:</strong></p>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">llm_providers</span><span class="p">:</span>
 </span><span id="line-2"><span class="w">  </span><span class="c1"># Latest models (examples - use any Anthropic chat model)</span>
-</span><span id="line-3"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-3-5-sonnet-20241022</span>
+</span><span id="line-3"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-opus-4-5</span>
 </span><span id="line-4"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
 </span><span id="line-5">
-</span><span id="line-6"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-3-5-haiku-20241022</span>
+</span><span id="line-6"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-sonnet-4-5</span>
 </span><span id="line-7"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
 </span><span id="line-8">
 </span><span id="line-9"><span class="w">  </span><span class="c1"># Use any model name from Anthropic's API</span>
-</span><span id="line-10"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-sonnet-4</span>
+</span><span id="line-10"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-haiku-4-5</span>
 </span><span id="line-11"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
 </span></code></pre></div>
 </div>
@@ -450,7 +429,7 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <p><strong>Provider Prefix:</strong> <code class="docutils literal notranslate"><span class="pre">groq/</span></code></p>
 <p><strong>API Endpoint:</strong> <code class="docutils literal notranslate"><span class="pre">/openai/v1/chat/completions</span></code> (transformed internally)</p>
 <p><strong>Authentication:</strong> API Key - Get your Groq API key from <a class="reference external" href="https://console.groq.com/keys" rel="nofollow noopener">Groq Console<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.</p>
-<p><strong>Supported Chat Models:</strong> All Groq chat models including Llama 3, Mixtral, Gemma, and all future releases.</p>
+<p><strong>Supported Chat Models:</strong> All Groq chat models including Llama 4, GPT OSS, Mixtral, Gemma, and all future releases.</p>
 <table class="docutils align-default">
 <colgroup>
 <col style="width: 30.0%"/>
@@ -464,27 +443,30 @@ Any provider that implements the OpenAI API interface can be configured using cu
 </tr>
 </thead>
 <tbody>
-<tr class="row-even"><td><p>Llama 3.1 8B</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">groq/llama3-8b-8192</span></code></p></td>
+<tr class="row-even"><td><p>Llama 4 Maverick 17B</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">groq/llama-4-maverick-17b-128e-instruct</span></code></p></td>
 <td><p>Fast inference Llama model</p></td>
 </tr>
-<tr class="row-odd"><td><p>Llama 3.1 70B</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">groq/llama3-70b-8192</span></code></p></td>
-<td><p>Larger Llama model</p></td>
+<tr class="row-odd"><td><p>Llama 4 Scout 8B</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">groq/llama-4-scout-8b-128e-instruct</span></code></p></td>
+<td><p>Smaller Llama model</p></td>
 </tr>
-<tr class="row-even"><td><p>Mixtral 8x7B</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">groq/mixtral-8x7b-32768</span></code></p></td>
-<td><p>Mixture of experts model</p></td>
+<tr class="row-even"><td><p>GPT OSS 20B</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">groq/gpt-oss-20b</span></code></p></td>
+<td><p>Open source GPT model</p></td>
 </tr>
 </tbody>
 </table>
 <p><strong>Configuration Examples:</strong></p>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">groq/llama3-8b-8192</span>
+</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">groq/llama-4-maverick-17b-128e-instruct</span>
 </span><span id="line-3"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$GROQ_API_KEY</span>
 </span><span id="line-4">
-</span><span id="line-5"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">groq/mixtral-8x7b-32768</span>
+</span><span id="line-5"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">groq/llama-4-scout-8b-128e-instruct</span>
 </span><span id="line-6"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$GROQ_API_KEY</span>
+</span><span id="line-7">
+</span><span id="line-8"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">groq/gpt-oss-20b</span>
+</span><span id="line-9"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$GROQ_API_KEY</span>
 </span></code></pre></div>
 </div>
 </section>
@@ -493,7 +475,7 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <p><strong>Provider Prefix:</strong> <code class="docutils literal notranslate"><span class="pre">gemini/</span></code></p>
 <p><strong>API Endpoint:</strong> <code class="docutils literal notranslate"><span class="pre">/v1beta/openai/chat/completions</span></code> (transformed internally)</p>
 <p><strong>Authentication:</strong> API Key - Get your Google AI API key from <a class="reference external" href="https://aistudio.google.com/app/apikey" rel="nofollow noopener">Google AI Studio<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.</p>
-<p><strong>Supported Chat Models:</strong> All Google Gemini chat models including Gemini 1.5 Pro, Gemini 1.5 Flash, and all future releases.</p>
+<p><strong>Supported Chat Models:</strong> All Google Gemini chat models including Gemini 3 Pro, Gemini 3 Flash, and all future releases.</p>
 <table class="docutils align-default">
 <colgroup>
 <col style="width: 30.0%"/>
@@ -507,22 +489,22 @@ Any provider that implements the OpenAI API interface can be configured using cu
 </tr>
 </thead>
 <tbody>
-<tr class="row-even"><td><p>Gemini 1.5 Pro</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">gemini/gemini-1.5-pro</span></code></p></td>
+<tr class="row-even"><td><p>Gemini 3 Pro</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">gemini/gemini-3-pro</span></code></p></td>
 <td><p>Advanced reasoning and creativity</p></td>
 </tr>
-<tr class="row-odd"><td><p>Gemini 1.5 Flash</p></td>
-<td><p><code class="docutils literal notranslate"><span class="pre">gemini/gemini-1.5-flash</span></code></p></td>
+<tr class="row-odd"><td><p>Gemini 3 Flash</p></td>
+<td><p><code class="docutils literal notranslate"><span class="pre">gemini/gemini-3-flash</span></code></p></td>
 <td><p>Fast and efficient model</p></td>
 </tr>
 </tbody>
 </table>
 <p><strong>Configuration Examples:</strong></p>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gemini/gemini-1.5-pro</span>
+</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gemini/gemini-3-pro</span>
 </span><span id="line-3"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$GOOGLE_API_KEY</span>
 </span><span id="line-4">
-</span><span id="line-5"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gemini/gemini-1.5-flash</span>
+</span><span id="line-5"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gemini/gemini-3-flash</span>
 </span><span id="line-6"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$GOOGLE_API_KEY</span>
 </span></code></pre></div>
 </div>
@@ -724,7 +706,7 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <h3>Amazon Bedrock<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#amazon-bedrock" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#amazon-bedrock'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <p><strong>Provider Prefix:</strong> <code class="docutils literal notranslate"><span class="pre">amazon_bedrock/</span></code></p>
 <dl class="simple">
-<dt><strong>API Endpoint:</strong> Arch automatically constructs the endpoint as:</dt><dd><ul class="simple">
+<dt><strong>API Endpoint:</strong> Plano automatically constructs the endpoint as:</dt><dd><ul class="simple">
 <li><p>Non-streaming: <code class="docutils literal notranslate"><span class="pre">/model/{model-id}/converse</span></code></p></li>
 <li><p>Streaming: <code class="docutils literal notranslate"><span class="pre">/model/{model-id}/converse-stream</span></code></p></li>
 </ul>
@@ -894,7 +876,7 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <h3>Routing Preferences<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#routing-preferences" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#routing-preferences'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <p>Configure routing preferences for dynamic model selection:</p>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
+</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5.2</span>
 </span><span id="line-3"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
 </span><span id="line-4"><span class="w">    </span><span class="nt">routing_preferences</span><span class="p">:</span>
 </span><span id="line-5"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">complex_reasoning</span>
@@ -902,7 +884,7 @@ Any provider that implements the OpenAI API interface can be configured using cu
 </span><span id="line-7"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">code_review</span>
 </span><span id="line-8"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">reviewing and analyzing existing code for bugs and improvements</span>
 </span><span id="line-9">
-</span><span id="line-10"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-3-5-sonnet-20241022</span>
+</span><span id="line-10"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-sonnet-4-5</span>
 </span><span id="line-11"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
 </span><span id="line-12"><span class="w">    </span><span class="nt">routing_preferences</span><span class="p">:</span>
 </span><span id="line-13"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">creative_writing</span>
@@ -914,14 +896,14 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <section id="model-selection-guidelines">
 <h2>Model Selection Guidelines<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#model-selection-guidelines" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#model-selection-guidelines'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <p><strong>For Production Applications:</strong>
-- <strong>High Performance</strong>: OpenAI GPT-4o, Anthropic Claude 3.5 Sonnet
-- <strong>Cost-Effective</strong>: OpenAI GPT-4o mini, Anthropic Claude 3.5 Haiku
+- <strong>High Performance</strong>: OpenAI GPT-5.2, Anthropic Claude Sonnet 4.5
+- <strong>Cost-Effective</strong>: OpenAI GPT-5, Anthropic Claude Haiku 4.5
 - <strong>Code Tasks</strong>: DeepSeek Coder, Together AI Code Llama
 - <strong>Local Deployment</strong>: Ollama with Llama 3.1 or Code Llama</p>
 <p><strong>For Development/Testing:</strong>
 - <strong>Fast Iteration</strong>: Groq models (optimized inference)
 - <strong>Local Testing</strong>: Ollama models
-- <strong>Cost Control</strong>: Smaller models like GPT-4o mini or Mistral Small</p>
+- <strong>Cost Control</strong>: Smaller models like GPT-4o or Mistral Small</p>
 </section>
 <section id="see-also">
 <h2>See Also<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#see-also" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#see-also'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
@@ -940,7 +922,7 @@ Any provider that implements the OpenAI API interface can be configured using cu
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        LLM Providers
+        Model (LLM) Providers
       </a>
 </div>
 <div class="ml-auto">
@@ -995,12 +977,12 @@ Any provider that implements the OpenAI API interface can be configured using cu
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/concepts/prompt_target.html b/concepts/prompt_target.html
index 3f476c5e..4ef2e4ef 100755
--- a/concepts/prompt_target.html
+++ b/concepts/prompt_target.html
@@ -7,18 +7,18 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Prompt Target | Arch Docs v0.3.22</title>
-<meta content="Prompt Target | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Prompt Target | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Prompt Target | Plano Docs v0.4</title>
+<meta content="Prompt Target | Plano Docs v0.4" property="og:title"/>
+<meta content="Prompt Target | Plano Docs v0.4" name="twitter:title"/>
 <link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/concepts/prompt_target.html" rel="canonical"/>
 <link href="../_static/favicon.ico" rel="icon"/>
 <link href="../search.html" rel="search" title="Search"/>
-<link href="../guides/prompt_guard.html" rel="next" title="Prompt Guard"/>
+<link href="../guides/orchestration.html" rel="next" title="Orchestration"/>
 <link href="llm_providers/model_aliases.html" rel="prev" title="Model Aliases"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul class="current">
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -163,11 +158,15 @@
 <div id="content" role="main">
 <section id="prompt-target">
 <span id="id1"></span><h1>Prompt Target<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#prompt-target"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p><strong>Prompt Targets</strong> are a core concept in Arch, empowering developers to clearly define how user prompts are interpreted, processed, and routed within their generative AI applications. Prompts can seamlessly be routed either to specialized AI agents capable of handling sophisticated, context-driven tasks or to targeted tools provided by your application, offering users a fast, precise, and personalized experience.</p>
-<p>This section covers the essentials of prompt targets—what they are, how to configure them, their practical uses, and recommended best practices—to help you fully utilize this feature in your applications.</p>
-<section id="what-are-prompt-targets">
-<h2>What Are Prompt Targets?<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#what-are-prompt-targets" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#what-are-prompt-targets'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Prompt targets are endpoints within Arch that handle specific types of user prompts. They act as the bridge between user inputs and your backend agents or tools (APIs), enabling Arch to route, process, and manage prompts efficiently. Defining prompt targets helps you decouple your application’s core logic from processing and handling complexities, leading to clearer code organization, better scalability, and easier maintenance.</p>
+<p>A Prompt Target is a deterministic, task-specific backend function or API endpoint that your application calls via Plano.
+Unlike agents (which handle wide-ranging, open-ended tasks), prompt targets are designed for focused, specific workloads where Plano can add value through input clarification and validation.</p>
+<p>Plano helps by:</p>
+<ul class="simple">
+<li><p><strong>Clarifying and validating input</strong>: Plano enriches incoming prompts with metadata (e.g., detecting follow-ups or clarifying requests) and can extract structured parameters from natural language before passing them to your backend.</p></li>
+<li><p><strong>Enabling high determinism</strong>: Since the task is specific and well-defined, Plano can reliably extract the information your backend needs without ambiguity.</p></li>
+<li><p><strong>Reducing backend work</strong>: Your backend receives clean, validated, structured inputs—so you can focus on business logic instead of parsing and validation.</p></li>
+</ul>
+<p>For example, a prompt target might be “schedule a meeting” (specific task, deterministic inputs like date, time, attendees) or “retrieve documents” (well-defined RAG query with clear intent). Prompt targets are typically called from your application code via Plano’s internal listener.</p>
 <table class="docutils align-default" style="width: 100%">
 <thead>
 <tr class="row-odd"><th class="head"><p><strong>Capability</strong></p></th>
@@ -190,23 +189,19 @@
 </tbody>
 </table>
 <section id="key-features">
-<h3>Key Features<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#key-features" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#key-features'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<h2>Key Features<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#key-features" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#key-features'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <p>Below are the key features of prompt targets that empower developers to build efficient, scalable, and personalized GenAI solutions:</p>
 <ul class="simple">
 <li><p><strong>Design Scenarios</strong>: Define prompt targets to effectively handle specific agentic scenarios.</p></li>
 <li><p><strong>Input Management</strong>: Specify required and optional parameters for each target.</p></li>
 <li><p><strong>Tools Integration</strong>: Seamlessly connect prompts to backend APIs or functions.</p></li>
 <li><p><strong>Error Handling</strong>: Direct errors to designated handlers for streamlined troubleshooting.</p></li>
-<li><p><strong>Metadata Enrichment</strong>: Attach additional context to prompts for enhanced processing.</p></li>
+<li><p><strong>Multi-Turn Support</strong>: Manage follow-up prompts and clarifications in conversational flows.</p></li>
 </ul>
 </section>
-</section>
-<section id="configuring-prompt-targets">
-<h2>Configuring Prompt Targets<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuring-prompt-targets" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configuring-prompt-targets'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Configuring prompt targets involves defining them in Arch’s configuration file. Each Prompt target specifies how a particular type of prompt should be handled, including the endpoint to invoke and any parameters required.</p>
 <section id="basic-configuration">
-<h3>Basic Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#basic-configuration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#basic-configuration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>A prompt target configuration includes the following elements:</p>
+<h2>Basic Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#basic-configuration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#basic-configuration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Configuring prompt targets involves defining them in Plano’s configuration file. Each Prompt target specifies how a particular type of prompt should be handled, including the endpoint to invoke and any parameters required. A prompt target configuration includes the following elements:</p>
 <ul class="simple">
 <li><p><code class="docutils literal notranslate"><span class="pre">name</span></code>: A unique identifier for the prompt target.</p></li>
 <li><p><code class="docutils literal notranslate"><span class="pre">description</span></code>: A brief explanation of what the prompt target does.</p></li>
@@ -215,9 +210,9 @@
 </ul>
 </section>
 <section id="defining-parameters">
-<span id="defining-prompt-target-parameters"></span><h3>Defining Parameters<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#defining-parameters" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#defining-parameters'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Parameters are the pieces of information that Arch needs to extract from the user’s prompt to perform the desired action.
-Each parameter can be marked as required or optional. Here is a full list of parameter attributes that Arch can support:</p>
+<span id="defining-prompt-target-parameters"></span><h2>Defining Parameters<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#defining-parameters" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#defining-parameters'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Parameters are the pieces of information that Plano needs to extract from the user’s prompt to perform the desired action.
+Each parameter can be marked as required or optional. Here is a full list of parameter attributes that Plano can support:</p>
 <table class="docutils align-default" style="width: 100%">
 <thead>
 <tr class="row-odd"><th class="head"><p><strong>Attribute</strong></p></th>
@@ -256,9 +251,9 @@ Each parameter can be marked as required or optional. Here is a full list of par
 </table>
 </section>
 <section id="example-configuration-for-tools">
-<h3>Example Configuration For Tools<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-configuration-for-tools" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-configuration-for-tools'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<div class="literal-block-wrapper docutils container" id="id2">
-<div class="code-block-caption"><span class="caption-text">Tools and Function Calling Configuration Example</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<h2>Example Configuration For Tools<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-configuration-for-tools" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-configuration-for-tools'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<div class="literal-block-wrapper docutils container" id="id3">
+<div class="code-block-caption"><span class="caption-text">Tools and Function Calling Configuration Example</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id3"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">prompt_targets</span><span class="p">:</span>
 </span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">get_weather</span>
 </span><span id="line-3"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Get the current weather for a location</span>
@@ -279,52 +274,158 @@ Each parameter can be marked as required or optional. Here is a full list of par
 </div>
 </div>
 </section>
-<section id="example-configuration-for-agents">
-<h3>Example Configuration For Agents<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-configuration-for-agents" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-configuration-for-agents'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<div class="literal-block-wrapper docutils container" id="id3">
-<div class="code-block-caption"><span class="caption-text">Agent Orchestration Configuration Example</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id3"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">overrides</span><span class="p">:</span>
-</span><span id="line-2"><span class="w">  </span><span class="nt">use_agent_orchestrator</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-3">
-</span><span id="line-4"><span class="nt">prompt_targets</span><span class="p">:</span>
-</span><span id="line-5"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">sales_agent</span>
-</span><span id="line-6"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">handles queries related to sales and purchases</span>
-</span><span id="line-7">
-</span><span id="line-8"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">issues_and_repairs</span>
-</span><span id="line-9"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">handles issues, repairs, or refunds</span>
-</span><span id="line-10">
-</span><span id="line-11"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">escalate_to_human</span>
-</span><span id="line-12"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">escalates to human agent</span>
+<section id="multi-turn">
+<span id="plano-multi-turn-guide"></span><h2>Multi-Turn<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#multi-turn" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#multi-turn'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Developers often <a class="reference external" href="https://www.reddit.com/r/LocalLLaMA/comments/18mqwg6/best_practice_for_rag_with_followup_chat/" rel="nofollow noopener">struggle<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> to efficiently handle
+<code class="docutils literal notranslate"><span class="pre">follow-up</span></code> or <code class="docutils literal notranslate"><span class="pre">clarification</span></code> questions. Specifically, when users ask for changes or additions to previous responses, it requires developers to
+re-write prompts using LLMs with precise prompt engineering techniques. This process is slow, manual, error prone and adds latency and token cost for
+common scenarios that can be managed more efficiently.</p>
+<p>Plano is highly capable of accurately detecting and processing prompts in multi-turn scenarios so that you can buil fast and accurate agents in minutes.
+Below are some cnversational examples that you can build via Plano. Each example is enriched with annotations (via ** [Plano] ** ) that illustrates how Plano
+processess conversational messages on your behalf.</p>
+<p>Example 1: Adjusting Retrieval</p>
+<div class="highlight-text notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">User: What are the benefits of renewable energy?
+</span><span id="line-2">**[Plano]**: Check if there is an available &lt;prompt_target&gt; that can handle this user query.
+</span><span id="line-3">**[Plano]**: Found "get_info_for_energy_source" prompt_target in arch_config.yaml. Forward prompt to the endpoint configured in "get_info_for_energy_source"
+</span><span id="line-4">...
+</span><span id="line-5">Assistant: Renewable energy reduces greenhouse gas emissions, lowers air pollution, and provides sustainable power sources like solar and wind.
+</span><span id="line-6">
+</span><span id="line-7">User: Include cost considerations in the response.
+</span><span id="line-8">**[Plano]**: Follow-up detected. Forward prompt history to the "get_info_for_energy_source" prompt_target and post the following parameters consideration="cost"
+</span><span id="line-9">...
+</span><span id="line-10">Assistant: Renewable energy reduces greenhouse gas emissions, lowers air pollution, and provides sustainable power sources like solar and wind. While the initial setup costs can be high, long-term savings from reduced fuel expenses and government incentives make it cost-effective.
+</span></code></pre></div>
+</div>
+<section id="example-2-switching-intent">
+<h3>Example 2: Switching Intent<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-2-switching-intent" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-2-switching-intent'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<div class="highlight-text notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">User: What are the symptoms of diabetes?
+</span><span id="line-2">**[Plano]**: Check if there is an available &lt;prompt_target&gt; that can handle this user query.
+</span><span id="line-3">**[Plano]**: Found "diseases_symptoms" prompt_target in arch_config.yaml. Forward disease=diabeteres to "diseases_symptoms" prompt target
+</span><span id="line-4">...
+</span><span id="line-5">Assistant: Common symptoms include frequent urination, excessive thirst, fatigue, and blurry vision.
+</span><span id="line-6">
+</span><span id="line-7">User: How is it diagnosed?
+</span><span id="line-8">**[Plano]**: New intent detected.
+</span><span id="line-9">**[Plano]**: Found "disease_diagnoses" prompt_target in arch_config.yaml. Forward disease=diabeteres to "disease_diagnoses" prompt target
+</span><span id="line-10">...
+</span><span id="line-11">Assistant: Diabetes is diagnosed through blood tests like fasting blood sugar, A1C, or an oral glucose tolerance test.
+</span></code></pre></div>
+</div>
+</section>
+<section id="build-multi-turn-rag-apps">
+<h3>Build Multi-Turn RAG Apps<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#build-multi-turn-rag-apps" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#build-multi-turn-rag-apps'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>The following section describes how you can easilly add support for multi-turn scenarios via Plano. You process and manage multi-turn prompts
+just like you manage single-turn ones. Plano handles the conpleixity of detecting the correct intent based on the last user prompt and
+the covnersational history, extracts relevant parameters needed by downstream APIs, and dipatches calls to any upstream LLMs to summarize the
+response from your APIs.</p>
+</section>
+<section id="step-1-define-plano-config">
+<span id="multi-turn-subsection-prompt-target"></span><h3>Step 1: Define Plano Config<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-1-define-plano-config" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-1-define-plano-config'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<div class="literal-block-wrapper docutils container" id="id4">
+<div class="code-block-caption"><span class="caption-text">Plano Config</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id4"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1</span>
+</span><span id="line-2"><span class="linenos"> 2</span><span class="nt">listener</span><span class="p">:</span>
+</span><span id="line-3"><span class="linenos"> 3</span><span class="w">  </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1</span>
+</span><span id="line-4"><span class="linenos"> 4</span><span class="w">  </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">8080</span><span class="w"> </span><span class="c1">#If you configure port 443, you'll need to update the listener with tls_certificates</span>
+</span><span id="line-5"><span class="linenos"> 5</span><span class="w">  </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">huggingface</span>
+</span><span id="line-6"><span class="linenos"> 6</span>
+</span><span id="line-7"><span class="linenos"> 7</span><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
+</span><span id="line-8"><span class="linenos"> 8</span><span class="nt">llm_providers</span><span class="p">:</span>
+</span><span id="line-9"><span class="linenos"> 9</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">OpenAI</span>
+</span><span id="line-10"><span class="linenos">10</span><span class="w">    </span><span class="nt">provider</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
+</span><span id="line-11"><span class="linenos">11</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-12"><span class="linenos">12</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-3.5-turbo</span>
+</span><span id="line-13"><span class="linenos">13</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-14"><span class="linenos">14</span>
+</span><span id="line-15"><span class="linenos">15</span><span class="c1"># default system prompt used by all prompt targets</span>
+</span><span id="line-16"><span class="linenos">16</span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
+</span><span id="line-17"><span class="linenos">17</span><span class="w">   </span><span class="no">You are a helpful assistant and can offer information about energy sources. You will get a JSON object with energy_source and consideration fields. Focus on answering using those fields</span>
+</span><span id="line-18"><span class="linenos">18</span>
+</span><span id="line-19"><span class="linenos">19</span><span class="nt">prompt_targets</span><span class="p">:</span>
+</span><span id="line-20"><span class="linenos">20</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">get_info_for_energy_source</span>
+</span><span id="line-21"><span class="linenos">21</span><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">get information about an energy source</span>
+</span><span id="line-22"><span class="linenos">22</span><span class="w">    </span><span class="nt">parameters</span><span class="p">:</span>
+</span><span id="line-23"><span class="linenos">23</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">energy_source</span>
+</span><span id="line-24"><span class="linenos">24</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
+</span><span id="line-25"><span class="linenos">25</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">a source of energy</span>
+</span><span id="line-26"><span class="linenos">26</span><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-27"><span class="linenos">27</span><span class="w">        </span><span class="nt">enum</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">[</span><span class="nv">renewable</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">fossil</span><span class="p p-Indicator">]</span>
+</span><span id="line-28"><span class="linenos">28</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">consideration</span>
+</span><span id="line-29"><span class="linenos">29</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
+</span><span id="line-30"><span class="linenos">30</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">a specific type of consideration for an energy source</span>
+</span><span id="line-31"><span class="linenos">31</span><span class="w">        </span><span class="nt">enum</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">[</span><span class="nv">cost</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">economic</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">technology</span><span class="p p-Indicator">]</span>
+</span><span id="line-32"><span class="linenos">32</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
+</span><span id="line-33"><span class="linenos">33</span><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">rag_energy_source_agent</span>
+</span><span id="line-34"><span class="linenos">34</span><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/energy_source_info</span>
+</span><span id="line-35"><span class="linenos">35</span><span class="w">      </span><span class="nt">http_method</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">POST</span>
 </span></code></pre></div>
 </div>
 </div>
-<div class="admonition note">
-<p class="admonition-title">Note</p>
-<p>Today, you can use Arch to coordinate more specific agentic scenarios via tools and function calling, or use it for high-level agent routing and hand off scenarios. In the future, we plan to offer you the ability to combine these two approaches for more complex scenarios. Please see <a class="reference external" href="https://github.com/katanemo/archgw/issues/442" rel="nofollow noopener">github issues<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> for more details.</p>
-</div>
 </section>
-</section>
-<section id="routing-logic">
-<h2>Routing Logic<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#routing-logic" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#routing-logic'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Prompt targets determine where and how user prompts are processed. Arch uses intelligent routing logic to ensure that prompts are directed to the appropriate targets based on their intent and context.</p>
-<section id="default-targets">
-<h3>Default Targets<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#default-targets" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#default-targets'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>For general-purpose prompts that do not match any specific prompt target, Arch routes them to a designated default target. This is useful for handling open-ended queries like document summarization or information extraction.</p>
-</section>
-<section id="intent-matching">
-<h3>Intent Matching<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#intent-matching" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#intent-matching'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Arch analyzes the user’s prompt to determine its intent and matches it with the most suitable prompt target based on the name and description defined in the configuration.</p>
-<p>For example:</p>
-<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">Prompt:<span class="w"> </span><span class="s2">"Can you reboot the router?"</span>
-</span><span id="line-2">Matching<span class="w"> </span>Target:<span class="w"> </span>reboot_device<span class="w"> </span><span class="o">(</span>based<span class="w"> </span>on<span class="w"> </span>description<span class="w"> </span>matching<span class="w"> </span><span class="s2">"reboot devices"</span><span class="o">)</span>
+<section id="step-2-process-request-in-flask">
+<h3>Step 2: Process Request in Flask<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-process-request-in-flask" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-2-process-request-in-flask'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Once the prompt targets are configured as above, handle parameters across multi-turn as if its a single-turn request</p>
+<div class="literal-block-wrapper docutils container" id="id5">
+<div class="code-block-caption"><span class="caption-text">Parameter handling with Flask</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id5"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="kn">import</span><span class="w"> </span><span class="nn">os</span>
+</span><span id="line-2"><span class="linenos"> 2</span><span class="kn">import</span><span class="w"> </span><span class="nn">gradio</span><span class="w"> </span><span class="k">as</span><span class="w"> </span><span class="nn">gr</span>
+</span><span id="line-3"><span class="linenos"> 3</span>
+</span><span id="line-4"><span class="linenos"> 4</span><span class="kn">from</span><span class="w"> </span><span class="nn">fastapi</span><span class="w"> </span><span class="kn">import</span> <span class="n">FastAPI</span><span class="p">,</span> <span class="n">HTTPException</span>
+</span><span id="line-5"><span class="linenos"> 5</span><span class="kn">from</span><span class="w"> </span><span class="nn">pydantic</span><span class="w"> </span><span class="kn">import</span> <span class="n">BaseModel</span>
+</span><span id="line-6"><span class="linenos"> 6</span><span class="kn">from</span><span class="w"> </span><span class="nn">typing</span><span class="w"> </span><span class="kn">import</span> <span class="n">Optional</span>
+</span><span id="line-7"><span class="linenos"> 7</span><span class="kn">from</span><span class="w"> </span><span class="nn">openai</span><span class="w"> </span><span class="kn">import</span> <span class="n">OpenAI</span>
+</span><span id="line-8"><span class="linenos"> 8</span><span class="kn">from</span><span class="w"> </span><span class="nn">common</span><span class="w"> </span><span class="kn">import</span> <span class="n">create_gradio_app</span>
+</span><span id="line-9"><span class="linenos"> 9</span>
+</span><span id="line-10"><span class="linenos">10</span><span class="n">app</span> <span class="o">=</span> <span class="n">FastAPI</span><span class="p">()</span>
+</span><span id="line-11"><span class="linenos">11</span>
+</span><span id="line-12"><span class="linenos">12</span>
+</span><span id="line-13"><span class="linenos">13</span><span class="c1"># Define the request model</span>
+</span><span id="line-14"><span class="linenos">14</span><span class="k">class</span><span class="w"> </span><span class="nc">EnergySourceRequest</span><span class="p">(</span><span class="n">BaseModel</span><span class="p">):</span>
+</span><span id="line-15"><span class="linenos">15</span>    <span class="n">energy_source</span><span class="p">:</span> <span class="nb">str</span>
+</span><span id="line-16"><span class="linenos">16</span>    <span class="n">consideration</span><span class="p">:</span> <span class="n">Optional</span><span class="p">[</span><span class="nb">str</span><span class="p">]</span> <span class="o">=</span> <span class="kc">None</span>
+</span><span id="line-17"><span class="linenos">17</span>
+</span><span id="line-18"><span class="linenos">18</span>
+</span><span id="line-19"><span class="linenos">19</span><span class="k">class</span><span class="w"> </span><span class="nc">EnergySourceResponse</span><span class="p">(</span><span class="n">BaseModel</span><span class="p">):</span>
+</span><span id="line-20"><span class="linenos">20</span>    <span class="n">energy_source</span><span class="p">:</span> <span class="nb">str</span>
+</span><span id="line-21"><span class="linenos">21</span>    <span class="n">consideration</span><span class="p">:</span> <span class="n">Optional</span><span class="p">[</span><span class="nb">str</span><span class="p">]</span> <span class="o">=</span> <span class="kc">None</span>
+</span><span id="line-22"><span class="linenos">22</span>
+</span><span id="line-23"><span class="linenos">23</span>
+</span><span id="line-24"><span class="linenos">24</span><span class="c1"># Post method for device summary</span>
+</span><span id="line-25"><span class="linenos">25</span><span class="nd">@app</span><span class="o">.</span><span class="n">post</span><span class="p">(</span><span class="s2">"/agent/energy_source_info"</span><span class="p">)</span>
+</span><span id="line-26"><span class="linenos">26</span><span class="k">def</span><span class="w"> </span><span class="nf">get_workforce</span><span class="p">(</span><span class="n">request</span><span class="p">:</span> <span class="n">EnergySourceRequest</span><span class="p">):</span>
+</span><span id="line-27"><span class="linenos">27</span><span class="w">    </span><span class="sd">"""</span>
+</span><span id="line-28"><span class="linenos">28</span><span class="sd">    Endpoint to get details about energy source</span>
+</span><span id="line-29"><span class="linenos">29</span><span class="sd">    """</span>
+</span><span id="line-30"><span class="linenos">30</span>    <span class="n">considertion</span> <span class="o">=</span> <span class="s2">"You don't have any specific consideration. Feel free to talk in a more open ended fashion"</span>
+</span><span id="line-31"><span class="linenos">31</span>
+</span><span id="line-32"><span class="linenos">32</span>    <span class="k">if</span> <span class="n">request</span><span class="o">.</span><span class="n">consideration</span> <span class="ow">is</span> <span class="ow">not</span> <span class="kc">None</span><span class="p">:</span>
+</span><span id="line-33"><span class="linenos">33</span>        <span class="n">considertion</span> <span class="o">=</span> <span class="sa">f</span><span class="s2">"Add specific focus on the following consideration when you summarize the content for the energy source: </span><span class="si">{</span><span class="n">request</span><span class="o">.</span><span class="n">consideration</span><span class="si">}</span><span class="s2">"</span>
+</span><span id="line-34"><span class="linenos">34</span>
+</span><span id="line-35"><span class="linenos">35</span>    <span class="n">response</span> <span class="o">=</span> <span class="p">{</span>
+</span><span id="line-36"><span class="linenos">36</span>        <span class="s2">"energy_source"</span><span class="p">:</span> <span class="n">request</span><span class="o">.</span><span class="n">energy_source</span><span class="p">,</span>
+</span><span id="line-37"><span class="linenos">37</span>        <span class="s2">"consideration"</span><span class="p">:</span> <span class="n">considertion</span><span class="p">,</span>
+</span><span id="line-38"><span class="linenos">38</span>    <span class="p">}</span>
+</span><span id="line-39"><span class="linenos">39</span>    <span class="k">return</span> <span class="n">response</span>
 </span></code></pre></div>
 </div>
+</div>
+</section>
+<section id="demo-app">
+<h3>Demo App<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#demo-app" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#demo-app'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>For your convenience, we’ve built a <a class="reference external" href="https://github.com/katanemo/archgw/tree/main/demos/samples_python/multi_turn_rag_agent" rel="nofollow noopener">demo app<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>
+that you can test and modify locally for multi-turn RAG scenarios.</p>
+<figure class="align-center" id="id6">
+<a class="reference internal image-reference" href="../_images/mutli-turn-example.png"><img alt="../_images/mutli-turn-example.png" src="../_images/mutli-turn-example.png" style="width: 100%;"/>
+</a>
+<figcaption>
+<p><span class="caption-text">Example multi-turn user conversation showing adjusting retrieval</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id6"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></p>
+</figcaption>
+</figure>
 </section>
 </section>
 <section id="summary">
 <h2>Summary<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#summary" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#summary'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Prompt targets are essential for defining how user prompts are handled within your generative AI applications using Arch.</p>
-<p>By carefully configuring prompt targets, you can ensure that prompts are accurately routed, necessary parameters are extracted, and backend services are invoked seamlessly. This modular approach not only simplifies your application’s architecture but also enhances scalability, maintainability, and overall user experience.</p>
+<p>By carefully designing prompt targets as deterministic, task-specific entry points, you ensure that prompts are routed to the right workload, necessary parameters are cleanly extracted and validated, and backend services are invoked with structured inputs. This clear separation between prompt handling and business logic simplifies your architecture, makes behavior more predictable and testable, and improves the scalability and maintainability of your agentic applications.</p>
 </section>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
@@ -337,8 +438,8 @@ Each parameter can be marked as required or optional. Here is a full list of par
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../guides/prompt_guard.html">
-        Prompt Guard
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../guides/orchestration.html">
+        Orchestration
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -347,20 +448,16 @@ Each parameter can be marked as required or optional. Here is a full list of par
 </div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
 <div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
 <ul>
-<li><a :data-current="activeSection === '#what-are-prompt-targets'" class="reference internal" href="#what-are-prompt-targets">What Are Prompt Targets?</a><ul>
 <li><a :data-current="activeSection === '#key-features'" class="reference internal" href="#key-features">Key Features</a></li>
-</ul>
-</li>
-<li><a :data-current="activeSection === '#configuring-prompt-targets'" class="reference internal" href="#configuring-prompt-targets">Configuring Prompt Targets</a><ul>
 <li><a :data-current="activeSection === '#basic-configuration'" class="reference internal" href="#basic-configuration">Basic Configuration</a></li>
 <li><a :data-current="activeSection === '#defining-parameters'" class="reference internal" href="#defining-parameters">Defining Parameters</a></li>
 <li><a :data-current="activeSection === '#example-configuration-for-tools'" class="reference internal" href="#example-configuration-for-tools">Example Configuration For Tools</a></li>
-<li><a :data-current="activeSection === '#example-configuration-for-agents'" class="reference internal" href="#example-configuration-for-agents">Example Configuration For Agents</a></li>
-</ul>
-</li>
-<li><a :data-current="activeSection === '#routing-logic'" class="reference internal" href="#routing-logic">Routing Logic</a><ul>
-<li><a :data-current="activeSection === '#default-targets'" class="reference internal" href="#default-targets">Default Targets</a></li>
-<li><a :data-current="activeSection === '#intent-matching'" class="reference internal" href="#intent-matching">Intent Matching</a></li>
+<li><a :data-current="activeSection === '#multi-turn'" class="reference internal" href="#multi-turn">Multi-Turn</a><ul>
+<li><a :data-current="activeSection === '#example-2-switching-intent'" class="reference internal" href="#example-2-switching-intent">Example 2: Switching Intent</a></li>
+<li><a :data-current="activeSection === '#build-multi-turn-rag-apps'" class="reference internal" href="#build-multi-turn-rag-apps">Build Multi-Turn RAG Apps</a></li>
+<li><a :data-current="activeSection === '#step-1-define-plano-config'" class="reference internal" href="#step-1-define-plano-config">Step 1: Define Plano Config</a></li>
+<li><a :data-current="activeSection === '#step-2-process-request-in-flask'" class="reference internal" href="#step-2-process-request-in-flask">Step 2: Process Request in Flask</a></li>
+<li><a :data-current="activeSection === '#demo-app'" class="reference internal" href="#demo-app">Demo App</a></li>
 </ul>
 </li>
 <li><a :data-current="activeSection === '#summary'" class="reference internal" href="#summary">Summary</a></li>
@@ -372,12 +469,12 @@ Each parameter can be marked as required or optional. Here is a full list of par
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../_static/doctools.js?v=9bcbadda"></script>
 <script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
diff --git a/concepts/tech_overview/error_target.html b/concepts/tech_overview/error_target.html
deleted file mode 100755
index 7bbb245b..00000000
--- a/concepts/tech_overview/error_target.html
+++ /dev/null
@@ -1,253 +0,0 @@
-<!DOCTYPE html>
-
-<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
-<head>
-<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
-<meta charset="utf-8"/>
-<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
-<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
-<meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Error Target | Arch Docs v0.3.22</title>
-<meta content="Error Target | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Error Target | Arch Docs v0.3.22" name="twitter:title"/>
-<link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
-<link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
-<link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
-<link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/concepts/tech_overview/error_target.html" rel="canonical"/>
-<link href="../../_static/favicon.ico" rel="icon"/>
-<link href="../../search.html" rel="search" title="Search"/>
-<link href="../llm_providers/llm_providers.html" rel="next" title="LLM Providers"/>
-<link href="request_lifecycle.html" rel="prev" title="Request Lifecycle"/>
-<script>
-    <!-- Prevent Flash of wrong theme -->
-      const userPreference = localStorage.getItem('darkMode');
-      let mode;
-      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
-        mode = 'dark';
-        document.documentElement.classList.add('dark');
-      } else {
-        mode = 'light';
-      }
-      if (!userPreference) {localStorage.setItem('darkMode', mode)}
-    </script>
-</head>
-<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
-<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
-      Skip to content
-    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
-<div class="hidden mr-4 md:flex">
-<a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
-<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
-</svg>
-<span class="sr-only">Toggle navigation menu</span>
-</button>
-<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
-<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../../search.html" class="relative flex items-center group" id="searchbox" method="get">
-<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
-<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
-<span class="text-xs">⌘</span>
-    K
-  </kbd>
-</form>
-</div>
-<nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
-<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
-<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
-</div>
-</a>
-<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
-<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
-</svg>
-<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
-</svg>
-</button>
-</nav>
-</div>
-</div>
-</header>
-<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
-<a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a>
-<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
-<div class="overflow-y-auto h-full w-full relative pr-6">
-
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
-<script>
-  window.dataLayer = window.dataLayer || [];
-  function gtag(){dataLayer.push(arguments);}
-  gtag('js', new Date());
-
-  gtag('config', 'G-K2LXXSX6HB');
-</script>
-<nav class="table w-full min-w-full my-6 lg:my-8">
-<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul class="current">
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2 current"><a class="current reference internal" href="#">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../prompt_target.html">Prompt Target</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
-</ul>
-</li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
-</ul>
-</nav>
-</div>
-</div>
-<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
-<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
-</svg>
-</button>
-</aside>
-<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
-<div class="w-full min-w-0 mx-auto">
-<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
-<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
-<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
-<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
-</svg>
-</a>
-<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="tech_overview.html">Tech Overview</a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Error Target</span>
-</nav>
-<div id="content" role="main">
-<section id="error-target">
-<span id="id1"></span><h1>Error Target<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#error-target"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p><strong>Error targets</strong> are designed to capture and manage specific issues or exceptions that occur during Arch’s function or system’s execution.</p>
-<p>These endpoints receive errors forwarded from Arch when issues arise, such as improper function/API calls, guardrail violations, or other processing errors.
-The errors are communicated to the application via headers like <code class="docutils literal notranslate"><span class="pre">X-Arch-[ERROR-TYPE]</span></code>, enabling you to respond appropriately and handle errors gracefully.</p>
-<section id="key-concepts">
-<h2>Key Concepts<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#key-concepts" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#key-concepts'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<ul class="simple">
-<li><p><strong>Error Type</strong>: Categorizes the nature of the error, such as “ValidationError” or “RuntimeError.” These error types help in identifying what kind of issue occurred and provide context for troubleshooting.</p></li>
-<li><p><strong>Error Message</strong>: A clear, human-readable message describing the error. This should provide enough detail to inform users or developers of the root cause or required action.</p></li>
-<li><p><strong>Parameter-Specific Errors</strong>: Errors that arise due to invalid or missing parameters when invoking a function. These errors are critical for ensuring the correctness of inputs.</p></li>
-</ul>
-</section>
-<section id="error-header-example">
-<h2>Error Header Example<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#error-header-example" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#error-header-example'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<div class="literal-block-wrapper docutils container" id="id2">
-<div class="code-block-caption"><span class="caption-text">Error Header Example</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="w">  </span>HTTP/1.1<span class="w"> </span><span class="m">400</span><span class="w"> </span>Bad<span class="w"> </span>Request
-</span><span id="line-2"><span class="w">  </span>X-Arch-Error-Type:<span class="w"> </span>FunctionValidationError
-</span><span id="line-3"><span class="w">  </span>X-Arch-Error-Message:<span class="w"> </span>Tools<span class="w"> </span>call<span class="w"> </span>parsing<span class="w"> </span>failure
-</span><span id="line-4"><span class="w">  </span>X-Arch-Target-Prompt:<span class="w"> </span>createUser
-</span><span id="line-5"><span class="w">  </span>Content-Type:<span class="w"> </span>application/json
-</span><span id="line-6">
-</span><span id="line-7"><span class="w">  </span><span class="s2">"messages"</span>:<span class="w"> </span><span class="o">[</span>
-</span><span id="line-8"><span class="w">      </span><span class="o">{</span>
-</span><span id="line-9"><span class="w">        </span><span class="s2">"role"</span>:<span class="w"> </span><span class="s2">"user"</span>,
-</span><span id="line-10"><span class="w">        </span><span class="s2">"content"</span>:<span class="w"> </span><span class="s2">"Please create a user with the following ID: 1234"</span>
-</span><span id="line-11"><span class="w">      </span><span class="o">}</span>,
-</span><span id="line-12"><span class="w">      </span><span class="o">{</span>
-</span><span id="line-13"><span class="w">        </span><span class="s2">"role"</span>:<span class="w"> </span><span class="s2">"system"</span>,
-</span><span id="line-14"><span class="w">        </span><span class="s2">"content"</span>:<span class="w"> </span><span class="s2">"Expected a string for 'user_id', but got an integer."</span>
-</span><span id="line-15"><span class="w">      </span><span class="o">}</span>
-</span><span id="line-16"><span class="w">  </span><span class="o">]</span>
-</span></code></pre></div>
-</div>
-</div>
-</section>
-<section id="best-practices-and-tips">
-<h2>Best Practices and Tips<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#best-practices-and-tips" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#best-practices-and-tips'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<ul class="simple">
-<li><p><strong>Graceful Degradation</strong>: If an error occurs, fail gracefully by providing fallback logic or alternative flows when possible.</p></li>
-<li><p><strong>Log Errors</strong>: Always log errors on the server side for later analysis.</p></li>
-<li><p><strong>Client-Side Handling</strong>: Make sure the client can interpret error responses and provide meaningful feedback to the user. Clients should not display raw error codes or stack traces but rather handle them gracefully.</p></li>
-</ul>
-</section>
-</section>
-</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
-<div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="request_lifecycle.html">
-<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="15 18 9 12 15 6"></polyline>
-</svg>
-        Request Lifecycle
-      </a>
-</div>
-<div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../llm_providers/llm_providers.html">
-        LLM Providers
-        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="9 18 15 12 9 6"></polyline>
-</svg>
-</a>
-</div>
-</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
-<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
-<ul>
-<li><a :data-current="activeSection === '#key-concepts'" class="reference internal" href="#key-concepts">Key Concepts</a></li>
-<li><a :data-current="activeSection === '#error-header-example'" class="reference internal" href="#error-header-example">Error Header Example</a></li>
-<li><a :data-current="activeSection === '#best-practices-and-tips'" class="reference internal" href="#best-practices-and-tips">Best Practices and Tips</a></li>
-</ul>
-</div>
-</aside>
-</main>
-</div>
-</div><footer class="py-6 border-t border-border md:py-0">
-<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
-<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
-</div>
-</div>
-</footer>
-</div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
-<script src="../../_static/doctools.js?v=9bcbadda"></script>
-<script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
-<script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
-<script src="../../_static/design-tabs.js?v=f930bc37"></script>
-</body>
-</html>
\ No newline at end of file
diff --git a/concepts/tech_overview/prompt.html b/concepts/tech_overview/prompt.html
deleted file mode 100755
index 34666684..00000000
--- a/concepts/tech_overview/prompt.html
+++ /dev/null
@@ -1,416 +0,0 @@
-<!DOCTYPE html>
-
-<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
-<head>
-<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
-<meta charset="utf-8"/>
-<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
-<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
-<meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Prompts | Arch Docs v0.3.22</title>
-<meta content="Prompts | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Prompts | Arch Docs v0.3.22" name="twitter:title"/>
-<link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
-<link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
-<link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
-<link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/concepts/tech_overview/prompt.html" rel="canonical"/>
-<link href="../../_static/favicon.ico" rel="icon"/>
-<link href="../../search.html" rel="search" title="Search"/>
-<link href="model_serving.html" rel="next" title="Model Serving"/>
-<link href="listener.html" rel="prev" title="Listener"/>
-<script>
-    <!-- Prevent Flash of wrong theme -->
-      const userPreference = localStorage.getItem('darkMode');
-      let mode;
-      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
-        mode = 'dark';
-        document.documentElement.classList.add('dark');
-      } else {
-        mode = 'light';
-      }
-      if (!userPreference) {localStorage.setItem('darkMode', mode)}
-    </script>
-</head>
-<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
-<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
-      Skip to content
-    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
-<div class="hidden mr-4 md:flex">
-<a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
-<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
-</svg>
-<span class="sr-only">Toggle navigation menu</span>
-</button>
-<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
-<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../../search.html" class="relative flex items-center group" id="searchbox" method="get">
-<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
-<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
-<span class="text-xs">⌘</span>
-    K
-  </kbd>
-</form>
-</div>
-<nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
-<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
-<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
-</div>
-</a>
-<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
-<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
-</svg>
-<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
-</svg>
-</button>
-</nav>
-</div>
-</div>
-</header>
-<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
-<a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a>
-<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
-<div class="overflow-y-auto h-full w-full relative pr-6">
-
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
-<script>
-  window.dataLayer = window.dataLayer || [];
-  function gtag(){dataLayer.push(arguments);}
-  gtag('js', new Date());
-
-  gtag('config', 'G-K2LXXSX6HB');
-</script>
-<nav class="table w-full min-w-full my-6 lg:my-8">
-<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul class="current">
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="listener.html">Listener</a></li>
-<li class="toctree-l2 current"><a class="current reference internal" href="#">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../prompt_target.html">Prompt Target</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
-</ul>
-</li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
-</ul>
-</nav>
-</div>
-</div>
-<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
-<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
-</svg>
-</button>
-</aside>
-<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
-<div class="w-full min-w-0 mx-auto">
-<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
-<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
-<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
-<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
-</svg>
-</a>
-<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="tech_overview.html">Tech Overview</a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Prompts</span>
-</nav>
-<div id="content" role="main">
-<section id="prompts">
-<span id="arch-overview-prompt-handling"></span><h1>Prompts<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#prompts"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Arch’s primary design point is to securely accept, process and handle prompts. To do that effectively,
-Arch relies on Envoy’s HTTP <a class="reference external" href="https://www.envoyproxy.io/docs/envoy/v1.31.2/intro/arch_overview/http/http_connection_management" rel="nofollow noopener">connection management<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>,
-subsystem and its <strong>prompt handler</strong> subsystem engineered with purpose-built LLMs to
-implement critical functionality on behalf of developers so that you can stay focused on business logic.</p>
-<p>Arch’s <strong>prompt handler</strong> subsystem interacts with the <strong>model subsystem</strong> through Envoy’s cluster manager system to ensure robust, resilient and fault-tolerant experience in managing incoming prompts.</p>
-<div class="admonition seealso">
-<p class="admonition-title">See also</p>
-<p>Read more about the <a class="reference internal" href="model_serving.html#model-serving"><span class="std std-ref">model subsystem</span></a> and how the LLMs are hosted in Arch.</p>
-</div>
-<section id="messages">
-<h2>Messages<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#messages" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#messages'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch accepts messages directly from the body of the HTTP request in a format that follows the <a class="reference external" href="https://huggingface.co/docs/text-generation-inference/en/messages_api" rel="nofollow noopener">Hugging Face Messages API<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.
-This design allows developers to pass a list of messages, where each message is represented as a dictionary
-containing two key-value pairs:</p>
-<blockquote>
-<div><ul class="simple">
-<li><p><strong>Role</strong>: Defines the role of the message sender, such as “user” or “assistant”.</p></li>
-<li><p><strong>Content</strong>: Contains the actual text of the message.</p></li>
-</ul>
-</div></blockquote>
-</section>
-<section id="prompt-guard">
-<h2>Prompt Guard<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#prompt-guard" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#prompt-guard'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch is engineered with <a class="reference external" href="https://huggingface.co/collections/katanemo/arch-guard-6702bdc08b889e4bce8f446d" rel="nofollow noopener">Arch-Guard<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, an industry leading safety layer, powered by a
-compact and high-performing LLM that monitors incoming prompts to detect and reject jailbreak attempts -
-ensuring that unauthorized or harmful behaviors are intercepted early in the process.</p>
-<p>To add jailbreak guardrails, see example below:</p>
-<div class="literal-block-wrapper docutils container" id="id1">
-<div class="code-block-caption"><span class="caption-text">Example Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id1"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1.0</span>
-</span><span id="line-2"><span class="linenos"> 2</span>
-</span><span id="line-3"><span class="linenos"> 3</span><span class="nt">listeners</span><span class="p">:</span>
-</span><span id="line-4"><span class="linenos"> 4</span><span class="w">  </span><span class="nt">ingress_traffic</span><span class="p">:</span>
-</span><span id="line-5"><span class="linenos"> 5</span><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
-</span><span id="line-6"><span class="linenos"> 6</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">10000</span>
-</span><span id="line-7"><span class="linenos"> 7</span><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-8"><span class="linenos"> 8</span><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
-</span><span id="line-9"><span class="linenos"> 9</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
-</span><span id="line-11"><span class="linenos">11</span><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-12"><span class="linenos">12</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-13"><span class="linenos">13</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
-</span><span id="line-14"><span class="linenos">14</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-15"><span class="linenos">15</span>
-</span><span id="line-16"><span class="linenos">16</span><span class="c1"># default system prompt used by all prompt targets</span>
-</span><span id="line-17"><span class="linenos">17</span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">You are a network assistant that just offers facts; not advice on manufacturers or purchasing decisions.</span>
-</span><span id="line-18"><span class="linenos">18</span>
-</span><span id="line-19"><span class="linenos">19</span><span class="nt">prompt_guards</span><span class="p">:</span>
-</span><span id="line-20"><span class="linenos">20</span><span class="w">  </span><span class="nt">input_guards</span><span class="p">:</span>
-</span><span id="line-21"><mark><span class="linenos">21</span><span class="w">    </span><span class="nt">jailbreak</span><span class="p">:</span>
-</mark></span><span id="line-22"><mark><span class="linenos">22</span><span class="w">      </span><span class="nt">on_exception</span><span class="p">:</span>
-</mark></span><span id="line-23"><mark><span class="linenos">23</span><span class="w">        </span><span class="nt">message</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Looks like you're curious about my abilities, but I can only provide assistance within my programmed parameters.</span>
-</mark></span><span id="line-24"><mark><span class="linenos">24</span>
-</mark></span><span id="line-25"><mark><span class="linenos">25</span><span class="nt">prompt_targets</span><span class="p">:</span>
-</mark></span></code></pre></div>
-</div>
-</div>
-<div class="admonition note">
-<p class="admonition-title">Note</p>
-<p>As a roadmap item, Arch will expose the ability for developers to define custom guardrails via Arch-Guard,
-and add support for additional safety checks defined by developers and hazardous categories like, violent crimes, privacy, hate,
-etc. To offer feedback on our roadmap, please visit our <a class="reference external" href="https://github.com/orgs/katanemo/projects/1" rel="nofollow noopener">github page<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a></p>
-</div>
-</section>
-<section id="prompt-targets">
-<h2>Prompt Targets<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#prompt-targets" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#prompt-targets'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Once a prompt passes any configured guardrail checks, Arch processes the contents of the incoming conversation
-and identifies where to forward the conversation to via its <code class="docutils literal notranslate"><span class="pre">prompt</span> <span class="pre">target</span></code> primitive. Prompt targets are endpoints
-that receive prompts that are processed by Arch. For example, Arch enriches incoming prompts with metadata like knowing
-when a user’s intent has changed so that you can build faster, more accurate RAG apps.</p>
-<p>Configuring <code class="docutils literal notranslate"><span class="pre">prompt_targets</span></code> is simple. See example below:</p>
-<div class="literal-block-wrapper docutils container" id="id2">
-<div class="code-block-caption"><span class="caption-text">Example Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1.0</span>
-</span><span id="line-2"><span class="linenos"> 2</span>
-</span><span id="line-3"><span class="linenos"> 3</span><span class="nt">listeners</span><span class="p">:</span>
-</span><span id="line-4"><span class="linenos"> 4</span><span class="w">  </span><span class="nt">ingress_traffic</span><span class="p">:</span>
-</span><span id="line-5"><span class="linenos"> 5</span><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
-</span><span id="line-6"><span class="linenos"> 6</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">10000</span>
-</span><span id="line-7"><span class="linenos"> 7</span><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-8"><span class="linenos"> 8</span><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
-</span><span id="line-9"><span class="linenos"> 9</span>
-</span><span id="line-10"><span class="linenos">10</span><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
-</span><span id="line-11"><span class="linenos">11</span><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-12"><span class="linenos">12</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-13"><span class="linenos">13</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
-</span><span id="line-14"><span class="linenos">14</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-15"><span class="linenos">15</span>
-</span><span id="line-16"><span class="linenos">16</span><span class="c1"># default system prompt used by all prompt targets</span>
-</span><span id="line-17"><span class="linenos">17</span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">You are a network assistant that just offers facts; not advice on manufacturers or purchasing decisions.</span>
-</span><span id="line-18"><span class="linenos">18</span>
-</span><span id="line-19"><span class="linenos">19</span><span class="nt">prompt_guards</span><span class="p">:</span>
-</span><span id="line-20"><span class="linenos">20</span><span class="w">  </span><span class="nt">input_guards</span><span class="p">:</span>
-</span><span id="line-21"><span class="linenos">21</span><span class="w">    </span><span class="nt">jailbreak</span><span class="p">:</span>
-</span><span id="line-22"><span class="linenos">22</span><span class="w">      </span><span class="nt">on_exception</span><span class="p">:</span>
-</span><span id="line-23"><span class="linenos">23</span><span class="w">        </span><span class="nt">message</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Looks like you're curious about my abilities, but I can only provide assistance within my programmed parameters.</span>
-</span><span id="line-24"><span class="linenos">24</span>
-</span><span id="line-25"><span class="linenos">25</span><span class="nt">prompt_targets</span><span class="p">:</span>
-</span><span id="line-26"><span class="linenos">26</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">information_extraction</span>
-</span><span id="line-27"><span class="linenos">27</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-28"><span class="linenos">28</span><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">handel all scenarios that are question and answer in nature. Like summarization, information extraction, etc.</span>
-</span><span id="line-29"><span class="linenos">29</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
-</span><span id="line-30"><span class="linenos">30</span><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</span><span id="line-31"><span class="linenos">31</span><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/summary</span>
-</span><span id="line-32"><span class="linenos">32</span><span class="w">    </span><span class="c1"># Arch uses the default LLM and treats the response from the endpoint as the prompt to send to the LLM</span>
-</span><span id="line-33"><span class="linenos">33</span><span class="w">    </span><span class="nt">auto_llm_dispatch_on_response</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-34"><span class="linenos">34</span><span class="w">    </span><span class="c1"># override system prompt for this prompt target</span>
-</span><span id="line-35"><span class="linenos">35</span><span class="w">    </span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">You are a helpful information extraction assistant. Use the information that is provided to you.</span>
-</span><span id="line-36"><span class="linenos">36</span>
-</span><span id="line-37"><span class="linenos">37</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">reboot_network_device</span>
-</span><span id="line-38"><span class="linenos">38</span><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Reboot a specific network device</span>
-</span><span id="line-39"><mark><span class="linenos">39</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
-</mark></span><span id="line-40"><mark><span class="linenos">40</span><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</mark></span><span id="line-41"><mark><span class="linenos">41</span><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/action</span>
-</mark></span><span id="line-42"><mark><span class="linenos">42</span><span class="w">    </span><span class="nt">parameters</span><span class="p">:</span>
-</mark></span><span id="line-43"><mark><span class="linenos">43</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_id</span>
-</mark></span><span id="line-44"><mark><span class="linenos">44</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
-</mark></span><span id="line-45"><mark><span class="linenos">45</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Identifier of the network device to reboot.</span>
-</mark></span><span id="line-46"><mark><span class="linenos">46</span><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</mark></span><span id="line-47"><mark><span class="linenos">47</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">confirmation</span>
-</mark></span><span id="line-48"><mark><span class="linenos">48</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">bool</span>
-</mark></span><span id="line-49"><mark><span class="linenos">49</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Confirmation flag to proceed with reboot.</span>
-</mark></span><span id="line-50"><mark><span class="linenos">50</span><span class="w">        </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">false</span>
-</mark></span><span id="line-51"><mark><span class="linenos">51</span><span class="w">        </span><span class="nt">enum</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">[</span><span class="nv">true</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">false</span><span class="p p-Indicator">]</span>
-</mark></span><span id="line-52"><mark><span class="linenos">52</span>
-</mark></span><span id="line-53"><mark><span class="linenos">53</span><span class="c1"># Arch creates a round-robin load balancing between different endpoints, managed via the cluster subsystem.</span>
-</mark></span><span id="line-54"><span class="linenos">54</span><span class="nt">endpoints</span><span class="p">:</span>
-</span><span id="line-55"><span class="linenos">55</span><span class="w">  </span><span class="nt">app_server</span><span class="p">:</span>
-</span><span id="line-56"><span class="linenos">56</span><span class="w">    </span><span class="c1"># value could be ip address or a hostname with port</span>
-</span><span id="line-57"><span class="linenos">57</span><span class="w">    </span><span class="c1"># this could also be a list of endpoints for load balancing</span>
-</span><span id="line-58"><span class="linenos">58</span><span class="w">    </span><span class="c1"># for example endpoint: [ ip1:port, ip2:port ]</span>
-</span><span id="line-59"><span class="linenos">59</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1:80</span>
-</span><span id="line-60"><span class="linenos">60</span><span class="w">    </span><span class="c1"># max time to wait for a connection to be established</span>
-</span><span id="line-61"><span class="linenos">61</span><span class="w">    </span><span class="nt">connect_timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.005s</span>
-</span></code></pre></div>
-</div>
-</div>
-<div class="admonition seealso">
-<p class="admonition-title">See also</p>
-<p>Check <a class="reference internal" href="../prompt_target.html#prompt-target"><span class="std std-ref">Prompt Target</span></a> for more details!</p>
-</div>
-<section id="intent-matching">
-<h3>Intent Matching<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#intent-matching" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#intent-matching'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Arch uses fast text embedding and intent recognition approaches to first detect the intent of each incoming prompt.
-This intent matching phase analyzes the prompt’s content and matches it against predefined prompt targets, ensuring that each prompt is forwarded to the most appropriate endpoint.
-Arch’s intent matching framework considers both the name and description of each prompt target, and uses a composite matching score between embedding similarity and intent classification scores to enhance accuracy in forwarding decisions.</p>
-<ul class="simple">
-<li><p><strong>Intent Recognition</strong>: NLI techniques further refine the matching process by evaluating the semantic alignment between the prompt and potential targets.</p></li>
-<li><p><strong>Text Embedding</strong>: By embedding the prompt and comparing it to known target vectors, Arch effectively identifies the closest match, ensuring that the prompt is handled by the correct downstream service.</p></li>
-</ul>
-</section>
-<section id="agentic-apps-via-prompt-targets">
-<h3>Agentic Apps via Prompt Targets<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#agentic-apps-via-prompt-targets" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#agentic-apps-via-prompt-targets'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>To support agentic apps, like scheduling travel plans or sharing comments on a document - via prompts, Arch uses its function calling abilities to extract critical information from the incoming prompt (or a set of prompts) needed by a downstream backend API or function call before calling it directly.
-For more details on how you can build agentic applications using Arch, see our full guide <a class="reference internal" href="../../build_with_arch/agent.html#arch-agent-guide"><span class="std std-ref">here</span></a>:</p>
-<div class="admonition note">
-<p class="admonition-title">Note</p>
-<p><a class="reference external" href="https://huggingface.co/collections/katanemo/arch-function-66f209a693ea8df14317ad68" rel="nofollow noopener">Arch-Function<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> is a collection of dedicated agentic models engineered in Arch to extract information from a (set of) prompts and executes necessary backend API calls.
-This allows for efficient handling of agentic tasks, such as scheduling data retrieval, by dynamically interacting with backend services.
-Arch-Function achieves state-of-the-art performance, comparable with frontier models like Claude Sonnet 3.5 ang GPT-4, while being 44x cheaper ($0.10M/token hosted) and 10x faster (p50 latencies of 200ms).</p>
-</div>
-</section>
-</section>
-<section id="prompting-llms">
-<h2>Prompting LLMs<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#prompting-llms" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#prompting-llms'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch is a single piece of software that is designed to manage both ingress and egress prompt traffic, drawing its distributed proxy nature from the robust <a class="reference external" href="https://envoyproxy.io" rel="nofollow noopener">Envoy<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.
-This makes it extremely efficient and capable of handling upstream connections to LLMs.
-If your application is originating code to an API-based LLM, simply use the OpenAI client and configure it with Arch.
-By sending traffic through Arch, you can propagate traces, manage and monitor traffic, apply rate limits, and utilize a large set of traffic management capabilities in a centralized way.</p>
-<div class="admonition attention">
-<p class="admonition-title">Attention</p>
-<p>When you start Arch, it automatically creates a listener port for egress calls to upstream LLMs. This is based on the
-<code class="docutils literal notranslate"><span class="pre">llm_providers</span></code> configuration section in the <code class="docutils literal notranslate"><span class="pre">arch_config.yml</span></code> file. Arch binds itself to a local address such as
-<code class="docutils literal notranslate"><span class="pre">127.0.0.1:12000</span></code>.</p>
-</div>
-<section id="example-using-openai-client-with-arch-as-an-egress-gateway">
-<h3>Example: Using OpenAI Client with Arch as an Egress Gateway<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-using-openai-client-with-arch-as-an-egress-gateway" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-using-openai-client-with-arch-as-an-egress-gateway'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="kn">import</span><span class="w"> </span><span class="nn">openai</span>
-</span><span id="line-2">
-</span><span id="line-3"><span class="c1"># Set the OpenAI API base URL to the Arch gateway endpoint</span>
-</span><span id="line-4"><span class="n">openai</span><span class="o">.</span><span class="n">api_base</span> <span class="o">=</span> <span class="s2">"http://127.0.0.1:12000"</span>
-</span><span id="line-5">
-</span><span id="line-6"><span class="c1"># No need to set openai.api_key since it's configured in Arch's gateway</span>
-</span><span id="line-7">
-</span><span id="line-8"><span class="c1"># Use the OpenAI client as usual</span>
-</span><span id="line-9"><span class="n">response</span> <span class="o">=</span> <span class="n">openai</span><span class="o">.</span><span class="n">Completion</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
-</span><span id="line-10">   <span class="n">model</span><span class="o">=</span><span class="s2">"text-davinci-003"</span><span class="p">,</span>
-</span><span id="line-11">   <span class="n">prompt</span><span class="o">=</span><span class="s2">"What is the capital of France?"</span>
-</span><span id="line-12"><span class="p">)</span>
-</span><span id="line-13">
-</span><span id="line-14"><span class="nb">print</span><span class="p">(</span><span class="s2">"OpenAI Response:"</span><span class="p">,</span> <span class="n">response</span><span class="o">.</span><span class="n">choices</span><span class="p">[</span><span class="mi">0</span><span class="p">]</span><span class="o">.</span><span class="n">text</span><span class="o">.</span><span class="n">strip</span><span class="p">())</span>
-</span></code></pre></div>
-</div>
-<p>In these examples, the OpenAI client is used to send traffic directly through the Arch egress proxy to the LLM of your choice, such as OpenAI.
-The OpenAI client is configured to route traffic via Arch by setting the proxy to <code class="docutils literal notranslate"><span class="pre">127.0.0.1:12000</span></code>, assuming Arch is running locally and bound to that address and port.
-This setup allows you to take advantage of Arch’s advanced traffic management features while interacting with LLM APIs like OpenAI.</p>
-</section>
-</section>
-</section>
-</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
-<div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="listener.html">
-<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="15 18 9 12 15 6"></polyline>
-</svg>
-        Listener
-      </a>
-</div>
-<div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="model_serving.html">
-        Model Serving
-        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="9 18 15 12 9 6"></polyline>
-</svg>
-</a>
-</div>
-</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
-<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
-<ul>
-<li><a :data-current="activeSection === '#messages'" class="reference internal" href="#messages">Messages</a></li>
-<li><a :data-current="activeSection === '#prompt-guard'" class="reference internal" href="#prompt-guard">Prompt Guard</a></li>
-<li><a :data-current="activeSection === '#prompt-targets'" class="reference internal" href="#prompt-targets">Prompt Targets</a><ul>
-<li><a :data-current="activeSection === '#intent-matching'" class="reference internal" href="#intent-matching">Intent Matching</a></li>
-<li><a :data-current="activeSection === '#agentic-apps-via-prompt-targets'" class="reference internal" href="#agentic-apps-via-prompt-targets">Agentic Apps via Prompt Targets</a></li>
-</ul>
-</li>
-<li><a :data-current="activeSection === '#prompting-llms'" class="reference internal" href="#prompting-llms">Prompting LLMs</a><ul>
-<li><a :data-current="activeSection === '#example-using-openai-client-with-arch-as-an-egress-gateway'" class="reference internal" href="#example-using-openai-client-with-arch-as-an-egress-gateway">Example: Using OpenAI Client with Arch as an Egress Gateway</a></li>
-</ul>
-</li>
-</ul>
-</div>
-</aside>
-</main>
-</div>
-</div><footer class="py-6 border-t border-border md:py-0">
-<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
-<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
-</div>
-</div>
-</footer>
-</div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
-<script src="../../_static/doctools.js?v=9bcbadda"></script>
-<script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
-<script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
-<script src="../../_static/design-tabs.js?v=f930bc37"></script>
-</body>
-</html>
\ No newline at end of file
diff --git a/concepts/tech_overview/terminology.html b/concepts/tech_overview/terminology.html
deleted file mode 100755
index bd95fdaf..00000000
--- a/concepts/tech_overview/terminology.html
+++ /dev/null
@@ -1,237 +0,0 @@
-<!DOCTYPE html>
-
-<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
-<head>
-<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
-<meta charset="utf-8"/>
-<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
-<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
-<meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Terminology | Arch Docs v0.3.22</title>
-<meta content="Terminology | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Terminology | Arch Docs v0.3.22" name="twitter:title"/>
-<link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
-<link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
-<link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
-<link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/concepts/tech_overview/terminology.html" rel="canonical"/>
-<link href="../../_static/favicon.ico" rel="icon"/>
-<link href="../../search.html" rel="search" title="Search"/>
-<link href="threading_model.html" rel="next" title="Threading Model"/>
-<link href="tech_overview.html" rel="prev" title="Tech Overview"/>
-<script>
-    <!-- Prevent Flash of wrong theme -->
-      const userPreference = localStorage.getItem('darkMode');
-      let mode;
-      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
-        mode = 'dark';
-        document.documentElement.classList.add('dark');
-      } else {
-        mode = 'light';
-      }
-      if (!userPreference) {localStorage.setItem('darkMode', mode)}
-    </script>
-</head>
-<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
-<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
-      Skip to content
-    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
-<div class="hidden mr-4 md:flex">
-<a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
-<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
-</svg>
-<span class="sr-only">Toggle navigation menu</span>
-</button>
-<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
-<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../../search.html" class="relative flex items-center group" id="searchbox" method="get">
-<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
-<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
-<span class="text-xs">⌘</span>
-    K
-  </kbd>
-</form>
-</div>
-<nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
-<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
-<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
-</div>
-</a>
-<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
-<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
-</svg>
-<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
-</svg>
-</button>
-</nav>
-</div>
-</div>
-</header>
-<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
-<a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a>
-<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
-<div class="overflow-y-auto h-full w-full relative pr-6">
-
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
-<script>
-  window.dataLayer = window.dataLayer || [];
-  function gtag(){dataLayer.push(arguments);}
-  gtag('js', new Date());
-
-  gtag('config', 'G-K2LXXSX6HB');
-</script>
-<nav class="table w-full min-w-full my-6 lg:my-8">
-<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul class="current">
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
-<li class="toctree-l2 current"><a class="current reference internal" href="#">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../prompt_target.html">Prompt Target</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
-</ul>
-</li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
-</ul>
-</nav>
-</div>
-</div>
-<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
-<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
-</svg>
-</button>
-</aside>
-<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
-<div class="w-full min-w-0 mx-auto">
-<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
-<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
-<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
-<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
-</svg>
-</a>
-<div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="tech_overview.html">Tech Overview</a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Terminology</span>
-</nav>
-<div id="content" role="main">
-<section id="terminology">
-<span id="arch-terminology"></span><h1>Terminology<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#terminology"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>A few definitions before we dive into the main architecture documentation. Also note, Arch borrows from Envoy’s terminology
-to keep things consistent in logs and traces, and introduces and clarifies concepts are is relates to LLM applications.</p>
-<p><strong>Agent</strong>: An application that uses LLMs to handle wide-ranging tasks from users via prompts. This could be as simple
-as retrieving or summarizing data from an API, or being able to trigger complex actions like adjusting ad campaigns, or
-changing travel plans via prompts.</p>
-<p><strong>Arch Config</strong>: Arch operates based on a configuration that controls the behavior of a single instance of the Arch gateway.
-This where you enable capabilities like LLM routing, fast function calling (via prompt_targets), applying guardrails, and enabling critical
-features like metrics and tracing. For the full configuration reference of <cite>arch_config.yaml</cite> see <a class="reference internal" href="../../resources/configuration_reference.html#configuration-reference"><span class="std std-ref">here</span></a>.</p>
-<p><strong>Downstream(Ingress)</strong>: An downstream client (web application, etc.) connects to Arch, sends prompts, and receives responses.</p>
-<p><strong>Upstream(Egress)</strong>: An upstream host that receives connections and prompts from Arch, and returns context or responses for a prompt</p>
-<a class="reference internal image-reference" href="../../_images/network-topology-ingress-egress.jpg"><img alt="../../_images/network-topology-ingress-egress.jpg" class="align-center" src="../../_images/network-topology-ingress-egress.jpg" style="width: 100%;"/>
-</a>
-<p><strong>Listener</strong>: A <a class="reference internal" href="listener.html#arch-overview-listeners"><span class="std std-ref">listener</span></a> is a named network location (e.g., port, address, path etc.) that Arch
-listens on to process prompts before forwarding them to your application server endpoints. rch enables you to configure one listener
-for downstream connections (like port 80, 443) and creates a separate internal listener for calls that initiate from your application
-code to LLMs.</p>
-<div class="admonition note">
-<p class="admonition-title">Note</p>
-<p>When you start Arch, you specify a listener address/port that you want to bind downstream. But, Arch uses are predefined port
-that you can use (<code class="docutils literal notranslate"><span class="pre">127.0.0.1:12000</span></code>) to proxy egress calls originating from your application to LLMs (API-based or hosted).
-For more details, check out <a class="reference internal" href="../llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">LLM providers</span></a>.</p>
-</div>
-<p><strong>Prompt Target</strong>: Arch offers a primitive called <a class="reference internal" href="../prompt_target.html#prompt-target"><span class="std std-ref">prompt target</span></a> to help separate business logic from
-undifferentiated work in building generative AI apps. Prompt targets are endpoints that receive prompts that are processed by Arch.
-For example, Arch enriches incoming prompts with metadata like knowing when a request is a follow-up or clarifying prompt so that you
-can build faster, more accurate retrieval (RAG) apps. To support agentic apps, like scheduling travel plans or sharing comments on a
-document - via prompts, Arch uses its function calling abilities to extract critical information from the incoming prompt (or a set of
-prompts) needed by a downstream backend API or function call before calling it directly.</p>
-<p><strong>Model Serving</strong>: Arch is a set of <cite>two</cite> self-contained processes that are designed to run alongside your application servers
-(or on a separate host connected via a network).The <a class="reference internal" href="model_serving.html#model-serving"><span class="std std-ref">model serving</span></a> process helps Arch make intelligent decisions
-about the incoming prompts. The model server is designed to call the (fast) purpose-built LLMs in Arch.</p>
-<p><strong>Error Target</strong>: <a class="reference internal" href="error_target.html#error-target"><span class="std std-ref">Error targets</span></a> are those endpoints that receive forwarded errors from Arch when issues arise,
-such as failing to properly call a function/API, detecting violations of guardrails, or encountering other processing errors.
-These errors are communicated to the application via headers <code class="docutils literal notranslate"><span class="pre">X-Arch-[ERROR-TYPE]</span></code>, allowing it to handle the errors gracefully
-and take appropriate actions.</p>
-</section>
-</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
-<div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="tech_overview.html">
-<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="15 18 9 12 15 6"></polyline>
-</svg>
-        Tech Overview
-      </a>
-</div>
-<div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="threading_model.html">
-        Threading Model
-        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="9 18 15 12 9 6"></polyline>
-</svg>
-</a>
-</div>
-</div></div>
-</main>
-</div>
-</div><footer class="py-6 border-t border-border md:py-0">
-<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
-<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
-</div>
-</div>
-</footer>
-</div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
-<script src="../../_static/doctools.js?v=9bcbadda"></script>
-<script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
-<script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
-<script src="../../_static/design-tabs.js?v=f930bc37"></script>
-</body>
-</html>
\ No newline at end of file
diff --git a/get_started/intro_to_arch.html b/get_started/intro_to_plano.html
similarity index 53%
rename from get_started/intro_to_arch.html
rename to get_started/intro_to_plano.html
index f061c329..fe374caa 100755
--- a/get_started/intro_to_arch.html
+++ b/get_started/intro_to_plano.html
@@ -7,15 +7,15 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Intro to Arch | Arch Docs v0.3.22</title>
-<meta content="Intro to Arch | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Intro to Arch | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Intro to Plano | Plano Docs v0.4</title>
+<meta content="Intro to Plano | Plano Docs v0.4" property="og:title"/>
+<meta content="Intro to Plano | Plano Docs v0.4" name="twitter:title"/>
 <link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/get_started/intro_to_arch.html" rel="canonical"/>
+<link href="./docs/get_started/intro_to_plano.html" rel="canonical"/>
 <link href="../_static/favicon.ico" rel="icon"/>
 <link href="../search.html" rel="search" title="Search"/>
 <link href="quickstart.html" rel="next" title="Quickstart"/>
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul class="current">
 <li class="toctree-l1"><a class="reference internal" href="overview.html">Overview</a></li>
-<li class="toctree-l1 current"><a class="current reference internal" href="#">Intro to Arch</a></li>
+<li class="toctree-l1 current"><a class="current reference internal" href="#">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,74 +148,56 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
 </a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Intro to Arch</span>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Intro to Plano</span>
 </nav>
 <div id="content" role="main">
-<section id="intro-to-arch">
-<span id="id1"></span><h1>Intro to Arch<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#intro-to-arch"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>AI demos are easy to build. But past the thrill of a quick hack, you are left building, maintaining and scaling low-level plumbing code for agents that slows down AI innovation.
-For example:</p>
+<section id="intro-to-plano">
+<span id="id1"></span><h1>Intro to Plano<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#intro-to-plano"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p>Building agentic demos is easy. Delivering agentic applications safely, reliably, and repeatably to production is hard. After a quick hack, you end up building the “hidden AI middleware” to reach production: routing logic to reach the right agent, guardrail hooks for safety and moderation, evaluation and observability glue for continuous learning, and model/provider quirks — scattered across frameworks and application code.</p>
+<p>Plano solves this by moving core delivery concerns into a unified, out-of-process dataplane. Core capabilities:</p>
 <ul class="simple">
-<li><p>You want to build specialized agents, but get stuck writing <strong>routing and handoff</strong> code.</p></li>
-<li><p>You bogged down with prompt engineering work to <strong>clarify user intent and validate inputs</strong>.</p></li>
-<li><p>You want to <strong>quickly and safely use new LLMs</strong> but get stuck writing integration code.</p></li>
-<li><p>You waste cycles writing and maintaining <strong>observability</strong> code, when it can be transparent.</p></li>
-<li><p>You want to <strong>apply guardrails</strong>, but have to write custom code for each prompt and LLM.</p></li>
+<li><p><strong>🚦 Orchestration:</strong> Low-latency orchestration between agents, and add new agents without changing app code. When routing lives inside app code, it becomes hard to evolve and easy to duplicate. Moving orchestration into a centrally managed dataplane lets you change strategies without touching your agents, improving performance and reducing maintenance burden while avoiding tight coupling.</p></li>
+<li><p><strong>🛡️ Guardrails &amp; Memory Hooks:</strong> Apply jailbreak protection, content policies, and context workflows (e.g., rewriting, retrieval, redaction) once via <a class="reference internal" href="../concepts/filter_chain.html#filter-chain"><span class="std std-ref">Filter Chains</span></a> at the dataplane. Instead of re-implementing these in every agentic service, you get centralized governance, reduced code duplication, and consistent behavior across your stack.</p></li>
+<li><p><strong>🔗 Model Agility:</strong> Route by model, alias (semantic names), or automatically via preferences so agents stay decoupled from specific providers. Swap or add models without refactoring prompts, tool-calling, or streaming handlers throughout your codebase by using Plano’s smart routing and unified API.</p></li>
+<li><p><strong>🕵 Agentic Signals™:</strong> Zero-code capture of behavior signals, traces, and metrics consistently across every agent. Rather than stitching together logging and metrics per framework, Plano surfaces traces, token usage, and learning signals in one place so you can iterate safely.</p></li>
 </ul>
-<p>Arch is designed to solve these problems by providing a unified, out-of-process architecture that integrates with your existing application stack, enabling you to focus on building high-level features rather than plumbing — all without locking you into a framework.</p>
+<p>Built by core contributors to the widely adopted Envoy Proxy &lt;<a class="reference external" href="https://www.envoyproxy.io/" rel="nofollow noopener">https://www.envoyproxy.io/<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>&gt;_, Plano gives you a production‑grade foundation for agentic applications. It helps <strong>developers</strong> stay focused on the core logic of their agents, helps <strong>product teams</strong> shorten feedback loops for learning, and helps <strong>engineering teams</strong>  standardize policy and safety across agents and LLMs. Plano is grounded in open protocols (de facto: OpenAI‑style v1/responses, de jure: MCP) and proven patterns like sidecar deployments, so it plugs in cleanly while remaining robust, scalable, and flexible.</p>
+<p>In practice, achieving the above goal is incredibly difficult. Plano attempts to do so by providing the following high level features:</p>
 <figure class="align-center" id="id2">
-<a class="reference internal image-reference" href="../_images/arch_network_diagram_high_level.png"><img alt="../_images/arch_network_diagram_high_level.png" src="../_images/arch_network_diagram_high_level.png" style="width: 100%;"/>
+<a class="reference internal image-reference" href="../_images/plano_network_diagram_high_level.png"><img alt="../_images/plano_network_diagram_high_level.png" src="../_images/plano_network_diagram_high_level.png" style="width: 100%;"/>
 </a>
 <figcaption>
-<p><span class="caption-text">High-level network flow of where Arch Gateway sits in your agentic stack. Designed for both ingress and egress prompt traffic.</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></p>
+<p><span class="caption-text">High-level network flow of where Plano sits in your agentic stack. Designed for both ingress and egress prompt traffic.</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></p>
 </figcaption>
 </figure>
-<p><a class="reference external" href="https://github.com/katanemo/arch" rel="nofollow noopener">Arch<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> is a smart edge and AI gateway for AI-native apps - built by the contributors of Envoy Proxy with the belief that:</p>
-<blockquote>
-<div><p><em>Prompts are nuanced and opaque user requests, which require the same capabilities as traditional HTTP requests
-including secure handling, intelligent routing, robust observability, and integration with backend (API)
-systems for personalization - all outside business logic.</em></p>
-</div></blockquote>
-<p>In practice, achieving the above goal is incredibly difficult. Arch attempts to do so by providing the following high level features:</p>
-<p><strong>Out-of-process architecture, built on</strong> <a class="reference external" href="http://envoyproxy.io/" rel="nofollow noopener">Envoy<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>:
-Arch takes a dependency on Envoy and is a self-contained process that is designed to run alongside your application servers.
-Arch uses Envoy’s HTTP connection management subsystem, HTTP L7 filtering and telemetry capabilities to extend the functionality exclusively for prompts and LLMs.
-This gives Arch several advantages:</p>
-<ul class="simple">
-<li><p>Arch builds on Envoy’s proven success. Envoy is used at massive scale by the leading technology companies of our time including <a class="reference external" href="https://www.airbnb.com" rel="nofollow noopener">AirBnB<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, <a class="reference external" href="https://www.dropbox.com" rel="nofollow noopener">Dropbox<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, <a class="reference external" href="https://www.google.com" rel="nofollow noopener">Google<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, <a class="reference external" href="https://www.reddit.com" rel="nofollow noopener">Reddit<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, <a class="reference external" href="https://www.stripe.com" rel="nofollow noopener">Stripe<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, etc. Its battle tested and scales linearly with usage and enables developers to focus on what really matters: application features and business logic.</p></li>
-<li><p>Arch works with any application language. A single Arch deployment can act as gateway for AI applications written in Python, Java, C++, Go, Php, etc.</p></li>
-<li><p>Arch can be deployed and upgraded quickly across your infrastructure transparently without the horrid pain of deploying library upgrades in your applications.</p></li>
-</ul>
-<p><strong>Engineered with Fast Task-Specific LLMs (TLMs):</strong> Arch is engineered with specialized LLMs that are designed for the fast, cost-effective and accurate handling of prompts.
+<p><strong>Engineered with Task-Specific LLMs (TLMs):</strong> Plano is engineered with specialized LLMs that are designed for fast, cost-effective and accurate handling of prompts.
 These LLMs are designed to be best-in-class for critical tasks like:</p>
 <ul class="simple">
-<li><p><strong>Function Calling:</strong> Arch helps you easily personalize your applications by enabling calls to application-specific (API) operations via user prompts.
-This involves any predefined functions or APIs you want to expose to users to perform tasks, gather information, or manipulate data.
-With function calling, you have flexibility to support “agentic” experiences tailored to specific use cases - from updating insurance claims to creating ad campaigns - via prompts.
-Arch analyzes prompts, extracts critical information from prompts, engages in lightweight conversation to gather any missing parameters and makes API calls so that you can focus on writing business logic.
-For more details, read <a class="reference internal" href="../guides/function_calling.html#function-calling"><span class="std std-ref">Function Calling</span></a>.</p></li>
-<li><p><strong>Prompt Guard:</strong> Arch helps you improve the safety of your application by applying prompt guardrails in a centralized way for better governance hygiene.
+<li><p><strong>Agent Orchestration:</strong> <a class="reference external" href="https://huggingface.co/collections/katanemo/plano-orchestrator" rel="nofollow noopener">Plano-Orchestrator<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> is a family of state-of-the-art routing and orchestration models that decide which agent(s) or LLM(s) should handle each request, and in what sequence. Built for real-world multi-agent deployments, it analyzes user intent and conversation context to make precise routing and orchestration decisions while remaining efficient enough for low-latency production use across general chat, coding, and long-context multi-turn conversations.</p></li>
+<li><p><strong>Function Calling:</strong> Plano lets you expose application-specific (API) operations as tools so that your agents can update records, fetch data, or trigger determininistic workflows via prompts. Under the hood this is backed by Arch-Function-Chat; for more details, read <a class="reference internal" href="../guides/function_calling.html#function-calling"><span class="std std-ref">Function Calling</span></a>.</p></li>
+<li><p><strong>Guardrails:</strong> Plano helps you improve the safety of your application by applying prompt guardrails in a centralized way for better governance hygiene.
 With prompt guardrails you can prevent <code class="docutils literal notranslate"><span class="pre">jailbreak</span> <span class="pre">attempts</span></code> present in user’s prompts without having to write a single line of code.
-To learn more about how to configure guardrails available in Arch, read <a class="reference internal" href="../guides/prompt_guard.html#prompt-guard"><span class="std std-ref">Prompt Guard</span></a>.</p></li>
+To learn more about how to configure guardrails available in Plano, read <a class="reference internal" href="../guides/prompt_guard.html#prompt-guard"><span class="std std-ref">Prompt Guard</span></a>.</p></li>
+</ul>
+<p><strong>Model Proxy:</strong> Plano offers several capabilities for LLM calls originating from your applications, including smart retries on errors from upstream LLMs and automatic cut-over to other LLMs configured in Plano for continuous availability and disaster recovery scenarios. From your application’s perspective you keep using an OpenAI-compatible API, while Plano owns resiliency and failover policies in one place.
+Plano extends Envoy’s <a class="reference external" href="https://www.envoyproxy.io/docs/envoy/latest/intro/arch_overview/upstream/cluster_manager" rel="nofollow noopener">cluster subsystem<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> to manage upstream connections to LLMs so that you can build resilient, provider-agnostic AI applications.</p>
+<p><strong>Edge Proxy:</strong> There is substantial benefit in using the same software at the edge (observability, traffic shaping algorithms, applying guardrails, etc.) as for outbound LLM inference use cases. Plano has the feature set that makes it exceptionally well suited as an edge gateway for AI applications.
+This includes TLS termination, applying guardrails early in the request flow, and intelligently deciding which agent(s) or LLM(s) should handle each request and in what sequence. In practice, you configure listeners and policies once, and every inbound and outbound call flows through the same hardened gateway.</p>
+<p><strong>Zero-Code Agent Signals™ &amp; Tracing:</strong> Zero-code capture of behavior signals, traces, and metrics consistently across every agent. Plano propagates trace context using the W3C Trace Context standard, specifically through the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header. This allows each component in the system to record its part of the request flow, enabling end-to-end tracing across the entire application. By using OpenTelemetry, Plano ensures that developers can capture this trace data consistently and in a format compatible with various observability tools.</p>
+<p><strong>Best-In Class Monitoring:</strong> Plano offers several monitoring metrics that help you understand three critical aspects of your application: latency, token usage, and error rates by an upstream LLM provider. Latency measures the speed at which your application is responding to users, which includes metrics like time to first token (TFT), time per output token (TOT) metrics, and the total latency as perceived by users.</p>
+<p><strong>Out-of-process architecture, built on</strong> <a class="reference external" href="http://envoyproxy.io/" rel="nofollow noopener">Envoy<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>:
+Plano takes a dependency on Envoy and is a self-contained process that is designed to run alongside your application servers. Plano uses Envoy’s HTTP connection management subsystem, HTTP L7 filtering and telemetry capabilities to extend the functionality exclusively for prompts and LLMs.
+This gives Plano several advantages:</p>
+<ul class="simple">
+<li><p>Plano builds on Envoy’s proven success. Envoy is used at massive scale by the leading technology companies of our time including <a class="reference external" href="https://www.airbnb.com" rel="nofollow noopener">AirBnB<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, <a class="reference external" href="https://www.dropbox.com" rel="nofollow noopener">Dropbox<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, <a class="reference external" href="https://www.google.com" rel="nofollow noopener">Google<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, <a class="reference external" href="https://www.reddit.com" rel="nofollow noopener">Reddit<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, <a class="reference external" href="https://www.stripe.com" rel="nofollow noopener">Stripe<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, etc. Its battle tested and scales linearly with usage and enables developers to focus on what really matters: application features and business logic.</p></li>
+<li><p>Plano works with any application language. A single Plano deployment can act as gateway for AI applications written in Python, Java, C++, Go, Php, etc.</p></li>
+<li><p>Plano can be deployed and upgraded quickly across your infrastructure transparently without the horrid pain of deploying library upgrades in your applications.</p></li>
 </ul>
-<p><strong>Traffic Management:</strong> Arch offers several capabilities for LLM calls originating from your applications, including smart retries on errors from upstream LLMs, and automatic cut-over to other LLMs configured in Arch for continuous availability and disaster recovery scenarios.
-Arch extends Envoy’s <a class="reference external" href="https://www.envoyproxy.io/docs/envoy/latest/intro/arch_overview/upstream/cluster_manager" rel="nofollow noopener">cluster subsystem<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> to manage upstream connections to LLMs so that you can build resilient AI applications.</p>
-<p><strong>Front/edge Gateway:</strong> There is substantial benefit in using the same software at the edge (observability, traffic shaping algorithms, applying guardrails, etc.) as for outbound LLM inference use cases.
-Arch has the feature set that makes it exceptionally well suited as an edge gateway for AI applications.
-This includes TLS termination, applying guardrail early in the process, intelligent parameter gathering from prompts, and prompt-based routing to backend APIs.</p>
-<p><strong>Best-In Class Monitoring:</strong> Arch offers several monitoring metrics that help you understand three critical aspects of
-your application: latency, token usage, and error rates by an upstream LLM provider. Latency measures the speed at which
-your application is responding to users, which includes metrics like time to first token (TFT), time per output token (TOT)
-metrics, and the total latency as perceived by users.</p>
-<p><strong>End-to-End Tracing:</strong> Arch propagates trace context using the W3C Trace Context standard, specifically through the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header.
-This allows each component in the system to record its part of the request flow, enabling end-to-end tracing across the entire application.
-By using OpenTelemetry, Arch ensures that developers can capture this trace data consistently and in a format compatible with various observability tools.
-For more details, read <a class="reference internal" href="../guides/observability/tracing.html#arch-overview-tracing"><span class="std std-ref">Tracing</span></a>.</p>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
@@ -245,12 +222,12 @@ For more details, read <a class="reference internal" href="../guides/observabili
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../_static/doctools.js?v=9bcbadda"></script>
 <script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
diff --git a/get_started/overview.html b/get_started/overview.html
index 56ee6e12..1a9c08e4 100755
--- a/get_started/overview.html
+++ b/get_started/overview.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Overview | Arch Docs v0.3.22</title>
-<meta content="Overview | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Overview | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Overview | Plano Docs v0.4</title>
+<meta content="Overview | Plano Docs v0.4" property="og:title"/>
+<meta content="Overview | Plano Docs v0.4" name="twitter:title"/>
 <link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/get_started/overview.html" rel="canonical"/>
 <link href="../_static/favicon.ico" rel="icon"/>
 <link href="../search.html" rel="search" title="Search"/>
-<link href="intro_to_arch.html" rel="next" title="Intro to Arch"/>
-<link href="../index.html" rel="prev" title="Welcome to Arch!"/>
+<link href="intro_to_plano.html" rel="next" title="Intro to Plano"/>
+<link href="../index.html" rel="prev" title="Welcome to Plano!"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul class="current">
 <li class="toctree-l1 current"><a class="current reference internal" href="#">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -163,20 +158,20 @@
 <div id="content" role="main">
 <section id="overview">
 <span id="id1"></span><h1>Overview<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#overview"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p><a class="reference external" href="https://github.com/katanemo/arch" rel="nofollow noopener">Arch<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> is a smart edge and AI gateway for AI agents - one that is natively designed to handle and process prompts, not just network traffic.</p>
-<p>Built by contributors to the widely adopted <a class="reference external" href="https://www.envoyproxy.io/" rel="nofollow noopener">Envoy Proxy<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, Arch handles the <em>pesky low-level work</em> in building agentic apps — like applying guardrails, clarifying vague user input, routing prompts to the right agent, and unifying access to any LLM. It’s a protocol-friendly and framework-agnostic infrastructure layer designed to help you build and ship agentic apps faster.</p>
-<p>In this documentation, you will learn how to quickly set up Arch to trigger API calls via prompts, apply prompt guardrails without writing any application-level logic,
-simplify the interaction with upstream LLMs, and improve observability all while simplifying your application development process.</p>
+<p><a class="reference external" href="https://github.com/katanemo/plano" rel="nofollow noopener">Plano<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> is delivery infrastructure for agentic apps. A models-native proxy server and data plane designed to help you build agents faster, and deliver them reliably to production.</p>
+<p>Plano pulls out the rote plumbing work (the “hidden AI middleware”) and decouples you from brittle, ever‑changing framework abstractions. It centralizes what shouldn’t be bespoke in every codebase like agent routing and orchestration, rich agentic signals and traces for continuous improvement, guardrail filters for safety and moderation, and smart LLM routing APIs for UX and DX agility. Use any language or AI framework, and ship agents to production faster with Plano.</p>
+<p>Built by core contributors to the widely adopted <a class="reference external" href="https://www.envoyproxy.io/" rel="nofollow noopener">Envoy Proxy<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, Plano gives you a production‑grade foundation for agentic applications. It helps <strong>developers</strong> stay focused on the core logic of their agents, helps <strong>product teams</strong> shorten feedback loops for learning, and helps <strong>engineering teams</strong>  standardize policy and safety across agents and LLMs. Plano is grounded in open protocols (de facto: OpenAI‑style v1/responses, de jure: MCP) and proven patterns like sidecar deployments, so it plugs in cleanly while remaining robust, scalable, and flexible.</p>
+<p>In this documentation, you’ll learn how to set up Plano quickly, trigger API calls via prompts, apply guardrails without tight coupling with application code, simplify model and provider integration, and improve observability — so that you can focus on what matters most: the core product logic of your agents.</p>
 <figure class="align-center" id="id2">
-<a class="reference internal image-reference" href="../_images/arch_network_diagram_high_level.png"><img alt="../_images/arch_network_diagram_high_level.png" src="../_images/arch_network_diagram_high_level.png" style="width: 100%;"/>
+<a class="reference internal image-reference" href="../_images/plano_network_diagram_high_level.png"><img alt="../_images/plano_network_diagram_high_level.png" src="../_images/plano_network_diagram_high_level.png" style="width: 100%;"/>
 </a>
 <figcaption>
-<p><span class="caption-text">High-level network flow of where Arch Gateway sits in your agentic stack. Designed for both ingress and egress prompt traffic.</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></p>
+<p><span class="caption-text">High-level network flow of where Plano sits in your agentic stack. Designed for both ingress and egress traffic.</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></p>
 </figcaption>
 </figure>
 <section id="get-started">
 <h2>Get Started<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#get-started" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#get-started'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>This section introduces you to Arch and helps you get set up quickly:</p>
+<p>This section introduces you to Plano and helps you get set up quickly:</p>
 <div class="sd-container-fluid sd-sphinx-override sd-mb-4 docutils">
 <div class="sd-row sd-row-cols-3 sd-row-cols-xs-3 sd-row-cols-sm-3 sd-row-cols-md-3 sd-row-cols-lg-3 docutils">
 <div class="sd-col sd-d-flex-row docutils">
@@ -184,7 +179,7 @@ simplify the interaction with upstream LLMs, and improve observability all while
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
 <svg aria-hidden="true" class="sd-octicon sd-octicon-apps" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M1.5 3.25c0-.966.784-1.75 1.75-1.75h2.5c.966 0 1.75.784 1.75 1.75v2.5A1.75 1.75 0 0 1 5.75 7.5h-2.5A1.75 1.75 0 0 1 1.5 5.75Zm7 0c0-.966.784-1.75 1.75-1.75h2.5c.966 0 1.75.784 1.75 1.75v2.5a1.75 1.75 0 0 1-1.75 1.75h-2.5A1.75 1.75 0 0 1 8.5 5.75Zm-7 7c0-.966.784-1.75 1.75-1.75h2.5c.966 0 1.75.784 1.75 1.75v2.5a1.75 1.75 0 0 1-1.75 1.75h-2.5a1.75 1.75 0 0 1-1.75-1.75Zm7 0c0-.966.784-1.75 1.75-1.75h2.5c.966 0 1.75.784 1.75 1.75v2.5a1.75 1.75 0 0 1-1.75 1.75h-2.5a1.75 1.75 0 0 1-1.75-1.75ZM3.25 3a.25.25 0 0 0-.25.25v2.5c0 .138.112.25.25.25h2.5A.25.25 0 0 0 6 5.75v-2.5A.25.25 0 0 0 5.75 3Zm7 0a.25.25 0 0 0-.25.25v2.5c0 .138.112.25.25.25h2.5a.25.25 0 0 0 .25-.25v-2.5a.25.25 0 0 0-.25-.25Zm-7 7a.25.25 0 0 0-.25.25v2.5c0 .138.112.25.25.25h2.5a.25.25 0 0 0 .25-.25v-2.5a.25.25 0 0 0-.25-.25Zm7 0a.25.25 0 0 0-.25.25v2.5c0 .138.112.25.25.25h2.5a.25.25 0 0 0 .25-.25v-2.5a.25.25 0 0 0-.25-.25Z"></path></svg> Overview</div>
-<p class="sd-card-text">Overview of Arch and Doc navigation</p>
+<p class="sd-card-text">Overview of Plano and Doc navigation</p>
 </div>
 <a class="sd-stretched-link sd-hide-link-text reference external" href="overview.html"><span>overview.html</span></a></div>
 </div>
@@ -192,10 +187,10 @@ simplify the interaction with upstream LLMs, and improve observability all while
 <div class="sd-card sd-sphinx-override sd-w-100 sd-shadow-sm sd-card-hover docutils">
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
-<svg aria-hidden="true" class="sd-octicon sd-octicon-book" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M0 1.75A.75.75 0 0 1 .75 1h4.253c1.227 0 2.317.59 3 1.501A3.743 3.743 0 0 1 11.006 1h4.245a.75.75 0 0 1 .75.75v10.5a.75.75 0 0 1-.75.75h-4.507a2.25 2.25 0 0 0-1.591.659l-.622.621a.75.75 0 0 1-1.06 0l-.622-.621A2.25 2.25 0 0 0 5.258 13H.75a.75.75 0 0 1-.75-.75Zm7.251 10.324.004-5.073-.002-2.253A2.25 2.25 0 0 0 5.003 2.5H1.5v9h3.757a3.75 3.75 0 0 1 1.994.574ZM8.755 4.75l-.004 7.322a3.752 3.752 0 0 1 1.992-.572H14.5v-9h-3.495a2.25 2.25 0 0 0-2.25 2.25Z"></path></svg> Intro to Arch</div>
-<p class="sd-card-text">Explore Arch’s features and developer workflow</p>
+<svg aria-hidden="true" class="sd-octicon sd-octicon-book" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M0 1.75A.75.75 0 0 1 .75 1h4.253c1.227 0 2.317.59 3 1.501A3.743 3.743 0 0 1 11.006 1h4.245a.75.75 0 0 1 .75.75v10.5a.75.75 0 0 1-.75.75h-4.507a2.25 2.25 0 0 0-1.591.659l-.622.621a.75.75 0 0 1-1.06 0l-.622-.621A2.25 2.25 0 0 0 5.258 13H.75a.75.75 0 0 1-.75-.75Zm7.251 10.324.004-5.073-.002-2.253A2.25 2.25 0 0 0 5.003 2.5H1.5v9h3.757a3.75 3.75 0 0 1 1.994.574ZM8.755 4.75l-.004 7.322a3.752 3.752 0 0 1 1.992-.572H14.5v-9h-3.495a2.25 2.25 0 0 0-2.25 2.25Z"></path></svg> Intro to Plano</div>
+<p class="sd-card-text">Explore Plano’s features and developer workflow</p>
 </div>
-<a class="sd-stretched-link sd-hide-link-text reference external" href="intro_to_arch.html"><span>intro_to_arch.html</span></a></div>
+<a class="sd-stretched-link sd-hide-link-text reference external" href="intro_to_plano.html"><span>intro_to_plano.html</span></a></div>
 </div>
 <div class="sd-col sd-d-flex-row docutils">
 <div class="sd-card sd-sphinx-override sd-w-100 sd-shadow-sm sd-card-hover docutils">
@@ -211,24 +206,24 @@ simplify the interaction with upstream LLMs, and improve observability all while
 </section>
 <section id="concepts">
 <h2>Concepts<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#concepts" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#concepts'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Deep dive into essential ideas and mechanisms behind Arch:</p>
+<p>Deep dive into essential ideas and mechanisms behind Plano:</p>
 <div class="sd-container-fluid sd-sphinx-override sd-mb-4 docutils">
 <div class="sd-row sd-row-cols-3 sd-row-cols-xs-3 sd-row-cols-sm-3 sd-row-cols-md-3 sd-row-cols-lg-3 docutils">
 <div class="sd-col sd-d-flex-row docutils">
 <div class="sd-card sd-sphinx-override sd-w-100 sd-shadow-sm sd-card-hover docutils">
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
-<svg aria-hidden="true" class="sd-octicon sd-octicon-package" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="m8.878.392 5.25 3.045c.54.314.872.89.872 1.514v6.098a1.75 1.75 0 0 1-.872 1.514l-5.25 3.045a1.75 1.75 0 0 1-1.756 0l-5.25-3.045A1.75 1.75 0 0 1 1 11.049V4.951c0-.624.332-1.201.872-1.514L7.122.392a1.75 1.75 0 0 1 1.756 0ZM7.875 1.69l-4.63 2.685L8 7.133l4.755-2.758-4.63-2.685a.248.248 0 0 0-.25 0ZM2.5 5.677v5.372c0 .09.047.171.125.216l4.625 2.683V8.432Zm6.25 8.271 4.625-2.683a.25.25 0 0 0 .125-.216V5.677L8.75 8.432Z"></path></svg> Tech Overview</div>
-<p class="sd-card-text">Learn about the technology stack</p>
+<svg aria-hidden="true" class="sd-octicon sd-octicon-package" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="m8.878.392 5.25 3.045c.54.314.872.89.872 1.514v6.098a1.75 1.75 0 0 1-.872 1.514l-5.25 3.045a1.75 1.75 0 0 1-1.756 0l-5.25-3.045A1.75 1.75 0 0 1 1 11.049V4.951c0-.624.332-1.201.872-1.514L7.122.392a1.75 1.75 0 0 1 1.756 0ZM7.875 1.69l-4.63 2.685L8 7.133l4.755-2.758-4.63-2.685a.248.248 0 0 0-.25 0ZM2.5 5.677v5.372c0 .09.047.171.125.216l4.625 2.683V8.432Zm6.25 8.271 4.625-2.683a.25.25 0 0 0 .125-.216V5.677L8.75 8.432Z"></path></svg> Agents</div>
+<p class="sd-card-text">Learn about how to build and scale agents with Plano</p>
 </div>
-<a class="sd-stretched-link sd-hide-link-text reference external" href="../concepts/tech_overview/tech_overview.html"><span>../concepts/tech_overview/tech_overview.html</span></a></div>
+<a class="sd-stretched-link sd-hide-link-text reference external" href="../concepts/agents.html"><span>../concepts/agents.html</span></a></div>
 </div>
 <div class="sd-col sd-d-flex-row docutils">
 <div class="sd-card sd-sphinx-override sd-w-100 sd-shadow-sm sd-card-hover docutils">
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
-<svg aria-hidden="true" class="sd-octicon sd-octicon-webhook" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M5.5 4.25a2.25 2.25 0 0 1 4.5 0 .75.75 0 0 0 1.5 0 3.75 3.75 0 1 0-6.14 2.889l-2.272 4.258a.75.75 0 0 0 1.324.706L7 7.25a.75.75 0 0 0-.309-1.015A2.25 2.25 0 0 1 5.5 4.25Z"></path><path d="M7.364 3.607a.75.75 0 0 1 1.03.257l2.608 4.349a3.75 3.75 0 1 1-.628 6.785.75.75 0 0 1 .752-1.299 2.25 2.25 0 1 0-.033-3.88.75.75 0 0 1-1.03-.256L7.107 4.636a.75.75 0 0 1 .257-1.03Z"></path><path d="M2.9 8.776A.75.75 0 0 1 2.625 9.8 2.25 2.25 0 1 0 6 11.75a.75.75 0 0 1 .75-.751h5.5a.75.75 0 0 1 0 1.5H7.425a3.751 3.751 0 1 1-5.55-3.998.75.75 0 0 1 1.024.274Z"></path></svg> LLM Providers</div>
-<p class="sd-card-text">Explore Arch’s LLM integration options</p>
+<svg aria-hidden="true" class="sd-octicon sd-octicon-webhook" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M5.5 4.25a2.25 2.25 0 0 1 4.5 0 .75.75 0 0 0 1.5 0 3.75 3.75 0 1 0-6.14 2.889l-2.272 4.258a.75.75 0 0 0 1.324.706L7 7.25a.75.75 0 0 0-.309-1.015A2.25 2.25 0 0 1 5.5 4.25Z"></path><path d="M7.364 3.607a.75.75 0 0 1 1.03.257l2.608 4.349a3.75 3.75 0 1 1-.628 6.785.75.75 0 0 1 .752-1.299 2.25 2.25 0 1 0-.033-3.88.75.75 0 0 1-1.03-.256L7.107 4.636a.75.75 0 0 1 .257-1.03Z"></path><path d="M2.9 8.776A.75.75 0 0 1 2.625 9.8 2.25 2.25 0 1 0 6 11.75a.75.75 0 0 1 .75-.751h5.5a.75.75 0 0 1 0 1.5H7.425a3.751 3.751 0 1 1-5.55-3.998.75.75 0 0 1 1.024.274Z"></path></svg> Model Providers</div>
+<p class="sd-card-text">Explore Plano’s LLM integration options</p>
 </div>
 <a class="sd-stretched-link sd-hide-link-text reference external" href="../concepts/llm_providers/llm_providers.html"><span>../concepts/llm_providers/llm_providers.html</span></a></div>
 </div>
@@ -237,7 +232,7 @@ simplify the interaction with upstream LLMs, and improve observability all while
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
 <svg aria-hidden="true" class="sd-octicon sd-octicon-workflow" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M0 1.75C0 .784.784 0 1.75 0h3.5C6.216 0 7 .784 7 1.75v3.5A1.75 1.75 0 0 1 5.25 7H4v4a1 1 0 0 0 1 1h4v-1.25C9 9.784 9.784 9 10.75 9h3.5c.966 0 1.75.784 1.75 1.75v3.5A1.75 1.75 0 0 1 14.25 16h-3.5A1.75 1.75 0 0 1 9 14.25v-.75H5A2.5 2.5 0 0 1 2.5 11V7h-.75A1.75 1.75 0 0 1 0 5.25Zm1.75-.25a.25.25 0 0 0-.25.25v3.5c0 .138.112.25.25.25h3.5a.25.25 0 0 0 .25-.25v-3.5a.25.25 0 0 0-.25-.25Zm9 9a.25.25 0 0 0-.25.25v3.5c0 .138.112.25.25.25h3.5a.25.25 0 0 0 .25-.25v-3.5a.25.25 0 0 0-.25-.25Z"></path></svg> Prompt Target</div>
-<p class="sd-card-text">Understand how Arch handles prompts</p>
+<p class="sd-card-text">Understand how Plano handles prompts</p>
 </div>
 <a class="sd-stretched-link sd-hide-link-text reference external" href="../concepts/prompt_target.html"><span>../concepts/prompt_target.html</span></a></div>
 </div>
@@ -246,14 +241,14 @@ simplify the interaction with upstream LLMs, and improve observability all while
 </section>
 <section id="guides">
 <h2>Guides<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#guides" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#guides'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Step-by-step tutorials for practical Arch use cases and scenarios:</p>
+<p>Step-by-step tutorials for practical Plano use cases and scenarios:</p>
 <div class="sd-container-fluid sd-sphinx-override sd-mb-4 docutils">
 <div class="sd-row sd-row-cols-3 sd-row-cols-xs-3 sd-row-cols-sm-3 sd-row-cols-md-3 sd-row-cols-lg-3 docutils">
 <div class="sd-col sd-d-flex-row docutils">
 <div class="sd-card sd-sphinx-override sd-w-100 sd-shadow-sm sd-card-hover docutils">
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
-<svg aria-hidden="true" class="sd-octicon sd-octicon-shield-check" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="m8.533.133 5.25 1.68A1.75 1.75 0 0 1 15 3.48V7c0 1.566-.32 3.182-1.303 4.682-.983 1.498-2.585 2.813-5.032 3.855a1.697 1.697 0 0 1-1.33 0c-2.447-1.042-4.049-2.357-5.032-3.855C1.32 10.182 1 8.566 1 7V3.48a1.75 1.75 0 0 1 1.217-1.667l5.25-1.68a1.748 1.748 0 0 1 1.066 0Zm-.61 1.429.001.001-5.25 1.68a.251.251 0 0 0-.174.237V7c0 1.36.275 2.666 1.057 3.859.784 1.194 2.121 2.342 4.366 3.298a.196.196 0 0 0 .154 0c2.245-.957 3.582-2.103 4.366-3.297C13.225 9.666 13.5 8.358 13.5 7V3.48a.25.25 0 0 0-.174-.238l-5.25-1.68a.25.25 0 0 0-.153 0ZM11.28 6.28l-3.5 3.5a.75.75 0 0 1-1.06 0l-1.5-1.5a.749.749 0 0 1 .326-1.275.749.749 0 0 1 .734.215l.97.97 2.97-2.97a.751.751 0 0 1 1.042.018.751.751 0 0 1 .018 1.042Z"></path></svg> Prompt Guard</div>
+<svg aria-hidden="true" class="sd-octicon sd-octicon-shield-check" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="m8.533.133 5.25 1.68A1.75 1.75 0 0 1 15 3.48V7c0 1.566-.32 3.182-1.303 4.682-.983 1.498-2.585 2.813-5.032 3.855a1.697 1.697 0 0 1-1.33 0c-2.447-1.042-4.049-2.357-5.032-3.855C1.32 10.182 1 8.566 1 7V3.48a1.75 1.75 0 0 1 1.217-1.667l5.25-1.68a1.748 1.748 0 0 1 1.066 0Zm-.61 1.429.001.001-5.25 1.68a.251.251 0 0 0-.174.237V7c0 1.36.275 2.666 1.057 3.859.784 1.194 2.121 2.342 4.366 3.298a.196.196 0 0 0 .154 0c2.245-.957 3.582-2.103 4.366-3.297C13.225 9.666 13.5 8.358 13.5 7V3.48a.25.25 0 0 0-.174-.238l-5.25-1.68a.25.25 0 0 0-.153 0ZM11.28 6.28l-3.5 3.5a.75.75 0 0 1-1.06 0l-1.5-1.5a.749.749 0 0 1 .326-1.275.749.749 0 0 1 .734.215l.97.97 2.97-2.97a.751.751 0 0 1 1.042.018.751.751 0 0 1 .018 1.042Z"></path></svg> Guardrails</div>
 <p class="sd-card-text">Instructions on securing and validating prompts</p>
 </div>
 <a class="sd-stretched-link sd-hide-link-text reference external" href="../guides/prompt_guard.html"><span>../guides/prompt_guard.html</span></a></div>
@@ -262,45 +257,45 @@ simplify the interaction with upstream LLMs, and improve observability all while
 <div class="sd-card sd-sphinx-override sd-w-100 sd-shadow-sm sd-card-hover docutils">
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
-<svg aria-hidden="true" class="sd-octicon sd-octicon-code-square" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M0 1.75C0 .784.784 0 1.75 0h12.5C15.216 0 16 .784 16 1.75v12.5A1.75 1.75 0 0 1 14.25 16H1.75A1.75 1.75 0 0 1 0 14.25Zm1.75-.25a.25.25 0 0 0-.25.25v12.5c0 .138.112.25.25.25h12.5a.25.25 0 0 0 .25-.25V1.75a.25.25 0 0 0-.25-.25Zm7.47 3.97a.75.75 0 0 1 1.06 0l2 2a.75.75 0 0 1 0 1.06l-2 2a.749.749 0 0 1-1.275-.326.749.749 0 0 1 .215-.734L10.69 8 9.22 6.53a.75.75 0 0 1 0-1.06ZM6.78 6.53 5.31 8l1.47 1.47a.749.749 0 0 1-.326 1.275.749.749 0 0 1-.734-.215l-2-2a.75.75 0 0 1 0-1.06l2-2a.751.751 0 0 1 1.042.018.751.751 0 0 1 .018 1.042Z"></path></svg> Function Calling</div>
-<p class="sd-card-text">A guide to effective function calling</p>
+<svg aria-hidden="true" class="sd-octicon sd-octicon-code-square" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M0 1.75C0 .784.784 0 1.75 0h12.5C15.216 0 16 .784 16 1.75v12.5A1.75 1.75 0 0 1 14.25 16H1.75A1.75 1.75 0 0 1 0 14.25Zm1.75-.25a.25.25 0 0 0-.25.25v12.5c0 .138.112.25.25.25h12.5a.25.25 0 0 0 .25-.25V1.75a.25.25 0 0 0-.25-.25Zm7.47 3.97a.75.75 0 0 1 1.06 0l2 2a.75.75 0 0 1 0 1.06l-2 2a.749.749 0 0 1-1.275-.326.749.749 0 0 1 .215-.734L10.69 8 9.22 6.53a.75.75 0 0 1 0-1.06ZM6.78 6.53 5.31 8l1.47 1.47a.749.749 0 0 1-.326 1.275.749.749 0 0 1-.734-.215l-2-2a.75.75 0 0 1 0-1.06l2-2a.751.751 0 0 1 1.042.018.751.751 0 0 1 .018 1.042Z"></path></svg> LLM Routing</div>
+<p class="sd-card-text">A guide to effective model selection strategies</p>
 </div>
-<a class="sd-stretched-link sd-hide-link-text reference external" href="../guides/function_calling.html"><span>../guides/function_calling.html</span></a></div>
+<a class="sd-stretched-link sd-hide-link-text reference external" href="../guides/llm_router.html"><span>../guides/llm_router.html</span></a></div>
 </div>
 <div class="sd-col sd-d-flex-row docutils">
 <div class="sd-card sd-sphinx-override sd-w-100 sd-shadow-sm sd-card-hover docutils">
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
-<svg aria-hidden="true" class="sd-octicon sd-octicon-issue-opened" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M8 9.5a1.5 1.5 0 1 0 0-3 1.5 1.5 0 0 0 0 3Z"></path><path d="M8 0a8 8 0 1 1 0 16A8 8 0 0 1 8 0ZM1.5 8a6.5 6.5 0 1 0 13 0 6.5 6.5 0 0 0-13 0Z"></path></svg> Observability</div>
-<p class="sd-card-text">Learn to monitor and troubleshoot Arch</p>
+<svg aria-hidden="true" class="sd-octicon sd-octicon-issue-opened" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M8 9.5a1.5 1.5 0 1 0 0-3 1.5 1.5 0 0 0 0 3Z"></path><path d="M8 0a8 8 0 1 1 0 16A8 8 0 0 1 8 0ZM1.5 8a6.5 6.5 0 1 0 13 0 6.5 6.5 0 0 0-13 0Z"></path></svg> State Management</div>
+<p class="sd-card-text">Learn to manage conversation and application state</p>
 </div>
-<a class="sd-stretched-link sd-hide-link-text reference external" href="../guides/observability/observability.html"><span>../guides/observability/observability.html</span></a></div>
+<a class="sd-stretched-link sd-hide-link-text reference external" href="../guides/state.html"><span>../guides/state.html</span></a></div>
 </div>
 </div>
 </div>
 </section>
-<section id="build-with-arch">
-<h2>Build with Arch<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#build-with-arch" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#build-with-arch'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>For developers extending and customizing Arch for specialized needs:</p>
+<section id="build-with-plano">
+<h2>Build with Plano<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#build-with-plano" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#build-with-plano'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>End to end examples demonstrating how to build agentic applications using Plano:</p>
 <div class="sd-container-fluid sd-sphinx-override sd-mb-4 docutils">
 <div class="sd-row sd-row-cols-2 sd-row-cols-xs-2 sd-row-cols-sm-2 sd-row-cols-md-2 sd-row-cols-lg-2 docutils">
 <div class="sd-col sd-d-flex-row docutils">
 <div class="sd-card sd-sphinx-override sd-w-100 sd-shadow-sm sd-card-hover docutils">
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
-<svg aria-hidden="true" class="sd-octicon sd-octicon-dependabot" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M5.75 7.5a.75.75 0 0 1 .75.75v1.5a.75.75 0 0 1-1.5 0v-1.5a.75.75 0 0 1 .75-.75Zm5.25.75a.75.75 0 0 0-1.5 0v1.5a.75.75 0 0 0 1.5 0v-1.5Z"></path><path d="M6.25 0h2A.75.75 0 0 1 9 .75V3.5h3.25a2.25 2.25 0 0 1 2.25 2.25V8h.75a.75.75 0 0 1 0 1.5h-.75v2.75a2.25 2.25 0 0 1-2.25 2.25h-8.5a2.25 2.25 0 0 1-2.25-2.25V9.5H.75a.75.75 0 0 1 0-1.5h.75V5.75A2.25 2.25 0 0 1 3.75 3.5H7.5v-2H6.25a.75.75 0 0 1 0-1.5ZM3 5.75v6.5c0 .414.336.75.75.75h8.5a.75.75 0 0 0 .75-.75v-6.5a.75.75 0 0 0-.75-.75h-8.5a.75.75 0 0 0-.75.75Z"></path></svg> Agentic Workflow</div>
-<p class="sd-card-text">Discover how to create and manage custom agents within Arch</p>
+<svg aria-hidden="true" class="sd-octicon sd-octicon-dependabot" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M5.75 7.5a.75.75 0 0 1 .75.75v1.5a.75.75 0 0 1-1.5 0v-1.5a.75.75 0 0 1 .75-.75Zm5.25.75a.75.75 0 0 0-1.5 0v1.5a.75.75 0 0 0 1.5 0v-1.5Z"></path><path d="M6.25 0h2A.75.75 0 0 1 9 .75V3.5h3.25a2.25 2.25 0 0 1 2.25 2.25V8h.75a.75.75 0 0 1 0 1.5h-.75v2.75a2.25 2.25 0 0 1-2.25 2.25h-8.5a2.25 2.25 0 0 1-2.25-2.25V9.5H.75a.75.75 0 0 1 0-1.5h.75V5.75A2.25 2.25 0 0 1 3.75 3.5H7.5v-2H6.25a.75.75 0 0 1 0-1.5ZM3 5.75v6.5c0 .414.336.75.75.75h8.5a.75.75 0 0 0 .75-.75v-6.5a.75.75 0 0 0-.75-.75h-8.5a.75.75 0 0 0-.75.75Z"></path></svg> Build Agentic Apps</div>
+<p class="sd-card-text">Discover how to create and manage custom agents within Plano</p>
 </div>
-<a class="sd-stretched-link sd-hide-link-text reference external" href="../build_with_arch/agent.html"><span>../build_with_arch/agent.html</span></a></div>
+<a class="sd-stretched-link sd-hide-link-text reference external" href="../get_started/quickstart.html#build-agentic-apps-with-plano"><span>../get_started/quickstart.html#build-agentic-apps-with-plano</span></a></div>
 </div>
 <div class="sd-col sd-d-flex-row docutils">
 <div class="sd-card sd-sphinx-override sd-w-100 sd-shadow-sm sd-card-hover docutils">
 <div class="sd-card-body docutils">
 <div class="sd-card-title sd-font-weight-bold docutils">
-<svg aria-hidden="true" class="sd-octicon sd-octicon-stack" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M7.122.392a1.75 1.75 0 0 1 1.756 0l5.003 2.902c.83.481.83 1.68 0 2.162L8.878 8.358a1.75 1.75 0 0 1-1.756 0L2.119 5.456a1.251 1.251 0 0 1 0-2.162ZM8.125 1.69a.248.248 0 0 0-.25 0l-4.63 2.685 4.63 2.685a.248.248 0 0 0 .25 0l4.63-2.685ZM1.601 7.789a.75.75 0 0 1 1.025-.273l5.249 3.044a.248.248 0 0 0 .25 0l5.249-3.044a.75.75 0 0 1 .752 1.298l-5.248 3.044a1.75 1.75 0 0 1-1.756 0L1.874 8.814A.75.75 0 0 1 1.6 7.789Zm0 3.5a.75.75 0 0 1 1.025-.273l5.249 3.044a.248.248 0 0 0 .25 0l5.249-3.044a.75.75 0 0 1 .752 1.298l-5.248 3.044a1.75 1.75 0 0 1-1.756 0l-5.248-3.044a.75.75 0 0 1-.273-1.025Z"></path></svg> RAG Application</div>
-<p class="sd-card-text">Integrate RAG for knowledge-driven responses</p>
+<svg aria-hidden="true" class="sd-octicon sd-octicon-stack" height="1.0em" version="1.1" viewbox="0 0 16 16" width="1.0em"><path d="M7.122.392a1.75 1.75 0 0 1 1.756 0l5.003 2.902c.83.481.83 1.68 0 2.162L8.878 8.358a1.75 1.75 0 0 1-1.756 0L2.119 5.456a1.251 1.251 0 0 1 0-2.162ZM8.125 1.69a.248.248 0 0 0-.25 0l-4.63 2.685 4.63 2.685a.248.248 0 0 0 .25 0l4.63-2.685ZM1.601 7.789a.75.75 0 0 1 1.025-.273l5.249 3.044a.248.248 0 0 0 .25 0l5.249-3.044a.75.75 0 0 1 .752 1.298l-5.248 3.044a1.75 1.75 0 0 1-1.756 0L1.874 8.814A.75.75 0 0 1 1.6 7.789Zm0 3.5a.75.75 0 0 1 1.025-.273l5.249 3.044a.248.248 0 0 0 .25 0l5.249-3.044a.75.75 0 0 1 .752 1.298l-5.248 3.044a1.75 1.75 0 0 1-1.756 0l-5.248-3.044a.75.75 0 0 1-.273-1.025Z"></path></svg> Build Multi-LLM Apps</div>
+<p class="sd-card-text">Learn how to route LLM calls through Plano for enhanced control and observability</p>
 </div>
-<a class="sd-stretched-link sd-hide-link-text reference external" href="../build_with_arch/rag.html"><span>../build_with_arch/rag.html</span></a></div>
+<a class="sd-stretched-link sd-hide-link-text reference external" href="../get_started/quickstart.html#use-plano-as-a-model-proxy-gateway"><span>../get_started/quickstart.html#use-plano-as-a-model-proxy-gateway</span></a></div>
 </div>
 </div>
 </div>
@@ -312,12 +307,12 @@ simplify the interaction with upstream LLMs, and improve observability all while
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Welcome to Arch!
+        Welcome to Plano!
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="intro_to_arch.html">
-        Intro to Arch
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="intro_to_plano.html">
+        Intro to Plano
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -329,7 +324,7 @@ simplify the interaction with upstream LLMs, and improve observability all while
 <li><a :data-current="activeSection === '#get-started'" class="reference internal" href="#get-started">Get Started</a></li>
 <li><a :data-current="activeSection === '#concepts'" class="reference internal" href="#concepts">Concepts</a></li>
 <li><a :data-current="activeSection === '#guides'" class="reference internal" href="#guides">Guides</a></li>
-<li><a :data-current="activeSection === '#build-with-arch'" class="reference internal" href="#build-with-arch">Build with Arch</a></li>
+<li><a :data-current="activeSection === '#build-with-plano'" class="reference internal" href="#build-with-plano">Build with Plano</a></li>
 </ul>
 </div>
 </aside>
@@ -338,12 +333,12 @@ simplify the interaction with upstream LLMs, and improve observability all while
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../_static/doctools.js?v=9bcbadda"></script>
 <script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
diff --git a/get_started/quickstart.html b/get_started/quickstart.html
index 778883fa..93325747 100755
--- a/get_started/quickstart.html
+++ b/get_started/quickstart.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Quickstart | Arch Docs v0.3.22</title>
-<meta content="Quickstart | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Quickstart | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Quickstart | Plano Docs v0.4</title>
+<meta content="Quickstart | Plano Docs v0.4" property="og:title"/>
+<meta content="Quickstart | Plano Docs v0.4" name="twitter:title"/>
 <link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/get_started/quickstart.html" rel="canonical"/>
 <link href="../_static/favicon.ico" rel="icon"/>
 <link href="../search.html" rel="search" title="Search"/>
-<link href="../concepts/tech_overview/tech_overview.html" rel="next" title="Tech Overview"/>
-<link href="intro_to_arch.html" rel="prev" title="Intro to Arch"/>
+<link href="../concepts/listeners.html" rel="next" title="Listeners"/>
+<link href="intro_to_plano.html" rel="prev" title="Intro to Plano"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul class="current">
 <li class="toctree-l1"><a class="reference internal" href="overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1 current"><a class="current reference internal" href="#">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -163,7 +158,17 @@
 <div id="content" role="main">
 <section id="quickstart">
 <span id="id1"></span><h1>Quickstart<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#quickstart"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Follow this guide to learn how to quickly set up Arch and integrate it into your generative AI applications.</p>
+<p>Follow this guide to learn how to quickly set up Plano and integrate it into your generative AI applications. You can:</p>
+<ul class="simple">
+<li><p><a class="reference internal" href="#quickstart-agents"><span class="std std-ref">Build agents</span></a> for multi-step workflows (e.g., travel assistants with flights and hotels).</p></li>
+<li><p><a class="reference internal" href="#quickstart-prompt-targets"><span class="std std-ref">Call deterministic APIs via prompt targets</span></a> to turn instructions directly into function calls.</p></li>
+<li><p><a class="reference internal" href="#llm-routing-quickstart"><span class="std std-ref">Use Plano as a model proxy (Gateway)</span></a> to standardize access to multiple LLM providers.</p></li>
+</ul>
+<div class="admonition note">
+<p class="admonition-title">Note</p>
+<p>This quickstart assumes basic familiarity with agents and prompt targets from the Concepts section. For background, see <a class="reference internal" href="../concepts/agents.html#agents"><span class="std std-ref">Agents</span></a> and <a class="reference internal" href="../concepts/prompt_target.html#prompt-target"><span class="std std-ref">Prompt Target</span></a>.</p>
+<p>The full agent and backend API implementations used here are available in the <a class="reference external" href="https://github.com/plano-ai/plano-quickstart" rel="nofollow noopener">plano-quickstart repository<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>. This guide focuses on wiring and configuring Plano (orchestration, prompt targets, and the model proxy), not application code.</p>
+</div>
 <section id="prerequisites">
 <h2>Prerequisites<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#prerequisites" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#prerequisites'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <p>Before you begin, ensure you have the following:</p>
@@ -172,24 +177,91 @@
 <li><p><a class="reference external" href="https://docs.docker.com/compose/install/" rel="nofollow noopener">Docker Compose<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> (v2.29)</p></li>
 <li><p><a class="reference external" href="https://www.python.org/downloads/" rel="nofollow noopener">Python<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> (v3.10+)</p></li>
 </ol>
-<p>Arch’s CLI allows you to manage and interact with the Arch gateway efficiently. To install the CLI, simply run the following command:</p>
+<p>Plano’s CLI allows you to manage and interact with the Plano efficiently. To install the CLI, simply run the following command:</p>
 <div class="admonition tip">
 <p class="admonition-title">Tip</p>
-<p>We recommend that developers create a new Python virtual environment to isolate dependencies before installing Arch. This ensures that <code class="docutils literal notranslate"><span class="pre">archgw</span></code> and its dependencies do not interfere with other packages on your system.</p>
+<p>We recommend that developers create a new Python virtual environment to isolate dependencies before installing Plano. This ensures that <code class="docutils literal notranslate"><span class="pre">plano</span></code> and its dependencies do not interfere with other packages on your system.</p>
 </div>
 <div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="gp">$ </span>python<span class="w"> </span>-m<span class="w"> </span>venv<span class="w"> </span>venv
 </span><span id="line-2"><span class="gp">$ </span><span class="nb">source</span><span class="w"> </span>venv/bin/activate<span class="w">   </span><span class="c1"># On Windows, use: venv\Scripts\activate</span>
-</span><span id="line-3"><span class="gp">$ </span>pip<span class="w"> </span>install<span class="w"> </span><span class="nv">archgw</span><span class="o">==</span><span class="m">0</span>.3.22
+</span><span id="line-3"><span class="gp">$ </span>pip<span class="w"> </span>install<span class="w"> </span><span class="nv">plano</span><span class="o">==</span><span class="m">0</span>.4.0
 </span></code></pre></div>
 </div>
 </section>
-<section id="build-ai-agent-with-arch-gateway">
-<h2>Build AI Agent with Arch Gateway<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#build-ai-agent-with-arch-gateway" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#build-ai-agent-with-arch-gateway'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>In the following quickstart, we will show you how easy it is to build an AI agent with the Arch gateway. We will build a currency exchange agent using the following simple steps. For this demo, we will use <cite>https://api.frankfurter.dev/</cite> to fetch the latest prices for currencies and assume USD as the base currency.</p>
-<section id="step-1-create-arch-config-file">
-<h3>Step 1. Create arch config file<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-1-create-arch-config-file" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-1-create-arch-config-file'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Create <code class="docutils literal notranslate"><span class="pre">arch_config.yaml</span></code> file with the following content:</p>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="w"> </span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1.0</span>
+<section id="build-agentic-apps-with-plano">
+<h2>Build Agentic Apps with Plano<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#build-agentic-apps-with-plano" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#build-agentic-apps-with-plano'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Plano helps you build agentic applications in two complementary ways:</p>
+<ul class="simple">
+<li><p><strong>Orchestrate agents</strong>: Let Plano decide which agent or LLM should handle each request and in what sequence.</p></li>
+<li><p><strong>Call deterministic backends</strong>: Use prompt targets to turn natural-language prompts into structured, validated API calls.</p></li>
+</ul>
+<section id="building-agents-with-plano-orchestration">
+<span id="quickstart-agents"></span><h3>Building agents with Plano orchestration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#building-agents-with-plano-orchestration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#building-agents-with-plano-orchestration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Agents are where your business logic lives (the “inner loop”). Plano takes care of the “outer loop”—routing, sequencing, and managing calls across agents and LLMs.</p>
+<p>At a high level, building agents with Plano looks like this:</p>
+<ol class="arabic simple">
+<li><p><strong>Implement your agent</strong> in your framework of choice (Python, JS/TS, etc.), exposing it as an HTTP service.</p></li>
+<li><p><strong>Route LLM calls through Plano’s Model Proxy</strong>, so all models share a consistent interface and observability.</p></li>
+<li><p><strong>Configure Plano to orchestrate</strong>: define which agent(s) can handle which kinds of prompts, and let Plano decide when to call an agent vs. an LLM.</p></li>
+</ol>
+<p>This quickstart uses a simplified version of the Travel Booking Assistant; for the full multi-agent walkthrough, see <a class="reference internal" href="../guides/orchestration.html#agent-routing"><span class="std std-ref">Orchestration</span></a>.</p>
+<section id="step-1-minimal-orchestration-config">
+<h4>Step 1. Minimal orchestration config<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-1-minimal-orchestration-config"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>Here is a minimal configuration that wires Plano-Orchestrator to two HTTP services: one for flights and one for hotels.</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1.0</span>
+</span><span id="line-2">
+</span><span id="line-3"><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-4"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">flight_agent</span>
+</span><span id="line-5"><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10520</span><span class="w">  </span><span class="c1"># your flights service</span>
+</span><span id="line-6"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">hotel_agent</span>
+</span><span id="line-7"><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10530</span><span class="w">  </span><span class="c1"># your hotels service</span>
+</span><span id="line-8">
+</span><span id="line-9"><span class="nt">model_providers</span><span class="p">:</span>
+</span><span id="line-10"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
+</span><span id="line-11"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-12">
+</span><span id="line-13"><span class="nt">listeners</span><span class="p">:</span>
+</span><span id="line-14"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">agent</span>
+</span><span id="line-15"><span class="w">    </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">travel_assistant</span>
+</span><span id="line-16"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">8001</span>
+</span><span id="line-17"><span class="w">    </span><span class="nt">router</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">plano_orchestrator_v1</span>
+</span><span id="line-18"><span class="w">    </span><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-19"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">flight_agent</span>
+</span><span id="line-20"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Search for flights and provide flight status.</span>
+</span><span id="line-21"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">hotel_agent</span>
+</span><span id="line-22"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Find hotels and check availability.</span>
+</span><span id="line-23">
+</span><span id="line-24"><span class="nt">tracing</span><span class="p">:</span>
+</span><span id="line-25"><span class="w">  </span><span class="nt">random_sampling</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">100</span>
+</span></code></pre></div>
+</div>
+</section>
+<section id="step-2-start-your-agents-and-plano">
+<h4>Step 2. Start your agents and Plano<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-start-your-agents-and-plano"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>Run your <code class="docutils literal notranslate"><span class="pre">flight_agent</span></code> and <code class="docutils literal notranslate"><span class="pre">hotel_agent</span></code> services (see <a class="reference internal" href="../guides/orchestration.html#agent-routing"><span class="std std-ref">Orchestration</span></a> for a full Travel Booking example), then start Plano with the config above:</p>
+<div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="gp">$ </span>plano<span class="w"> </span>up<span class="w"> </span>plano_config.yaml
+</span></code></pre></div>
+</div>
+<p>Plano will start the orchestrator and expose an agent listener on port <code class="docutils literal notranslate"><span class="pre">8001</span></code>.</p>
+</section>
+<section id="step-3-send-a-prompt-and-let-plano-route">
+<h4>Step 3. Send a prompt and let Plano route<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-3-send-a-prompt-and-let-plano-route"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>Now send a request to Plano using the OpenAI-compatible chat completions API—the orchestrator will analyze the prompt and route it to the right agent based on intent:</p>
+<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">$<span class="w"> </span>curl<span class="w"> </span>--header<span class="w"> </span><span class="s1">'Content-Type: application/json'</span><span class="w"> </span><span class="se">\</span>
+</span><span id="line-2"><span class="w">  </span>--data<span class="w"> </span><span class="s1">'{"messages": [{"role": "user","content": "Find me flights from SFO to JFK tomorrow"}], "model": "openai/gpt-4o"}'</span><span class="w"> </span><span class="se">\</span>
+</span><span id="line-3"><span class="w">  </span>http://localhost:8001/v1/chat/completions
+</span></code></pre></div>
+</div>
+<p>You can then ask a follow-up like “Also book me a hotel near JFK” and Plano-Orchestrator will route to <code class="docutils literal notranslate"><span class="pre">hotel_agent</span></code>—your agents stay focused on business logic while Plano handles routing.</p>
+</section>
+</section>
+<section id="deterministic-api-calls-with-prompt-targets">
+<span id="quickstart-prompt-targets"></span><h3>Deterministic API calls with prompt targets<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#deterministic-api-calls-with-prompt-targets" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#deterministic-api-calls-with-prompt-targets'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Next, we’ll show Plano’s deterministic API calling using a single prompt target. We’ll build a currency exchange backend powered by <cite>https://api.frankfurter.dev/</cite>, assuming USD as the base currency.</p>
+<section id="step-1-create-plano-config-file">
+<h4>Step 1. Create plano config file<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-1-create-plano-config-file"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>Create <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code> file with the following content:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1.0</span>
 </span><span id="line-2">
 </span><span id="line-3"><span class="nt">listeners</span><span class="p">:</span>
 </span><span id="line-4"><span class="w">  </span><span class="nt">ingress_traffic</span><span class="p">:</span>
@@ -198,54 +270,48 @@
 </span><span id="line-7"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
 </span><span id="line-8"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
 </span><span id="line-9">
-</span><span id="line-10"><span class="w w-Error"> </span><span class="nt">llm_providers</span><span class="p">:</span>
+</span><span id="line-10"><span class="w w-Error"> </span><span class="nt">model_providers</span><span class="p">:</span>
 </span><span id="line-11"><span class="w">   </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
 </span><span id="line-12"><span class="w">     </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
 </span><span id="line-13">
 </span><span id="line-14"><span class="w"> </span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
 </span><span id="line-15"><span class="w">   </span><span class="no">You are a helpful assistant.</span>
 </span><span id="line-16">
-</span><span id="line-17"><span class="w"> </span><span class="nt">prompt_guards</span><span class="p">:</span>
-</span><span id="line-18"><span class="w">   </span><span class="nt">input_guards</span><span class="p">:</span>
-</span><span id="line-19"><span class="w">     </span><span class="nt">jailbreak</span><span class="p">:</span>
-</span><span id="line-20"><span class="w">       </span><span class="nt">on_exception</span><span class="p">:</span>
-</span><span id="line-21"><span class="w">         </span><span class="nt">message</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Looks like you're curious about my abilities, but I can only provide assistance for currency exchange.</span>
-</span><span id="line-22">
-</span><span id="line-23"><span class="w"> </span><span class="nt">prompt_targets</span><span class="p">:</span>
-</span><span id="line-24"><span class="w">   </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">currency_exchange</span>
-</span><span id="line-25"><span class="w">     </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Get currency exchange rate from USD to other currencies</span>
-</span><span id="line-26"><span class="w">     </span><span class="nt">parameters</span><span class="p">:</span>
-</span><span id="line-27"><span class="w">       </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">currency_symbol</span>
-</span><span id="line-28"><span class="w">         </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">the currency that needs conversion</span>
-</span><span id="line-29"><span class="w">         </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-30"><span class="w">         </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
-</span><span id="line-31"><span class="w">         </span><span class="nt">in_path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-32"><span class="w">     </span><span class="nt">endpoint</span><span class="p">:</span>
-</span><span id="line-33"><span class="w">       </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">frankfurther_api</span>
-</span><span id="line-34"><span class="w">       </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/v1/latest?base=USD&amp;symbols={currency_symbol}</span>
-</span><span id="line-35"><span class="w">     </span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
-</span><span id="line-36"><span class="w">       </span><span class="no">You are a helpful assistant. Show me the currency symbol you want to convert from USD.</span>
+</span><span id="line-17"><span class="w"> </span><span class="nt">prompt_targets</span><span class="p">:</span>
+</span><span id="line-18"><span class="w">   </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">currency_exchange</span>
+</span><span id="line-19"><span class="w">     </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Get currency exchange rate from USD to other currencies</span>
+</span><span id="line-20"><span class="w">     </span><span class="nt">parameters</span><span class="p">:</span>
+</span><span id="line-21"><span class="w">       </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">currency_symbol</span>
+</span><span id="line-22"><span class="w">         </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">the currency that needs conversion</span>
+</span><span id="line-23"><span class="w">         </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-24"><span class="w">         </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
+</span><span id="line-25"><span class="w">         </span><span class="nt">in_path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-26"><span class="w">     </span><span class="nt">endpoint</span><span class="p">:</span>
+</span><span id="line-27"><span class="w">       </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">frankfurther_api</span>
+</span><span id="line-28"><span class="w">       </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/v1/latest?base=USD&amp;symbols={currency_symbol}</span>
+</span><span id="line-29"><span class="w">     </span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
+</span><span id="line-30"><span class="w">       </span><span class="no">You are a helpful assistant. Show me the currency symbol you want to convert from USD.</span>
+</span><span id="line-31">
+</span><span id="line-32"><span class="w">   </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">get_supported_currencies</span>
+</span><span id="line-33"><span class="w">     </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Get list of supported currencies for conversion</span>
+</span><span id="line-34"><span class="w">     </span><span class="nt">endpoint</span><span class="p">:</span>
+</span><span id="line-35"><span class="w">       </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">frankfurther_api</span>
+</span><span id="line-36"><span class="w">       </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/v1/currencies</span>
 </span><span id="line-37">
-</span><span id="line-38"><span class="w">   </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">get_supported_currencies</span>
-</span><span id="line-39"><span class="w">     </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Get list of supported currencies for conversion</span>
-</span><span id="line-40"><span class="w">     </span><span class="nt">endpoint</span><span class="p">:</span>
-</span><span id="line-41"><span class="w">       </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">frankfurther_api</span>
-</span><span id="line-42"><span class="w">       </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/v1/currencies</span>
-</span><span id="line-43">
-</span><span id="line-44"><span class="w"> </span><span class="nt">endpoints</span><span class="p">:</span>
-</span><span id="line-45"><span class="w">   </span><span class="nt">frankfurther_api</span><span class="p">:</span>
-</span><span id="line-46"><span class="w">     </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">api.frankfurter.dev:443</span>
-</span><span id="line-47"><span class="w">     </span><span class="nt">protocol</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">https</span>
+</span><span id="line-38"><span class="w"> </span><span class="nt">endpoints</span><span class="p">:</span>
+</span><span id="line-39"><span class="w">   </span><span class="nt">frankfurther_api</span><span class="p">:</span>
+</span><span id="line-40"><span class="w">     </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">api.frankfurter.dev:443</span>
+</span><span id="line-41"><span class="w">     </span><span class="nt">protocol</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">https</span>
 </span></code></pre></div>
 </div>
 </section>
-<section id="step-2-start-arch-gateway-with-currency-conversion-config">
-<h3>Step 2. Start arch gateway with currency conversion config<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-start-arch-gateway-with-currency-conversion-config" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-2-start-arch-gateway-with-currency-conversion-config'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<div class="highlight-sh notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">$<span class="w"> </span>archgw<span class="w"> </span>up<span class="w"> </span>arch_config.yaml
-</span><span id="line-2"><span class="m">2024</span>-12-05<span class="w"> </span><span class="m">16</span>:56:27,979<span class="w"> </span>-<span class="w"> </span>cli.main<span class="w"> </span>-<span class="w"> </span>INFO<span class="w"> </span>-<span class="w"> </span>Starting<span class="w"> </span>archgw<span class="w"> </span>cli<span class="w"> </span>version:<span class="w"> </span><span class="m">0</span>.1.5
+<section id="step-2-start-plano-with-currency-conversion-config">
+<h4>Step 2. Start plano with currency conversion config<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-start-plano-with-currency-conversion-config"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<div class="highlight-sh notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">$<span class="w"> </span>plano<span class="w"> </span>up<span class="w"> </span>plano_config.yaml
+</span><span id="line-2"><span class="m">2024</span>-12-05<span class="w"> </span><span class="m">16</span>:56:27,979<span class="w"> </span>-<span class="w"> </span>cli.main<span class="w"> </span>-<span class="w"> </span>INFO<span class="w"> </span>-<span class="w"> </span>Starting<span class="w"> </span>plano<span class="w"> </span>cli<span class="w"> </span>version:<span class="w"> </span><span class="m">0</span>.1.5
 </span><span id="line-3">...
 </span><span id="line-4"><span class="m">2024</span>-12-05<span class="w"> </span><span class="m">16</span>:56:28,485<span class="w"> </span>-<span class="w"> </span>cli.utils<span class="w"> </span>-<span class="w"> </span>INFO<span class="w"> </span>-<span class="w"> </span>Schema<span class="w"> </span>validation<span class="w"> </span>successful!
-</span><span id="line-5"><span class="m">2024</span>-12-05<span class="w"> </span><span class="m">16</span>:56:28,485<span class="w"> </span>-<span class="w"> </span>cli.main<span class="w"> </span>-<span class="w"> </span>INFO<span class="w"> </span>-<span class="w"> </span>Starting<span class="w"> </span>arch<span class="w"> </span>model<span class="w"> </span>server<span class="w"> </span>and<span class="w"> </span>arch<span class="w"> </span>gateway
+</span><span id="line-5"><span class="m">2024</span>-12-05<span class="w"> </span><span class="m">16</span>:56:28,485<span class="w"> </span>-<span class="w"> </span>cli.main<span class="w"> </span>-<span class="w"> </span>INFO<span class="w"> </span>-<span class="w"> </span>Starting<span class="w"> </span>plano<span class="w"> </span>model<span class="w"> </span>server<span class="w"> </span>and<span class="w"> </span>plano<span class="w"> </span>gateway
 </span><span id="line-6">...
 </span><span id="line-7"><span class="m">2024</span>-12-05<span class="w"> </span><span class="m">16</span>:56:51,647<span class="w"> </span>-<span class="w"> </span>cli.core<span class="w"> </span>-<span class="w"> </span>INFO<span class="w"> </span>-<span class="w"> </span>Container<span class="w"> </span>is<span class="w"> </span>healthy!
 </span></code></pre></div>
@@ -254,7 +320,7 @@
 <p>Some sample queries you can ask include: <code class="docutils literal notranslate"><span class="pre">what</span> <span class="pre">is</span> <span class="pre">currency</span> <span class="pre">rate</span> <span class="pre">for</span> <span class="pre">gbp?</span></code> or <code class="docutils literal notranslate"><span class="pre">show</span> <span class="pre">me</span> <span class="pre">list</span> <span class="pre">of</span> <span class="pre">currencies</span> <span class="pre">for</span> <span class="pre">conversion</span></code>.</p>
 </section>
 <section id="step-3-interacting-with-gateway-using-curl-command">
-<h3>Step 3. Interacting with gateway using curl command<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-3-interacting-with-gateway-using-curl-command" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-3-interacting-with-gateway-using-curl-command'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<h4>Step 3. Interacting with gateway using curl command<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-3-interacting-with-gateway-using-curl-command"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
 <p>Here is a sample curl command you can use to interact:</p>
 <div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">$<span class="w"> </span>curl<span class="w"> </span>--header<span class="w"> </span><span class="s1">'Content-Type: application/json'</span><span class="w"> </span><span class="se">\</span>
 </span><span id="line-2"><span class="w">  </span>--data<span class="w"> </span><span class="s1">'{"messages": [{"role": "user","content": "what is exchange rate for gbp"}], "model": "none"}'</span><span class="w"> </span><span class="se">\</span>
@@ -273,12 +339,13 @@
 </div>
 </section>
 </section>
-<section id="use-arch-gateway-as-llm-router">
-<h2>Use Arch Gateway as LLM Router<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#use-arch-gateway-as-llm-router" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#use-arch-gateway-as-llm-router'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+</section>
+<section id="use-plano-as-a-model-proxy-gateway">
+<span id="llm-routing-quickstart"></span><h2>Use Plano as a Model Proxy (Gateway)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#use-plano-as-a-model-proxy-gateway" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#use-plano-as-a-model-proxy-gateway'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <section id="id2">
-<h3>Step 1. Create arch config file<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#id2'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Arch operates based on a configuration file where you can define LLM providers, prompt targets, guardrails, etc. Below is an example configuration that defines OpenAI and Mistral LLM providers.</p>
-<p>Create <code class="docutils literal notranslate"><span class="pre">arch_config.yaml</span></code> file with the following content:</p>
+<h3>Step 1. Create plano config file<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#id2'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Plano operates based on a configuration file where you can define LLM providers, prompt targets, guardrails, etc. Below is an example configuration that defines OpenAI and Mistral LLM providers.</p>
+<p>Create <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code> file with the following content:</p>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="w"> </span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1.0</span>
 </span><span id="line-2">
 </span><span id="line-3"><span class="nt">listeners</span><span class="p">:</span>
@@ -288,7 +355,7 @@
 </span><span id="line-7"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
 </span><span id="line-8"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
 </span><span id="line-9">
-</span><span id="line-10"><span class="w w-Error"> </span><span class="nt">llm_providers</span><span class="p">:</span>
+</span><span id="line-10"><span class="w w-Error"> </span><span class="nt">model_providers</span><span class="p">:</span>
 </span><span id="line-11"><span class="w">   </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
 </span><span id="line-12"><span class="w">     </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
 </span><span id="line-13"><span class="w">     </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
@@ -298,14 +365,14 @@
 </span></code></pre></div>
 </div>
 </section>
-<section id="step-2-start-arch-gateway">
-<h3>Step 2. Start arch gateway<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-start-arch-gateway" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-2-start-arch-gateway'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<section id="step-2-start-plano">
+<h3>Step 2. Start plano<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-start-plano" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-2-start-plano'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <p>Once the config file is created, ensure that you have environment variables set up for <code class="docutils literal notranslate"><span class="pre">MISTRAL_API_KEY</span></code> and <code class="docutils literal notranslate"><span class="pre">OPENAI_API_KEY</span></code> (or these are defined in a <code class="docutils literal notranslate"><span class="pre">.env</span></code> file).</p>
-<p>Start the Arch gateway:</p>
-<div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="gp">$ </span>archgw<span class="w"> </span>up<span class="w"> </span>arch_config.yaml
-</span><span id="line-2"><span class="go">2024-12-05 11:24:51,288 - cli.main - INFO - Starting archgw cli version: 0.1.5</span>
+<p>Start Plano:</p>
+<div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="gp">$ </span>plano<span class="w"> </span>up<span class="w"> </span>plano_config.yaml
+</span><span id="line-2"><span class="go">2024-12-05 11:24:51,288 - cli.main - INFO - Starting plano cli version: 0.1.5</span>
 </span><span id="line-3"><span class="go">2024-12-05 11:24:51,825 - cli.utils - INFO - Schema validation successful!</span>
-</span><span id="line-4"><span class="go">2024-12-05 11:24:51,825 - cli.main - INFO - Starting arch model server and arch gateway</span>
+</span><span id="line-4"><span class="go">2024-12-05 11:24:51,825 - cli.main - INFO - Starting plano</span>
 </span><span id="line-5"><span class="go">...</span>
 </span><span id="line-6"><span class="go">2024-12-05 11:25:16,131 - cli.core - INFO - Container is healthy!</span>
 </span></code></pre></div>
@@ -315,19 +382,19 @@
 <h3>Step 3: Interact with LLM<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-3-interact-with-llm" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-3-interact-with-llm'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <section id="step-3-1-using-openai-python-client">
 <h4>Step 3.1: Using OpenAI Python client<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-3-1-using-openai-python-client"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
-<p>Make outbound calls via the Arch gateway:</p>
+<p>Make outbound calls via the Plano gateway:</p>
 <div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="kn">from</span><span class="w"> </span><span class="nn">openai</span><span class="w"> </span><span class="kn">import</span> <span class="n">OpenAI</span>
 </span><span id="line-2">
 </span><span id="line-3"><span class="c1"># Use the OpenAI client as usual</span>
 </span><span id="line-4"><span class="n">client</span> <span class="o">=</span> <span class="n">OpenAI</span><span class="p">(</span>
-</span><span id="line-5">  <span class="c1"># No need to set a specific openai.api_key since it's configured in Arch's gateway</span>
+</span><span id="line-5">  <span class="c1"># No need to set a specific openai.api_key since it's configured in Plano's gateway</span>
 </span><span id="line-6">  <span class="n">api_key</span><span class="o">=</span><span class="s1">'--'</span><span class="p">,</span>
-</span><span id="line-7">  <span class="c1"># Set the OpenAI API base URL to the Arch gateway endpoint</span>
+</span><span id="line-7">  <span class="c1"># Set the OpenAI API base URL to the Plano gateway endpoint</span>
 </span><span id="line-8">  <span class="n">base_url</span><span class="o">=</span><span class="s2">"http://127.0.0.1:12000/v1"</span>
 </span><span id="line-9"><span class="p">)</span>
 </span><span id="line-10">
 </span><span id="line-11"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
-</span><span id="line-12">    <span class="c1"># we select model from arch_config file</span>
+</span><span id="line-12">    <span class="c1"># we select model from plano_config file</span>
 </span><span id="line-13">    <span class="n">model</span><span class="o">=</span><span class="s2">"--"</span><span class="p">,</span>
 </span><span id="line-14">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"What is the capital of France?"</span><span class="p">}],</span>
 </span><span id="line-15"><span class="p">)</span>
@@ -357,54 +424,32 @@
 </span><span id="line-17"><span class="o">}</span>
 </span></code></pre></div>
 </div>
-<p>You can override model selection using the <code class="docutils literal notranslate"><span class="pre">x-arch-llm-provider-hint</span></code> header. For example, to use Mistral, use the following curl command:</p>
-<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">$<span class="w"> </span>curl<span class="w"> </span>--header<span class="w"> </span><span class="s1">'Content-Type: application/json'</span><span class="w"> </span><span class="se">\</span>
-</span><span id="line-2"><span class="w">  </span>--header<span class="w"> </span><span class="s1">'x-arch-llm-provider-hint: ministral-3b'</span><span class="w"> </span><span class="se">\</span>
-</span><span id="line-3"><span class="w">  </span>--data<span class="w"> </span><span class="s1">'{"messages": [{"role": "user","content": "What is the capital of France?"}], "model": "none"}'</span><span class="w"> </span><span class="se">\</span>
-</span><span id="line-4"><span class="w">  </span>http://localhost:12000/v1/chat/completions
-</span><span id="line-5">
-</span><span id="line-6"><span class="o">{</span>
-</span><span id="line-7"><span class="w">  </span>...
-</span><span id="line-8"><span class="w">  </span><span class="s2">"model"</span>:<span class="w"> </span><span class="s2">"ministral-3b-latest"</span>,
-</span><span id="line-9"><span class="w">  </span><span class="s2">"choices"</span>:<span class="w"> </span><span class="o">[</span>
-</span><span id="line-10"><span class="w">    </span><span class="o">{</span>
-</span><span id="line-11"><span class="w">      </span><span class="s2">"messages"</span>:<span class="w"> </span><span class="o">{</span>
-</span><span id="line-12"><span class="w">        </span><span class="s2">"role"</span>:<span class="w"> </span><span class="s2">"assistant"</span>,
-</span><span id="line-13"><span class="w">        </span><span class="s2">"content"</span>:<span class="w"> </span><span class="s2">"The capital of France is Paris. It is the most populous city in France and is known for its iconic landmarks such as the Eiffel Tower, the Louvre Museum, and Notre-Dame Cathedral. Paris is also a major global center for art, fashion, gastronomy, and culture."</span>,
-</span><span id="line-14"><span class="w">      </span><span class="o">}</span>,
-</span><span id="line-15"><span class="w">      </span>...
-</span><span id="line-16"><span class="w">    </span><span class="o">}</span>
-</span><span id="line-17"><span class="w">  </span><span class="o">]</span>,
-</span><span id="line-18"><span class="w">  </span>...
-</span><span id="line-19"><span class="o">}</span>
-</span></code></pre></div>
-</div>
 </section>
 </section>
 </section>
 </section>
 <section id="next-steps">
 <h1>Next Steps<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#next-steps"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Congratulations! You’ve successfully set up Arch and made your first prompt-based request. To further enhance your GenAI applications, explore the following resources:</p>
+<p>Congratulations! You’ve successfully set up Plano and made your first prompt-based request. To further enhance your GenAI applications, explore the following resources:</p>
 <ul class="simple">
 <li><p><a class="reference internal" href="overview.html#overview"><span class="std std-ref">Full Documentation</span></a>: Comprehensive guides and references.</p></li>
-<li><p><a class="reference external" href="https://github.com/katanemo/arch" rel="nofollow noopener">GitHub Repository<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>: Access the source code, contribute, and track updates.</p></li>
-<li><p><a class="reference external" href="https://github.com/katanemo/arch#contact" rel="nofollow noopener">Support<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>: Get help and connect with the Arch community .</p></li>
+<li><p><a class="reference external" href="https://github.com/katanemo/plano" rel="nofollow noopener">GitHub Repository<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>: Access the source code, contribute, and track updates.</p></li>
+<li><p><a class="reference external" href="https://github.com/katanemo/plano#contact" rel="nofollow noopener">Support<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>: Get help and connect with the Plano community .</p></li>
 </ul>
-<p>With Arch, building scalable, fast, and personalized GenAI applications has never been easier. Dive deeper into Arch’s capabilities and start creating innovative AI-driven experiences today!</p>
+<p>With Plano, building scalable, fast, and personalized GenAI applications has never been easier. Dive deeper into Plano’s capabilities and start creating innovative AI-driven experiences today!</p>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="intro_to_arch.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="intro_to_plano.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Intro to Arch
+        Intro to Plano
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../concepts/tech_overview/tech_overview.html">
-        Tech Overview
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../concepts/listeners.html">
+        Listeners
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -415,15 +460,24 @@
 <ul>
 <li><ul>
 <li><a :data-current="activeSection === '#prerequisites'" class="reference internal" href="#prerequisites">Prerequisites</a></li>
-<li><a :data-current="activeSection === '#build-ai-agent-with-arch-gateway'" class="reference internal" href="#build-ai-agent-with-arch-gateway">Build AI Agent with Arch Gateway</a><ul>
-<li><a :data-current="activeSection === '#step-1-create-arch-config-file'" class="reference internal" href="#step-1-create-arch-config-file">Step 1. Create arch config file</a></li>
-<li><a :data-current="activeSection === '#step-2-start-arch-gateway-with-currency-conversion-config'" class="reference internal" href="#step-2-start-arch-gateway-with-currency-conversion-config">Step 2. Start arch gateway with currency conversion config</a></li>
+<li><a :data-current="activeSection === '#build-agentic-apps-with-plano'" class="reference internal" href="#build-agentic-apps-with-plano">Build Agentic Apps with Plano</a><ul>
+<li><a :data-current="activeSection === '#building-agents-with-plano-orchestration'" class="reference internal" href="#building-agents-with-plano-orchestration">Building agents with Plano orchestration</a><ul>
+<li><a :data-current="activeSection === '#step-1-minimal-orchestration-config'" class="reference internal" href="#step-1-minimal-orchestration-config">Step 1. Minimal orchestration config</a></li>
+<li><a :data-current="activeSection === '#step-2-start-your-agents-and-plano'" class="reference internal" href="#step-2-start-your-agents-and-plano">Step 2. Start your agents and Plano</a></li>
+<li><a :data-current="activeSection === '#step-3-send-a-prompt-and-let-plano-route'" class="reference internal" href="#step-3-send-a-prompt-and-let-plano-route">Step 3. Send a prompt and let Plano route</a></li>
+</ul>
+</li>
+<li><a :data-current="activeSection === '#deterministic-api-calls-with-prompt-targets'" class="reference internal" href="#deterministic-api-calls-with-prompt-targets">Deterministic API calls with prompt targets</a><ul>
+<li><a :data-current="activeSection === '#step-1-create-plano-config-file'" class="reference internal" href="#step-1-create-plano-config-file">Step 1. Create plano config file</a></li>
+<li><a :data-current="activeSection === '#step-2-start-plano-with-currency-conversion-config'" class="reference internal" href="#step-2-start-plano-with-currency-conversion-config">Step 2. Start plano with currency conversion config</a></li>
 <li><a :data-current="activeSection === '#step-3-interacting-with-gateway-using-curl-command'" class="reference internal" href="#step-3-interacting-with-gateway-using-curl-command">Step 3. Interacting with gateway using curl command</a></li>
 </ul>
 </li>
-<li><a :data-current="activeSection === '#use-arch-gateway-as-llm-router'" class="reference internal" href="#use-arch-gateway-as-llm-router">Use Arch Gateway as LLM Router</a><ul>
-<li><a :data-current="activeSection === '#id2'" class="reference internal" href="#id2">Step 1. Create arch config file</a></li>
-<li><a :data-current="activeSection === '#step-2-start-arch-gateway'" class="reference internal" href="#step-2-start-arch-gateway">Step 2. Start arch gateway</a></li>
+</ul>
+</li>
+<li><a :data-current="activeSection === '#use-plano-as-a-model-proxy-gateway'" class="reference internal" href="#use-plano-as-a-model-proxy-gateway">Use Plano as a Model Proxy (Gateway)</a><ul>
+<li><a :data-current="activeSection === '#id2'" class="reference internal" href="#id2">Step 1. Create plano config file</a></li>
+<li><a :data-current="activeSection === '#step-2-start-plano'" class="reference internal" href="#step-2-start-plano">Step 2. Start plano</a></li>
 <li><a :data-current="activeSection === '#step-3-interact-with-llm'" class="reference internal" href="#step-3-interact-with-llm">Step 3: Interact with LLM</a><ul>
 <li><a :data-current="activeSection === '#step-3-1-using-openai-python-client'" class="reference internal" href="#step-3-1-using-openai-python-client">Step 3.1: Using OpenAI Python client</a></li>
 <li><a :data-current="activeSection === '#step-3-2-using-curl-command'" class="reference internal" href="#step-3-2-using-curl-command">Step 3.2: Using curl command</a></li>
@@ -442,12 +496,12 @@
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../_static/doctools.js?v=9bcbadda"></script>
 <script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
diff --git a/guides/agent_routing.html b/guides/agent_routing.html
deleted file mode 100755
index 30ccff5c..00000000
--- a/guides/agent_routing.html
+++ /dev/null
@@ -1,311 +0,0 @@
-<!DOCTYPE html>
-
-<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
-<head>
-<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
-<meta charset="utf-8"/>
-<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
-<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
-<meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Agent Routing and Hand Off | Arch Docs v0.3.22</title>
-<meta content="Agent Routing and Hand Off | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Agent Routing and Hand Off | Arch Docs v0.3.22" name="twitter:title"/>
-<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
-<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
-<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
-<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/guides/agent_routing.html" rel="canonical"/>
-<link href="../_static/favicon.ico" rel="icon"/>
-<link href="../search.html" rel="search" title="Search"/>
-<link href="function_calling.html" rel="next" title="Function Calling"/>
-<link href="prompt_guard.html" rel="prev" title="Prompt Guard"/>
-<script>
-    <!-- Prevent Flash of wrong theme -->
-      const userPreference = localStorage.getItem('darkMode');
-      let mode;
-      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
-        mode = 'dark';
-        document.documentElement.classList.add('dark');
-      } else {
-        mode = 'light';
-      }
-      if (!userPreference) {localStorage.setItem('darkMode', mode)}
-    </script>
-</head>
-<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
-<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
-      Skip to content
-    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
-<div class="hidden mr-4 md:flex">
-<a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
-<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
-</svg>
-<span class="sr-only">Toggle navigation menu</span>
-</button>
-<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
-<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
-<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
-<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
-<span class="text-xs">⌘</span>
-    K
-  </kbd>
-</form>
-</div>
-<nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
-<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
-<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
-</div>
-</a>
-<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
-<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
-</svg>
-<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
-</svg>
-</button>
-</nav>
-</div>
-</div>
-</header>
-<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
-<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
-</a>
-<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
-<div class="overflow-y-auto h-full w-full relative pr-6">
-
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
-<script>
-  window.dataLayer = window.dataLayer || [];
-  function gtag(){dataLayer.push(arguments);}
-  gtag('js', new Date());
-
-  gtag('config', 'G-K2LXXSX6HB');
-</script>
-<nav class="table w-full min-w-full my-6 lg:my-8">
-<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../concepts/prompt_target.html">Prompt Target</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
-<ul class="current">
-<li class="toctree-l1"><a class="reference internal" href="prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1 current"><a class="current reference internal" href="#">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="function_calling.html">Function Calling</a></li>
-<li class="toctree-l1"><a class="reference internal" href="llm_router.html">LLM Routing</a></li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="observability/tracing.html">Tracing</a></li>
-<li class="toctree-l2"><a class="reference internal" href="observability/monitoring.html">Monitoring</a></li>
-<li class="toctree-l2"><a class="reference internal" href="observability/access_logging.html">Access Logging</a></li>
-</ul>
-</li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
-</ul>
-</nav>
-</div>
-</div>
-<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
-<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
-<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
-</svg>
-</button>
-</aside>
-<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
-<div class="w-full min-w-0 mx-auto">
-<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
-<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
-<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
-<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
-</svg>
-</a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Agent Routing and Hand Off</span>
-</nav>
-<div id="content" role="main">
-<section id="agent-routing-and-hand-off">
-<span id="agent-routing"></span><h1>Agent Routing and Hand Off<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#agent-routing-and-hand-off"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Agent Routing and Hand Off is a key feature in Arch that enables intelligent routing of user prompts to specialized AI agents or human agents based on the nature and complexity of the user’s request.</p>
-<p>This capability significantly enhances the efficiency and personalization of interactions, ensuring each prompt receives the most appropriate and effective handling. The following section describes
-the workflow, configuration, and implementation of Agent routing and hand off in Arch.</p>
-<ol class="arabic simple">
-<li><p><strong>Agent Selection</strong>
-When a user submits a prompt, Arch analyzes the input to determine the intent and complexity. Based on the analysis, Arch selects the most suitable agent configured within your application to handle the specific category of the user’s request—such as sales inquiries, technical issues, or complex scenarios requiring human attention.</p></li>
-<li><p><strong>Prompt Routing</strong>
-After selecting the appropriate agent, Arch routes the user’s prompt to the designated agent’s endpoint and waits for the agent to respond back with the processed output or further instructions.</p></li>
-<li><p><strong>Hand Off</strong>
-Based on follow-up queries from the user, Arch repeats the process of analysis, agent selection, and routing to ensure a seamless hand off between AI agents as needed.</p></li>
-</ol>
-<div class="literal-block-wrapper docutils container" id="id1">
-<div class="code-block-caption"><span class="caption-text">Agent Routing and Hand Off Configuration Example</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id1"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">prompt_targets</span><span class="p">:</span>
-</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">sales_agent</span>
-</span><span id="line-3"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Handles queries related to sales and purchases</span>
-</span><span id="line-4">
-</span><span id="line-5"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">issues_and_repairs</span>
-</span><span id="line-6"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">handles issues, repairs, or refunds</span>
-</span><span id="line-7">
-</span><span id="line-8"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">escalate_to_human</span>
-</span><span id="line-9"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">escalates to human agent</span>
-</span></code></pre></div>
-</div>
-</div>
-<div class="literal-block-wrapper docutils container" id="id2">
-<div class="code-block-caption"><span class="caption-text">Agent Routing and Hand Off Implementation Example via FastAPI</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="k">class</span><span class="w"> </span><span class="nc">Agent</span><span class="p">:</span>
-</span><span id="line-2">    <span class="k">def</span><span class="w"> </span><span class="fm">__init__</span><span class="p">(</span><span class="bp">self</span><span class="p">,</span> <span class="n">role</span><span class="p">:</span> <span class="nb">str</span><span class="p">,</span> <span class="n">instructions</span><span class="p">:</span> <span class="nb">str</span><span class="p">):</span>
-</span><span id="line-3">        <span class="bp">self</span><span class="o">.</span><span class="n">system_prompt</span> <span class="o">=</span> <span class="sa">f</span><span class="s2">"You are a </span><span class="si">{</span><span class="n">role</span><span class="si">}</span><span class="s2">.</span><span class="se">\n</span><span class="si">{</span><span class="n">instructions</span><span class="si">}</span><span class="s2">"</span>
-</span><span id="line-4">
-</span><span id="line-5">    <span class="k">def</span><span class="w"> </span><span class="nf">handle</span><span class="p">(</span><span class="bp">self</span><span class="p">,</span> <span class="n">req</span><span class="p">:</span> <span class="n">ChatCompletionsRequest</span><span class="p">):</span>
-</span><span id="line-6">        <span class="n">messages</span> <span class="o">=</span> <span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"system"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="bp">self</span><span class="o">.</span><span class="n">get_system_prompt</span><span class="p">()}]</span> <span class="o">+</span> <span class="p">[</span>
-</span><span id="line-7">            <span class="n">message</span><span class="o">.</span><span class="n">model_dump</span><span class="p">()</span> <span class="k">for</span> <span class="n">message</span> <span class="ow">in</span> <span class="n">req</span><span class="o">.</span><span class="n">messages</span>
-</span><span id="line-8">        <span class="p">]</span>
-</span><span id="line-9">        <span class="k">return</span> <span class="n">call_openai</span><span class="p">(</span><span class="n">messages</span><span class="p">,</span> <span class="n">req</span><span class="o">.</span><span class="n">stream</span><span class="p">)</span> <span class="c1">#call_openai is a placeholder for the actual API call</span>
-</span><span id="line-10">
-</span><span id="line-11">    <span class="k">def</span><span class="w"> </span><span class="nf">get_system_prompt</span><span class="p">(</span><span class="bp">self</span><span class="p">)</span> <span class="o">-&gt;</span> <span class="nb">str</span><span class="p">:</span>
-</span><span id="line-12">        <span class="k">return</span> <span class="bp">self</span><span class="o">.</span><span class="n">system_prompt</span>
-</span><span id="line-13">
-</span><span id="line-14"><span class="c1"># Define your agents</span>
-</span><span id="line-15"><span class="n">AGENTS</span> <span class="o">=</span> <span class="p">{</span>
-</span><span id="line-16">    <span class="s2">"sales_agent"</span><span class="p">:</span> <span class="n">Agent</span><span class="p">(</span>
-</span><span id="line-17">        <span class="n">role</span><span class="o">=</span><span class="s2">"sales agent"</span><span class="p">,</span>
-</span><span id="line-18">        <span class="n">instructions</span><span class="o">=</span><span class="p">(</span>
-</span><span id="line-19">            <span class="s2">"Always answer in a sentence or less.</span><span class="se">\n</span><span class="s2">"</span>
-</span><span id="line-20">            <span class="s2">"Follow the following routine with the user:</span><span class="se">\n</span><span class="s2">"</span>
-</span><span id="line-21">            <span class="s2">"1. Engage</span><span class="se">\n</span><span class="s2">"</span>
-</span><span id="line-22">            <span class="s2">"2. Quote ridiculous price</span><span class="se">\n</span><span class="s2">"</span>
-</span><span id="line-23">            <span class="s2">"3. Reveal caveat if user agrees."</span>
-</span><span id="line-24">        <span class="p">),</span>
-</span><span id="line-25">    <span class="p">),</span>
-</span><span id="line-26">    <span class="s2">"issues_and_repairs"</span><span class="p">:</span> <span class="n">Agent</span><span class="p">(</span>
-</span><span id="line-27">        <span class="n">role</span><span class="o">=</span><span class="s2">"issues and repairs agent"</span><span class="p">,</span>
-</span><span id="line-28">        <span class="n">instructions</span><span class="o">=</span><span class="s2">"Propose a solution, offer refund if necessary."</span><span class="p">,</span>
-</span><span id="line-29">    <span class="p">),</span>
-</span><span id="line-30">    <span class="s2">"escalate_to_human"</span><span class="p">:</span> <span class="n">Agent</span><span class="p">(</span>
-</span><span id="line-31">        <span class="n">role</span><span class="o">=</span><span class="s2">"human escalation agent"</span><span class="p">,</span> <span class="n">instructions</span><span class="o">=</span><span class="s2">"Escalate issues to a human."</span>
-</span><span id="line-32">    <span class="p">),</span>
-</span><span id="line-33">    <span class="s2">"unknown_agent"</span><span class="p">:</span> <span class="n">Agent</span><span class="p">(</span>
-</span><span id="line-34">        <span class="n">role</span><span class="o">=</span><span class="s2">"general assistant"</span><span class="p">,</span> <span class="n">instructions</span><span class="o">=</span><span class="s2">"Assist the user in general queries."</span>
-</span><span id="line-35">    <span class="p">),</span>
-</span><span id="line-36"><span class="p">}</span>
-</span><span id="line-37">
-</span><span id="line-38"><span class="c1">#handle the request from arch gateway</span>
-</span><span id="line-39"><span class="nd">@app</span><span class="o">.</span><span class="n">post</span><span class="p">(</span><span class="s2">"/v1/chat/completions"</span><span class="p">)</span>
-</span><span id="line-40"><span class="k">def</span><span class="w"> </span><span class="nf">completion_api</span><span class="p">(</span><span class="n">req</span><span class="p">:</span> <span class="n">ChatCompletionsRequest</span><span class="p">,</span> <span class="n">request</span><span class="p">:</span> <span class="n">Request</span><span class="p">):</span>
-</span><span id="line-41">
-</span><span id="line-42">    <span class="n">agent_name</span> <span class="o">=</span> <span class="n">req</span><span class="o">.</span><span class="n">metadata</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"agent-name"</span><span class="p">,</span> <span class="s2">"unknown_agent"</span><span class="p">)</span>
-</span><span id="line-43">    <span class="n">agent</span> <span class="o">=</span> <span class="n">AGENTS</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="n">agent_name</span><span class="p">)</span>
-</span><span id="line-44">    <span class="n">logger</span><span class="o">.</span><span class="n">info</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Routing to agent: </span><span class="si">{</span><span class="n">agent_name</span><span class="si">}</span><span class="s2">"</span><span class="p">)</span>
-</span><span id="line-45">
-</span><span id="line-46">    <span class="k">return</span> <span class="n">agent</span><span class="o">.</span><span class="n">handle</span><span class="p">(</span><span class="n">req</span><span class="p">)</span>
-</span></code></pre></div>
-</div>
-</div>
-<div class="admonition note">
-<p class="admonition-title">Note</p>
-<p>The above example demonstrates a simple implementation of Agent Routing and Hand Off using FastAPI. For the full implementation of this example
-please see our <a class="reference external" href="https://github.com/katanemo/archgw/tree/main/demos/use_cases/orchestrating_agents" rel="nofollow noopener">GitHub demo<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.</p>
-</div>
-<section id="example-use-cases">
-<h2>Example Use Cases<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-use-cases" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-use-cases'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Agent Routing and Hand Off is particularly beneficial in scenarios such as:</p>
-<ul class="simple">
-<li><p><strong>Customer Support</strong>: Routing common customer queries to automated support agents, while escalating complex or sensitive issues to human support staff.</p></li>
-<li><p><strong>Sales and Marketing</strong>: Automatically directing potential leads and sales inquiries to specialized sales agents for timely and targeted follow-ups.</p></li>
-<li><p><strong>Technical Assistance</strong>: Managing user-reported issues, repairs, or refunds by assigning them to the correct technical or support agent efficiently.</p></li>
-</ul>
-</section>
-<section id="best-practices-and-tips">
-<h2>Best Practices and Tips<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#best-practices-and-tips" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#best-practices-and-tips'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>When implementing Agent Routing and Hand Off in your applications, consider these best practices:</p>
-<ul class="simple">
-<li><p>Clearly define agent responsibilities: Ensure each agent or human endpoint has a clear, specific description of the prompts they handle, reducing mis-routing.</p></li>
-<li><p>Monitor and optimize routes: Regularly review how prompts are routed to adjust and optimize agent definitions and configurations.</p></li>
-</ul>
-<div class="admonition note">
-<p class="admonition-title">Note</p>
-<p>To observe traffic to and from agents, please read more about <a class="reference internal" href="observability/observability.html#observability"><span class="std std-ref">observability</span></a> in Arch.</p>
-</div>
-<p>By carefully configuring and managing your Agent routing and hand off, you can significantly improve your application’s responsiveness, performance, and overall user satisfaction.</p>
-</section>
-</section>
-</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
-<div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="prompt_guard.html">
-<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="15 18 9 12 15 6"></polyline>
-</svg>
-        Prompt Guard
-      </a>
-</div>
-<div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="function_calling.html">
-        Function Calling
-        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
-<polyline points="9 18 15 12 9 6"></polyline>
-</svg>
-</a>
-</div>
-</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
-<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
-<ul>
-<li><a :data-current="activeSection === '#example-use-cases'" class="reference internal" href="#example-use-cases">Example Use Cases</a></li>
-<li><a :data-current="activeSection === '#best-practices-and-tips'" class="reference internal" href="#best-practices-and-tips">Best Practices and Tips</a></li>
-</ul>
-</div>
-</aside>
-</main>
-</div>
-</div><footer class="py-6 border-t border-border md:py-0">
-<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
-<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
-</div>
-</div>
-</footer>
-</div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
-<script src="../_static/doctools.js?v=9bcbadda"></script>
-<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
-<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
-<script src="../_static/design-tabs.js?v=f930bc37"></script>
-</body>
-</html>
\ No newline at end of file
diff --git a/guides/function_calling.html b/guides/function_calling.html
index 9522ad99..0db4a5fa 100755
--- a/guides/function_calling.html
+++ b/guides/function_calling.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Function Calling | Arch Docs v0.3.22</title>
-<meta content="Function Calling | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Function Calling | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Function Calling | Plano Docs v0.4</title>
+<meta content="Function Calling | Plano Docs v0.4" property="og:title"/>
+<meta content="Function Calling | Plano Docs v0.4" name="twitter:title"/>
 <link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/guides/function_calling.html" rel="canonical"/>
 <link href="../_static/favicon.ico" rel="icon"/>
 <link href="../search.html" rel="search" title="Search"/>
-<link href="llm_router.html" rel="next" title="LLM Routing"/>
-<link href="agent_routing.html" rel="prev" title="Agent Routing and Hand Off"/>
+<link href="observability/observability.html" rel="next" title="Observability"/>
+<link href="llm_router.html" rel="prev" title="LLM Routing"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul class="current">
-<li class="toctree-l1"><a class="reference internal" href="prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1 current"><a class="current reference internal" href="#">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1 current"><a class="current reference internal" href="#">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -163,7 +158,7 @@
 <div id="content" role="main">
 <section id="function-calling">
 <span id="id1"></span><h1>Function Calling<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#function-calling"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p><strong>Function Calling</strong> is a powerful feature in Arch that allows your application to dynamically execute backend functions or services based on user prompts.
+<p><strong>Function Calling</strong> is a powerful feature in Plano that allows your application to dynamically execute backend functions or services based on user prompts.
 This enables seamless integration between natural language interactions and backend operations, turning user inputs into actionable results.</p>
 <section id="what-is-function-calling">
 <h2>What is Function Calling?<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#what-is-function-calling" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#what-is-function-calling'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
@@ -175,17 +170,17 @@ This feature bridges the gap between generative AI systems and functional busine
 <ol class="arabic">
 <li><p><strong>Prompt Parsing</strong></p>
 <blockquote>
-<div><p>When a user submits a prompt, Arch analyzes it to determine the intent. Based on this intent, the system identifies whether a function needs to be invoked and which parameters should be extracted.</p>
+<div><p>When a user submits a prompt, Plano analyzes it to determine the intent. Based on this intent, the system identifies whether a function needs to be invoked and which parameters should be extracted.</p>
 </div></blockquote>
 </li>
 <li><p><strong>Parameter Extraction</strong></p>
 <blockquote>
-<div><p>Arch’s advanced natural language processing capabilities automatically extract parameters from the prompt that are necessary for executing the function. These parameters can include text, numbers, dates, locations, or other relevant data points.</p>
+<div><p>Plano’s advanced natural language processing capabilities automatically extract parameters from the prompt that are necessary for executing the function. These parameters can include text, numbers, dates, locations, or other relevant data points.</p>
 </div></blockquote>
 </li>
 <li><p><strong>Function Invocation</strong></p>
 <blockquote>
-<div><p>Once the necessary parameters have been extracted, Arch invokes the relevant backend function. This function could be an API, a database query, or any other form of backend logic. The function is executed with the extracted parameters to produce the desired output.</p>
+<div><p>Once the necessary parameters have been extracted, Plano invokes the relevant backend function. This function could be an API, a database query, or any other form of backend logic. The function is executed with the extracted parameters to produce the desired output.</p>
 </div></blockquote>
 </li>
 <li><p><strong>Response Handling</strong></p>
@@ -234,10 +229,10 @@ Achieving performance on par with GPT-4, these models set a new benchmark in the
 </section>
 <section id="implementing-function-calling">
 <h2>Implementing Function Calling<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#implementing-function-calling" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#implementing-function-calling'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Here’s a step-by-step guide to configuring function calling within your Arch setup:</p>
+<p>Here’s a step-by-step guide to configuring function calling within your Plano setup:</p>
 <section id="step-1-define-the-function">
 <h3>Step 1: Define the Function<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-1-define-the-function" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-1-define-the-function'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>First, create or identify the backend function you want Arch to call. This could be an API endpoint, a script, or any other executable backend logic.</p>
+<p>First, create or identify the backend function you want Plano to call. This could be an API endpoint, a script, or any other executable backend logic.</p>
 <div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="kn">import</span><span class="w"> </span><span class="nn">requests</span>
 </span><span id="line-2">
 </span><span id="line-3"><span class="k">def</span><span class="w"> </span><span class="nf">get_weather</span><span class="p">(</span><span class="n">location</span><span class="p">:</span> <span class="nb">str</span><span class="p">,</span> <span class="n">unit</span><span class="p">:</span> <span class="nb">str</span> <span class="o">=</span> <span class="s2">"fahrenheit"</span><span class="p">):</span>
@@ -263,8 +258,8 @@ Achieving performance on par with GPT-4, these models set a new benchmark in the
 </section>
 <section id="step-2-configure-prompt-targets">
 <h3>Step 2: Configure Prompt Targets<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-2-configure-prompt-targets" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-2-configure-prompt-targets'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Next, map the function to a prompt target, defining the intent and parameters that Arch will extract from the user’s prompt.
-Specify the parameters your function needs and how Arch should interpret these.</p>
+<p>Next, map the function to a prompt target, defining the intent and parameters that Plano will extract from the user’s prompt.
+Specify the parameters your function needs and how Plano should interpret these.</p>
 <div class="literal-block-wrapper docutils container" id="id3">
 <div class="code-block-caption"><span class="caption-text">Prompt Target Example Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id3"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">prompt_targets</span><span class="p">:</span>
@@ -290,20 +285,20 @@ Specify the parameters your function needs and how Arch should interpret these.<
 <p>For a complete refernce of attributes that you can configure in a prompt target, see <a class="reference internal" href="../concepts/prompt_target.html#defining-prompt-target-parameters"><span class="std std-ref">here</span></a>.</p>
 </div>
 </section>
-<section id="step-3-arch-takes-over">
-<h3>Step 3: Arch Takes Over<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-3-arch-takes-over" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-3-arch-takes-over'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Once you have defined the functions and configured the prompt targets, Arch Gateway takes care of the remaining work.
+<section id="step-3-plano-takes-over">
+<h3>Step 3: Plano Takes Over<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#step-3-plano-takes-over" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#step-3-plano-takes-over'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Once you have defined the functions and configured the prompt targets, Plano takes care of the remaining work.
 It will automatically validate parameters, and ensure that the required parameters (e.g., location) are present in the prompt, and add validation rules if necessary.</p>
 <figure class="align-center" id="id4">
-<a class="reference internal image-reference" href="../_images/arch_network_diagram_high_level.png"><img alt="../_images/arch_network_diagram_high_level.png" src="../_images/arch_network_diagram_high_level.png" style="width: 100%;"/>
+<a class="reference internal image-reference" href="../_images/plano_network_diagram_high_level.png"><img alt="../_images/plano_network_diagram_high_level.png" src="../_images/plano_network_diagram_high_level.png" style="width: 100%;"/>
 </a>
 <figcaption>
-<p><span class="caption-text">High-level network flow of where Arch Gateway sits in your agentic stack. Managing incoming and outgoing prompt traffic</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id4"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></p>
+<p><span class="caption-text">High-level network flow of where Plano sits in your agentic stack. Managing incoming and outgoing prompt traffic</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id4"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></p>
 </figcaption>
 </figure>
-<p>Once a downstream function (API) is called, Arch Gateway takes the response and sends it an upstream LLM to complete the request (for summarization, Q/A, text generation tasks).
-For more details on how Arch Gateway enables you to centralize usage of LLMs, please read <a class="reference internal" href="../concepts/llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">LLM providers</span></a>.</p>
-<p>By completing these steps, you enable Arch to manage the process from validation to response, ensuring users receive consistent, reliable results - and that you are focused
+<p>Once a downstream function (API) is called, Plano  takes the response and sends it an upstream LLM to complete the request (for summarization, Q/A, text generation tasks).
+For more details on how Plano  enables you to centralize usage of LLMs, please read <a class="reference internal" href="../concepts/llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">LLM providers</span></a>.</p>
+<p>By completing these steps, you enable Plano to manage the process from validation to response, ensuring users receive consistent, reliable results - and that you are focused
 on the stuff that matters most.</p>
 </section>
 </section>
@@ -320,7 +315,7 @@ on the stuff that matters most.</p>
 </section>
 <section id="best-practices-and-tips">
 <h2>Best Practices and Tips<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#best-practices-and-tips" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#best-practices-and-tips'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>When integrating function calling into your generative AI applications, keep these tips in mind to get the most out of our Arch-Function models:</p>
+<p>When integrating function calling into your generative AI applications, keep these tips in mind to get the most out of our Plano-Function models:</p>
 <ul class="simple">
 <li><p><strong>Keep it clear and simple</strong>: Your function names and parameters should be straightforward and easy to understand. Think of it like explaining a task to a smart colleague - the clearer you are, the better the results.</p></li>
 <li><p><strong>Context is king</strong>: Don’t skimp on the descriptions for your functions and parameters. The more context you provide, the better the LLM can understand when and how to use each function.</p></li>
@@ -333,16 +328,16 @@ on the stuff that matters most.</p>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="agent_routing.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="llm_router.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Agent Routing and Hand Off
+        LLM Routing
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="llm_router.html">
-        LLM Routing
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="observability/observability.html">
+        Observability
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -360,7 +355,7 @@ on the stuff that matters most.</p>
 <li><a :data-current="activeSection === '#implementing-function-calling'" class="reference internal" href="#implementing-function-calling">Implementing Function Calling</a><ul>
 <li><a :data-current="activeSection === '#step-1-define-the-function'" class="reference internal" href="#step-1-define-the-function">Step 1: Define the Function</a></li>
 <li><a :data-current="activeSection === '#step-2-configure-prompt-targets'" class="reference internal" href="#step-2-configure-prompt-targets">Step 2: Configure Prompt Targets</a></li>
-<li><a :data-current="activeSection === '#step-3-arch-takes-over'" class="reference internal" href="#step-3-arch-takes-over">Step 3: Arch Takes Over</a></li>
+<li><a :data-current="activeSection === '#step-3-plano-takes-over'" class="reference internal" href="#step-3-plano-takes-over">Step 3: Plano Takes Over</a></li>
 </ul>
 </li>
 <li><a :data-current="activeSection === '#example-use-cases'" class="reference internal" href="#example-use-cases">Example Use Cases</a></li>
@@ -373,12 +368,12 @@ on the stuff that matters most.</p>
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../_static/doctools.js?v=9bcbadda"></script>
 <script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
diff --git a/guides/llm_router.html b/guides/llm_router.html
index c887e9a8..3a954251 100755
--- a/guides/llm_router.html
+++ b/guides/llm_router.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>LLM Routing | Arch Docs v0.3.22</title>
-<meta content="LLM Routing | Arch Docs v0.3.22" property="og:title"/>
-<meta content="LLM Routing | Arch Docs v0.3.22" name="twitter:title"/>
+<title>LLM Routing | Plano Docs v0.4</title>
+<meta content="LLM Routing | Plano Docs v0.4" property="og:title"/>
+<meta content="LLM Routing | Plano Docs v0.4" name="twitter:title"/>
 <link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/guides/llm_router.html" rel="canonical"/>
 <link href="../_static/favicon.ico" rel="icon"/>
 <link href="../search.html" rel="search" title="Search"/>
-<link href="observability/observability.html" rel="next" title="Observability"/>
-<link href="function_calling.html" rel="prev" title="Function Calling"/>
+<link href="function_calling.html" rel="next" title="Function Calling"/>
+<link href="orchestration.html" rel="prev" title="Orchestration"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul class="current">
-<li class="toctree-l1"><a class="reference internal" href="prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="orchestration.html">Orchestration</a></li>
 <li class="toctree-l1 current"><a class="current reference internal" href="#">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -163,133 +158,185 @@
 <div id="content" role="main">
 <section id="llm-routing">
 <span id="llm-router"></span><h1>LLM Routing<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#llm-routing"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>With the rapid proliferation of large language models (LLM) — each optimized for different strengths, style, or latency/cost profile — routing has become an essential technique to operationalize the use of different models.</p>
-<p>Arch provides three distinct routing approaches to meet different use cases:</p>
-<ol class="arabic simple">
-<li><p><strong>Model-based Routing</strong>: Direct routing to specific models using provider/model names</p></li>
-<li><p><strong>Alias-based Routing</strong>: Semantic routing using custom aliases that map to underlying models</p></li>
-<li><p><strong>Preference-aligned Routing</strong>: Intelligent routing using the Arch-Router model based on context and user-defined preferences</p></li>
-</ol>
-<p>This enables optimal performance, cost efficiency, and response quality by matching requests with the most suitable model from your available LLM fleet.</p>
+<p>With the rapid proliferation of large language models (LLMs) — each optimized for different strengths, style, or latency/cost profile — routing has become an essential technique to operationalize the use of different models. Plano provides three distinct routing approaches to meet different use cases: <a class="reference internal" href="#model-based-routing"><span class="std std-ref">Model-based routing</span></a>, <a class="reference internal" href="#alias-based-routing"><span class="std std-ref">Alias-based routing</span></a>, and <a class="reference internal" href="#preference-aligned-routing"><span class="std std-ref">Preference-aligned routing</span></a>. This enables optimal performance, cost efficiency, and response quality by matching requests with the most suitable model from your available LLM fleet.</p>
+<div class="admonition note">
+<p class="admonition-title">Note</p>
+<p>For details on supported model providers, configuration options, and client libraries, see <a class="reference internal" href="../concepts/llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">LLM Providers</span></a>.</p>
+</div>
 <section id="routing-methods">
 <h2>Routing Methods<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#routing-methods" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#routing-methods'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <section id="model-based-routing">
-<h3>Model-based Routing<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#model-based-routing" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#model-based-routing'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<span id="id1"></span><h3>Model-based routing<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#model-based-routing" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#model-based-routing'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <p>Direct routing allows you to specify exact provider and model combinations using the format <code class="docutils literal notranslate"><span class="pre">provider/model-name</span></code>:</p>
 <ul class="simple">
-<li><p>Use provider-specific names like <code class="docutils literal notranslate"><span class="pre">openai/gpt-4o</span></code> or <code class="docutils literal notranslate"><span class="pre">anthropic/claude-3-5-sonnet-20241022</span></code></p></li>
+<li><p>Use provider-specific names like <code class="docutils literal notranslate"><span class="pre">openai/gpt-5.2</span></code> or <code class="docutils literal notranslate"><span class="pre">anthropic/claude-sonnet-4-5</span></code></p></li>
 <li><p>Provides full control and transparency over which model handles each request</p></li>
 <li><p>Ideal for production workloads where you want predictable routing behavior</p></li>
 </ul>
+<section id="configuration">
+<h4>Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuration"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>Configure your LLM providers with specific provider/model names:</p>
+<div class="literal-block-wrapper docutils container" id="id9">
+<div class="code-block-caption"><span class="caption-text">Model-based Routing Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id9"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">listeners</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="nt">egress_traffic</span><span class="p">:</span>
+</span><span id="line-3"><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
+</span><span id="line-4"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">12000</span>
+</span><span id="line-5"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
+</span><span id="line-6"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
+</span><span id="line-7">
+</span><span id="line-8"><span class="nt">llm_providers</span><span class="p">:</span>
+</span><span id="line-9"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5.2</span>
+</span><span id="line-10"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-11"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-12">
+</span><span id="line-13"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5</span>
+</span><span id="line-14"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-15">
+</span><span id="line-16"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-sonnet-4-5</span>
+</span><span id="line-17"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
+</span></code></pre></div>
+</div>
+</div>
+</section>
+<section id="client-usage">
+<h4>Client usage<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#client-usage"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>Clients specify exact models:</p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Direct provider/model specification</span>
+</span><span id="line-2"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-3">    <span class="n">model</span><span class="o">=</span><span class="s2">"openai/gpt-5.2"</span><span class="p">,</span>
+</span><span id="line-4">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Hello!"</span><span class="p">}]</span>
+</span><span id="line-5"><span class="p">)</span>
+</span><span id="line-6">
+</span><span id="line-7"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-8">    <span class="n">model</span><span class="o">=</span><span class="s2">"anthropic/claude-sonnet-4-5"</span><span class="p">,</span>
+</span><span id="line-9">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Write a story"</span><span class="p">}]</span>
+</span><span id="line-10"><span class="p">)</span>
+</span></code></pre></div>
+</div>
+</section>
 </section>
 <section id="alias-based-routing">
-<h3>Alias-based Routing<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#alias-based-routing" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#alias-based-routing'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<span id="id2"></span><h3>Alias-based routing<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#alias-based-routing" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#alias-based-routing'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <p>Alias-based routing lets you create semantic model names that decouple your application from specific providers:</p>
 <ul class="simple">
-<li><p>Use meaningful names like <code class="docutils literal notranslate"><span class="pre">fast-model</span></code>, <code class="docutils literal notranslate"><span class="pre">reasoning-model</span></code>, or <code class="docutils literal notranslate"><span class="pre">arch.summarize.v1</span></code> (see <a class="reference internal" href="../concepts/llm_providers/model_aliases.html#model-aliases"><span class="std std-ref">Model Aliases</span></a>)</p></li>
+<li><p>Use meaningful names like <code class="docutils literal notranslate"><span class="pre">fast-model</span></code>, <code class="docutils literal notranslate"><span class="pre">reasoning-model</span></code>, or <code class="docutils literal notranslate"><span class="pre">plano.summarize.v1</span></code> (see <a class="reference internal" href="../concepts/llm_providers/model_aliases.html#model-aliases"><span class="std std-ref">Model Aliases</span></a>)</p></li>
 <li><p>Maps semantic names to underlying provider models for easier experimentation and provider switching</p></li>
 <li><p>Ideal for applications that want abstraction from specific model names while maintaining control</p></li>
 </ul>
+<section id="id3">
+<h4>Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id3"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>Configure semantic aliases that map to underlying models:</p>
+<div class="literal-block-wrapper docutils container" id="id10">
+<div class="code-block-caption"><span class="caption-text">Alias-based Routing Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id10"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">listeners</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="nt">egress_traffic</span><span class="p">:</span>
+</span><span id="line-3"><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
+</span><span id="line-4"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">12000</span>
+</span><span id="line-5"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
+</span><span id="line-6"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
+</span><span id="line-7">
+</span><span id="line-8"><span class="nt">llm_providers</span><span class="p">:</span>
+</span><span id="line-9"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5.2</span>
+</span><span id="line-10"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-11">
+</span><span id="line-12"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5</span>
+</span><span id="line-13"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-14">
+</span><span id="line-15"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-sonnet-4-5</span>
+</span><span id="line-16"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
+</span><span id="line-17">
+</span><span id="line-18"><span class="nt">model_aliases</span><span class="p">:</span>
+</span><span id="line-19"><span class="w">  </span><span class="c1"># Model aliases - friendly names that map to actual provider names</span>
+</span><span id="line-20"><span class="w">  </span><span class="nt">fast-model</span><span class="p">:</span>
+</span><span id="line-21"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-5.2</span>
+</span><span id="line-22">
+</span><span id="line-23"><span class="w">  </span><span class="nt">reasoning-model</span><span class="p">:</span>
+</span><span id="line-24"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-5</span>
+</span><span id="line-25">
+</span><span id="line-26"><span class="w">  </span><span class="nt">creative-model</span><span class="p">:</span>
+</span><span id="line-27"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">claude-sonnet-4-5</span>
+</span></code></pre></div>
+</div>
+</div>
+</section>
+<section id="id4">
+<h4>Client usage<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id4"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>Clients use semantic names:</p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Using semantic aliases</span>
+</span><span id="line-2"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-3">    <span class="n">model</span><span class="o">=</span><span class="s2">"fast-model"</span><span class="p">,</span>  <span class="c1"># Routes to best available fast model</span>
+</span><span id="line-4">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Quick summary please"</span><span class="p">}]</span>
+</span><span id="line-5"><span class="p">)</span>
+</span><span id="line-6">
+</span><span id="line-7"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-8">    <span class="n">model</span><span class="o">=</span><span class="s2">"reasoning-model"</span><span class="p">,</span>  <span class="c1"># Routes to best reasoning model</span>
+</span><span id="line-9">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Solve this complex problem"</span><span class="p">}]</span>
+</span><span id="line-10"><span class="p">)</span>
+</span></code></pre></div>
+</div>
+</section>
 </section>
 <section id="preference-aligned-routing-arch-router">
-<span id="preference-aligned-routing"></span><h3>Preference-aligned Routing (Arch-Router)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#preference-aligned-routing-arch-router" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#preference-aligned-routing-arch-router'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Traditional LLM routing approaches face significant limitations: they evaluate performance using benchmarks that often fail to capture human preferences, select from fixed model pools, and operate as “black boxes” without practical mechanisms for encoding user preferences.</p>
-<p>Arch’s preference-aligned routing addresses these challenges by applying a fundamental engineering principle: decoupling. The framework separates route selection (matching queries to human-readable policies) from model assignment (mapping policies to specific LLMs). This separation allows you to define routing policies using descriptive labels like <code class="docutils literal notranslate"><span class="pre">Domain:</span> <span class="pre">'finance',</span> <span class="pre">Action:</span> <span class="pre">'analyze_earnings_report'</span></code> rather than cryptic identifiers, while independently configuring which models handle each policy.</p>
-<p>The <a class="reference external" href="https://huggingface.co/katanemo/Arch-Router-1.5B" rel="nofollow noopener">Arch-Router<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> model automatically selects the most appropriate LLM based on:</p>
+<span id="preference-aligned-routing"></span><h3>Preference-aligned routing (Arch-Router)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#preference-aligned-routing-arch-router" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#preference-aligned-routing-arch-router'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Preference-aligned routing uses the <a class="reference external" href="https://huggingface.co/katanemo/Arch-Router-1.5B" rel="nofollow noopener">Arch-Router<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> model to pick the best LLM based on domain, action, and your configured preferences instead of hard-coding a model.</p>
 <ul class="simple">
-<li><p>Domain Analysis: Identifies the subject matter (e.g., legal, healthcare, programming)</p></li>
-<li><p>Action Classification: Determines the type of operation (e.g., summarization, code generation, translation)</p></li>
-<li><p>User-Defined Preferences: Maps domains and actions to preferred models using transparent, configurable routing decisions</p></li>
-<li><p>Human Preference Alignment: Uses domain-action mappings that capture subjective evaluation criteria, ensuring routing aligns with real-world user needs rather than just benchmark scores</p></li>
+<li><p><strong>Domain</strong>: High-level topic of the request (e.g., legal, healthcare, programming).</p></li>
+<li><p><strong>Action</strong>: What the user wants to do (e.g., summarize, generate code, translate).</p></li>
+<li><p><strong>Routing preferences</strong>: Your mapping from (domain, action) to preferred models.</p></li>
 </ul>
-<p>This approach supports seamlessly adding new models without retraining and is ideal for dynamic, context-aware routing that adapts to request content and intent.</p>
+<p>Arch-Router analyzes each prompt to infer domain and action, then applies your preferences to select a model. This decouples <strong>routing policy</strong> (how to choose) from <strong>model assignment</strong> (what to run), making routing transparent, controllable, and easy to extend as you add or swap models.</p>
+<section id="id5">
+<h4>Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id5"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>To configure preference-aligned dynamic routing, define routing preferences that map domains and actions to specific models:</p>
+<div class="literal-block-wrapper docutils container" id="id11">
+<div class="code-block-caption"><span class="caption-text">Preference-Aligned Dynamic Routing Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id11"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">listeners</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="nt">egress_traffic</span><span class="p">:</span>
+</span><span id="line-3"><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
+</span><span id="line-4"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">12000</span>
+</span><span id="line-5"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
+</span><span id="line-6"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
+</span><span id="line-7">
+</span><span id="line-8"><span class="nt">llm_providers</span><span class="p">:</span>
+</span><span id="line-9"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5.2</span>
+</span><span id="line-10"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-11"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-12">
+</span><span id="line-13"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5</span>
+</span><span id="line-14"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-15"><span class="w">    </span><span class="nt">routing_preferences</span><span class="p">:</span>
+</span><span id="line-16"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">code understanding</span>
+</span><span id="line-17"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">understand and explain existing code snippets, functions, or libraries</span>
+</span><span id="line-18"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">complex reasoning</span>
+</span><span id="line-19"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">deep analysis, mathematical problem solving, and logical reasoning</span>
+</span><span id="line-20">
+</span><span id="line-21"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-sonnet-4-5</span>
+</span><span id="line-22"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
+</span><span id="line-23"><span class="w">    </span><span class="nt">routing_preferences</span><span class="p">:</span>
+</span><span id="line-24"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">creative writing</span>
+</span><span id="line-25"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">creative content generation, storytelling, and writing assistance</span>
+</span><span id="line-26"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">code generation</span>
+</span><span id="line-27"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">generating new code snippets, functions, or boilerplate based on user prompts</span>
+</span></code></pre></div>
+</div>
+</div>
+</section>
+<section id="id6">
+<h4>Client usage<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id6"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
+<p>Clients can let the router decide or still specify aliases:</p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Let Arch-Router choose based on content</span>
+</span><span id="line-2"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-3">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Write a creative story about space exploration"</span><span class="p">}]</span>
+</span><span id="line-4">    <span class="c1"># No model specified - router will analyze and choose claude-sonnet-4-5</span>
+</span><span id="line-5"><span class="p">)</span>
+</span></code></pre></div>
+</div>
 </section>
 </section>
-<section id="model-based-routing-workflow">
-<h2>Model-based Routing Workflow<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#model-based-routing-workflow" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#model-based-routing-workflow'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>For direct model routing, the process is straightforward:</p>
-<ol class="arabic">
-<li><p><strong>Client Request</strong></p>
-<blockquote>
-<div><p>The client specifies the exact model using provider/model format (<code class="docutils literal notranslate"><span class="pre">openai/gpt-4o</span></code>).</p>
-</div></blockquote>
-</li>
-<li><p><strong>Provider Validation</strong></p>
-<blockquote>
-<div><p>Arch validates that the specified provider and model are configured and available.</p>
-</div></blockquote>
-</li>
-<li><p><strong>Direct Routing</strong></p>
-<blockquote>
-<div><p>The request is sent directly to the specified model without analysis or decision-making.</p>
-</div></blockquote>
-</li>
-<li><p><strong>Response Handling</strong></p>
-<blockquote>
-<div><p>The response is returned to the client with optional metadata about the routing decision.</p>
-</div></blockquote>
-</li>
-</ol>
 </section>
-<section id="alias-based-routing-workflow">
-<h2>Alias-based Routing Workflow<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#alias-based-routing-workflow" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#alias-based-routing-workflow'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>For alias-based routing, the process includes name resolution:</p>
-<ol class="arabic">
-<li><p><strong>Client Request</strong></p>
-<blockquote>
-<div><p>The client specifies a semantic alias name (<code class="docutils literal notranslate"><span class="pre">reasoning-model</span></code>).</p>
-</div></blockquote>
-</li>
-<li><p><strong>Alias Resolution</strong></p>
-<blockquote>
-<div><p>Arch resolves the alias to the actual provider/model name based on configuration.</p>
-</div></blockquote>
-</li>
-<li><p><strong>Model Selection</strong></p>
-<blockquote>
-<div><p>If the alias maps to multiple models, Arch selects one based on availability and load balancing.</p>
-</div></blockquote>
-</li>
-<li><p><strong>Request Forwarding</strong></p>
-<blockquote>
-<div><p>The request is forwarded to the resolved model.</p>
-</div></blockquote>
-</li>
-<li><p><strong>Response Handling</strong></p>
-<blockquote>
-<div><p>The response is returned with optional metadata about the alias resolution.</p>
-</div></blockquote>
-</li>
-</ol>
-</section>
-<section id="preference-aligned-routing-workflow-arch-router">
-<span id="preference-aligned-routing-workflow"></span><h2>Preference-aligned Routing Workflow (Arch-Router)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#preference-aligned-routing-workflow-arch-router" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#preference-aligned-routing-workflow-arch-router'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>For preference-aligned dynamic routing, the process involves intelligent analysis:</p>
-<ol class="arabic">
-<li><p><strong>Prompt Analysis</strong></p>
-<blockquote>
-<div><p>When a user submits a prompt without specifying a model, the Arch-Router analyzes it to determine the domain (subject matter) and action (type of operation requested).</p>
-</div></blockquote>
-</li>
-<li><p><strong>Model Selection</strong></p>
-<blockquote>
-<div><p>Based on the analyzed intent and your configured routing preferences, the Router selects the most appropriate model from your available LLM fleet.</p>
-</div></blockquote>
-</li>
-<li><p><strong>Request Forwarding</strong></p>
-<blockquote>
-<div><p>Once the optimal model is identified, our gateway forwards the original prompt to the selected LLM endpoint. The routing decision is transparent and can be logged for monitoring and optimization purposes.</p>
-</div></blockquote>
-</li>
-<li><p><strong>Response Handling</strong></p>
-<blockquote>
-<div><p>After the selected model processes the request, the response is returned through the gateway. The gateway can optionally add routing metadata or performance metrics to help you understand and optimize your routing decisions.</p>
-</div></blockquote>
-</li>
-</ol>
-</section>
-<section id="id1">
-<h2>Arch-Router<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id1" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#id1'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<section id="id7">
+<h2>Arch-Router<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id7" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#id7'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <p>The <a class="reference external" href="https://huggingface.co/katanemo/Arch-Router-1.5B" rel="nofollow noopener">Arch-Router<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> is a state-of-the-art <strong>preference-based routing model</strong> specifically designed to address the limitations of traditional LLM routing. This compact 1.5B model delivers production-ready performance with low latency and high accuracy while solving key routing challenges.</p>
 <p><strong>Addressing Traditional Routing Limitations:</strong></p>
 <p><strong>Human Preference Alignment</strong>
@@ -312,152 +359,23 @@ Provides a practical mechanism to encode user preferences through domain-action
 <li><p><strong>Production-Ready Performance</strong>: Optimized for low-latency, high-throughput applications in multi-model environments.</p></li>
 </ul>
 </section>
-<section id="implementing-routing">
-<h2>Implementing Routing<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#implementing-routing" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#implementing-routing'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p><strong>Model-based Routing</strong></p>
-<p>For direct model routing, configure your LLM providers with specific provider/model names:</p>
-<div class="literal-block-wrapper docutils container" id="id3">
-<div class="code-block-caption"><span class="caption-text">Model-based Routing Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id3"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">listeners</span><span class="p">:</span>
-</span><span id="line-2"><span class="w">  </span><span class="nt">egress_traffic</span><span class="p">:</span>
-</span><span id="line-3"><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
-</span><span id="line-4"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">12000</span>
-</span><span id="line-5"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-6"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
-</span><span id="line-7">
-</span><span id="line-8"><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-9"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o-mini</span>
-</span><span id="line-10"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-11"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-12">
-</span><span id="line-13"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
-</span><span id="line-14"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-15">
-</span><span id="line-16"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-3-5-sonnet-20241022</span>
-</span><span id="line-17"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
-</span></code></pre></div>
-</div>
-</div>
-<p>Clients specify exact models:</p>
-<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Direct provider/model specification</span>
-</span><span id="line-2"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
-</span><span id="line-3">    <span class="n">model</span><span class="o">=</span><span class="s2">"openai/gpt-4o-mini"</span><span class="p">,</span>
-</span><span id="line-4">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Hello!"</span><span class="p">}]</span>
-</span><span id="line-5"><span class="p">)</span>
-</span><span id="line-6">
-</span><span id="line-7"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
-</span><span id="line-8">    <span class="n">model</span><span class="o">=</span><span class="s2">"anthropic/claude-3-5-sonnet-20241022"</span><span class="p">,</span>
-</span><span id="line-9">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Write a story"</span><span class="p">}]</span>
-</span><span id="line-10"><span class="p">)</span>
-</span></code></pre></div>
-</div>
-<p><strong>Alias-based Routing</strong></p>
-<p>Configure semantic aliases that map to underlying models:</p>
-<div class="literal-block-wrapper docutils container" id="id4">
-<div class="code-block-caption"><span class="caption-text">Alias-based Routing Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id4"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">listeners</span><span class="p">:</span>
-</span><span id="line-2"><span class="w">  </span><span class="nt">egress_traffic</span><span class="p">:</span>
-</span><span id="line-3"><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
-</span><span id="line-4"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">12000</span>
-</span><span id="line-5"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-6"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
-</span><span id="line-7">
-</span><span id="line-8"><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-9"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o-mini</span>
-</span><span id="line-10"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-11">
-</span><span id="line-12"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
-</span><span id="line-13"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-14">
-</span><span id="line-15"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-3-5-sonnet-20241022</span>
-</span><span id="line-16"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
-</span><span id="line-17">
-</span><span id="line-18"><span class="nt">model_aliases</span><span class="p">:</span>
-</span><span id="line-19"><span class="w">  </span><span class="c1"># Model aliases - friendly names that map to actual provider names</span>
-</span><span id="line-20"><span class="w">  </span><span class="nt">fast-model</span><span class="p">:</span>
-</span><span id="line-21"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-4o-mini</span>
-</span><span id="line-22">
-</span><span id="line-23"><span class="w">  </span><span class="nt">reasoning-model</span><span class="p">:</span>
-</span><span id="line-24"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-4o</span>
-</span><span id="line-25">
-</span><span id="line-26"><span class="w">  </span><span class="nt">creative-model</span><span class="p">:</span>
-</span><span id="line-27"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">claude-3-5-sonnet-20241022</span>
-</span></code></pre></div>
-</div>
-</div>
-<p>Clients use semantic names:</p>
-<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Using semantic aliases</span>
-</span><span id="line-2"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
-</span><span id="line-3">    <span class="n">model</span><span class="o">=</span><span class="s2">"fast-model"</span><span class="p">,</span>  <span class="c1"># Routes to best available fast model</span>
-</span><span id="line-4">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Quick summary please"</span><span class="p">}]</span>
-</span><span id="line-5"><span class="p">)</span>
-</span><span id="line-6">
-</span><span id="line-7"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
-</span><span id="line-8">    <span class="n">model</span><span class="o">=</span><span class="s2">"reasoning-model"</span><span class="p">,</span>  <span class="c1"># Routes to best reasoning model</span>
-</span><span id="line-9">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Solve this complex problem"</span><span class="p">}]</span>
-</span><span id="line-10"><span class="p">)</span>
-</span></code></pre></div>
-</div>
-<p><strong>Preference-aligned Routing (Arch-Router)</strong></p>
-<p>To configure preference-aligned dynamic routing, you need to define routing preferences that map domains and actions to specific models:</p>
-<div class="literal-block-wrapper docutils container" id="id5">
-<div class="code-block-caption"><span class="caption-text">Preference-Aligned Dynamic Routing Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id5"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">listeners</span><span class="p">:</span>
-</span><span id="line-2"><span class="w">  </span><span class="nt">egress_traffic</span><span class="p">:</span>
-</span><span id="line-3"><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
-</span><span id="line-4"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">12000</span>
-</span><span id="line-5"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-6"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
-</span><span id="line-7">
-</span><span id="line-8"><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-9"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o-mini</span>
-</span><span id="line-10"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-11"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-12">
-</span><span id="line-13"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
-</span><span id="line-14"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-15"><span class="w">    </span><span class="nt">routing_preferences</span><span class="p">:</span>
-</span><span id="line-16"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">code understanding</span>
-</span><span id="line-17"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">understand and explain existing code snippets, functions, or libraries</span>
-</span><span id="line-18"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">complex reasoning</span>
-</span><span id="line-19"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">deep analysis, mathematical problem solving, and logical reasoning</span>
-</span><span id="line-20">
-</span><span id="line-21"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-3-5-sonnet-20241022</span>
-</span><span id="line-22"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
-</span><span id="line-23"><span class="w">    </span><span class="nt">routing_preferences</span><span class="p">:</span>
-</span><span id="line-24"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">creative writing</span>
-</span><span id="line-25"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">creative content generation, storytelling, and writing assistance</span>
-</span><span id="line-26"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">code generation</span>
-</span><span id="line-27"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">generating new code snippets, functions, or boilerplate based on user prompts</span>
-</span></code></pre></div>
-</div>
-</div>
-<p>Clients can let the router decide or use aliases:</p>
-<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Let Arch-Router choose based on content</span>
-</span><span id="line-2"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
-</span><span id="line-3">    <span class="n">messages</span><span class="o">=</span><span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="s2">"Write a creative story about space exploration"</span><span class="p">}]</span>
-</span><span id="line-4">    <span class="c1"># No model specified - router will analyze and choose claude-3-5-sonnet-20241022</span>
-</span><span id="line-5"><span class="p">)</span>
-</span></code></pre></div>
-</div>
-</section>
 <section id="combining-routing-methods">
 <h2>Combining Routing Methods<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#combining-routing-methods" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#combining-routing-methods'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <p>You can combine static model selection with dynamic routing preferences for maximum flexibility:</p>
-<div class="literal-block-wrapper docutils container" id="id6">
-<div class="code-block-caption"><span class="caption-text">Hybrid Routing Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id6"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="literal-block-wrapper docutils container" id="id12">
+<div class="code-block-caption"><span class="caption-text">Hybrid Routing Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id12"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o-mini</span>
+</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5.2</span>
 </span><span id="line-3"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
 </span><span id="line-4"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
 </span><span id="line-5">
-</span><span id="line-6"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
+</span><span id="line-6"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-5</span>
 </span><span id="line-7"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
 </span><span id="line-8"><span class="w">    </span><span class="nt">routing_preferences</span><span class="p">:</span>
 </span><span id="line-9"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">complex_reasoning</span>
 </span><span id="line-10"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">deep analysis and complex problem solving</span>
 </span><span id="line-11">
-</span><span id="line-12"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-3-5-sonnet-20241022</span>
+</span><span id="line-12"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-sonnet-4-5</span>
 </span><span id="line-13"><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
 </span><span id="line-14"><span class="w">    </span><span class="nt">routing_preferences</span><span class="p">:</span>
 </span><span id="line-15"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">creative_tasks</span>
@@ -466,14 +384,14 @@ Provides a practical mechanism to encode user preferences through domain-action
 </span><span id="line-18"><span class="nt">model_aliases</span><span class="p">:</span>
 </span><span id="line-19"><span class="w">  </span><span class="c1"># Model aliases - friendly names that map to actual provider names</span>
 </span><span id="line-20"><span class="w">  </span><span class="nt">fast-model</span><span class="p">:</span>
-</span><span id="line-21"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-4o-mini</span>
+</span><span id="line-21"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-5.2</span>
 </span><span id="line-22">
 </span><span id="line-23"><span class="w">  </span><span class="nt">reasoning-model</span><span class="p">:</span>
-</span><span id="line-24"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-4o</span>
+</span><span id="line-24"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-5</span>
 </span><span id="line-25">
 </span><span id="line-26"><span class="w">  </span><span class="c1"># Aliases that can also participate in dynamic routing</span>
 </span><span id="line-27"><span class="w">  </span><span class="nt">creative-model</span><span class="p">:</span>
-</span><span id="line-28"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">claude-3-5-sonnet-20241022</span>
+</span><span id="line-28"><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">claude-sonnet-4-5</span>
 </span></code></pre></div>
 </div>
 </div>
@@ -493,8 +411,8 @@ Provides a practical mechanism to encode user preferences through domain-action
 <li><p><strong>Conversational Routing</strong>: Track conversation context to identify when topics shift between domains or when the type of assistance needed changes mid-conversation.</p></li>
 </ul>
 </section>
-<section id="best-practicesm">
-<h2>Best practicesm<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#best-practicesm" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#best-practicesm'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<section id="best-practices">
+<h2>Best practices<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#best-practices" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#best-practices'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <ul class="simple">
 <li><p><strong>💡Consistent Naming:</strong>  Route names should align with their descriptions.</p>
 <ul>
@@ -521,22 +439,31 @@ Provides a practical mechanism to encode user preferences through domain-action
 </ul>
 </li>
 <li><p><strong>💡Nouns Descriptor:</strong> Preference-based routers perform better with noun-centric descriptors, as they offer more stable and semantically rich signals for matching.</p></li>
-<li><p><strong>💡Domain Inclusion:</strong> for best user experience, you should always include domain route. This help the router fall back to domain when action is not</p></li>
+<li><p><strong>💡Domain Inclusion:</strong> for best user experience, you should always include a domain route. This helps the router fall back to domain when action is not confidently inferred.</p></li>
+</ul>
+</section>
+<section id="unsupported-features">
+<h2>Unsupported Features<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#unsupported-features" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#unsupported-features'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>The following features are <strong>not supported</strong> by the Arch-Router model:</p>
+<ul class="simple">
+<li><p><strong>Multi-modality</strong>: The model is not trained to process raw image or audio inputs. It can handle textual queries <em>about</em> these modalities (e.g., “generate an image of a cat”), but cannot interpret encoded multimedia data directly.</p></li>
+<li><p><strong>Function calling</strong>: Arch-Router is designed for <strong>semantic preference matching</strong>, not exact intent classification or tool execution. For structured function invocation, use models in the Plano Function Calling collection instead.</p></li>
+<li><p><strong>System prompt dependency</strong>: Arch-Router routes based solely on the user’s conversation history. It does not use or rely on system prompts for routing decisions.</p></li>
 </ul>
 </section>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="function_calling.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="orchestration.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Function Calling
+        Orchestration
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="observability/observability.html">
-        Observability
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="function_calling.html">
+        Function Calling
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -546,19 +473,28 @@ Provides a practical mechanism to encode user preferences through domain-action
 <div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
 <ul>
 <li><a :data-current="activeSection === '#routing-methods'" class="reference internal" href="#routing-methods">Routing Methods</a><ul>
-<li><a :data-current="activeSection === '#model-based-routing'" class="reference internal" href="#model-based-routing">Model-based Routing</a></li>
-<li><a :data-current="activeSection === '#alias-based-routing'" class="reference internal" href="#alias-based-routing">Alias-based Routing</a></li>
-<li><a :data-current="activeSection === '#preference-aligned-routing-arch-router'" class="reference internal" href="#preference-aligned-routing-arch-router">Preference-aligned Routing (Arch-Router)</a></li>
+<li><a :data-current="activeSection === '#model-based-routing'" class="reference internal" href="#model-based-routing">Model-based routing</a><ul>
+<li><a :data-current="activeSection === '#configuration'" class="reference internal" href="#configuration">Configuration</a></li>
+<li><a :data-current="activeSection === '#client-usage'" class="reference internal" href="#client-usage">Client usage</a></li>
 </ul>
 </li>
-<li><a :data-current="activeSection === '#model-based-routing-workflow'" class="reference internal" href="#model-based-routing-workflow">Model-based Routing Workflow</a></li>
-<li><a :data-current="activeSection === '#alias-based-routing-workflow'" class="reference internal" href="#alias-based-routing-workflow">Alias-based Routing Workflow</a></li>
-<li><a :data-current="activeSection === '#preference-aligned-routing-workflow-arch-router'" class="reference internal" href="#preference-aligned-routing-workflow-arch-router">Preference-aligned Routing Workflow (Arch-Router)</a></li>
-<li><a :data-current="activeSection === '#id1'" class="reference internal" href="#id1">Arch-Router</a></li>
-<li><a :data-current="activeSection === '#implementing-routing'" class="reference internal" href="#implementing-routing">Implementing Routing</a></li>
+<li><a :data-current="activeSection === '#alias-based-routing'" class="reference internal" href="#alias-based-routing">Alias-based routing</a><ul>
+<li><a :data-current="activeSection === '#id3'" class="reference internal" href="#id3">Configuration</a></li>
+<li><a :data-current="activeSection === '#id4'" class="reference internal" href="#id4">Client usage</a></li>
+</ul>
+</li>
+<li><a :data-current="activeSection === '#preference-aligned-routing-arch-router'" class="reference internal" href="#preference-aligned-routing-arch-router">Preference-aligned routing (Arch-Router)</a><ul>
+<li><a :data-current="activeSection === '#id5'" class="reference internal" href="#id5">Configuration</a></li>
+<li><a :data-current="activeSection === '#id6'" class="reference internal" href="#id6">Client usage</a></li>
+</ul>
+</li>
+</ul>
+</li>
+<li><a :data-current="activeSection === '#id7'" class="reference internal" href="#id7">Arch-Router</a></li>
 <li><a :data-current="activeSection === '#combining-routing-methods'" class="reference internal" href="#combining-routing-methods">Combining Routing Methods</a></li>
 <li><a :data-current="activeSection === '#example-use-cases'" class="reference internal" href="#example-use-cases">Example Use Cases</a></li>
-<li><a :data-current="activeSection === '#best-practicesm'" class="reference internal" href="#best-practicesm">Best practicesm</a></li>
+<li><a :data-current="activeSection === '#best-practices'" class="reference internal" href="#best-practices">Best practices</a></li>
+<li><a :data-current="activeSection === '#unsupported-features'" class="reference internal" href="#unsupported-features">Unsupported Features</a></li>
 </ul>
 </div>
 </aside>
@@ -567,12 +503,12 @@ Provides a practical mechanism to encode user preferences through domain-action
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../_static/doctools.js?v=9bcbadda"></script>
 <script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
diff --git a/guides/observability/access_logging.html b/guides/observability/access_logging.html
index 452e359d..f54924d4 100755
--- a/guides/observability/access_logging.html
+++ b/guides/observability/access_logging.html
@@ -7,18 +7,18 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Access Logging | Arch Docs v0.3.22</title>
-<meta content="Access Logging | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Access Logging | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Access Logging | Plano Docs v0.4</title>
+<meta content="Access Logging | Plano Docs v0.4" property="og:title"/>
+<meta content="Access Logging | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/guides/observability/access_logging.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
 <link href="../../search.html" rel="search" title="Search"/>
-<link href="../../build_with_arch/agent.html" rel="next" title="Agentic Apps"/>
+<link href="../prompt_guard.html" rel="next" title="Guardrails"/>
 <link href="monitoring.html" rel="prev" title="Monitoring"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul class="current">
-<li class="toctree-l1"><a class="reference internal" href="../prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../function_calling.html">Function Calling</a></li>
 <li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="monitoring.html">Monitoring</a></li>
 <li class="toctree-l2 current"><a class="current reference internal" href="#">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -164,14 +159,14 @@
 <div id="content" role="main">
 <section id="access-logging">
 <span id="arch-access-logging"></span><h1>Access Logging<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#access-logging"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Access logging in Arch refers to the logging of detailed information about each request and response that flows through Arch.
-It provides visibility into the traffic passing through Arch, which is crucial for monitoring, debugging, and analyzing the
+<p>Access logging in Plano refers to the logging of detailed information about each request and response that flows through Plano.
+It provides visibility into the traffic passing through Plano, which is crucial for monitoring, debugging, and analyzing the
 behavior of AI applications and their interactions.</p>
 <section id="key-features">
 <h2>Key Features<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#key-features" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#key-features'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <ul class="simple">
 <li><p><strong>Per-Request Logging</strong>:
-Each request that passes through Arch is logged. This includes important metadata such as HTTP method,
+Each request that passes through Plano is logged. This includes important metadata such as HTTP method,
 path, response status code, request duration, upstream host, and more.</p></li>
 <li><p><strong>Integration with Monitoring Tools</strong>:
 Access logs can be exported to centralized logging systems (e.g., ELK stack or Fluentd) or used to feed monitoring and alerting systems.</p></li>
@@ -180,21 +175,21 @@ Access logs can be exported to centralized logging systems (e.g., ELK stack or F
 </section>
 <section id="how-it-works">
 <h2>How It Works<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#how-it-works" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#how-it-works'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch gateway exposes access logs for every call it manages on your behalf. By default these access logs can be found under <code class="docutils literal notranslate"><span class="pre">~/archgw_logs</span></code>. For example:</p>
-<div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="gp">$ </span>tail<span class="w"> </span>-F<span class="w"> </span>~/archgw_logs/access_*.log
+<p>Plano exposes access logs for every call it manages on your behalf. By default these access logs can be found under <code class="docutils literal notranslate"><span class="pre">~/plano_logs</span></code>. For example:</p>
+<div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="gp">$ </span>tail<span class="w"> </span>-F<span class="w"> </span>~/plano_logs/access_*.log
 </span><span id="line-2">
-</span><span id="line-3"><span class="go">==&gt; /Users/adilhafeez/archgw_logs/access_llm.log &lt;==</span>
+</span><span id="line-3"><span class="go">==&gt; /Users/username/plano_logs/access_llm.log &lt;==</span>
 </span><span id="line-4"><span class="go">[2024-10-10T03:55:49.537Z] "POST /v1/chat/completions HTTP/1.1" 0 DC 0 0 770 - "-" "OpenAI/Python 1.51.0" "469793af-b25f-9b57-b265-f376e8d8c586" "api.openai.com" "162.159.140.245:443"</span>
 </span><span id="line-5">
-</span><span id="line-6"><span class="go">==&gt; /Users/adilhafeez/archgw_logs/access_internal.log &lt;==</span>
+</span><span id="line-6"><span class="go">==&gt; /Users/username/plano_logs/access_internal.log &lt;==</span>
 </span><span id="line-7"><span class="go">[2024-10-10T03:56:03.906Z] "POST /embeddings HTTP/1.1" 200 - 52 21797 54 53 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "model_server" "192.168.65.254:51000"</span>
 </span><span id="line-8"><span class="go">[2024-10-10T03:56:03.961Z] "POST /zeroshot HTTP/1.1" 200 - 106 218 87 87 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "model_server" "192.168.65.254:51000"</span>
 </span><span id="line-9"><span class="go">[2024-10-10T03:56:04.050Z] "POST /v1/chat/completions HTTP/1.1" 200 - 1301 614 441 441 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "model_server" "192.168.65.254:51000"</span>
 </span><span id="line-10"><span class="go">[2024-10-10T03:56:04.492Z] "POST /hallucination HTTP/1.1" 200 - 556 127 104 104 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "model_server" "192.168.65.254:51000"</span>
 </span><span id="line-11"><span class="go">[2024-10-10T03:56:04.598Z] "POST /insurance_claim_details HTTP/1.1" 200 - 447 125 17 17 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "api_server" "192.168.65.254:18083"</span>
 </span><span id="line-12">
-</span><span id="line-13"><span class="go">==&gt; /Users/adilhafeez/archgw_logs/access_ingress.log &lt;==</span>
-</span><span id="line-14"><span class="go">[2024-10-10T03:56:03.905Z] "POST /v1/chat/completions HTTP/1.1" 200 - 463 1022 1695 984 "-" "OpenAI/Python 1.51.0" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "arch_llm_listener" "0.0.0.0:12000"</span>
+</span><span id="line-13"><span class="go">==&gt; /Users/username/plano_logs/access_ingress.log &lt;==</span>
+</span><span id="line-14"><span class="go">[2024-10-10T03:56:03.905Z] "POST /v1/chat/completions HTTP/1.1" 200 - 463 1022 1695 984 "-" "OpenAI/Python 1.51.0" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "plano_llm_listener" "0.0.0.0:12000"</span>
 </span></code></pre></div>
 </div>
 </section>
@@ -212,7 +207,7 @@ Access logs can be exported to centralized logging systems (e.g., ELK stack or F
 <li><p>DURATION: The total time taken to process the request.</p></li>
 </ul>
 <p>For example for following request:</p>
-<div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="go">[2024-10-10T03:56:03.905Z] "POST /v1/chat/completions HTTP/1.1" 200 - 463 1022 1695 984 "-" "OpenAI/Python 1.51.0" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "arch_llm_listener" "0.0.0.0:12000"</span>
+<div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="go">[2024-10-10T03:56:03.905Z] "POST /v1/chat/completions HTTP/1.1" 200 - 463 1022 1695 984 "-" "OpenAI/Python 1.51.0" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "plano_llm_listener" "0.0.0.0:12000"</span>
 </span></code></pre></div>
 </div>
 <p>Total duration was 1695ms, and the upstream service took 984ms to process the request. Bytes received and sent were 463 and 1022 respectively.</p>
@@ -228,8 +223,8 @@ Access logs can be exported to centralized logging systems (e.g., ELK stack or F
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../../build_with_arch/agent.html">
-        Agentic Apps
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../prompt_guard.html">
+        Guardrails
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -249,12 +244,12 @@ Access logs can be exported to centralized logging systems (e.g., ELK stack or F
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/guides/observability/monitoring.html b/guides/observability/monitoring.html
index f7372f18..0a316442 100755
--- a/guides/observability/monitoring.html
+++ b/guides/observability/monitoring.html
@@ -7,13 +7,13 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Monitoring | Arch Docs v0.3.22</title>
-<meta content="Monitoring | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Monitoring | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Monitoring | Plano Docs v0.4</title>
+<meta content="Monitoring | Plano Docs v0.4" property="og:title"/>
+<meta content="Monitoring | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/guides/observability/monitoring.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul class="current">
-<li class="toctree-l1"><a class="reference internal" href="../prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../function_calling.html">Function Calling</a></li>
 <li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="tracing.html">Tracing</a></li>
 <li class="toctree-l2 current"><a class="current reference internal" href="#">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -168,11 +163,11 @@
 and instrumentation for generating, collecting, processing, and exporting telemetry data, such as traces,
 metrics, and logs. Its flexible design supports a wide range of backends and seamlessly integrates with
 modern application tools.</p>
-<p>Arch acts a <em>source</em> for several monitoring metrics related to <strong>prompts</strong> and <strong>LLMs</strong> natively integrated
+<p>Plano acts a <em>source</em> for several monitoring metrics related to <strong>agents</strong> and <strong>LLMs</strong> natively integrated
 via <a class="reference external" href="https://opentelemetry.io/" rel="nofollow noopener">OpenTelemetry<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> to help you understand three critical aspects of your application:
 latency, token usage, and error rates by an upstream LLM provider. Latency measures the speed at which your application
 is responding to users, which includes metrics like time to first token (TFT), time per output token (TOT) metrics, and
-the total latency as perceived by users. Below are some screenshots how Arch integrates natively with tools like
+the total latency as perceived by users. Below are some screenshots how Plano integrates natively with tools like
 <a class="reference external" href="https://grafana.com/grafana/dashboards/" rel="nofollow noopener">Grafana<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> via <a class="reference external" href="https://prometheus.io/" rel="nofollow noopener">Promethus<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a></p>
 <section id="metrics-dashboard-via-grafana">
 <h2>Metrics Dashboard (via Grafana)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#metrics-dashboard-via-grafana" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#metrics-dashboard-via-grafana'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
@@ -185,7 +180,7 @@ the total latency as perceived by users. Below are some screenshots how Arch int
 </section>
 <section id="configure-monitoring">
 <h2>Configure Monitoring<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configure-monitoring" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configure-monitoring'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch gateway publishes stats endpoint at <a class="reference external" href="http://localhost:19901/stats" rel="nofollow noopener">http://localhost:19901/stats<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>. As noted above, Arch is a source for metrics. To view and manipulate dashbaords, you will
+<p>Plano publishes stats endpoint at <a class="reference external" href="http://localhost:19901/stats" rel="nofollow noopener">http://localhost:19901/stats<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>. As noted above, Plano is a source for metrics. To view and manipulate dashbaords, you will
 need to configiure <a class="reference external" href="https://prometheus.io/" rel="nofollow noopener">Promethus<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> (as a metrics store) and <a class="reference external" href="https://grafana.com/grafana/dashboards/" rel="nofollow noopener">Grafana<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> for dashboards. Below
 are some sample configuration files for both, respectively.</p>
 <div class="literal-block-wrapper docutils container" id="id5">
@@ -202,7 +197,7 @@ are some sample configuration files for both, respectively.</p>
 </span><span id="line-10"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">10s</span>
 </span><span id="line-11"><span class="w">    </span><span class="nt">api_version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v2</span>
 </span><span id="line-12"><span class="nt">scrape_configs</span><span class="p">:</span>
-</span><span id="line-13"><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">job_name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">archgw</span>
+</span><span id="line-13"><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">job_name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">plano</span>
 </span><span id="line-14"><span class="w">    </span><span class="l l-Scalar l-Scalar-Plain">honor_timestamps</span><span class="p p-Indicator">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
 </span><span id="line-15"><span class="w">    </span><span class="l l-Scalar l-Scalar-Plain">scrape_interval</span><span class="p p-Indicator">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">15s</span>
 </span><span id="line-16"><span class="w">    </span><span class="l l-Scalar l-Scalar-Plain">scrape_timeout</span><span class="p p-Indicator">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">10s</span>
@@ -261,12 +256,12 @@ are some sample configuration files for both, respectively.</p>
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/guides/observability/observability.html b/guides/observability/observability.html
index b3715f0b..33cca860 100755
--- a/guides/observability/observability.html
+++ b/guides/observability/observability.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Observability | Arch Docs v0.3.22</title>
-<meta content="Observability | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Observability | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Observability | Plano Docs v0.4</title>
+<meta content="Observability | Plano Docs v0.4" property="og:title"/>
+<meta content="Observability | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/guides/observability/observability.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
 <link href="../../search.html" rel="search" title="Search"/>
 <link href="tracing.html" rel="next" title="Tracing"/>
-<link href="../llm_router.html" rel="prev" title="LLM Routing"/>
+<link href="../function_calling.html" rel="prev" title="Function Calling"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul class="current">
-<li class="toctree-l1"><a class="reference internal" href="../prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../function_calling.html">Function Calling</a></li>
 <li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="current reference internal expandable" href="#">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -192,11 +187,11 @@
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../llm_router.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../function_calling.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        LLM Routing
+        Function Calling
       </a>
 </div>
 <div class="ml-auto">
@@ -213,12 +208,12 @@
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/guides/observability/tracing.html b/guides/observability/tracing.html
index b72f69f4..78a2db77 100755
--- a/guides/observability/tracing.html
+++ b/guides/observability/tracing.html
@@ -7,13 +7,13 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Tracing | Arch Docs v0.3.22</title>
-<meta content="Tracing | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Tracing | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Tracing | Plano Docs v0.4</title>
+<meta content="Tracing | Plano Docs v0.4" property="og:title"/>
+<meta content="Tracing | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/guides/observability/tracing.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul class="current">
-<li class="toctree-l1"><a class="reference internal" href="../prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../function_calling.html">Function Calling</a></li>
 <li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
 <li class="toctree-l2 current"><a class="current reference internal" href="#">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -175,9 +170,9 @@ modern application tools. A key feature of OpenTelemetry is its commitment to st
 requests in an AI application. With tracing, you can capture a detailed view of how requests propagate
 through various services and components, which is crucial for <strong>debugging</strong>, <strong>performance optimization</strong>,
 and understanding complex AI agent architectures like Co-pilots.</p>
-<p><strong>Arch</strong> propagates trace context using the W3C Trace Context standard, specifically through the
+<p><strong>Plano</strong> propagates trace context using the W3C Trace Context standard, specifically through the
 <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header. This allows each component in the system to record its part of the request
-flow, enabling <strong>end-to-end tracing</strong> across the entire application. By using OpenTelemetry, Arch ensures
+flow, enabling <strong>end-to-end tracing</strong> across the entire application. By using OpenTelemetry, Plano ensures
 that developers can capture this trace data consistently and in a format compatible with various observability
 tools.</p>
 <a class="reference internal image-reference" href="../../_images/tracing.png"><img alt="../../_images/tracing.png" class="align-center" src="../../_images/tracing.png" style="width: 100%;"/>
@@ -197,8 +192,8 @@ making it easy to visualize traces in the tools you’re already usi</p></li>
 <section id="how-to-initiate-a-trace">
 <h2>How to Initiate A Trace<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#how-to-initiate-a-trace" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#how-to-initiate-a-trace'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <ol class="arabic simple">
-<li><p><strong>Enable Tracing Configuration</strong>: Simply add the <code class="docutils literal notranslate"><span class="pre">random_sampling</span></code> in <code class="docutils literal notranslate"><span class="pre">tracing</span></code> section to 100`` flag to in the <a class="reference internal" href="../../concepts/tech_overview/listener.html#arch-overview-listeners"><span class="std std-ref">listener</span></a> config</p></li>
-<li><p><strong>Trace Context Propagation</strong>: Arch automatically propagates the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header. When a request is received, Arch will:</p>
+<li><p><strong>Enable Tracing Configuration</strong>: Simply add the <code class="docutils literal notranslate"><span class="pre">random_sampling</span></code> in <code class="docutils literal notranslate"><span class="pre">tracing</span></code> section to 100`` flag to in the <a class="reference internal" href="../../concepts/listeners.html#plano-overview-listeners"><span class="std std-ref">listener</span></a> config</p></li>
+<li><p><strong>Trace Context Propagation</strong>: Plano automatically propagates the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header. When a request is received, Plano will:</p>
 <ul class="simple">
 <li><p>Generate a new <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header if one is not present.</p></li>
 <li><p>Extract the trace context from the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header if it exists.</p></li>
@@ -212,7 +207,7 @@ You can adjust this value from 0-100.</p></li>
 </section>
 <section id="trace-propagation">
 <h2>Trace Propagation<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#trace-propagation" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#trace-propagation'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch uses the W3C Trace Context standard for trace propagation, which relies on the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header.
+<p>Plano uses the W3C Trace Context standard for trace propagation, which relies on the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header.
 This header carries tracing information in a standardized format, enabling interoperability between different
 tracing systems.</p>
 <section id="header-format">
@@ -231,7 +226,7 @@ tracing systems.</p>
 <section id="instrumentation">
 <h3>Instrumentation<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#instrumentation" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#instrumentation'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <p>To integrate AI tracing, your application needs to follow a few simple steps. The steps
-below are very common practice, and not unique to Arch, when you reading tracing headers and export
+below are very common practice, and not unique to Plano, when you reading tracing headers and export
 <a class="reference external" href="https://docs.lightstep.com/docs/understand-distributed-tracing" rel="nofollow noopener">spans<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> for distributed tracing.</p>
 <ul class="simple">
 <li><p>Read the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header from incoming requests.</p></li>
@@ -296,94 +291,6 @@ below are very common practice, and not unique to Arch, when you reading tracing
 </div>
 </section>
 </section>
-<section id="ai-agent-tracing-visualization-example">
-<h3>AI Agent Tracing Visualization Example<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#ai-agent-tracing-visualization-example" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#ai-agent-tracing-visualization-example'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>The following is an example of tracing for an AI-powered customer support system.
-A customer interacts with AI agents, which forward their requests through different
-specialized services and external systems.</p>
-<div class="highlight-default notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="o">+--------------------------+</span>
-</span><span id="line-2"><span class="o">|</span>   <span class="n">Customer</span> <span class="n">Interaction</span>   <span class="o">|</span>
-</span><span id="line-3"><span class="o">+--------------------------+</span>
-</span><span id="line-4">           <span class="o">|</span>
-</span><span id="line-5">           <span class="n">v</span>
-</span><span id="line-6"><span class="o">+--------------------------+</span>        <span class="o">+--------------------------+</span>
-</span><span id="line-7"><span class="o">|</span>  <span class="n">Agent</span> <span class="mi">1</span> <span class="p">(</span><span class="n">Main</span> <span class="o">-</span> <span class="n">Arch</span><span class="p">)</span>   <span class="o">|</span> <span class="o">----&gt;</span>  <span class="o">|</span> <span class="n">External</span> <span class="n">Payment</span> <span class="n">Service</span> <span class="o">|</span>
-</span><span id="line-8"><span class="o">+--------------------------+</span>        <span class="o">+--------------------------+</span>
-</span><span id="line-9">           <span class="o">|</span>                                  <span class="o">|</span>
-</span><span id="line-10">           <span class="n">v</span>                                  <span class="n">v</span>
-</span><span id="line-11"><span class="o">+--------------------------+</span>        <span class="o">+--------------------------+</span>
-</span><span id="line-12"><span class="o">|</span>  <span class="n">Agent</span> <span class="mi">2</span> <span class="p">(</span><span class="n">Support</span> <span class="o">-</span> <span class="n">Arch</span><span class="p">)</span><span class="o">|</span> <span class="o">----&gt;</span>  <span class="o">|</span>   <span class="n">Internal</span> <span class="n">Tech</span> <span class="n">Support</span>  <span class="o">|</span>
-</span><span id="line-13"><span class="o">+--------------------------+</span>        <span class="o">+--------------------------+</span>
-</span><span id="line-14">           <span class="o">|</span>                                  <span class="o">|</span>
-</span><span id="line-15">           <span class="n">v</span>                                  <span class="n">v</span>
-</span><span id="line-16"><span class="o">+--------------------------+</span>        <span class="o">+--------------------------+</span>
-</span><span id="line-17"><span class="o">|</span> <span class="n">Agent</span> <span class="mi">3</span> <span class="p">(</span><span class="n">Orders</span><span class="o">-</span> <span class="n">Arch</span><span class="p">)</span>   <span class="o">|</span> <span class="o">----&gt;</span>  <span class="o">|</span>   <span class="n">Inventory</span> <span class="n">Management</span>   <span class="o">|</span>
-</span><span id="line-18"><span class="o">+--------------------------+</span>        <span class="o">+--------------------------+</span>
-</span></code></pre></div>
-</div>
-<section id="trace-breakdown">
-<h4>Trace Breakdown:<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#trace-breakdown"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h4>
-<ul class="simple">
-<li><dl class="simple">
-<dt>Customer Interaction:</dt><dd><ul>
-<li><p>Span 1: Customer initiates a request via the AI-powered chatbot for billing support (e.g., asking for payment details).</p></li>
-</ul>
-</dd>
-</dl>
-</li>
-<li><dl class="simple">
-<dt>AI Agent 1 (Main - Arch):</dt><dd><ul>
-<li><p>Span 2: AI Agent 1 (Main) processes the request and identifies it as related to billing, forwarding the request
-to an external payment service.</p></li>
-<li><p>Span 3: AI Agent 1 determines that additional technical support is needed for processing and forwards the request
-to AI Agent 2.</p></li>
-</ul>
-</dd>
-</dl>
-</li>
-<li><dl class="simple">
-<dt>External Payment Service:</dt><dd><ul>
-<li><p>Span 4: The external payment service processes the payment-related request (e.g., verifying payment status) and sends
-the response back to AI Agent 1.</p></li>
-</ul>
-</dd>
-</dl>
-</li>
-<li><dl class="simple">
-<dt>AI Agent 2 (Tech - Arch):</dt><dd><ul>
-<li><p>Span 5: AI Agent 2, responsible for technical queries, processes a request forwarded from AI Agent 1 (e.g., checking for
-any account issues).</p></li>
-<li><p>Span 6: AI Agent 2 forwards the query to Internal Tech Support for further investigation.</p></li>
-</ul>
-</dd>
-</dl>
-</li>
-<li><dl class="simple">
-<dt>Internal Tech Support:</dt><dd><ul>
-<li><p>Span 7: Internal Tech Support processes the request (e.g., resolving account access issues) and responds to AI Agent 2.</p></li>
-</ul>
-</dd>
-</dl>
-</li>
-<li><dl class="simple">
-<dt>AI Agent 3 (Orders - Arch):</dt><dd><ul>
-<li><p>Span 8: AI Agent 3 handles order-related queries. AI Agent 1 forwards the request to AI Agent 3 after payment verification
-is completed.</p></li>
-<li><p>Span 9: AI Agent 3 forwards a request to the Inventory Management system to confirm product availability for a pending order.</p></li>
-</ul>
-</dd>
-</dl>
-</li>
-<li><dl class="simple">
-<dt>Inventory Management:</dt><dd><ul>
-<li><p>Span 10: The Inventory Management system checks stock and availability and returns the information to AI Agent 3.</p></li>
-</ul>
-</dd>
-</dl>
-</li>
-</ul>
-</section>
-</section>
 </section>
 <section id="integrating-with-tracing-tools">
 <h2>Integrating with Tracing Tools<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#integrating-with-tracing-tools" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#integrating-with-tracing-tools'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
@@ -461,10 +368,10 @@ is completed.</p></li>
 </section>
 <section id="langtrace">
 <h3>Langtrace<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#langtrace" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#langtrace'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Langtrace is an observability tool designed specifically for large language models (LLMs). It helps you capture, analyze, and understand how LLMs are used in your applications including those built using Arch.</p>
+<p>Langtrace is an observability tool designed specifically for large language models (LLMs). It helps you capture, analyze, and understand how LLMs are used in your applications including those built using Plano.</p>
 <p>To send tracing data to <a class="reference external" href="https://docs.langtrace.ai/supported-integrations/llm-tools/arch" rel="nofollow noopener">Langtrace<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>:</p>
 <ol class="arabic">
-<li><p><strong>Configure Arch</strong>: Make sure Arch is installed and setup correctly. For more information, refer to the <a class="reference external" href="https://github.com/katanemo/archgw?tab=readme-ov-file#prerequisites" rel="nofollow noopener">installation guide<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.</p></li>
+<li><p><strong>Configure Plano</strong>: Make sure Plano is installed and setup correctly. For more information, refer to the <a class="reference external" href="https://github.com/katanemo/archgw?tab=readme-ov-file#prerequisites" rel="nofollow noopener">installation guide<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.</p></li>
 <li><p><strong>Install Langtrace</strong>: Install the Langtrace SDK.:</p>
 <div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="gp">$ </span>pip<span class="w"> </span>install<span class="w"> </span>langtrace-python-sdk
 </span></code></pre></div>
@@ -512,7 +419,7 @@ is completed.</p></li>
 </section>
 <section id="summary">
 <h2>Summary<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#summary" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#summary'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>By leveraging the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header for trace context propagation, Arch enables developers to implement
+<p>By leveraging the <code class="docutils literal notranslate"><span class="pre">traceparent</span></code> header for trace context propagation, Plano enables developers to implement
 tracing efficiently. This approach simplifies the process of collecting and analyzing tracing data in common
 tools like AWS X-Ray and Datadog, enhancing observability and facilitating faster debugging and optimization.</p>
 </section>
@@ -560,10 +467,6 @@ tools like AWS X-Ray and Datadog, enhancing observability and facilitating faste
 <li><a :data-current="activeSection === '#example-with-opentelemetry-in-python'" class="reference internal" href="#example-with-opentelemetry-in-python">Example with OpenTelemetry in Python</a></li>
 </ul>
 </li>
-<li><a :data-current="activeSection === '#ai-agent-tracing-visualization-example'" class="reference internal" href="#ai-agent-tracing-visualization-example">AI Agent Tracing Visualization Example</a><ul>
-<li><a :data-current="activeSection === '#trace-breakdown'" class="reference internal" href="#trace-breakdown">Trace Breakdown:</a></li>
-</ul>
-</li>
 </ul>
 </li>
 <li><a :data-current="activeSection === '#integrating-with-tracing-tools'" class="reference internal" href="#integrating-with-tracing-tools">Integrating with Tracing Tools</a><ul>
@@ -583,12 +486,12 @@ tools like AWS X-Ray and Datadog, enhancing observability and facilitating faste
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/guides/orchestration.html b/guides/orchestration.html
new file mode 100755
index 00000000..da45cb1b
--- /dev/null
+++ b/guides/orchestration.html
@@ -0,0 +1,939 @@
+<!DOCTYPE html>
+
+<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
+<head>
+<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
+<meta charset="utf-8"/>
+<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
+<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
+<meta content="width=device-width, initial-scale=1" name="viewport"/>
+<title>Orchestration | Plano Docs v0.4</title>
+<meta content="Orchestration | Plano Docs v0.4" property="og:title"/>
+<meta content="Orchestration | Plano Docs v0.4" name="twitter:title"/>
+<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
+<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
+<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
+<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
+<link href="./docs/guides/orchestration.html" rel="canonical"/>
+<link href="../_static/favicon.ico" rel="icon"/>
+<link href="../search.html" rel="search" title="Search"/>
+<link href="llm_router.html" rel="next" title="LLM Routing"/>
+<link href="../concepts/prompt_target.html" rel="prev" title="Prompt Target"/>
+<script>
+    <!-- Prevent Flash of wrong theme -->
+      const userPreference = localStorage.getItem('darkMode');
+      let mode;
+      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
+        mode = 'dark';
+        document.documentElement.classList.add('dark');
+      } else {
+        mode = 'light';
+      }
+      if (!userPreference) {localStorage.setItem('darkMode', mode)}
+    </script>
+</head>
+<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
+<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
+      Skip to content
+    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
+<div class="hidden mr-4 md:flex">
+<a class="flex items-center mr-6" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
+<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
+</svg>
+<span class="sr-only">Toggle navigation menu</span>
+</button>
+<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
+<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
+<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
+<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
+<span class="text-xs">⌘</span>
+    K
+  </kbd>
+</form>
+</div>
+<nav class="flex items-center space-x-1">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
+<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
+<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
+</div>
+</a>
+<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
+<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
+</svg>
+<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
+</svg>
+</button>
+</nav>
+</div>
+</div>
+</header>
+<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
+<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a>
+<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
+<div class="overflow-y-auto h-full w-full relative pr-6">
+
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
+<script>
+  window.dataLayer = window.dataLayer || [];
+  function gtag(){dataLayer.push(arguments);}
+  gtag('js', new Date());
+
+  gtag('config', 'G-EH2VW19FXE');
+</script>
+<nav class="table w-full min-w-full my-6 lg:my-8">
+<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/prompt_target.html">Prompt Target</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
+<ul class="current">
+<li class="toctree-l1 current"><a class="current reference internal" href="#">Orchestration</a></li>
+<li class="toctree-l1"><a class="reference internal" href="llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="function_calling.html">Function Calling</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="observability/tracing.html">Tracing</a></li>
+<li class="toctree-l2"><a class="reference internal" href="observability/monitoring.html">Monitoring</a></li>
+<li class="toctree-l2"><a class="reference internal" href="observability/access_logging.html">Access Logging</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="state.html">Conversational State</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
+<ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
+</ul>
+</nav>
+</div>
+</div>
+<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
+<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
+</svg>
+</button>
+</aside>
+<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
+<div class="w-full min-w-0 mx-auto">
+<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
+<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
+<span class="hidden md:inline">Plano Docs v0.4</span>
+<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
+<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
+</svg>
+</a>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Orchestration</span>
+</nav>
+<div id="content" role="main">
+<section id="orchestration">
+<span id="agent-routing"></span><h1>Orchestration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#orchestration"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p>Building multi-agent systems allow you to route requests across multiple specialized agents, each designed to handle specific types of tasks.
+Plano makes it easy to build and scale these systems by managing the orchestration layer—deciding which agent(s) should handle each request—while you focus on implementing individual agent logic.</p>
+<p>This guide shows you how to configure and implement multi-agent orchestration in Plano using a real-world example: a <strong>Travel Booking Assistant</strong> that routes queries to specialized agents for weather and flights.</p>
+<section id="how-it-works">
+<h2>How It Works<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#how-it-works" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#how-it-works'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Plano’s orchestration layer analyzes incoming prompts and routes them to the most appropriate agent based on user intent and conversation context. The workflow is:</p>
+<ol class="arabic simple">
+<li><p><strong>User submits a prompt</strong>: The request arrives at Plano’s agent listener.</p></li>
+<li><p><strong>Agent selection</strong>: Plano uses an LLM to analyze the prompt and determine user intent and complexity. By default, this uses <a class="reference external" href="https://huggingface.co/collections/katanemo/plano-orchestrator" rel="nofollow noopener">Plano-Orchestrator-30B-A3B<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, which offers performance of foundation models at 1/10th the cost. The LLM routes the request to the most suitable agent configured in your system—such as a weather agent or flight agent.</p></li>
+<li><p><strong>Agent handles request</strong>: Once the selected agent receives the request object from Plano, it manages its own <a class="reference internal" href="../concepts/agents.html#agents"><span class="std std-ref">inner loop</span></a> until the task is complete. This means the agent autonomously calls models, invokes tools, processes data, and reasons about next steps—all within its specialized domain—before returning the final response.</p></li>
+<li><p><strong>Seamless handoffs</strong>: For multi-turn conversations, Plano repeats the intent analysis for each follow-up query, enabling smooth handoffs between agents as the conversation evolves.</p></li>
+</ol>
+</section>
+<section id="example-travel-booking-assistant">
+<h2>Example: Travel Booking Assistant<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-travel-booking-assistant" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-travel-booking-assistant'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Let’s walk through a complete multi-agent system: a Travel Booking Assistant that helps users plan trips by providing weather forecasts and flight information. This system uses two specialized agents:</p>
+<ul class="simple">
+<li><p><strong>Weather Agent</strong>: Provides real-time weather conditions and multi-day forecasts</p></li>
+<li><p><strong>Flight Agent</strong>: Searches for flights between airports with real-time tracking</p></li>
+</ul>
+</section>
+<section id="configuration">
+<h2>Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configuration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Configure your agents in the <code class="docutils literal notranslate"><span class="pre">listeners</span></code> section of your <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code>:</p>
+<div class="literal-block-wrapper docutils container" id="id1">
+<div class="code-block-caption"><span class="caption-text">Travel Booking Multi-Agent Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id1"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.3.0</span>
+</span><span id="line-2"><span class="linenos"> 2</span>
+</span><span id="line-3"><span class="linenos"> 3</span><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-4"><span class="linenos"> 4</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">weather_agent</span>
+</span><span id="line-5"><span class="linenos"> 5</span><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10510</span>
+</span><span id="line-6"><span class="linenos"> 6</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">flight_agent</span>
+</span><span id="line-7"><span class="linenos"> 7</span><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10520</span>
+</span><span id="line-8"><span class="linenos"> 8</span>
+</span><span id="line-9"><span class="linenos"> 9</span><span class="nt">model_providers</span><span class="p">:</span>
+</span><span id="line-10"><span class="linenos">10</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
+</span><span id="line-11"><span class="linenos">11</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-12"><span class="linenos">12</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-13"><span class="linenos">13</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o-mini</span>
+</span><span id="line-14"><span class="linenos">14</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span><span class="w"> </span><span class="c1"># smaller, faster, cheaper model for extracting entities like location</span>
+</span><span id="line-15"><span class="linenos">15</span>
+</span><span id="line-16"><span class="linenos">16</span><span class="nt">listeners</span><span class="p">:</span>
+</span><span id="line-17"><span class="linenos">17</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">agent</span>
+</span><span id="line-18"><span class="linenos">18</span><span class="w">    </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">travel_booking_service</span>
+</span><span id="line-19"><span class="linenos">19</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">8001</span>
+</span><span id="line-20"><span class="linenos">20</span><span class="w">    </span><span class="nt">router</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">plano_orchestrator_v1</span>
+</span><span id="line-21"><span class="linenos">21</span><span class="w">    </span><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-22"><span class="linenos">22</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">weather_agent</span>
+</span><span id="line-23"><span class="linenos">23</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
+</span><span id="line-24"><span class="linenos">24</span>
+</span><span id="line-25"><span class="linenos">25</span><span class="w">          </span><span class="no">WeatherAgent is a specialized AI assistant for real-time weather information and forecasts. It provides accurate weather data for any city worldwide using the Open-Meteo API, helping travelers plan their trips with up-to-date weather conditions.</span>
+</span><span id="line-26"><span class="linenos">26</span>
+</span><span id="line-27"><span class="linenos">27</span><span class="w">          </span><span class="no">Capabilities:</span>
+</span><span id="line-28"><span class="linenos">28</span><span class="w">            </span><span class="no">* Get real-time weather conditions and multi-day forecasts for any city worldwide using Open-Meteo API (free, no API key needed)</span>
+</span><span id="line-29"><span class="linenos">29</span><span class="w">            </span><span class="no">* Provides current temperature</span>
+</span><span id="line-30"><span class="linenos">30</span><span class="w">            </span><span class="no">* Provides multi-day forecasts</span>
+</span><span id="line-31"><span class="linenos">31</span><span class="w">            </span><span class="no">* Provides weather conditions</span>
+</span><span id="line-32"><span class="linenos">32</span><span class="w">            </span><span class="no">* Provides sunrise/sunset times</span>
+</span><span id="line-33"><span class="linenos">33</span><span class="w">            </span><span class="no">* Provides detailed weather information</span>
+</span><span id="line-34"><span class="linenos">34</span><span class="w">            </span><span class="no">* Understands conversation context to resolve location references from previous messages</span>
+</span><span id="line-35"><span class="linenos">35</span><span class="w">            </span><span class="no">* Handles weather-related questions including "What's the weather in [city]?", "What's the forecast for [city]?", "How's the weather in [city]?"</span>
+</span><span id="line-36"><span class="linenos">36</span><span class="w">            </span><span class="no">* When queries include both weather and other travel questions (e.g., flights, currency), this agent answers ONLY the weather part</span>
+</span><span id="line-37"><span class="linenos">37</span>
+</span><span id="line-38"><span class="linenos">38</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">flight_agent</span>
+</span><span id="line-39"><span class="linenos">39</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
+</span><span id="line-40"><span class="linenos">40</span>
+</span><span id="line-41"><span class="linenos">41</span><span class="w">          </span><span class="no">FlightAgent is an AI-powered tool specialized in providing live flight information between airports. It leverages the FlightAware AeroAPI to deliver real-time flight status, gate information, and delay updates.</span>
+</span><span id="line-42"><span class="linenos">42</span>
+</span><span id="line-43"><span class="linenos">43</span><span class="w">          </span><span class="no">Capabilities:</span>
+</span><span id="line-44"><span class="linenos">44</span><span class="w">            </span><span class="no">* Get live flight information between airports using FlightAware AeroAPI</span>
+</span><span id="line-45"><span class="linenos">45</span><span class="w">            </span><span class="no">* Shows real-time flight status</span>
+</span><span id="line-46"><span class="linenos">46</span><span class="w">            </span><span class="no">* Shows scheduled/estimated/actual departure and arrival times</span>
+</span><span id="line-47"><span class="linenos">47</span><span class="w">            </span><span class="no">* Shows gate and terminal information</span>
+</span><span id="line-48"><span class="linenos">48</span><span class="w">            </span><span class="no">* Shows delays</span>
+</span><span id="line-49"><span class="linenos">49</span><span class="w">            </span><span class="no">* Shows aircraft type</span>
+</span><span id="line-50"><span class="linenos">50</span><span class="w">            </span><span class="no">* Shows flight status</span>
+</span><span id="line-51"><span class="linenos">51</span><span class="w">            </span><span class="no">* Automatically resolves city names to airport codes (IATA/ICAO)</span>
+</span><span id="line-52"><span class="linenos">52</span><span class="w">            </span><span class="no">* Understands conversation context to infer origin/destination from follow-up questions</span>
+</span><span id="line-53"><span class="linenos">53</span><span class="w">            </span><span class="no">* Handles flight-related questions including "What flights go from [city] to [city]?", "Do flights go to [city]?", "Are there direct flights from [city]?"</span>
+</span><span id="line-54"><span class="linenos">54</span><span class="w">            </span><span class="no">* When queries include both flight and other travel questions (e.g., weather, currency), this agent answers ONLY the flight part</span>
+</span><span id="line-55"><span class="linenos">55</span>
+</span><span id="line-56"><span class="linenos">56</span><span class="nt">tracing</span><span class="p">:</span>
+</span><span id="line-57"><span class="linenos">57</span><span class="w">  </span><span class="nt">random_sampling</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">100</span>
+</span></code></pre></div>
+</div>
+</div>
+<p><strong>Key Configuration Elements:</strong></p>
+<ul class="simple">
+<li><p><strong>agent listener</strong>: A listener of <code class="docutils literal notranslate"><span class="pre">type:</span> <span class="pre">agent</span></code> tells Plano to perform intent analysis and routing for incoming requests.</p></li>
+<li><p><strong>agents list</strong>: Define each agent with an <code class="docutils literal notranslate"><span class="pre">id</span></code>, <code class="docutils literal notranslate"><span class="pre">description</span></code> (used for routing decisions)</p></li>
+<li><p><strong>router</strong>: The <code class="docutils literal notranslate"><span class="pre">plano_orchestrator_v1</span></code> router uses Plano-Orchestrator to analyze user intent and select the appropriate agent.</p></li>
+<li><p><strong>filter_chain</strong>: Optionally attach <a class="reference internal" href="../concepts/filter_chain.html#filter-chain"><span class="std std-ref">filter chains</span></a> to agents for guardrails, query rewriting, or context enrichment.</p></li>
+</ul>
+<p><strong>Writing Effective Agent Descriptions</strong></p>
+<p>Agent descriptions are critical—they’re used by Plano-Orchestrator to make routing decisions. Effective descriptions should include:</p>
+<ul class="simple">
+<li><p><strong>Clear introduction</strong>: A concise statement explaining what the agent is and its primary purpose</p></li>
+<li><p><strong>Capabilities section</strong>: A bulleted list of specific capabilities, including:</p>
+<ul>
+<li><p>What APIs or data sources it uses (e.g., “Open-Meteo API”, “FlightAware AeroAPI”)</p></li>
+<li><p>What information it provides (e.g., “current temperature”, “multi-day forecasts”, “gate information”)</p></li>
+<li><p>How it handles context (e.g., “Understands conversation context to resolve location references”)</p></li>
+<li><p>What question patterns it handles (e.g., “What’s the weather in [city]?”)</p></li>
+<li><p>How it handles multi-part queries (e.g., “When queries include both weather and flights, this agent answers ONLY the weather part”)</p></li>
+</ul>
+</li>
+</ul>
+<p>Here’s an example of a well-structured agent description:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">weather_agent</span>
+</span><span id="line-2"><span class="w">  </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">|</span>
+</span><span id="line-3">
+</span><span id="line-4"><span class="w">    </span><span class="no">WeatherAgent is a specialized AI assistant for real-time weather information</span>
+</span><span id="line-5"><span class="w">    </span><span class="no">and forecasts. It provides accurate weather data for any city worldwide using</span>
+</span><span id="line-6"><span class="w">    </span><span class="no">the Open-Meteo API, helping travelers plan their trips with up-to-date weather</span>
+</span><span id="line-7"><span class="w">    </span><span class="no">conditions.</span>
+</span><span id="line-8">
+</span><span id="line-9"><span class="w">    </span><span class="no">Capabilities:</span>
+</span><span id="line-10"><span class="w">      </span><span class="no">* Get real-time weather conditions and multi-day forecasts for any city worldwide</span>
+</span><span id="line-11"><span class="w">      </span><span class="no">* Provides current temperature, weather conditions, sunrise/sunset times</span>
+</span><span id="line-12"><span class="w">      </span><span class="no">* Provides detailed weather information including multi-day forecasts</span>
+</span><span id="line-13"><span class="w">      </span><span class="no">* Understands conversation context to resolve location references from previous messages</span>
+</span><span id="line-14"><span class="w">      </span><span class="no">* Handles weather-related questions including "What's the weather in [city]?"</span>
+</span><span id="line-15"><span class="w">      </span><span class="no">* When queries include both weather and other travel questions (e.g., flights),</span>
+</span><span id="line-16"><span class="w">        </span><span class="no">this agent answers ONLY the weather part</span>
+</span></code></pre></div>
+</div>
+<div class="admonition note">
+<p class="admonition-title">Note</p>
+<p>We will soon support “Agents as Tools” via Model Context Protocol (MCP), enabling agents to dynamically discover and invoke other agents as tools. Track progress on <a class="reference external" href="https://github.com/katanemo/archgw/issues/646" rel="nofollow noopener">GitHub Issue #646<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.</p>
+</div>
+</section>
+<section id="implementation">
+<h2>Implementation<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#implementation" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#implementation'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Agents are HTTP services that receive routed requests from Plano. Each agent implements the OpenAI Chat Completions API format, making them compatible with standard LLM clients.</p>
+<section id="agent-structure">
+<h3>Agent Structure<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#agent-structure" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#agent-structure'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Let’s examine the Weather Agent implementation:</p>
+<div class="literal-block-wrapper docutils container" id="id2">
+<div class="code-block-caption"><span class="caption-text">Weather Agent - Core Structure</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nd">@app</span><span class="o">.</span><span class="n">post</span><span class="p">(</span><span class="s2">"/v1/chat/completions"</span><span class="p">)</span>
+</span><span id="line-2"><span class="linenos"> 2</span><span class="k">async</span> <span class="k">def</span><span class="w"> </span><span class="nf">handle_request</span><span class="p">(</span><span class="n">request</span><span class="p">:</span> <span class="n">Request</span><span class="p">):</span>
+</span><span id="line-3"><span class="linenos"> 3</span><span class="w">    </span><span class="sd">"""HTTP endpoint for chat completions with streaming support."""</span>
+</span><span id="line-4"><span class="linenos"> 4</span>
+</span><span id="line-5"><span class="linenos"> 5</span>    <span class="n">request_body</span> <span class="o">=</span> <span class="k">await</span> <span class="n">request</span><span class="o">.</span><span class="n">json</span><span class="p">()</span>
+</span><span id="line-6"><span class="linenos"> 6</span>    <span class="n">messages</span> <span class="o">=</span> <span class="n">request_body</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"messages"</span><span class="p">,</span> <span class="p">[])</span>
+</span><span id="line-7"><span class="linenos"> 7</span>    <span class="n">logger</span><span class="o">.</span><span class="n">info</span><span class="p">(</span>
+</span><span id="line-8"><span class="linenos"> 8</span>        <span class="s2">"messages detail json dumps: </span><span class="si">%s</span><span class="s2">"</span><span class="p">,</span>
+</span><span id="line-9"><span class="linenos"> 9</span>        <span class="n">json</span><span class="o">.</span><span class="n">dumps</span><span class="p">(</span><span class="n">messages</span><span class="p">,</span> <span class="n">indent</span><span class="o">=</span><span class="mi">2</span><span class="p">),</span>
+</span><span id="line-10"><span class="linenos">10</span>    <span class="p">)</span>
+</span><span id="line-11"><span class="linenos">11</span>
+</span><span id="line-12"><span class="linenos">12</span>    <span class="n">traceparent_header</span> <span class="o">=</span> <span class="n">request</span><span class="o">.</span><span class="n">headers</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"traceparent"</span><span class="p">)</span>
+</span><span id="line-13"><span class="linenos">13</span>    <span class="k">return</span> <span class="n">StreamingResponse</span><span class="p">(</span>
+</span><span id="line-14"><span class="linenos">14</span>        <span class="n">invoke_weather_agent</span><span class="p">(</span><span class="n">request</span><span class="p">,</span> <span class="n">request_body</span><span class="p">,</span> <span class="n">traceparent_header</span><span class="p">),</span>
+</span><span id="line-15"><span class="linenos">15</span>        <span class="n">media_type</span><span class="o">=</span><span class="s2">"text/plain"</span><span class="p">,</span>
+</span><span id="line-16"><span class="linenos">16</span>        <span class="n">headers</span><span class="o">=</span><span class="p">{</span>
+</span><span id="line-17"><span class="linenos">17</span>            <span class="s2">"content-type"</span><span class="p">:</span> <span class="s2">"text/event-stream"</span><span class="p">,</span>
+</span><span id="line-18"><span class="linenos">18</span>        <span class="p">},</span>
+</span><span id="line-19"><span class="linenos">19</span>    <span class="p">)</span>
+</span><span id="line-20"><span class="linenos">20</span>
+</span><span id="line-21"><span class="linenos">21</span>
+</span><span id="line-22"><span class="linenos">22</span><span class="k">async</span> <span class="k">def</span><span class="w"> </span><span class="nf">invoke_weather_agent</span><span class="p">(</span>
+</span></code></pre></div>
+</div>
+</div>
+<p><strong>Key Points:</strong></p>
+<ul class="simple">
+<li><p>Agents expose a <code class="docutils literal notranslate"><span class="pre">/v1/chat/completions</span></code> endpoint that matches OpenAI’s API format</p></li>
+<li><p>They use Plano’s LLM gateway (via <code class="docutils literal notranslate"><span class="pre">LLM_GATEWAY_ENDPOINT</span></code>) for all LLM calls</p></li>
+<li><p>They receive the full conversation history in <code class="docutils literal notranslate"><span class="pre">request_body.messages</span></code></p></li>
+</ul>
+</section>
+<section id="information-extraction-with-llms">
+<h3>Information Extraction with LLMs<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#information-extraction-with-llms" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#information-extraction-with-llms'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Agents use LLMs to extract structured information from natural language queries. This enables them to understand user intent and extract parameters needed for API calls.</p>
+<p>The Weather Agent extracts location information:</p>
+<div class="literal-block-wrapper docutils container" id="id3">
+<div class="code-block-caption"><span class="caption-text">Weather Agent - Location Extraction</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id3"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span>
+</span><span id="line-2"><span class="linenos"> 2</span>    <span class="n">instructions</span> <span class="o">=</span> <span class="s2">"""Extract the location for WEATHER queries. Return just the city name.</span>
+</span><span id="line-3"><span class="linenos"> 3</span>
+</span><span id="line-4"><span class="linenos"> 4</span><span class="s2">            Rules:</span>
+</span><span id="line-5"><span class="linenos"> 5</span><span class="s2">            1. For multi-part queries, extract ONLY the location mentioned with weather keywords ("weather in [location]")</span>
+</span><span id="line-6"><span class="linenos"> 6</span><span class="s2">            2. If user says "there" or "that city", it typically refers to the DESTINATION city in travel contexts (not the origin)</span>
+</span><span id="line-7"><span class="linenos"> 7</span><span class="s2">            3. For flight queries with weather, "there" means the destination city where they're traveling TO</span>
+</span><span id="line-8"><span class="linenos"> 8</span><span class="s2">            4. Return plain text (e.g., "London", "New York", "Paris, France")</span>
+</span><span id="line-9"><span class="linenos"> 9</span><span class="s2">            5. If no weather location found, return "NOT_FOUND"</span>
+</span><span id="line-10"><span class="linenos">10</span>
+</span><span id="line-11"><span class="linenos">11</span><span class="s2">            Examples:</span>
+</span><span id="line-12"><span class="linenos">12</span><span class="s2">            - "What's the weather in London?" -&gt; "London"</span>
+</span><span id="line-13"><span class="linenos">13</span><span class="s2">            - "Flights from Seattle to Atlanta, and show me the weather there" -&gt; "Atlanta"</span>
+</span><span id="line-14"><span class="linenos">14</span><span class="s2">            - "Can you get me flights from Seattle to Atlanta tomorrow, and also please show me the weather there" -&gt; "Atlanta"</span>
+</span><span id="line-15"><span class="linenos">15</span><span class="s2">            - "What's the weather in Seattle, and what is one flight that goes direct to Atlanta?" -&gt; "Seattle"</span>
+</span><span id="line-16"><span class="linenos">16</span><span class="s2">            - User asked about flights to Atlanta, then "what's the weather like there?" -&gt; "Atlanta"</span>
+</span><span id="line-17"><span class="linenos">17</span><span class="s2">            - "I'm going to Seattle" -&gt; "Seattle"</span>
+</span><span id="line-18"><span class="linenos">18</span><span class="s2">            - "What's happening?" -&gt; "NOT_FOUND"</span>
+</span><span id="line-19"><span class="linenos">19</span>
+</span><span id="line-20"><span class="linenos">20</span><span class="s2">            Extract location:"""</span>
+</span><span id="line-21"><span class="linenos">21</span>
+</span><span id="line-22"><span class="linenos">22</span>    <span class="k">try</span><span class="p">:</span>
+</span><span id="line-23"><span class="linenos">23</span>        <span class="n">user_messages</span> <span class="o">=</span> <span class="p">[</span>
+</span><span id="line-24"><span class="linenos">24</span>            <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"content"</span><span class="p">)</span> <span class="k">for</span> <span class="n">msg</span> <span class="ow">in</span> <span class="n">messages</span> <span class="k">if</span> <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"role"</span><span class="p">)</span> <span class="o">==</span> <span class="s2">"user"</span>
+</span><span id="line-25"><span class="linenos">25</span>        <span class="p">]</span>
+</span><span id="line-26"><span class="linenos">26</span>
+</span><span id="line-27"><span class="linenos">27</span>        <span class="k">if</span> <span class="ow">not</span> <span class="n">user_messages</span><span class="p">:</span>
+</span><span id="line-28"><span class="linenos">28</span>            <span class="n">location</span> <span class="o">=</span> <span class="s2">"New York"</span>
+</span><span id="line-29"><span class="linenos">29</span>        <span class="k">else</span><span class="p">:</span>
+</span><span id="line-30"><span class="linenos">30</span>            <span class="n">ctx</span> <span class="o">=</span> <span class="n">extract</span><span class="p">(</span><span class="n">request</span><span class="o">.</span><span class="n">headers</span><span class="p">)</span>
+</span><span id="line-31"><span class="linenos">31</span>            <span class="n">extra_headers</span> <span class="o">=</span> <span class="p">{}</span>
+</span><span id="line-32"><span class="linenos">32</span>            <span class="n">inject</span><span class="p">(</span><span class="n">extra_headers</span><span class="p">,</span> <span class="n">context</span><span class="o">=</span><span class="n">ctx</span><span class="p">)</span>
+</span><span id="line-33"><span class="linenos">33</span>
+</span><span id="line-34"><span class="linenos">34</span>            <span class="c1"># For location extraction, pass full conversation for context (e.g., "there" = previous destination)</span>
+</span><span id="line-35"><span class="linenos">35</span>            <span class="n">response</span> <span class="o">=</span> <span class="k">await</span> <span class="n">openai_client_via_plano</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-36"><span class="linenos">36</span>                <span class="n">model</span><span class="o">=</span><span class="n">LOCATION_MODEL</span><span class="p">,</span>
+</span><span id="line-37"><span class="linenos">37</span>                <span class="n">messages</span><span class="o">=</span><span class="p">[</span>
+</span><span id="line-38"><span class="linenos">38</span>                    <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"system"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="n">instructions</span><span class="p">},</span>
+</span><span id="line-39"><span class="linenos">39</span>                    <span class="o">*</span><span class="p">[</span>
+</span><span id="line-40"><span class="linenos">40</span>                        <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"role"</span><span class="p">),</span> <span class="s2">"content"</span><span class="p">:</span> <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"content"</span><span class="p">)}</span>
+</span><span id="line-41"><span class="linenos">41</span>                        <span class="k">for</span> <span class="n">msg</span> <span class="ow">in</span> <span class="n">messages</span>
+</span><span id="line-42"><span class="linenos">42</span>                    <span class="p">],</span>
+</span><span id="line-43"><span class="linenos">43</span>                <span class="p">],</span>
+</span><span id="line-44"><span class="linenos">44</span>                <span class="n">temperature</span><span class="o">=</span><span class="mf">0.1</span><span class="p">,</span>
+</span><span id="line-45"><span class="linenos">45</span>                <span class="n">max_tokens</span><span class="o">=</span><span class="mi">50</span><span class="p">,</span>
+</span><span id="line-46"><span class="linenos">46</span>                <span class="n">extra_headers</span><span class="o">=</span><span class="n">extra_headers</span> <span class="k">if</span> <span class="n">extra_headers</span> <span class="k">else</span> <span class="kc">None</span><span class="p">,</span>
+</span><span id="line-47"><span class="linenos">47</span>            <span class="p">)</span>
+</span></code></pre></div>
+</div>
+</div>
+<p>The Flight Agent extracts more complex information—origin, destination, and dates:</p>
+<div class="literal-block-wrapper docutils container" id="id4">
+<div class="code-block-caption"><span class="caption-text">Flight Agent - Flight Information Extraction</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id4"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="k">async</span> <span class="k">def</span><span class="w"> </span><span class="nf">extract_flight_route</span><span class="p">(</span><span class="n">messages</span><span class="p">:</span> <span class="nb">list</span><span class="p">,</span> <span class="n">request</span><span class="p">:</span> <span class="n">Request</span><span class="p">)</span> <span class="o">-&gt;</span> <span class="nb">dict</span><span class="p">:</span>
+</span><span id="line-2"><span class="linenos"> 2</span><span class="w">    </span><span class="sd">"""Extract origin, destination, and date from conversation using LLM."""</span>
+</span><span id="line-3"><span class="linenos"> 3</span>
+</span><span id="line-4"><span class="linenos"> 4</span>    <span class="n">extraction_prompt</span> <span class="o">=</span> <span class="s2">"""Extract flight origin, destination cities, and travel date from the conversation.</span>
+</span><span id="line-5"><span class="linenos"> 5</span>
+</span><span id="line-6"><span class="linenos"> 6</span><span class="s2">    Rules:</span>
+</span><span id="line-7"><span class="linenos"> 7</span><span class="s2">    1. Look for patterns: "flight from X to Y", "flights to Y", "fly from X"</span>
+</span><span id="line-8"><span class="linenos"> 8</span><span class="s2">    2. Extract dates like "tomorrow", "next week", "December 25", "12/25", "on Monday"</span>
+</span><span id="line-9"><span class="linenos"> 9</span><span class="s2">    3. Use conversation context to fill in missing details</span>
+</span><span id="line-10"><span class="linenos">10</span><span class="s2">    4. Return JSON: {"origin": "City" or null, "destination": "City" or null, "date": "YYYY-MM-DD" or null}</span>
+</span><span id="line-11"><span class="linenos">11</span>
+</span><span id="line-12"><span class="linenos">12</span><span class="s2">    Examples:</span>
+</span><span id="line-13"><span class="linenos">13</span><span class="s2">    - "Flight from Seattle to Atlanta tomorrow" -&gt; {"origin": "Seattle", "destination": "Atlanta", "date": "2025-12-24"}</span>
+</span><span id="line-14"><span class="linenos">14</span><span class="s2">    - "What flights go to New York?" -&gt; {"origin": null, "destination": "New York", "date": null}</span>
+</span><span id="line-15"><span class="linenos">15</span><span class="s2">    - "Flights to Miami on Christmas" -&gt; {"origin": null, "destination": "Miami", "date": "2025-12-25"}</span>
+</span><span id="line-16"><span class="linenos">16</span><span class="s2">    - "Show me flights from LA to NYC next Monday" -&gt; {"origin": "LA", "destination": "NYC", "date": "2025-12-30"}</span>
+</span><span id="line-17"><span class="linenos">17</span>
+</span><span id="line-18"><span class="linenos">18</span><span class="s2">    Today is December 23, 2025. Extract flight route and date:"""</span>
+</span><span id="line-19"><span class="linenos">19</span>
+</span><span id="line-20"><span class="linenos">20</span>    <span class="k">try</span><span class="p">:</span>
+</span><span id="line-21"><span class="linenos">21</span>        <span class="n">ctx</span> <span class="o">=</span> <span class="n">extract</span><span class="p">(</span><span class="n">request</span><span class="o">.</span><span class="n">headers</span><span class="p">)</span>
+</span><span id="line-22"><span class="linenos">22</span>        <span class="n">extra_headers</span> <span class="o">=</span> <span class="p">{}</span>
+</span><span id="line-23"><span class="linenos">23</span>        <span class="n">inject</span><span class="p">(</span><span class="n">extra_headers</span><span class="p">,</span> <span class="n">context</span><span class="o">=</span><span class="n">ctx</span><span class="p">)</span>
+</span><span id="line-24"><span class="linenos">24</span>
+</span><span id="line-25"><span class="linenos">25</span>        <span class="n">response</span> <span class="o">=</span> <span class="k">await</span> <span class="n">openai_client_via_plano</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-26"><span class="linenos">26</span>            <span class="n">model</span><span class="o">=</span><span class="n">EXTRACTION_MODEL</span><span class="p">,</span>
+</span><span id="line-27"><span class="linenos">27</span>            <span class="n">messages</span><span class="o">=</span><span class="p">[</span>
+</span><span id="line-28"><span class="linenos">28</span>                <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"system"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="n">extraction_prompt</span><span class="p">},</span>
+</span><span id="line-29"><span class="linenos">29</span>                <span class="o">*</span><span class="p">[</span>
+</span><span id="line-30"><span class="linenos">30</span>                    <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"role"</span><span class="p">),</span> <span class="s2">"content"</span><span class="p">:</span> <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"content"</span><span class="p">)}</span>
+</span><span id="line-31"><span class="linenos">31</span>                    <span class="k">for</span> <span class="n">msg</span> <span class="ow">in</span> <span class="n">messages</span><span class="p">[</span><span class="o">-</span><span class="mi">5</span><span class="p">:]</span>
+</span><span id="line-32"><span class="linenos">32</span>                <span class="p">],</span>
+</span><span id="line-33"><span class="linenos">33</span>            <span class="p">],</span>
+</span><span id="line-34"><span class="linenos">34</span>            <span class="n">temperature</span><span class="o">=</span><span class="mf">0.1</span><span class="p">,</span>
+</span><span id="line-35"><span class="linenos">35</span>            <span class="n">max_tokens</span><span class="o">=</span><span class="mi">100</span><span class="p">,</span>
+</span><span id="line-36"><span class="linenos">36</span>            <span class="n">extra_headers</span><span class="o">=</span><span class="n">extra_headers</span> <span class="k">if</span> <span class="n">extra_headers</span> <span class="k">else</span> <span class="kc">None</span><span class="p">,</span>
+</span><span id="line-37"><span class="linenos">37</span>        <span class="p">)</span>
+</span><span id="line-38"><span class="linenos">38</span>
+</span><span id="line-39"><span class="linenos">39</span>        <span class="n">result</span> <span class="o">=</span> <span class="n">response</span><span class="o">.</span><span class="n">choices</span><span class="p">[</span><span class="mi">0</span><span class="p">]</span><span class="o">.</span><span class="n">message</span><span class="o">.</span><span class="n">content</span><span class="o">.</span><span class="n">strip</span><span class="p">()</span>
+</span><span id="line-40"><span class="linenos">40</span>        <span class="k">if</span> <span class="s2">"```json"</span> <span class="ow">in</span> <span class="n">result</span><span class="p">:</span>
+</span><span id="line-41"><span class="linenos">41</span>            <span class="n">result</span> <span class="o">=</span> <span class="n">result</span><span class="o">.</span><span class="n">split</span><span class="p">(</span><span class="s2">"```json"</span><span class="p">)[</span><span class="mi">1</span><span class="p">]</span><span class="o">.</span><span class="n">split</span><span class="p">(</span><span class="s2">"```"</span><span class="p">)[</span><span class="mi">0</span><span class="p">]</span><span class="o">.</span><span class="n">strip</span><span class="p">()</span>
+</span><span id="line-42"><span class="linenos">42</span>        <span class="k">elif</span> <span class="s2">"```"</span> <span class="ow">in</span> <span class="n">result</span><span class="p">:</span>
+</span><span id="line-43"><span class="linenos">43</span>            <span class="n">result</span> <span class="o">=</span> <span class="n">result</span><span class="o">.</span><span class="n">split</span><span class="p">(</span><span class="s2">"```"</span><span class="p">)[</span><span class="mi">1</span><span class="p">]</span><span class="o">.</span><span class="n">split</span><span class="p">(</span><span class="s2">"```"</span><span class="p">)[</span><span class="mi">0</span><span class="p">]</span><span class="o">.</span><span class="n">strip</span><span class="p">()</span>
+</span><span id="line-44"><span class="linenos">44</span>
+</span><span id="line-45"><span class="linenos">45</span>        <span class="n">route</span> <span class="o">=</span> <span class="n">json</span><span class="o">.</span><span class="n">loads</span><span class="p">(</span><span class="n">result</span><span class="p">)</span>
+</span><span id="line-46"><span class="linenos">46</span>        <span class="k">return</span> <span class="p">{</span>
+</span><span id="line-47"><span class="linenos">47</span>            <span class="s2">"origin"</span><span class="p">:</span> <span class="n">route</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"origin"</span><span class="p">),</span>
+</span><span id="line-48"><span class="linenos">48</span>            <span class="s2">"destination"</span><span class="p">:</span> <span class="n">route</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"destination"</span><span class="p">),</span>
+</span><span id="line-49"><span class="linenos">49</span>            <span class="s2">"date"</span><span class="p">:</span> <span class="n">route</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"date"</span><span class="p">),</span>
+</span><span id="line-50"><span class="linenos">50</span>        <span class="p">}</span>
+</span><span id="line-51"><span class="linenos">51</span>    <span class="k">except</span> <span class="ne">Exception</span> <span class="k">as</span> <span class="n">e</span><span class="p">:</span>
+</span><span id="line-52"><span class="linenos">52</span>        <span class="n">logger</span><span class="o">.</span><span class="n">error</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Error extracting flight route: </span><span class="si">{</span><span class="n">e</span><span class="si">}</span><span class="s2">"</span><span class="p">)</span>
+</span></code></pre></div>
+</div>
+</div>
+<p><strong>Key Points:</strong></p>
+<ul class="simple">
+<li><p>Use smaller, faster models (like <code class="docutils literal notranslate"><span class="pre">gpt-4o-mini</span></code>) for extraction tasks</p></li>
+<li><p>Include conversation context to handle follow-up questions and pronouns</p></li>
+<li><p>Use structured prompts with clear output formats (JSON)</p></li>
+<li><p>Handle edge cases with fallback values</p></li>
+</ul>
+</section>
+<section id="calling-external-apis">
+<h3>Calling External APIs<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#calling-external-apis" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#calling-external-apis'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>After extracting information, agents call external APIs to fetch real-time data:</p>
+<div class="literal-block-wrapper docutils container" id="id5">
+<div class="code-block-caption"><span class="caption-text">Weather Agent - External API Call</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id5"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span>        <span class="c1"># Geocode city to get coordinates</span>
+</span><span id="line-2"><span class="linenos"> 2</span>        <span class="n">geocode_url</span> <span class="o">=</span> <span class="sa">f</span><span class="s2">"https://geocoding-api.open-meteo.com/v1/search?name=</span><span class="si">{</span><span class="n">quote</span><span class="p">(</span><span class="n">location</span><span class="p">)</span><span class="si">}</span><span class="s2">&amp;count=1&amp;language=en&amp;format=json"</span>
+</span><span id="line-3"><span class="linenos"> 3</span>        <span class="n">geocode_response</span> <span class="o">=</span> <span class="k">await</span> <span class="n">http_client</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="n">geocode_url</span><span class="p">)</span>
+</span><span id="line-4"><span class="linenos"> 4</span>
+</span><span id="line-5"><span class="linenos"> 5</span>        <span class="k">if</span> <span class="n">geocode_response</span><span class="o">.</span><span class="n">status_code</span> <span class="o">!=</span> <span class="mi">200</span> <span class="ow">or</span> <span class="ow">not</span> <span class="n">geocode_response</span><span class="o">.</span><span class="n">json</span><span class="p">()</span><span class="o">.</span><span class="n">get</span><span class="p">(</span>
+</span><span id="line-6"><span class="linenos"> 6</span>            <span class="s2">"results"</span>
+</span><span id="line-7"><span class="linenos"> 7</span>        <span class="p">):</span>
+</span><span id="line-8"><span class="linenos"> 8</span>            <span class="n">logger</span><span class="o">.</span><span class="n">warning</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Could not geocode </span><span class="si">{</span><span class="n">location</span><span class="si">}</span><span class="s2">, using New York"</span><span class="p">)</span>
+</span><span id="line-9"><span class="linenos"> 9</span>            <span class="n">location</span> <span class="o">=</span> <span class="s2">"New York"</span>
+</span><span id="line-10"><span class="linenos">10</span>            <span class="n">geocode_url</span> <span class="o">=</span> <span class="sa">f</span><span class="s2">"https://geocoding-api.open-meteo.com/v1/search?name=</span><span class="si">{</span><span class="n">quote</span><span class="p">(</span><span class="n">location</span><span class="p">)</span><span class="si">}</span><span class="s2">&amp;count=1&amp;language=en&amp;format=json"</span>
+</span><span id="line-11"><span class="linenos">11</span>            <span class="n">geocode_response</span> <span class="o">=</span> <span class="k">await</span> <span class="n">http_client</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="n">geocode_url</span><span class="p">)</span>
+</span><span id="line-12"><span class="linenos">12</span>
+</span><span id="line-13"><span class="linenos">13</span>        <span class="n">geocode_data</span> <span class="o">=</span> <span class="n">geocode_response</span><span class="o">.</span><span class="n">json</span><span class="p">()</span>
+</span><span id="line-14"><span class="linenos">14</span>        <span class="k">if</span> <span class="ow">not</span> <span class="n">geocode_data</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"results"</span><span class="p">):</span>
+</span><span id="line-15"><span class="linenos">15</span>            <span class="k">return</span> <span class="p">{</span>
+</span><span id="line-16"><span class="linenos">16</span>                <span class="s2">"location"</span><span class="p">:</span> <span class="n">location</span><span class="p">,</span>
+</span><span id="line-17"><span class="linenos">17</span>                <span class="s2">"weather"</span><span class="p">:</span> <span class="p">{</span>
+</span><span id="line-18"><span class="linenos">18</span>                    <span class="s2">"date"</span><span class="p">:</span> <span class="n">datetime</span><span class="o">.</span><span class="n">now</span><span class="p">()</span><span class="o">.</span><span class="n">strftime</span><span class="p">(</span><span class="s2">"%Y-%m-</span><span class="si">%d</span><span class="s2">"</span><span class="p">),</span>
+</span><span id="line-19"><span class="linenos">19</span>                    <span class="s2">"day_name"</span><span class="p">:</span> <span class="n">datetime</span><span class="o">.</span><span class="n">now</span><span class="p">()</span><span class="o">.</span><span class="n">strftime</span><span class="p">(</span><span class="s2">"%A"</span><span class="p">),</span>
+</span><span id="line-20"><span class="linenos">20</span>                    <span class="s2">"temperature_c"</span><span class="p">:</span> <span class="kc">None</span><span class="p">,</span>
+</span><span id="line-21"><span class="linenos">21</span>                    <span class="s2">"temperature_f"</span><span class="p">:</span> <span class="kc">None</span><span class="p">,</span>
+</span><span id="line-22"><span class="linenos">22</span>                    <span class="s2">"weather_code"</span><span class="p">:</span> <span class="kc">None</span><span class="p">,</span>
+</span><span id="line-23"><span class="linenos">23</span>                    <span class="s2">"error"</span><span class="p">:</span> <span class="s2">"Could not retrieve weather data"</span><span class="p">,</span>
+</span><span id="line-24"><span class="linenos">24</span>                <span class="p">},</span>
+</span><span id="line-25"><span class="linenos">25</span>            <span class="p">}</span>
+</span><span id="line-26"><span class="linenos">26</span>
+</span><span id="line-27"><span class="linenos">27</span>        <span class="n">result</span> <span class="o">=</span> <span class="n">geocode_data</span><span class="p">[</span><span class="s2">"results"</span><span class="p">][</span><span class="mi">0</span><span class="p">]</span>
+</span><span id="line-28"><span class="linenos">28</span>        <span class="n">location_name</span> <span class="o">=</span> <span class="n">result</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"name"</span><span class="p">,</span> <span class="n">location</span><span class="p">)</span>
+</span><span id="line-29"><span class="linenos">29</span>        <span class="n">latitude</span> <span class="o">=</span> <span class="n">result</span><span class="p">[</span><span class="s2">"latitude"</span><span class="p">]</span>
+</span><span id="line-30"><span class="linenos">30</span>        <span class="n">longitude</span> <span class="o">=</span> <span class="n">result</span><span class="p">[</span><span class="s2">"longitude"</span><span class="p">]</span>
+</span><span id="line-31"><span class="linenos">31</span>
+</span><span id="line-32"><span class="linenos">32</span>        <span class="n">logger</span><span class="o">.</span><span class="n">info</span><span class="p">(</span>
+</span><span id="line-33"><span class="linenos">33</span>            <span class="sa">f</span><span class="s2">"Geocoded '</span><span class="si">{</span><span class="n">location</span><span class="si">}</span><span class="s2">' to </span><span class="si">{</span><span class="n">location_name</span><span class="si">}</span><span class="s2"> (</span><span class="si">{</span><span class="n">latitude</span><span class="si">}</span><span class="s2">, </span><span class="si">{</span><span class="n">longitude</span><span class="si">}</span><span class="s2">)"</span>
+</span><span id="line-34"><span class="linenos">34</span>        <span class="p">)</span>
+</span><span id="line-35"><span class="linenos">35</span>
+</span><span id="line-36"><span class="linenos">36</span>        <span class="c1"># Get weather forecast</span>
+</span><span id="line-37"><span class="linenos">37</span>        <span class="n">weather_url</span> <span class="o">=</span> <span class="p">(</span>
+</span><span id="line-38"><span class="linenos">38</span>            <span class="sa">f</span><span class="s2">"https://api.open-meteo.com/v1/forecast?"</span>
+</span><span id="line-39"><span class="linenos">39</span>            <span class="sa">f</span><span class="s2">"latitude=</span><span class="si">{</span><span class="n">latitude</span><span class="si">}</span><span class="s2">&amp;longitude=</span><span class="si">{</span><span class="n">longitude</span><span class="si">}</span><span class="s2">&amp;"</span>
+</span><span id="line-40"><span class="linenos">40</span>            <span class="sa">f</span><span class="s2">"current=temperature_2m&amp;"</span>
+</span><span id="line-41"><span class="linenos">41</span>            <span class="sa">f</span><span class="s2">"daily=sunrise,sunset,temperature_2m_max,temperature_2m_min,weather_code&amp;"</span>
+</span><span id="line-42"><span class="linenos">42</span>            <span class="sa">f</span><span class="s2">"forecast_days=</span><span class="si">{</span><span class="n">days</span><span class="si">}</span><span class="s2">&amp;timezone=auto"</span>
+</span><span id="line-43"><span class="linenos">43</span>        <span class="p">)</span>
+</span><span id="line-44"><span class="linenos">44</span>
+</span><span id="line-45"><span class="linenos">45</span>        <span class="n">weather_response</span> <span class="o">=</span> <span class="k">await</span> <span class="n">http_client</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="n">weather_url</span><span class="p">)</span>
+</span><span id="line-46"><span class="linenos">46</span>        <span class="k">if</span> <span class="n">weather_response</span><span class="o">.</span><span class="n">status_code</span> <span class="o">!=</span> <span class="mi">200</span><span class="p">:</span>
+</span><span id="line-47"><span class="linenos">47</span>            <span class="k">return</span> <span class="p">{</span>
+</span><span id="line-48"><span class="linenos">48</span>                <span class="s2">"location"</span><span class="p">:</span> <span class="n">location_name</span><span class="p">,</span>
+</span><span id="line-49"><span class="linenos">49</span>                <span class="s2">"weather"</span><span class="p">:</span> <span class="p">{</span>
+</span><span id="line-50"><span class="linenos">50</span>                    <span class="s2">"date"</span><span class="p">:</span> <span class="n">datetime</span><span class="o">.</span><span class="n">now</span><span class="p">()</span><span class="o">.</span><span class="n">strftime</span><span class="p">(</span><span class="s2">"%Y-%m-</span><span class="si">%d</span><span class="s2">"</span><span class="p">),</span>
+</span><span id="line-51"><span class="linenos">51</span>                    <span class="s2">"day_name"</span><span class="p">:</span> <span class="n">datetime</span><span class="o">.</span><span class="n">now</span><span class="p">()</span><span class="o">.</span><span class="n">strftime</span><span class="p">(</span><span class="s2">"%A"</span><span class="p">),</span>
+</span><span id="line-52"><span class="linenos">52</span>                    <span class="s2">"temperature_c"</span><span class="p">:</span> <span class="kc">None</span><span class="p">,</span>
+</span><span id="line-53"><span class="linenos">53</span>                    <span class="s2">"temperature_f"</span><span class="p">:</span> <span class="kc">None</span><span class="p">,</span>
+</span><span id="line-54"><span class="linenos">54</span>                    <span class="s2">"weather_code"</span><span class="p">:</span> <span class="kc">None</span><span class="p">,</span>
+</span><span id="line-55"><span class="linenos">55</span>                    <span class="s2">"error"</span><span class="p">:</span> <span class="s2">"Could not retrieve weather data"</span><span class="p">,</span>
+</span><span id="line-56"><span class="linenos">56</span>                <span class="p">},</span>
+</span><span id="line-57"><span class="linenos">57</span>            <span class="p">}</span>
+</span><span id="line-58"><span class="linenos">58</span>
+</span><span id="line-59"><span class="linenos">59</span>        <span class="n">weather_data</span> <span class="o">=</span> <span class="n">weather_response</span><span class="o">.</span><span class="n">json</span><span class="p">()</span>
+</span><span id="line-60"><span class="linenos">60</span>        <span class="n">current_temp</span> <span class="o">=</span> <span class="n">weather_data</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"current"</span><span class="p">,</span> <span class="p">{})</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"temperature_2m"</span><span class="p">)</span>
+</span><span id="line-61"><span class="linenos">61</span>        <span class="n">daily</span> <span class="o">=</span> <span class="n">weather_data</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"daily"</span><span class="p">,</span> <span class="p">{})</span>
+</span><span id="line-62"><span class="linenos">62</span>
+</span></code></pre></div>
+</div>
+</div>
+<p>The Flight Agent calls FlightAware’s AeroAPI:</p>
+<div class="literal-block-wrapper docutils container" id="id6">
+<div class="code-block-caption"><span class="caption-text">Flight Agent - External API Call</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id6"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos">  1</span><span class="k">async</span> <span class="k">def</span><span class="w"> </span><span class="nf">get_flights</span><span class="p">(</span>
+</span><span id="line-2"><span class="linenos">  2</span>    <span class="n">origin_code</span><span class="p">:</span> <span class="nb">str</span><span class="p">,</span> <span class="n">dest_code</span><span class="p">:</span> <span class="nb">str</span><span class="p">,</span> <span class="n">travel_date</span><span class="p">:</span> <span class="n">Optional</span><span class="p">[</span><span class="nb">str</span><span class="p">]</span> <span class="o">=</span> <span class="kc">None</span>
+</span><span id="line-3"><span class="linenos">  3</span><span class="p">)</span> <span class="o">-&gt;</span> <span class="n">Optional</span><span class="p">[</span><span class="nb">dict</span><span class="p">]:</span>
+</span><span id="line-4"><span class="linenos">  4</span><span class="w">    </span><span class="sd">"""Get flights between two airports using FlightAware API.</span>
+</span><span id="line-5"><span class="linenos">  5</span>
+</span><span id="line-6"><span class="linenos">  6</span><span class="sd">    Args:</span>
+</span><span id="line-7"><span class="linenos">  7</span><span class="sd">        origin_code: Origin airport IATA code</span>
+</span><span id="line-8"><span class="linenos">  8</span><span class="sd">        dest_code: Destination airport IATA code</span>
+</span><span id="line-9"><span class="linenos">  9</span><span class="sd">        travel_date: Travel date in YYYY-MM-DD format, defaults to today</span>
+</span><span id="line-10"><span class="linenos"> 10</span>
+</span><span id="line-11"><span class="linenos"> 11</span><span class="sd">    Note: FlightAware API limits searches to 2 days in the future.</span>
+</span><span id="line-12"><span class="linenos"> 12</span><span class="sd">    """</span>
+</span><span id="line-13"><span class="linenos"> 13</span>    <span class="k">try</span><span class="p">:</span>
+</span><span id="line-14"><span class="linenos"> 14</span>        <span class="c1"># Use provided date or default to today</span>
+</span><span id="line-15"><span class="linenos"> 15</span>        <span class="k">if</span> <span class="n">travel_date</span><span class="p">:</span>
+</span><span id="line-16"><span class="linenos"> 16</span>            <span class="n">search_date</span> <span class="o">=</span> <span class="n">travel_date</span>
+</span><span id="line-17"><span class="linenos"> 17</span>        <span class="k">else</span><span class="p">:</span>
+</span><span id="line-18"><span class="linenos"> 18</span>            <span class="n">search_date</span> <span class="o">=</span> <span class="n">datetime</span><span class="o">.</span><span class="n">now</span><span class="p">()</span><span class="o">.</span><span class="n">strftime</span><span class="p">(</span><span class="s2">"%Y-%m-</span><span class="si">%d</span><span class="s2">"</span><span class="p">)</span>
+</span><span id="line-19"><span class="linenos"> 19</span>
+</span><span id="line-20"><span class="linenos"> 20</span>        <span class="c1"># Validate date is not too far in the future (FlightAware limit: 2 days)</span>
+</span><span id="line-21"><span class="linenos"> 21</span>        <span class="n">search_date_obj</span> <span class="o">=</span> <span class="n">datetime</span><span class="o">.</span><span class="n">strptime</span><span class="p">(</span><span class="n">search_date</span><span class="p">,</span> <span class="s2">"%Y-%m-</span><span class="si">%d</span><span class="s2">"</span><span class="p">)</span>
+</span><span id="line-22"><span class="linenos"> 22</span>        <span class="n">today</span> <span class="o">=</span> <span class="n">datetime</span><span class="o">.</span><span class="n">now</span><span class="p">()</span><span class="o">.</span><span class="n">replace</span><span class="p">(</span><span class="n">hour</span><span class="o">=</span><span class="mi">0</span><span class="p">,</span> <span class="n">minute</span><span class="o">=</span><span class="mi">0</span><span class="p">,</span> <span class="n">second</span><span class="o">=</span><span class="mi">0</span><span class="p">,</span> <span class="n">microsecond</span><span class="o">=</span><span class="mi">0</span><span class="p">)</span>
+</span><span id="line-23"><span class="linenos"> 23</span>        <span class="n">days_ahead</span> <span class="o">=</span> <span class="p">(</span><span class="n">search_date_obj</span> <span class="o">-</span> <span class="n">today</span><span class="p">)</span><span class="o">.</span><span class="n">days</span>
+</span><span id="line-24"><span class="linenos"> 24</span>
+</span><span id="line-25"><span class="linenos"> 25</span>        <span class="k">if</span> <span class="n">days_ahead</span> <span class="o">&gt;</span> <span class="mi">2</span><span class="p">:</span>
+</span><span id="line-26"><span class="linenos"> 26</span>            <span class="n">logger</span><span class="o">.</span><span class="n">warning</span><span class="p">(</span>
+</span><span id="line-27"><span class="linenos"> 27</span>                <span class="sa">f</span><span class="s2">"Requested date </span><span class="si">{</span><span class="n">search_date</span><span class="si">}</span><span class="s2"> is </span><span class="si">{</span><span class="n">days_ahead</span><span class="si">}</span><span class="s2"> days ahead, exceeds FlightAware 2-day limit"</span>
+</span><span id="line-28"><span class="linenos"> 28</span>            <span class="p">)</span>
+</span><span id="line-29"><span class="linenos"> 29</span>            <span class="k">return</span> <span class="p">{</span>
+</span><span id="line-30"><span class="linenos"> 30</span>                <span class="s2">"origin_code"</span><span class="p">:</span> <span class="n">origin_code</span><span class="p">,</span>
+</span><span id="line-31"><span class="linenos"> 31</span>                <span class="s2">"destination_code"</span><span class="p">:</span> <span class="n">dest_code</span><span class="p">,</span>
+</span><span id="line-32"><span class="linenos"> 32</span>                <span class="s2">"flights"</span><span class="p">:</span> <span class="p">[],</span>
+</span><span id="line-33"><span class="linenos"> 33</span>                <span class="s2">"count"</span><span class="p">:</span> <span class="mi">0</span><span class="p">,</span>
+</span><span id="line-34"><span class="linenos"> 34</span>                <span class="s2">"error"</span><span class="p">:</span> <span class="sa">f</span><span class="s2">"FlightAware API only provides flight data up to 2 days in the future. The requested date (</span><span class="si">{</span><span class="n">search_date</span><span class="si">}</span><span class="s2">) is </span><span class="si">{</span><span class="n">days_ahead</span><span class="si">}</span><span class="s2"> days ahead. Please search for today, tomorrow, or the day after."</span><span class="p">,</span>
+</span><span id="line-35"><span class="linenos"> 35</span>            <span class="p">}</span>
+</span><span id="line-36"><span class="linenos"> 36</span>
+</span><span id="line-37"><span class="linenos"> 37</span>        <span class="n">url</span> <span class="o">=</span> <span class="sa">f</span><span class="s2">"</span><span class="si">{</span><span class="n">AEROAPI_BASE_URL</span><span class="si">}</span><span class="s2">/airports/</span><span class="si">{</span><span class="n">origin_code</span><span class="si">}</span><span class="s2">/flights/to/</span><span class="si">{</span><span class="n">dest_code</span><span class="si">}</span><span class="s2">"</span>
+</span><span id="line-38"><span class="linenos"> 38</span>        <span class="n">headers</span> <span class="o">=</span> <span class="p">{</span><span class="s2">"x-apikey"</span><span class="p">:</span> <span class="n">AEROAPI_KEY</span><span class="p">}</span>
+</span><span id="line-39"><span class="linenos"> 39</span>        <span class="n">params</span> <span class="o">=</span> <span class="p">{</span>
+</span><span id="line-40"><span class="linenos"> 40</span>            <span class="s2">"start"</span><span class="p">:</span> <span class="sa">f</span><span class="s2">"</span><span class="si">{</span><span class="n">search_date</span><span class="si">}</span><span class="s2">T00:00:00Z"</span><span class="p">,</span>
+</span><span id="line-41"><span class="linenos"> 41</span>            <span class="s2">"end"</span><span class="p">:</span> <span class="sa">f</span><span class="s2">"</span><span class="si">{</span><span class="n">search_date</span><span class="si">}</span><span class="s2">T23:59:59Z"</span><span class="p">,</span>
+</span><span id="line-42"><span class="linenos"> 42</span>            <span class="s2">"connection"</span><span class="p">:</span> <span class="s2">"nonstop"</span><span class="p">,</span>
+</span><span id="line-43"><span class="linenos"> 43</span>            <span class="s2">"max_pages"</span><span class="p">:</span> <span class="mi">1</span><span class="p">,</span>
+</span><span id="line-44"><span class="linenos"> 44</span>        <span class="p">}</span>
+</span><span id="line-45"><span class="linenos"> 45</span>
+</span><span id="line-46"><span class="linenos"> 46</span>        <span class="n">response</span> <span class="o">=</span> <span class="k">await</span> <span class="n">http_client</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="n">url</span><span class="p">,</span> <span class="n">headers</span><span class="o">=</span><span class="n">headers</span><span class="p">,</span> <span class="n">params</span><span class="o">=</span><span class="n">params</span><span class="p">)</span>
+</span><span id="line-47"><span class="linenos"> 47</span>
+</span><span id="line-48"><span class="linenos"> 48</span>        <span class="k">if</span> <span class="n">response</span><span class="o">.</span><span class="n">status_code</span> <span class="o">!=</span> <span class="mi">200</span><span class="p">:</span>
+</span><span id="line-49"><span class="linenos"> 49</span>            <span class="n">logger</span><span class="o">.</span><span class="n">error</span><span class="p">(</span>
+</span><span id="line-50"><span class="linenos"> 50</span>                <span class="sa">f</span><span class="s2">"FlightAware API error </span><span class="si">{</span><span class="n">response</span><span class="o">.</span><span class="n">status_code</span><span class="si">}</span><span class="s2">: </span><span class="si">{</span><span class="n">response</span><span class="o">.</span><span class="n">text</span><span class="si">}</span><span class="s2">"</span>
+</span><span id="line-51"><span class="linenos"> 51</span>            <span class="p">)</span>
+</span><span id="line-52"><span class="linenos"> 52</span>            <span class="k">return</span> <span class="kc">None</span>
+</span><span id="line-53"><span class="linenos"> 53</span>
+</span><span id="line-54"><span class="linenos"> 54</span>        <span class="n">data</span> <span class="o">=</span> <span class="n">response</span><span class="o">.</span><span class="n">json</span><span class="p">()</span>
+</span><span id="line-55"><span class="linenos"> 55</span>        <span class="n">flights</span> <span class="o">=</span> <span class="p">[]</span>
+</span><span id="line-56"><span class="linenos"> 56</span>
+</span><span id="line-57"><span class="linenos"> 57</span>        <span class="c1"># Log raw API response for debugging</span>
+</span><span id="line-58"><span class="linenos"> 58</span>        <span class="n">logger</span><span class="o">.</span><span class="n">info</span><span class="p">(</span><span class="sa">f</span><span class="s2">"FlightAware API returned </span><span class="si">{</span><span class="nb">len</span><span class="p">(</span><span class="n">data</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s1">'flights'</span><span class="p">,</span><span class="w"> </span><span class="p">[]))</span><span class="si">}</span><span class="s2"> flights"</span><span class="p">)</span>
+</span><span id="line-59"><span class="linenos"> 59</span>
+</span><span id="line-60"><span class="linenos"> 60</span>        <span class="k">for</span> <span class="n">idx</span><span class="p">,</span> <span class="n">flight_group</span> <span class="ow">in</span> <span class="nb">enumerate</span><span class="p">(</span>
+</span><span id="line-61"><span class="linenos"> 61</span>            <span class="n">data</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"flights"</span><span class="p">,</span> <span class="p">[])[:</span><span class="mi">5</span><span class="p">]</span>
+</span><span id="line-62"><span class="linenos"> 62</span>        <span class="p">):</span>  <span class="c1"># Limit to 5 flights</span>
+</span><span id="line-63"><span class="linenos"> 63</span>            <span class="c1"># FlightAware API nests data in segments array</span>
+</span><span id="line-64"><span class="linenos"> 64</span>            <span class="n">segments</span> <span class="o">=</span> <span class="n">flight_group</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"segments"</span><span class="p">,</span> <span class="p">[])</span>
+</span><span id="line-65"><span class="linenos"> 65</span>            <span class="k">if</span> <span class="ow">not</span> <span class="n">segments</span><span class="p">:</span>
+</span><span id="line-66"><span class="linenos"> 66</span>                <span class="k">continue</span>
+</span><span id="line-67"><span class="linenos"> 67</span>
+</span><span id="line-68"><span class="linenos"> 68</span>            <span class="n">flight</span> <span class="o">=</span> <span class="n">segments</span><span class="p">[</span><span class="mi">0</span><span class="p">]</span>  <span class="c1"># Get first segment (direct flights only have one)</span>
+</span><span id="line-69"><span class="linenos"> 69</span>
+</span><span id="line-70"><span class="linenos"> 70</span>            <span class="c1"># Extract airport codes from nested objects</span>
+</span><span id="line-71"><span class="linenos"> 71</span>            <span class="n">flight_origin</span> <span class="o">=</span> <span class="kc">None</span>
+</span><span id="line-72"><span class="linenos"> 72</span>            <span class="n">flight_dest</span> <span class="o">=</span> <span class="kc">None</span>
+</span><span id="line-73"><span class="linenos"> 73</span>
+</span><span id="line-74"><span class="linenos"> 74</span>            <span class="k">if</span> <span class="nb">isinstance</span><span class="p">(</span><span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"origin"</span><span class="p">),</span> <span class="nb">dict</span><span class="p">):</span>
+</span><span id="line-75"><span class="linenos"> 75</span>                <span class="n">flight_origin</span> <span class="o">=</span> <span class="n">flight</span><span class="p">[</span><span class="s2">"origin"</span><span class="p">]</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"code_iata"</span><span class="p">)</span>
+</span><span id="line-76"><span class="linenos"> 76</span>
+</span><span id="line-77"><span class="linenos"> 77</span>            <span class="k">if</span> <span class="nb">isinstance</span><span class="p">(</span><span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"destination"</span><span class="p">),</span> <span class="nb">dict</span><span class="p">):</span>
+</span><span id="line-78"><span class="linenos"> 78</span>                <span class="n">flight_dest</span> <span class="o">=</span> <span class="n">flight</span><span class="p">[</span><span class="s2">"destination"</span><span class="p">]</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"code_iata"</span><span class="p">)</span>
+</span><span id="line-79"><span class="linenos"> 79</span>
+</span><span id="line-80"><span class="linenos"> 80</span>            <span class="c1"># Build flight object</span>
+</span><span id="line-81"><span class="linenos"> 81</span>            <span class="n">flights</span><span class="o">.</span><span class="n">append</span><span class="p">(</span>
+</span><span id="line-82"><span class="linenos"> 82</span>                <span class="p">{</span>
+</span><span id="line-83"><span class="linenos"> 83</span>                    <span class="s2">"airline"</span><span class="p">:</span> <span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"operator"</span><span class="p">),</span>
+</span><span id="line-84"><span class="linenos"> 84</span>                    <span class="s2">"flight_number"</span><span class="p">:</span> <span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"ident_iata"</span><span class="p">)</span> <span class="ow">or</span> <span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"ident"</span><span class="p">),</span>
+</span><span id="line-85"><span class="linenos"> 85</span>                    <span class="s2">"departure_time"</span><span class="p">:</span> <span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"scheduled_out"</span><span class="p">),</span>
+</span><span id="line-86"><span class="linenos"> 86</span>                    <span class="s2">"arrival_time"</span><span class="p">:</span> <span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"scheduled_in"</span><span class="p">),</span>
+</span><span id="line-87"><span class="linenos"> 87</span>                    <span class="s2">"origin"</span><span class="p">:</span> <span class="n">flight_origin</span><span class="p">,</span>
+</span><span id="line-88"><span class="linenos"> 88</span>                    <span class="s2">"destination"</span><span class="p">:</span> <span class="n">flight_dest</span><span class="p">,</span>
+</span><span id="line-89"><span class="linenos"> 89</span>                    <span class="s2">"aircraft_type"</span><span class="p">:</span> <span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"aircraft_type"</span><span class="p">),</span>
+</span><span id="line-90"><span class="linenos"> 90</span>                    <span class="s2">"status"</span><span class="p">:</span> <span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"status"</span><span class="p">),</span>
+</span><span id="line-91"><span class="linenos"> 91</span>                    <span class="s2">"terminal_origin"</span><span class="p">:</span> <span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"terminal_origin"</span><span class="p">),</span>
+</span><span id="line-92"><span class="linenos"> 92</span>                    <span class="s2">"gate_origin"</span><span class="p">:</span> <span class="n">flight</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"gate_origin"</span><span class="p">),</span>
+</span><span id="line-93"><span class="linenos"> 93</span>                <span class="p">}</span>
+</span><span id="line-94"><span class="linenos"> 94</span>            <span class="p">)</span>
+</span><span id="line-95"><span class="linenos"> 95</span>
+</span><span id="line-96"><span class="linenos"> 96</span>        <span class="k">return</span> <span class="p">{</span>
+</span><span id="line-97"><span class="linenos"> 97</span>            <span class="s2">"origin_code"</span><span class="p">:</span> <span class="n">origin_code</span><span class="p">,</span>
+</span><span id="line-98"><span class="linenos"> 98</span>            <span class="s2">"destination_code"</span><span class="p">:</span> <span class="n">dest_code</span><span class="p">,</span>
+</span><span id="line-99"><span class="linenos"> 99</span>            <span class="s2">"flights"</span><span class="p">:</span> <span class="n">flights</span><span class="p">,</span>
+</span><span id="line-100"><span class="linenos">100</span>            <span class="s2">"count"</span><span class="p">:</span> <span class="nb">len</span><span class="p">(</span><span class="n">flights</span><span class="p">),</span>
+</span><span id="line-101"><span class="linenos">101</span>        <span class="p">}</span>
+</span><span id="line-102"><span class="linenos">102</span>    <span class="k">except</span> <span class="ne">Exception</span> <span class="k">as</span> <span class="n">e</span><span class="p">:</span>
+</span><span id="line-103"><span class="linenos">103</span>        <span class="n">logger</span><span class="o">.</span><span class="n">error</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Error fetching flights: </span><span class="si">{</span><span class="n">e</span><span class="si">}</span><span class="s2">"</span><span class="p">)</span>
+</span><span id="line-104"><span class="linenos">104</span>        <span class="k">return</span> <span class="kc">None</span>
+</span><span id="line-105"><span class="linenos">105</span>
+</span></code></pre></div>
+</div>
+</div>
+<p><strong>Key Points:</strong></p>
+<ul class="simple">
+<li><p>Use async HTTP clients (like <code class="docutils literal notranslate"><span class="pre">httpx.AsyncClient</span></code>) for non-blocking API calls</p></li>
+<li><p>Transform external API responses into consistent, structured formats</p></li>
+<li><p>Handle errors gracefully with fallback values</p></li>
+<li><p>Cache or validate data when appropriate (e.g., airport code validation)</p></li>
+</ul>
+</section>
+<section id="preparing-context-and-generating-responses">
+<h3>Preparing Context and Generating Responses<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#preparing-context-and-generating-responses" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#preparing-context-and-generating-responses'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Agents combine extracted information, API data, and conversation history to generate responses:</p>
+<div class="literal-block-wrapper docutils container" id="id7">
+<div class="code-block-caption"><span class="caption-text">Weather Agent - Context Preparation and Response Generation</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id7"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span>    <span class="n">last_user_msg</span> <span class="o">=</span> <span class="n">get_last_user_content</span><span class="p">(</span><span class="n">messages</span><span class="p">)</span>
+</span><span id="line-2"><span class="linenos"> 2</span>    <span class="n">days</span> <span class="o">=</span> <span class="mi">1</span>
+</span><span id="line-3"><span class="linenos"> 3</span>
+</span><span id="line-4"><span class="linenos"> 4</span>    <span class="k">if</span> <span class="s2">"forecast"</span> <span class="ow">in</span> <span class="n">last_user_msg</span> <span class="ow">or</span> <span class="s2">"week"</span> <span class="ow">in</span> <span class="n">last_user_msg</span><span class="p">:</span>
+</span><span id="line-5"><span class="linenos"> 5</span>        <span class="n">days</span> <span class="o">=</span> <span class="mi">7</span>
+</span><span id="line-6"><span class="linenos"> 6</span>    <span class="k">elif</span> <span class="s2">"tomorrow"</span> <span class="ow">in</span> <span class="n">last_user_msg</span><span class="p">:</span>
+</span><span id="line-7"><span class="linenos"> 7</span>        <span class="n">days</span> <span class="o">=</span> <span class="mi">2</span>
+</span><span id="line-8"><span class="linenos"> 8</span>
+</span><span id="line-9"><span class="linenos"> 9</span>    <span class="c1"># Extract specific number of days if mentioned (e.g., "5 day forecast")</span>
+</span><span id="line-10"><span class="linenos">10</span>    <span class="kn">import</span><span class="w"> </span><span class="nn">re</span>
+</span><span id="line-11"><span class="linenos">11</span>
+</span><span id="line-12"><span class="linenos">12</span>    <span class="n">day_match</span> <span class="o">=</span> <span class="n">re</span><span class="o">.</span><span class="n">search</span><span class="p">(</span><span class="sa">r</span><span class="s2">"(\d{1,2})\s+day"</span><span class="p">,</span> <span class="n">last_user_msg</span><span class="p">)</span>
+</span><span id="line-13"><span class="linenos">13</span>    <span class="k">if</span> <span class="n">day_match</span><span class="p">:</span>
+</span><span id="line-14"><span class="linenos">14</span>        <span class="n">requested_days</span> <span class="o">=</span> <span class="nb">int</span><span class="p">(</span><span class="n">day_match</span><span class="o">.</span><span class="n">group</span><span class="p">(</span><span class="mi">1</span><span class="p">))</span>
+</span><span id="line-15"><span class="linenos">15</span>        <span class="n">days</span> <span class="o">=</span> <span class="nb">min</span><span class="p">(</span><span class="n">requested_days</span><span class="p">,</span> <span class="mi">16</span><span class="p">)</span>  <span class="c1"># API supports max 16 days</span>
+</span><span id="line-16"><span class="linenos">16</span>
+</span><span id="line-17"><span class="linenos">17</span>    <span class="c1"># Get live weather data (location extraction happens inside this function)</span>
+</span><span id="line-18"><span class="linenos">18</span>    <span class="n">weather_data</span> <span class="o">=</span> <span class="k">await</span> <span class="n">get_weather_data</span><span class="p">(</span><span class="n">request</span><span class="p">,</span> <span class="n">messages</span><span class="p">,</span> <span class="n">days</span><span class="p">)</span>
+</span><span id="line-19"><span class="linenos">19</span>
+</span><span id="line-20"><span class="linenos">20</span>    <span class="c1"># Create weather context to append to user message</span>
+</span><span id="line-21"><span class="linenos">21</span>    <span class="n">forecast_type</span> <span class="o">=</span> <span class="s2">"forecast"</span> <span class="k">if</span> <span class="n">days</span> <span class="o">&gt;</span> <span class="mi">1</span> <span class="k">else</span> <span class="s2">"current weather"</span>
+</span><span id="line-22"><span class="linenos">22</span>    <span class="n">weather_context</span> <span class="o">=</span> <span class="sa">f</span><span class="s2">"""</span>
+</span><span id="line-23"><span class="linenos">23</span>
+</span><span id="line-24"><span class="linenos">24</span><span class="s2">Weather data for </span><span class="si">{</span><span class="n">weather_data</span><span class="p">[</span><span class="s1">'location'</span><span class="p">]</span><span class="si">}</span><span class="s2"> (</span><span class="si">{</span><span class="n">forecast_type</span><span class="si">}</span><span class="s2">):</span>
+</span><span id="line-25"><span class="linenos">25</span><span class="si">{</span><span class="n">json</span><span class="o">.</span><span class="n">dumps</span><span class="p">(</span><span class="n">weather_data</span><span class="p">,</span><span class="w"> </span><span class="n">indent</span><span class="o">=</span><span class="mi">2</span><span class="p">)</span><span class="si">}</span><span class="s2">"""</span>
+</span><span id="line-26"><span class="linenos">26</span>
+</span><span id="line-27"><span class="linenos">27</span>    <span class="c1"># System prompt for weather agent</span>
+</span><span id="line-28"><span class="linenos">28</span>    <span class="n">instructions</span> <span class="o">=</span> <span class="s2">"""You are a weather assistant in a multi-agent system. You will receive weather data in JSON format with these fields:</span>
+</span><span id="line-29"><span class="linenos">29</span>
+</span><span id="line-30"><span class="linenos">30</span><span class="s2">    - "location": City name</span>
+</span><span id="line-31"><span class="linenos">31</span><span class="s2">    - "forecast": Array of weather objects, each with date, day_name, temperature_c, temperature_f, temperature_max_c, temperature_min_c, weather_code, sunrise, sunset</span>
+</span><span id="line-32"><span class="linenos">32</span><span class="s2">    - weather_code: WMO code (0=clear, 1-3=partly cloudy, 45-48=fog, 51-67=rain, 71-86=snow, 95-99=thunderstorm)</span>
+</span><span id="line-33"><span class="linenos">33</span>
+</span><span id="line-34"><span class="linenos">34</span><span class="s2">    Your task:</span>
+</span><span id="line-35"><span class="linenos">35</span><span class="s2">    1. Present the weather/forecast clearly for the location</span>
+</span><span id="line-36"><span class="linenos">36</span><span class="s2">    2. For single day: show current conditions</span>
+</span><span id="line-37"><span class="linenos">37</span><span class="s2">    3. For multi-day: show each day with date and conditions</span>
+</span><span id="line-38"><span class="linenos">38</span><span class="s2">    4. Include temperature in both Celsius and Fahrenheit</span>
+</span><span id="line-39"><span class="linenos">39</span><span class="s2">    5. Describe conditions naturally based on weather_code</span>
+</span><span id="line-40"><span class="linenos">40</span><span class="s2">    6. Use conversational language</span>
+</span><span id="line-41"><span class="linenos">41</span>
+</span><span id="line-42"><span class="linenos">42</span><span class="s2">    Important: If the conversation includes information from other agents (like flight details), acknowledge and build upon that context naturally. Your primary focus is weather, but maintain awareness of the full conversation.</span>
+</span><span id="line-43"><span class="linenos">43</span>
+</span><span id="line-44"><span class="linenos">44</span><span class="s2">    Remember: Only use the provided data. If fields are null, mention data is unavailable."""</span>
+</span><span id="line-45"><span class="linenos">45</span>
+</span><span id="line-46"><span class="linenos">46</span>    <span class="c1"># Build message history with weather data appended to the last user message</span>
+</span><span id="line-47"><span class="linenos">47</span>    <span class="n">response_messages</span> <span class="o">=</span> <span class="p">[{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"system"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="n">instructions</span><span class="p">}]</span>
+</span><span id="line-48"><span class="linenos">48</span>
+</span><span id="line-49"><span class="linenos">49</span>    <span class="k">for</span> <span class="n">i</span><span class="p">,</span> <span class="n">msg</span> <span class="ow">in</span> <span class="nb">enumerate</span><span class="p">(</span><span class="n">messages</span><span class="p">):</span>
+</span><span id="line-50"><span class="linenos">50</span>        <span class="c1"># Append weather data to the last user message</span>
+</span><span id="line-51"><span class="linenos">51</span>        <span class="k">if</span> <span class="n">i</span> <span class="o">==</span> <span class="nb">len</span><span class="p">(</span><span class="n">messages</span><span class="p">)</span> <span class="o">-</span> <span class="mi">1</span> <span class="ow">and</span> <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"role"</span><span class="p">)</span> <span class="o">==</span> <span class="s2">"user"</span><span class="p">:</span>
+</span><span id="line-52"><span class="linenos">52</span>            <span class="n">response_messages</span><span class="o">.</span><span class="n">append</span><span class="p">(</span>
+</span><span id="line-53"><span class="linenos">53</span>                <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="s2">"user"</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"content"</span><span class="p">)</span> <span class="o">+</span> <span class="n">weather_context</span><span class="p">}</span>
+</span><span id="line-54"><span class="linenos">54</span>            <span class="p">)</span>
+</span><span id="line-55"><span class="linenos">55</span>        <span class="k">else</span><span class="p">:</span>
+</span><span id="line-56"><span class="linenos">56</span>            <span class="n">response_messages</span><span class="o">.</span><span class="n">append</span><span class="p">(</span>
+</span><span id="line-57"><span class="linenos">57</span>                <span class="p">{</span><span class="s2">"role"</span><span class="p">:</span> <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"role"</span><span class="p">),</span> <span class="s2">"content"</span><span class="p">:</span> <span class="n">msg</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"content"</span><span class="p">)}</span>
+</span><span id="line-58"><span class="linenos">58</span>            <span class="p">)</span>
+</span><span id="line-59"><span class="linenos">59</span>
+</span><span id="line-60"><span class="linenos">60</span>    <span class="k">try</span><span class="p">:</span>
+</span><span id="line-61"><span class="linenos">61</span>        <span class="n">ctx</span> <span class="o">=</span> <span class="n">extract</span><span class="p">(</span><span class="n">request</span><span class="o">.</span><span class="n">headers</span><span class="p">)</span>
+</span><span id="line-62"><span class="linenos">62</span>        <span class="n">extra_headers</span> <span class="o">=</span> <span class="p">{</span><span class="s2">"x-envoy-max-retries"</span><span class="p">:</span> <span class="s2">"3"</span><span class="p">}</span>
+</span><span id="line-63"><span class="linenos">63</span>        <span class="n">inject</span><span class="p">(</span><span class="n">extra_headers</span><span class="p">,</span> <span class="n">context</span><span class="o">=</span><span class="n">ctx</span><span class="p">)</span>
+</span><span id="line-64"><span class="linenos">64</span>
+</span><span id="line-65"><span class="linenos">65</span>        <span class="n">stream</span> <span class="o">=</span> <span class="k">await</span> <span class="n">openai_client_via_plano</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-66"><span class="linenos">66</span>            <span class="n">model</span><span class="o">=</span><span class="n">WEATHER_MODEL</span><span class="p">,</span>
+</span><span id="line-67"><span class="linenos">67</span>            <span class="n">messages</span><span class="o">=</span><span class="n">response_messages</span><span class="p">,</span>
+</span><span id="line-68"><span class="linenos">68</span>            <span class="n">temperature</span><span class="o">=</span><span class="n">request_body</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"temperature"</span><span class="p">,</span> <span class="mf">0.7</span><span class="p">),</span>
+</span><span id="line-69"><span class="linenos">69</span>            <span class="n">max_tokens</span><span class="o">=</span><span class="n">request_body</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"max_tokens"</span><span class="p">,</span> <span class="mi">1000</span><span class="p">),</span>
+</span><span id="line-70"><span class="linenos">70</span>            <span class="n">stream</span><span class="o">=</span><span class="kc">True</span><span class="p">,</span>
+</span><span id="line-71"><span class="linenos">71</span>            <span class="n">extra_headers</span><span class="o">=</span><span class="n">extra_headers</span><span class="p">,</span>
+</span><span id="line-72"><span class="linenos">72</span>        <span class="p">)</span>
+</span><span id="line-73"><span class="linenos">73</span>
+</span><span id="line-74"><span class="linenos">74</span>        <span class="k">async</span> <span class="k">for</span> <span class="n">chunk</span> <span class="ow">in</span> <span class="n">stream</span><span class="p">:</span>
+</span><span id="line-75"><span class="linenos">75</span>            <span class="k">if</span> <span class="n">chunk</span><span class="o">.</span><span class="n">choices</span><span class="p">:</span>
+</span><span id="line-76"><span class="linenos">76</span>                <span class="k">yield</span> <span class="sa">f</span><span class="s2">"data: </span><span class="si">{</span><span class="n">chunk</span><span class="o">.</span><span class="n">model_dump_json</span><span class="p">()</span><span class="si">}</span><span class="se">\n\n</span><span class="s2">"</span>
+</span><span id="line-77"><span class="linenos">77</span>
+</span><span id="line-78"><span class="linenos">78</span>        <span class="k">yield</span> <span class="s2">"data: [DONE]</span><span class="se">\n\n</span><span class="s2">"</span>
+</span><span id="line-79"><span class="linenos">79</span>
+</span><span id="line-80"><span class="linenos">80</span>    <span class="k">except</span> <span class="ne">Exception</span> <span class="k">as</span> <span class="n">e</span><span class="p">:</span>
+</span><span id="line-81"><span class="linenos">81</span>        <span class="n">logger</span><span class="o">.</span><span class="n">error</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Error generating weather response: </span><span class="si">{</span><span class="n">e</span><span class="si">}</span><span class="s2">"</span><span class="p">)</span>
+</span></code></pre></div>
+</div>
+</div>
+<p><strong>Key Points:</strong></p>
+<ul class="simple">
+<li><p>Use system messages to provide structured data to the LLM</p></li>
+<li><p>Include full conversation history for context-aware responses</p></li>
+<li><p>Stream responses for better user experience</p></li>
+<li><p>Route all LLM calls through Plano’s gateway for consistent behavior and observability</p></li>
+</ul>
+</section>
+</section>
+<section id="best-practices">
+<h2>Best Practices<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#best-practices" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#best-practices'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p><strong>Write Clear Agent Descriptions</strong></p>
+<p>Agent descriptions are used by Plano-Orchestrator to make routing decisions. Be specific about what each agent handles:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Good - specific and actionable</span>
+</span><span id="line-2"><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">flight_agent</span>
+</span><span id="line-3"><span class="w">  </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Get live flight information between airports using FlightAware AeroAPI. Shows real-time flight status, scheduled/estimated/actual departure and arrival times, gate and terminal information, delays, aircraft type, and flight status. Automatically resolves city names to airport codes (IATA/ICAO). Understands conversation context to infer origin/destination from follow-up questions.</span>
+</span><span id="line-4">
+</span><span id="line-5"><span class="c1"># Less ideal - too vague</span>
+</span><span id="line-6"><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">flight_agent</span>
+</span><span id="line-7"><span class="w">  </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Handles flight queries</span>
+</span></code></pre></div>
+</div>
+<p><strong>Use Conversation Context Effectively</strong></p>
+<p>Include conversation history in your extraction and response generation:</p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Include conversation context for extraction</span>
+</span><span id="line-2"><span class="n">conversation_context</span> <span class="o">=</span> <span class="p">[]</span>
+</span><span id="line-3"><span class="k">for</span> <span class="n">msg</span> <span class="ow">in</span> <span class="n">messages</span><span class="p">:</span>
+</span><span id="line-4">    <span class="n">conversation_context</span><span class="o">.</span><span class="n">append</span><span class="p">({</span><span class="s2">"role"</span><span class="p">:</span> <span class="n">msg</span><span class="o">.</span><span class="n">role</span><span class="p">,</span> <span class="s2">"content"</span><span class="p">:</span> <span class="n">msg</span><span class="o">.</span><span class="n">content</span><span class="p">})</span>
+</span><span id="line-5">
+</span><span id="line-6"><span class="c1"># Use recent context (last 10 messages)</span>
+</span><span id="line-7"><span class="n">context_messages</span> <span class="o">=</span> <span class="n">conversation_context</span><span class="p">[</span><span class="o">-</span><span class="mi">10</span><span class="p">:]</span> <span class="k">if</span> <span class="nb">len</span><span class="p">(</span><span class="n">conversation_context</span><span class="p">)</span> <span class="o">&gt;</span> <span class="mi">10</span> <span class="k">else</span> <span class="n">conversation_context</span>
+</span></code></pre></div>
+</div>
+<p><strong>Route LLM Calls Through Plano’s Model Proxy</strong></p>
+<p>Always route LLM calls through Plano’s <a class="reference internal" href="../concepts/llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">Model Proxy</span></a> for consistent responses, smart routing, and rich observability:</p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="n">openai_client_via_plano</span> <span class="o">=</span> <span class="n">AsyncOpenAI</span><span class="p">(</span>
+</span><span id="line-2">    <span class="n">base_url</span><span class="o">=</span><span class="n">LLM_GATEWAY_ENDPOINT</span><span class="p">,</span>  <span class="c1"># Plano's LLM gateway</span>
+</span><span id="line-3">    <span class="n">api_key</span><span class="o">=</span><span class="s2">"EMPTY"</span><span class="p">,</span>
+</span><span id="line-4"><span class="p">)</span>
+</span><span id="line-5">
+</span><span id="line-6"><span class="n">response</span> <span class="o">=</span> <span class="k">await</span> <span class="n">openai_client_via_plano</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-7">    <span class="n">model</span><span class="o">=</span><span class="s2">"openai/gpt-4o"</span><span class="p">,</span>
+</span><span id="line-8">    <span class="n">messages</span><span class="o">=</span><span class="n">messages</span><span class="p">,</span>
+</span><span id="line-9">    <span class="n">stream</span><span class="o">=</span><span class="kc">True</span><span class="p">,</span>
+</span><span id="line-10"><span class="p">)</span>
+</span></code></pre></div>
+</div>
+<p><strong>Handle Errors Gracefully</strong></p>
+<p>Provide fallback values and clear error messages:</p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="k">async</span> <span class="k">def</span><span class="w"> </span><span class="nf">get_weather_data</span><span class="p">(</span><span class="n">request</span><span class="p">:</span> <span class="n">Request</span><span class="p">,</span> <span class="n">messages</span><span class="p">:</span> <span class="nb">list</span><span class="p">,</span> <span class="n">days</span><span class="p">:</span> <span class="nb">int</span> <span class="o">=</span> <span class="mi">1</span><span class="p">):</span>
+</span><span id="line-2">    <span class="k">try</span><span class="p">:</span>
+</span><span id="line-3">        <span class="c1"># ... extraction and API logic ...</span>
+</span><span id="line-4">        <span class="n">location</span> <span class="o">=</span> <span class="n">response</span><span class="o">.</span><span class="n">choices</span><span class="p">[</span><span class="mi">0</span><span class="p">]</span><span class="o">.</span><span class="n">message</span><span class="o">.</span><span class="n">content</span><span class="o">.</span><span class="n">strip</span><span class="p">()</span><span class="o">.</span><span class="n">strip</span><span class="p">(</span><span class="s2">"</span><span class="se">\"</span><span class="s2">'`.,!?"</span><span class="p">)</span>
+</span><span id="line-5">        <span class="k">if</span> <span class="ow">not</span> <span class="n">location</span> <span class="ow">or</span> <span class="n">location</span><span class="o">.</span><span class="n">upper</span><span class="p">()</span> <span class="o">==</span> <span class="s2">"NOT_FOUND"</span><span class="p">:</span>
+</span><span id="line-6">            <span class="n">location</span> <span class="o">=</span> <span class="s2">"New York"</span>  <span class="c1"># Fallback to default</span>
+</span><span id="line-7">        <span class="k">return</span> <span class="n">weather_data</span>
+</span><span id="line-8">    <span class="k">except</span> <span class="ne">Exception</span> <span class="k">as</span> <span class="n">e</span><span class="p">:</span>
+</span><span id="line-9">        <span class="n">logger</span><span class="o">.</span><span class="n">error</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Error getting weather data: </span><span class="si">{</span><span class="n">e</span><span class="si">}</span><span class="s2">"</span><span class="p">)</span>
+</span><span id="line-10">        <span class="k">return</span> <span class="p">{</span><span class="s2">"location"</span><span class="p">:</span> <span class="s2">"New York"</span><span class="p">,</span> <span class="s2">"weather"</span><span class="p">:</span> <span class="p">{</span><span class="s2">"error"</span><span class="p">:</span> <span class="s2">"Could not retrieve weather data"</span><span class="p">}}</span>
+</span></code></pre></div>
+</div>
+<p><strong>Use Appropriate Models for Tasks</strong></p>
+<p>Use smaller, faster models for extraction tasks and larger models for final responses:</p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Extraction: Use smaller, faster model</span>
+</span><span id="line-2"><span class="n">LOCATION_MODEL</span> <span class="o">=</span> <span class="s2">"openai/gpt-4o-mini"</span>
+</span><span id="line-3">
+</span><span id="line-4"><span class="c1"># Final response: Use larger, more capable model</span>
+</span><span id="line-5"><span class="n">WEATHER_MODEL</span> <span class="o">=</span> <span class="s2">"openai/gpt-4o"</span>
+</span></code></pre></div>
+</div>
+<p><strong>Stream Responses</strong></p>
+<p>Stream responses for better user experience:</p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="k">async</span> <span class="k">def</span><span class="w"> </span><span class="nf">invoke_weather_agent</span><span class="p">(</span><span class="n">request</span><span class="p">:</span> <span class="n">Request</span><span class="p">,</span> <span class="n">request_body</span><span class="p">:</span> <span class="nb">dict</span><span class="p">,</span> <span class="n">traceparent_header</span><span class="p">:</span> <span class="nb">str</span> <span class="o">=</span> <span class="kc">None</span><span class="p">):</span>
+</span><span id="line-2">    <span class="c1"># ... prepare messages with weather data ...</span>
+</span><span id="line-3">
+</span><span id="line-4">    <span class="n">stream</span> <span class="o">=</span> <span class="k">await</span> <span class="n">openai_client_via_plano</span><span class="o">.</span><span class="n">chat</span><span class="o">.</span><span class="n">completions</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-5">        <span class="n">model</span><span class="o">=</span><span class="n">WEATHER_MODEL</span><span class="p">,</span>
+</span><span id="line-6">        <span class="n">messages</span><span class="o">=</span><span class="n">response_messages</span><span class="p">,</span>
+</span><span id="line-7">        <span class="n">temperature</span><span class="o">=</span><span class="n">request_body</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"temperature"</span><span class="p">,</span> <span class="mf">0.7</span><span class="p">),</span>
+</span><span id="line-8">        <span class="n">max_tokens</span><span class="o">=</span><span class="n">request_body</span><span class="o">.</span><span class="n">get</span><span class="p">(</span><span class="s2">"max_tokens"</span><span class="p">,</span> <span class="mi">1000</span><span class="p">),</span>
+</span><span id="line-9">        <span class="n">stream</span><span class="o">=</span><span class="kc">True</span><span class="p">,</span>
+</span><span id="line-10">        <span class="n">extra_headers</span><span class="o">=</span><span class="n">extra_headers</span><span class="p">,</span>
+</span><span id="line-11">    <span class="p">)</span>
+</span><span id="line-12">
+</span><span id="line-13">    <span class="k">async</span> <span class="k">for</span> <span class="n">chunk</span> <span class="ow">in</span> <span class="n">stream</span><span class="p">:</span>
+</span><span id="line-14">        <span class="k">if</span> <span class="n">chunk</span><span class="o">.</span><span class="n">choices</span><span class="p">:</span>
+</span><span id="line-15">            <span class="k">yield</span> <span class="sa">f</span><span class="s2">"data: </span><span class="si">{</span><span class="n">chunk</span><span class="o">.</span><span class="n">model_dump_json</span><span class="p">()</span><span class="si">}</span><span class="se">\n\n</span><span class="s2">"</span>
+</span><span id="line-16">
+</span><span id="line-17">    <span class="k">yield</span> <span class="s2">"data: [DONE]</span><span class="se">\n\n</span><span class="s2">"</span>
+</span></code></pre></div>
+</div>
+</section>
+<section id="common-use-cases">
+<h2>Common Use Cases<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#common-use-cases" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#common-use-cases'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Multi-agent orchestration is particularly powerful for:</p>
+<p><strong>Travel and Booking Systems</strong></p>
+<p>Route queries to specialized agents for weather and flights:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">weather_agent</span>
+</span><span id="line-3"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Get real-time weather conditions and forecasts</span>
+</span><span id="line-4"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">flight_agent</span>
+</span><span id="line-5"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Search for flights and provide flight status</span>
+</span></code></pre></div>
+</div>
+<p><strong>Customer Support</strong></p>
+<p>Route common queries to automated support agents while escalating complex issues:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">tier1_support</span>
+</span><span id="line-3"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Handles common FAQs, password resets, and basic troubleshooting</span>
+</span><span id="line-4"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">tier2_support</span>
+</span><span id="line-5"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Handles complex technical issues requiring deep product knowledge</span>
+</span><span id="line-6"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">human_escalation</span>
+</span><span id="line-7"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Escalates sensitive issues or unresolved problems to human agents</span>
+</span></code></pre></div>
+</div>
+<p><strong>Sales and Marketing</strong></p>
+<p>Direct leads and inquiries to specialized sales agents:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">product_recommendation</span>
+</span><span id="line-3"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Recommends products based on user needs and preferences</span>
+</span><span id="line-4"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">pricing_agent</span>
+</span><span id="line-5"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Provides pricing information and quotes</span>
+</span><span id="line-6"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">sales_closer</span>
+</span><span id="line-7"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Handles final negotiations and closes deals</span>
+</span></code></pre></div>
+</div>
+<p><strong>Technical Documentation and Support</strong></p>
+<p>Combine RAG agents for documentation lookup with specialized troubleshooting agents:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">docs_agent</span>
+</span><span id="line-3"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Retrieves relevant documentation and guides</span>
+</span><span id="line-4"><span class="w">    </span><span class="nt">filter_chain</span><span class="p">:</span>
+</span><span id="line-5"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">query_rewriter</span>
+</span><span id="line-6"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">context_builder</span>
+</span><span id="line-7"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">troubleshoot_agent</span>
+</span><span id="line-8"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Diagnoses and resolves technical issues step by step</span>
+</span></code></pre></div>
+</div>
+</section>
+<section id="next-steps">
+<h2>Next Steps<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#next-steps" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#next-steps'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<ul class="simple">
+<li><p>Learn more about <a class="reference internal" href="../concepts/agents.html#agents"><span class="std std-ref">agents</span></a> and the inner vs. outer loop model</p></li>
+<li><p>Explore <a class="reference internal" href="../concepts/filter_chain.html#filter-chain"><span class="std std-ref">filter chains</span></a> for adding guardrails and context enrichment</p></li>
+<li><p>See <a class="reference internal" href="observability/observability.html#observability"><span class="std std-ref">observability</span></a> for monitoring multi-agent workflows</p></li>
+<li><p>Review the <a class="reference internal" href="../concepts/llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">LLM Providers</span></a> guide for model routing within agents</p></li>
+<li><p>Check out the complete <a class="reference external" href="https://github.com/katanemo/plano/tree/main/demos/use_cases/travel_booking" rel="nofollow noopener">Travel Booking demo<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> on GitHub</p></li>
+</ul>
+<div class="admonition note">
+<p class="admonition-title">Note</p>
+<p>To observe traffic to and from agents, please read more about <a class="reference internal" href="observability/observability.html#observability"><span class="std std-ref">observability</span></a> in Plano.</p>
+</div>
+<p>By carefully configuring and managing your Agent routing and hand off, you can significantly improve your application’s responsiveness, performance, and overall user satisfaction.</p>
+</section>
+</section>
+</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
+<div class="mr-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../concepts/prompt_target.html">
+<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="15 18 9 12 15 6"></polyline>
+</svg>
+        Prompt Target
+      </a>
+</div>
+<div class="ml-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="llm_router.html">
+        LLM Routing
+        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="9 18 15 12 9 6"></polyline>
+</svg>
+</a>
+</div>
+</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
+<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
+<ul>
+<li><a :data-current="activeSection === '#how-it-works'" class="reference internal" href="#how-it-works">How It Works</a></li>
+<li><a :data-current="activeSection === '#example-travel-booking-assistant'" class="reference internal" href="#example-travel-booking-assistant">Example: Travel Booking Assistant</a></li>
+<li><a :data-current="activeSection === '#configuration'" class="reference internal" href="#configuration">Configuration</a></li>
+<li><a :data-current="activeSection === '#implementation'" class="reference internal" href="#implementation">Implementation</a><ul>
+<li><a :data-current="activeSection === '#agent-structure'" class="reference internal" href="#agent-structure">Agent Structure</a></li>
+<li><a :data-current="activeSection === '#information-extraction-with-llms'" class="reference internal" href="#information-extraction-with-llms">Information Extraction with LLMs</a></li>
+<li><a :data-current="activeSection === '#calling-external-apis'" class="reference internal" href="#calling-external-apis">Calling External APIs</a></li>
+<li><a :data-current="activeSection === '#preparing-context-and-generating-responses'" class="reference internal" href="#preparing-context-and-generating-responses">Preparing Context and Generating Responses</a></li>
+</ul>
+</li>
+<li><a :data-current="activeSection === '#best-practices'" class="reference internal" href="#best-practices">Best Practices</a></li>
+<li><a :data-current="activeSection === '#common-use-cases'" class="reference internal" href="#common-use-cases">Common Use Cases</a></li>
+<li><a :data-current="activeSection === '#next-steps'" class="reference internal" href="#next-steps">Next Steps</a></li>
+</ul>
+</div>
+</aside>
+</main>
+</div>
+</div><footer class="py-6 border-t border-border md:py-0">
+<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
+<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
+</div>
+</div>
+</footer>
+</div>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
+<script src="../_static/doctools.js?v=9bcbadda"></script>
+<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
+<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
+<script src="../_static/design-tabs.js?v=f930bc37"></script>
+</body>
+</html>
\ No newline at end of file
diff --git a/guides/prompt_guard.html b/guides/prompt_guard.html
index a765aebb..4d5c77c4 100755
--- a/guides/prompt_guard.html
+++ b/guides/prompt_guard.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Prompt Guard | Arch Docs v0.3.22</title>
-<meta content="Prompt Guard | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Prompt Guard | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Guardrails | Plano Docs v0.4</title>
+<meta content="Guardrails | Plano Docs v0.4" property="og:title"/>
+<meta content="Guardrails | Plano Docs v0.4" name="twitter:title"/>
 <link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/guides/prompt_guard.html" rel="canonical"/>
 <link href="../_static/favicon.ico" rel="icon"/>
 <link href="../search.html" rel="search" title="Search"/>
-<link href="agent_routing.html" rel="next" title="Agent Routing and Hand Off"/>
-<link href="../concepts/prompt_target.html" rel="prev" title="Prompt Target"/>
+<link href="state.html" rel="next" title="Conversational State"/>
+<link href="observability/access_logging.html" rel="prev" title="Access Logging"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul class="current">
-<li class="toctree-l1 current"><a class="current reference internal" href="#">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1 current"><a class="current reference internal" href="#">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,103 +148,133 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
 </a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Prompt Guard</span>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Guardrails</span>
 </nav>
 <div id="content" role="main">
-<section id="prompt-guard">
-<span id="id1"></span><h1>Prompt Guard<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#prompt-guard"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p><strong>Prompt guard</strong> is a security and validation feature offered in Arch to protect agents, by filtering and analyzing prompts before they reach your application logic.
-In applications where prompts generate responses or execute specific actions based on user inputs, prompt guard minimizes risks like malicious inputs (or misaligned outputs).
-By adding a layer of input scrutiny, prompt guards ensures safer, more reliable, and accurate interactions with agents.</p>
-<section id="why-prompt-guard">
-<h2>Why Prompt Guard<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#why-prompt-guard" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#why-prompt-guard'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<section id="guardrails">
+<span id="prompt-guard"></span><h1>Guardrails<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#guardrails"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p><strong>Guardrails</strong> are Plano’s way of applying safety and validation checks to prompts before they reach your application logic. They are typically implemented as
+filters in a <a class="reference internal" href="../concepts/filter_chain.html#filter-chain"><span class="std std-ref">Filter Chain</span></a> attached to an agent, so every request passes through a consistent processing layer.</p>
+<section id="why-guardrails">
+<h2>Why Guardrails<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#why-guardrails" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#why-guardrails'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Guardrails are essential for maintaining control over AI-driven applications. They help enforce organizational policies, ensure compliance with regulations
+(like GDPR or HIPAA), and protect users from harmful or inappropriate content. In applications where prompts generate responses or trigger actions, guardrails
+minimize risks like malicious inputs, off-topic queries, or misaligned outputs—adding a consistent layer of input scrutiny that makes interactions safer,
+more reliable, and easier to reason about.</p>
 <ul class="simple">
-<li><dl class="simple">
-<dt><strong>Prompt Sanitization via Arch-Guard</strong></dt><dd><ul>
-<li><p><strong>Jailbreak Prevention</strong>: Detects and filters inputs that might attempt jailbreak attacks, like alternating LLM intended behavior, exposing the system prompt, or bypassing ethnics safety.</p></li>
+<li><p><strong>Jailbreak Prevention</strong>: Detect and filter inputs that attempt to change LLM behavior, expose system prompts, or bypass safety policies.</p></li>
+<li><p><strong>Domain and Topicality Enforcement</strong>: Ensure that agents only respond to prompts within an approved domain (for example, finance-only or healthcare-only use cases) and reject unrelated queries.</p></li>
+<li><p><strong>Dynamic Error Handling</strong>: Provide clear error messages when requests violate policy, helping users correct their inputs.</p></li>
 </ul>
-</dd>
-</dl>
-</li>
-<li><dl class="simple">
-<dt><strong>Dynamic Error Handling</strong></dt><dd><ul>
-<li><p><strong>Automatic Correction</strong>: Applies error-handling techniques to suggest corrections for minor input errors, such as typos or misformatted data.</p></li>
-<li><p><strong>Feedback Mechanism</strong>: Provides informative error messages to users, helping them understand how to correct input mistakes or adhere to guidelines.</p></li>
-</ul>
-</dd>
-</dl>
-</li>
-</ul>
-<div class="admonition note">
-<p class="admonition-title">Note</p>
-<p>Today, Arch offers support for jailbreak via Arch-Guard. We will be adding support for additional guards in Q1, 2025 (including response guardrails)</p>
-</div>
-<section id="what-is-arch-guard">
-<h3>What Is Arch-Guard<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#what-is-arch-guard" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#what-is-arch-guard'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p><a class="reference external" href="https://huggingface.co/collections/katanemo/arch-guard-6702bdc08b889e4bce8f446d" rel="nofollow noopener">Arch-Guard<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> is a robust classifier model specifically trained on a diverse corpus of prompt attacks.
-It excels at detecting explicitly malicious prompts, providing an essential layer of security for LLM applications.</p>
-<p>By embedding Arch-Guard within the Arch architecture, we empower developers to build robust, LLM-powered applications while prioritizing security and safety. With Arch-Guard, you can navigate the complexities of prompt management with confidence, knowing you have a reliable defense against malicious input.</p>
 </section>
-<section id="example-configuration">
-<h3>Example Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#example-configuration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#example-configuration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Here is an example of using Arch-Guard in Arch:</p>
-<div class="literal-block-wrapper docutils container" id="id2">
-<div class="code-block-caption"><span class="caption-text">Arch-Guard Example Configuration</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos">1</span><span class="w">      </span><span class="nt">on_exception</span><span class="p">:</span>
-</span><span id="line-2"><span class="linenos">2</span><span class="w">        </span><span class="nt">message</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Looks like you're curious about my abilities, but I can only provide assistance within my programmed parameters.</span>
-</span><span id="line-3"><span class="linenos">3</span>
-</span><span id="line-4"><span class="linenos">4</span><span class="nt">prompt_targets</span><span class="p">:</span>
-</span><span id="line-5"><span class="linenos">5</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">information_extraction</span>
+<section id="how-guardrails-work">
+<h2>How Guardrails Work<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#how-guardrails-work" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#how-guardrails-work'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Guardrails can be implemented as either in-process MCP filters or as HTTP-based filters. HTTP filters are external services that receive the request over HTTP, validate it, and return a response to allow or reject the request. This makes it easy to use filters written in any language or run them as independent services.</p>
+<p>Each filter receives the chat messages, evaluates them against policy, and either lets the request continue or raises a <code class="docutils literal notranslate"><span class="pre">ToolError</span></code> (or returns an error response) to reject it with a helpful error message.</p>
+<p>The example below shows an input guard for TechCorp’s customer support system that validates queries are within the company’s domain:</p>
+<div class="literal-block-wrapper docutils container" id="id1">
+<div class="code-block-caption"><span class="caption-text">Example domain validation guard using FastMCP</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id1"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="kn">from</span><span class="w"> </span><span class="nn">typing</span><span class="w"> </span><span class="kn">import</span> <span class="n">List</span>
+</span><span id="line-2"><span class="kn">from</span><span class="w"> </span><span class="nn">fastmcp.exceptions</span><span class="w"> </span><span class="kn">import</span> <span class="n">ToolError</span>
+</span><span id="line-3"><span class="kn">from</span><span class="w"> </span><span class="nn">.</span><span class="w"> </span><span class="kn">import</span> <span class="n">mcp</span>
+</span><span id="line-4">
+</span><span id="line-5"><span class="nd">@mcp</span><span class="o">.</span><span class="n">tool</span>
+</span><span id="line-6"><span class="k">async</span> <span class="k">def</span><span class="w"> </span><span class="nf">input_guards</span><span class="p">(</span><span class="n">messages</span><span class="p">:</span> <span class="n">List</span><span class="p">[</span><span class="n">ChatMessage</span><span class="p">])</span> <span class="o">-&gt;</span> <span class="n">List</span><span class="p">[</span><span class="n">ChatMessage</span><span class="p">]:</span>
+</span><span id="line-7"><span class="w">    </span><span class="sd">"""Validates queries are within TechCorp's domain."""</span>
+</span><span id="line-8">
+</span><span id="line-9">    <span class="c1"># Get the user's query</span>
+</span><span id="line-10">    <span class="n">user_query</span> <span class="o">=</span> <span class="nb">next</span><span class="p">(</span>
+</span><span id="line-11">        <span class="p">(</span><span class="n">msg</span><span class="o">.</span><span class="n">content</span> <span class="k">for</span> <span class="n">msg</span> <span class="ow">in</span> <span class="nb">reversed</span><span class="p">(</span><span class="n">messages</span><span class="p">)</span> <span class="k">if</span> <span class="n">msg</span><span class="o">.</span><span class="n">role</span> <span class="o">==</span> <span class="s2">"user"</span><span class="p">),</span>
+</span><span id="line-12">        <span class="s2">""</span>
+</span><span id="line-13">    <span class="p">)</span>
+</span><span id="line-14">
+</span><span id="line-15">    <span class="c1"># Use an LLM to validate the query scope (simplified)</span>
+</span><span id="line-16">    <span class="n">is_valid</span> <span class="o">=</span> <span class="k">await</span> <span class="n">validate_with_llm</span><span class="p">(</span><span class="n">user_query</span><span class="p">)</span>
+</span><span id="line-17">
+</span><span id="line-18">    <span class="k">if</span> <span class="ow">not</span> <span class="n">is_valid</span><span class="p">:</span>
+</span><span id="line-19">        <span class="k">raise</span> <span class="n">ToolError</span><span class="p">(</span>
+</span><span id="line-20">            <span class="s2">"I can only assist with questions related to TechCorp and its services. "</span>
+</span><span id="line-21">            <span class="s2">"Please ask about TechCorp's products, pricing, SLAs, or technical support."</span>
+</span><span id="line-22">        <span class="p">)</span>
+</span><span id="line-23">
+</span><span id="line-24">    <span class="k">return</span> <span class="n">messages</span>
 </span></code></pre></div>
 </div>
 </div>
+<p>To wire this guardrail into Plano, define the filter and add it to your agent’s filter chain:</p>
+<div class="literal-block-wrapper docutils container" id="id2">
+<div class="code-block-caption"><span class="caption-text">Plano configuration with input guard filter</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">filters</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">input_guards</span>
+</span><span id="line-3"><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://localhost:10500</span>
+</span><span id="line-4">
+</span><span id="line-5"><span class="nt">listeners</span><span class="p">:</span>
+</span><span id="line-6"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">agent</span>
+</span><span id="line-7"><span class="w">    </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">agent_1</span>
+</span><span id="line-8"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">8001</span>
+</span><span id="line-9"><span class="w">    </span><span class="nt">router</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">plano_orchestrator_v1</span>
+</span><span id="line-10"><span class="w">    </span><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-11"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">rag_agent</span>
+</span><span id="line-12"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">virtual assistant for retrieval augmented generation tasks</span>
+</span><span id="line-13"><span class="w">        </span><span class="nt">filter_chain</span><span class="p">:</span>
+</span><span id="line-14"><span class="w">          </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">input_guards</span>
+</span></code></pre></div>
+</div>
+</div>
+<p>When a request arrives at <code class="docutils literal notranslate"><span class="pre">agent_1</span></code>, Plano invokes the <code class="docutils literal notranslate"><span class="pre">input_guards</span></code> filter first. If validation passes, the request continues to
+the agent. If validation fails (<code class="docutils literal notranslate"><span class="pre">ToolError</span></code> raised), Plano returns an error response to the caller.</p>
 </section>
-</section>
-<section id="how-arch-guard-works">
-<h2>How Arch-Guard Works<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#how-arch-guard-works" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#how-arch-guard-works'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<ol class="arabic">
-<li><p><strong>Pre-Processing Stage</strong></p>
-<blockquote>
-<div><p>As a request or prompt is received, Arch Guard first performs validation. If any violations are detected, the input is flagged, and a tailored error message may be returned.</p>
-</div></blockquote>
-</li>
-<li><p><strong>Error Handling and Feedback</strong></p>
-<blockquote>
-<div><p>If the prompt contains errors or does not meet certain criteria, the user receives immediate feedback or correction suggestions, enhancing usability and reducing the chance of repeated input mistakes.</p>
-</div></blockquote>
-</li>
-</ol>
-</section>
-<section id="benefits-of-using-arch-guard">
-<h2>Benefits of Using Arch Guard<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#benefits-of-using-arch-guard" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#benefits-of-using-arch-guard'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<ul class="simple">
-<li><p><strong>Enhanced Security</strong>: Protects against injection attacks, harmful content, and misuse, securing both system and user data.</p></li>
-<li><p><strong>Better User Experience</strong>: Clear feedback and error correction improve user interactions by guiding them to correct input formats and constraints.</p></li>
-</ul>
-</section>
-<section id="summary">
-<h2>Summary<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#summary" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#summary'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Prompt guard is an essential tool for any prompt-based system that values security, accuracy, and compliance.
-By implementing Prompt Guard, developers can provide a robust layer of input validation and security, leading to better-performing, reliable, and safer applications.</p>
+<section id="testing-the-guardrail">
+<h2>Testing the Guardrail<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#testing-the-guardrail" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#testing-the-guardrail'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Here’s an example of the guardrail in action, rejecting a query about Apple Corporation (outside TechCorp’s domain):</p>
+<div class="literal-block-wrapper docutils container" id="id3">
+<div class="code-block-caption"><span class="caption-text">Request that violates the guardrail policy</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id3"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">curl<span class="w"> </span>-X<span class="w"> </span>POST<span class="w"> </span>http://localhost:8001/v1/chat/completions<span class="w"> </span><span class="se">\</span>
+</span><span id="line-2"><span class="w">  </span>-H<span class="w"> </span><span class="s2">"Content-Type: application/json"</span><span class="w"> </span><span class="se">\</span>
+</span><span id="line-3"><span class="w">  </span>-d<span class="w"> </span><span class="s1">'{</span>
+</span><span id="line-4"><span class="s1">    "model": "gpt-4",</span>
+</span><span id="line-5"><span class="s1">    "messages": [</span>
+</span><span id="line-6"><span class="s1">      {</span>
+</span><span id="line-7"><span class="s1">        "role": "user",</span>
+</span><span id="line-8"><span class="s1">        "content": "what is sla for apple corporation?"</span>
+</span><span id="line-9"><span class="s1">      }</span>
+</span><span id="line-10"><span class="s1">    ],</span>
+</span><span id="line-11"><span class="s1">    "stream": false</span>
+</span><span id="line-12"><span class="s1">  }'</span>
+</span></code></pre></div>
+</div>
+</div>
+<div class="literal-block-wrapper docutils container" id="id4">
+<div class="code-block-caption"><span class="caption-text">Error response from the guardrail</span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id4"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-json notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="p">{</span>
+</span><span id="line-2"><span class="w">  </span><span class="nt">"error"</span><span class="p">:</span><span class="w"> </span><span class="s2">"ClientError"</span><span class="p">,</span>
+</span><span id="line-3"><span class="w">  </span><span class="nt">"agent"</span><span class="p">:</span><span class="w"> </span><span class="s2">"input_guards"</span><span class="p">,</span>
+</span><span id="line-4"><span class="w">  </span><span class="nt">"status"</span><span class="p">:</span><span class="w"> </span><span class="mi">400</span><span class="p">,</span>
+</span><span id="line-5"><span class="w">  </span><span class="nt">"agent_response"</span><span class="p">:</span><span class="w"> </span><span class="s2">"I apologize, but I can only assist with questions related to TechCorp and its services. Your query appears to be outside this scope. The query is about SLA for Apple Corporation, which is unrelated to TechCorp.\n\nPlease ask me about TechCorp's products, services, pricing, SLAs, or technical support."</span>
+</span><span id="line-6"><span class="p">}</span>
+</span></code></pre></div>
+</div>
+</div>
+<p>This prevents out-of-scope queries from reaching your agent while providing clear feedback to users about why their request was rejected.</p>
 </section>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../concepts/prompt_target.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="observability/access_logging.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Prompt Target
+        Access Logging
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="agent_routing.html">
-        Agent Routing and Hand Off
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="state.html">
+        Conversational State
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -258,14 +283,9 @@ By implementing Prompt Guard, developers can provide a robust layer of input val
 </div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
 <div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
 <ul>
-<li><a :data-current="activeSection === '#why-prompt-guard'" class="reference internal" href="#why-prompt-guard">Why Prompt Guard</a><ul>
-<li><a :data-current="activeSection === '#what-is-arch-guard'" class="reference internal" href="#what-is-arch-guard">What Is Arch-Guard</a></li>
-<li><a :data-current="activeSection === '#example-configuration'" class="reference internal" href="#example-configuration">Example Configuration</a></li>
-</ul>
-</li>
-<li><a :data-current="activeSection === '#how-arch-guard-works'" class="reference internal" href="#how-arch-guard-works">How Arch-Guard Works</a></li>
-<li><a :data-current="activeSection === '#benefits-of-using-arch-guard'" class="reference internal" href="#benefits-of-using-arch-guard">Benefits of Using Arch Guard</a></li>
-<li><a :data-current="activeSection === '#summary'" class="reference internal" href="#summary">Summary</a></li>
+<li><a :data-current="activeSection === '#why-guardrails'" class="reference internal" href="#why-guardrails">Why Guardrails</a></li>
+<li><a :data-current="activeSection === '#how-guardrails-work'" class="reference internal" href="#how-guardrails-work">How Guardrails Work</a></li>
+<li><a :data-current="activeSection === '#testing-the-guardrail'" class="reference internal" href="#testing-the-guardrail">Testing the Guardrail</a></li>
 </ul>
 </div>
 </aside>
@@ -274,12 +294,12 @@ By implementing Prompt Guard, developers can provide a robust layer of input val
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../_static/doctools.js?v=9bcbadda"></script>
 <script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
diff --git a/guides/state.html b/guides/state.html
new file mode 100755
index 00000000..2b9b6b1c
--- /dev/null
+++ b/guides/state.html
@@ -0,0 +1,463 @@
+<!DOCTYPE html>
+
+<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
+<head>
+<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
+<meta charset="utf-8"/>
+<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
+<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
+<meta content="width=device-width, initial-scale=1" name="viewport"/>
+<title>Conversational State | Plano Docs v0.4</title>
+<meta content="Conversational State | Plano Docs v0.4" property="og:title"/>
+<meta content="Conversational State | Plano Docs v0.4" name="twitter:title"/>
+<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
+<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
+<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
+<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
+<link href="./docs/guides/state.html" rel="canonical"/>
+<link href="../_static/favicon.ico" rel="icon"/>
+<link href="../search.html" rel="search" title="Search"/>
+<link href="../resources/tech_overview/tech_overview.html" rel="next" title="Tech Overview"/>
+<link href="prompt_guard.html" rel="prev" title="Guardrails"/>
+<script>
+    <!-- Prevent Flash of wrong theme -->
+      const userPreference = localStorage.getItem('darkMode');
+      let mode;
+      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
+        mode = 'dark';
+        document.documentElement.classList.add('dark');
+      } else {
+        mode = 'light';
+      }
+      if (!userPreference) {localStorage.setItem('darkMode', mode)}
+    </script>
+</head>
+<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
+<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
+      Skip to content
+    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
+<div class="hidden mr-4 md:flex">
+<a class="flex items-center mr-6" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
+<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
+</svg>
+<span class="sr-only">Toggle navigation menu</span>
+</button>
+<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
+<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
+<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
+<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
+<span class="text-xs">⌘</span>
+    K
+  </kbd>
+</form>
+</div>
+<nav class="flex items-center space-x-1">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
+<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
+<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
+</div>
+</a>
+<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
+<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
+</svg>
+<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
+</svg>
+</button>
+</nav>
+</div>
+</div>
+</header>
+<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
+<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a>
+<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
+<div class="overflow-y-auto h-full w-full relative pr-6">
+
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
+<script>
+  window.dataLayer = window.dataLayer || [];
+  function gtag(){dataLayer.push(arguments);}
+  gtag('js', new Date());
+
+  gtag('config', 'G-EH2VW19FXE');
+</script>
+<nav class="table w-full min-w-full my-6 lg:my-8">
+<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/prompt_target.html">Prompt Target</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
+<ul class="current">
+<li class="toctree-l1"><a class="reference internal" href="orchestration.html">Orchestration</a></li>
+<li class="toctree-l1"><a class="reference internal" href="llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="function_calling.html">Function Calling</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="observability/tracing.html">Tracing</a></li>
+<li class="toctree-l2"><a class="reference internal" href="observability/monitoring.html">Monitoring</a></li>
+<li class="toctree-l2"><a class="reference internal" href="observability/access_logging.html">Access Logging</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1 current"><a class="current reference internal" href="#">Conversational State</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
+<ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../resources/llms_txt.html">llms.txt</a></li>
+</ul>
+</nav>
+</div>
+</div>
+<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
+<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
+</svg>
+</button>
+</aside>
+<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
+<div class="w-full min-w-0 mx-auto">
+<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
+<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
+<span class="hidden md:inline">Plano Docs v0.4</span>
+<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
+<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
+</svg>
+</a>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Conversational State</span>
+</nav>
+<div id="content" role="main">
+<section id="conversational-state">
+<span id="managing-conversational-state"></span><h1>Conversational State<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#conversational-state"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p>The OpenAI Responses API (<code class="docutils literal notranslate"><span class="pre">v1/responses</span></code>) is designed for multi-turn conversations where context needs to persist across requests. Plano provides a unified <code class="docutils literal notranslate"><span class="pre">v1/responses</span></code> API that works with <strong>any LLM provider</strong>—OpenAI, Anthropic, Azure OpenAI, DeepSeek, or any OpenAI-compatible provider—while automatically managing conversational state for you.</p>
+<p>Unlike the traditional Chat Completions API where you manually manage conversation history by including all previous messages in each request, Plano handles state management behind the scenes. This means you can use the Responses API with any model provider, and Plano will persist conversation context across requests—making it ideal for building conversational agents that remember context without bloating every request with full message history.</p>
+<section id="how-it-works">
+<h2>How It Works<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#how-it-works" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#how-it-works'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>When a client calls the Responses API:</p>
+<ol class="arabic simple">
+<li><p><strong>First request</strong>: Plano generates a unique <code class="docutils literal notranslate"><span class="pre">resp_id</span></code> and stores the conversation state (messages, model, provider, timestamp).</p></li>
+<li><p><strong>Subsequent requests</strong>: The client includes the <code class="docutils literal notranslate"><span class="pre">previous_resp_id</span></code> from the previous response. Plano retrieves the stored conversation state, merges it with the new input, and sends the combined context to the LLM.</p></li>
+<li><p><strong>Response</strong>: The LLM sees the full conversation history without the client needing to resend all previous messages.</p></li>
+</ol>
+<p>This pattern dramatically reduces bandwidth and makes it easier to build multi-turn agents—Plano handles the state plumbing so you can focus on agent logic.</p>
+<p><strong>Example Using OpenAI Python SDK:</strong></p>
+<div class="highlight-python notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="kn">from</span><span class="w"> </span><span class="nn">openai</span><span class="w"> </span><span class="kn">import</span> <span class="n">OpenAI</span>
+</span><span id="line-2">
+</span><span id="line-3"><span class="c1"># Point to Plano's Model Proxy endpoint</span>
+</span><span id="line-4"><span class="n">client</span> <span class="o">=</span> <span class="n">OpenAI</span><span class="p">(</span>
+</span><span id="line-5">    <span class="n">api_key</span><span class="o">=</span><span class="s2">"test-key"</span><span class="p">,</span>
+</span><span id="line-6">    <span class="n">base_url</span><span class="o">=</span><span class="s2">"http://127.0.0.1:12000/v1"</span>
+</span><span id="line-7"><span class="p">)</span>
+</span><span id="line-8">
+</span><span id="line-9"><span class="c1"># First turn - Plano creates a new conversation state</span>
+</span><span id="line-10"><span class="n">response</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">responses</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-11">    <span class="n">model</span><span class="o">=</span><span class="s2">"claude-sonnet-4-5"</span><span class="p">,</span>  <span class="c1"># Works with any configured provider</span>
+</span><span id="line-12">    <span class="nb">input</span><span class="o">=</span><span class="s2">"My name is Alice and I like Python"</span>
+</span><span id="line-13"><span class="p">)</span>
+</span><span id="line-14">
+</span><span id="line-15"><span class="c1"># Save the response_id for conversation continuity</span>
+</span><span id="line-16"><span class="n">resp_id</span> <span class="o">=</span> <span class="n">response</span><span class="o">.</span><span class="n">id</span>
+</span><span id="line-17"><span class="nb">print</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Assistant: </span><span class="si">{</span><span class="n">response</span><span class="o">.</span><span class="n">output_text</span><span class="si">}</span><span class="s2">"</span><span class="p">)</span>
+</span><span id="line-18">
+</span><span id="line-19"><span class="c1"># Second turn - Plano automatically retrieves previous context</span>
+</span><span id="line-20"><span class="n">resp2</span> <span class="o">=</span> <span class="n">client</span><span class="o">.</span><span class="n">responses</span><span class="o">.</span><span class="n">create</span><span class="p">(</span>
+</span><span id="line-21">    <span class="n">model</span><span class="o">=</span><span class="s2">"claude-sonnet-4-5"</span><span class="p">,</span> <span class="c1"># Make sure its configured in plano_config.yaml</span>
+</span><span id="line-22">    <span class="nb">input</span><span class="o">=</span><span class="s2">"Please list all the messages you have received in our conversation, numbering each one."</span><span class="p">,</span>
+</span><span id="line-23">    <span class="n">previous_response_id</span><span class="o">=</span><span class="n">resp_id</span><span class="p">,</span>
+</span><span id="line-24"><span class="p">)</span>
+</span><span id="line-25">
+</span><span id="line-26"><span class="nb">print</span><span class="p">(</span><span class="sa">f</span><span class="s2">"Assistant: </span><span class="si">{</span><span class="n">resp2</span><span class="o">.</span><span class="n">output_text</span><span class="si">}</span><span class="s2">"</span><span class="p">)</span>
+</span><span id="line-27"><span class="c1"># Output: "Your name is Alice and your favorite language is Python"</span>
+</span></code></pre></div>
+</div>
+<p>Notice how the second request only includes the new user message—Plano automatically merges it with the stored conversation history before sending to the LLM.</p>
+</section>
+<section id="configuration-overview">
+<h2>Configuration Overview<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuration-overview" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configuration-overview'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>State storage is configured in the <code class="docutils literal notranslate"><span class="pre">state_storage</span></code> section of your <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code>:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="nt">state_storage</span><span class="p">:</span>
+</span><span id="line-2"><span class="linenos"> 2</span><span class="w">  </span><span class="c1"># Type: memory | postgres</span>
+</span><span id="line-3"><mark><span class="linenos"> 3</span><span class="w">  </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">postgres</span>
+</mark></span><span id="line-4"><span class="linenos"> 4</span>
+</span><span id="line-5"><span class="linenos"> 5</span><span class="w">  </span><span class="c1"># Connection string for postgres type</span>
+</span><span id="line-6"><mark><span class="linenos"> 6</span><span class="w">  </span><span class="c1"># Environment variables are supported using $VAR_NAME or ${VAR_NAME} syntax</span>
+</mark></span><span id="line-7"><mark><span class="linenos"> 7</span><span class="w">  </span><span class="c1"># Replace [USER] and [HOST] with your actual database credentials</span>
+</mark></span><span id="line-8"><mark><span class="linenos"> 8</span><span class="w">  </span><span class="c1"># Variables like $DB_PASSWORD MUST be set before running config validation/rendering</span>
+</mark></span><span id="line-9"><mark><span class="linenos"> 9</span><span class="w">  </span><span class="c1"># Example: Replace [USER] with 'myuser' and [HOST] with 'db.example.com:5432'</span>
+</mark></span><span id="line-10"><mark><span class="linenos">10</span><span class="w">  </span><span class="nt">connection_string</span><span class="p">:</span><span class="w"> </span><span class="s">"postgresql://[USER]:$DB_PASSWORD@[HOST]:5432/postgres"</span>
+</mark></span></code></pre></div>
+</div>
+<p>Plano supports two storage backends:</p>
+<ul class="simple">
+<li><p><strong>Memory</strong>: Fast, ephemeral storage for development and testing. State is lost when Plano restarts.</p></li>
+<li><p><strong>PostgreSQL</strong>: Durable, production-ready storage with support for Supabase and self-hosted PostgreSQL instances.</p></li>
+</ul>
+<div class="admonition note">
+<p class="admonition-title">Note</p>
+<p>If you don’t configure <code class="docutils literal notranslate"><span class="pre">state_storage</span></code>, conversation state management is <strong>disabled</strong>. The Responses API will still work, but clients must manually include full conversation history in each request (similar to the Chat Completions API behavior).</p>
+</div>
+</section>
+<section id="memory-storage-development">
+<h2>Memory Storage (Development)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#memory-storage-development" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#memory-storage-development'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Memory storage keeps conversation state in-memory using a thread-safe <code class="docutils literal notranslate"><span class="pre">HashMap</span></code>. It’s perfect for local development, demos, and testing, but all state is lost when Plano restarts.</p>
+<p><strong>Configuration</strong></p>
+<p>Add this to your <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code>:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">state_storage</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">memory</span>
+</span></code></pre></div>
+</div>
+<p>That’s it. No additional setup required.</p>
+<p><strong>When to Use Memory Storage</strong></p>
+<ul class="simple">
+<li><p>Local development and debugging</p></li>
+<li><p>Demos and proof-of-concepts</p></li>
+<li><p>Automated testing environments</p></li>
+<li><p>Single-instance deployments where persistence isn’t critical</p></li>
+</ul>
+<p><strong>Limitations</strong></p>
+<ul class="simple">
+<li><p>State is lost on restart</p></li>
+<li><p>Not suitable for production workloads</p></li>
+<li><p>Cannot scale across multiple Plano instances</p></li>
+</ul>
+</section>
+<section id="postgresql-storage-production">
+<h2>PostgreSQL Storage (Production)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#postgresql-storage-production" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#postgresql-storage-production'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>PostgreSQL storage provides durable, production-grade conversation state management. It works with both self-hosted PostgreSQL and Supabase (PostgreSQL-as-a-service), making it ideal for scaling multi-agent systems in production.</p>
+<section id="prerequisites">
+<h3>Prerequisites<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#prerequisites" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#prerequisites'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Before configuring PostgreSQL storage, you need:</p>
+<ol class="arabic simple">
+<li><p>A PostgreSQL database (version 12 or later)</p></li>
+<li><p>Database credentials (host, user, password)</p></li>
+<li><p>The <code class="docutils literal notranslate"><span class="pre">conversation_states</span></code> table created in your database</p></li>
+</ol>
+<p><strong>Setting Up the Database</strong></p>
+<p>Run the SQL schema to create the required table:</p>
+<div class="highlight-sql notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos"> 1</span><span class="c1">-- Conversation State Storage Table</span>
+</span><span id="line-2"><span class="linenos"> 2</span><span class="c1">-- This table stores conversational context for the OpenAI Responses API</span>
+</span><span id="line-3"><span class="linenos"> 3</span><span class="c1">-- Run this SQL against your PostgreSQL/Supabase database before enabling conversation state storage</span>
+</span><span id="line-4"><span class="linenos"> 4</span>
+</span><span id="line-5"><span class="linenos"> 5</span><span class="k">CREATE</span><span class="w"> </span><span class="k">TABLE</span><span class="w"> </span><span class="k">IF</span><span class="w"> </span><span class="k">NOT</span><span class="w"> </span><span class="k">EXISTS</span><span class="w"> </span><span class="n">conversation_states</span><span class="w"> </span><span class="p">(</span>
+</span><span id="line-6"><span class="linenos"> 6</span><span class="w">    </span><span class="n">response_id</span><span class="w"> </span><span class="nb">TEXT</span><span class="w"> </span><span class="k">PRIMARY</span><span class="w"> </span><span class="k">KEY</span><span class="p">,</span>
+</span><span id="line-7"><span class="linenos"> 7</span><span class="w">    </span><span class="n">input_items</span><span class="w"> </span><span class="n">JSONB</span><span class="w"> </span><span class="k">NOT</span><span class="w"> </span><span class="k">NULL</span><span class="p">,</span>
+</span><span id="line-8"><span class="linenos"> 8</span><span class="w">    </span><span class="n">created_at</span><span class="w"> </span><span class="nb">BIGINT</span><span class="w"> </span><span class="k">NOT</span><span class="w"> </span><span class="k">NULL</span><span class="p">,</span>
+</span><span id="line-9"><span class="linenos"> 9</span><span class="w">    </span><span class="n">model</span><span class="w"> </span><span class="nb">TEXT</span><span class="w"> </span><span class="k">NOT</span><span class="w"> </span><span class="k">NULL</span><span class="p">,</span>
+</span><span id="line-10"><span class="linenos">10</span><span class="w">    </span><span class="n">provider</span><span class="w"> </span><span class="nb">TEXT</span><span class="w"> </span><span class="k">NOT</span><span class="w"> </span><span class="k">NULL</span><span class="p">,</span>
+</span><span id="line-11"><span class="linenos">11</span><span class="w">    </span><span class="n">updated_at</span><span class="w"> </span><span class="k">TIMESTAMP</span><span class="w"> </span><span class="k">DEFAULT</span><span class="w"> </span><span class="k">CURRENT_TIMESTAMP</span>
+</span><span id="line-12"><span class="linenos">12</span><span class="p">);</span>
+</span><span id="line-13"><span class="linenos">13</span>
+</span><span id="line-14"><span class="linenos">14</span><span class="c1">-- Indexes for common query patterns</span>
+</span><span id="line-15"><span class="linenos">15</span><span class="k">CREATE</span><span class="w"> </span><span class="k">INDEX</span><span class="w"> </span><span class="k">IF</span><span class="w"> </span><span class="k">NOT</span><span class="w"> </span><span class="k">EXISTS</span><span class="w"> </span><span class="n">idx_conversation_states_created_at</span>
+</span><span id="line-16"><span class="linenos">16</span><span class="w">    </span><span class="k">ON</span><span class="w"> </span><span class="n">conversation_states</span><span class="p">(</span><span class="n">created_at</span><span class="p">);</span>
+</span><span id="line-17"><span class="linenos">17</span>
+</span><span id="line-18"><span class="linenos">18</span><span class="k">CREATE</span><span class="w"> </span><span class="k">INDEX</span><span class="w"> </span><span class="k">IF</span><span class="w"> </span><span class="k">NOT</span><span class="w"> </span><span class="k">EXISTS</span><span class="w"> </span><span class="n">idx_conversation_states_provider</span>
+</span><span id="line-19"><span class="linenos">19</span><span class="w">    </span><span class="k">ON</span><span class="w"> </span><span class="n">conversation_states</span><span class="p">(</span><span class="n">provider</span><span class="p">);</span>
+</span><span id="line-20"><span class="linenos">20</span>
+</span><span id="line-21"><span class="linenos">21</span><span class="c1">-- Optional: Add a policy for automatic cleanup of old conversations</span>
+</span><span id="line-22"><span class="linenos">22</span><span class="c1">-- Uncomment and adjust the retention period as needed</span>
+</span><span id="line-23"><span class="linenos">23</span><span class="c1">-- CREATE INDEX IF NOT EXISTS idx_conversation_states_updated_at</span>
+</span><span id="line-24"><span class="linenos">24</span><span class="c1">--     ON conversation_states(updated_at);</span>
+</span><span id="line-25"><span class="linenos">25</span>
+</span><span id="line-26"><span class="linenos">26</span><span class="k">COMMENT</span><span class="w"> </span><span class="k">ON</span><span class="w"> </span><span class="k">TABLE</span><span class="w"> </span><span class="n">conversation_states</span><span class="w"> </span><span class="k">IS</span><span class="w"> </span><span class="s1">'Stores conversation history for OpenAI Responses API continuity'</span><span class="p">;</span>
+</span><span id="line-27"><span class="linenos">27</span><span class="k">COMMENT</span><span class="w"> </span><span class="k">ON</span><span class="w"> </span><span class="k">COLUMN</span><span class="w"> </span><span class="n">conversation_states</span><span class="p">.</span><span class="n">response_id</span><span class="w"> </span><span class="k">IS</span><span class="w"> </span><span class="s1">'Unique identifier for the conversation state'</span><span class="p">;</span>
+</span><span id="line-28"><span class="linenos">28</span><span class="k">COMMENT</span><span class="w"> </span><span class="k">ON</span><span class="w"> </span><span class="k">COLUMN</span><span class="w"> </span><span class="n">conversation_states</span><span class="p">.</span><span class="n">input_items</span><span class="w"> </span><span class="k">IS</span><span class="w"> </span><span class="s1">'JSONB array of conversation messages and context'</span><span class="p">;</span>
+</span><span id="line-29"><span class="linenos">29</span><span class="k">COMMENT</span><span class="w"> </span><span class="k">ON</span><span class="w"> </span><span class="k">COLUMN</span><span class="w"> </span><span class="n">conversation_states</span><span class="p">.</span><span class="n">created_at</span><span class="w"> </span><span class="k">IS</span><span class="w"> </span><span class="s1">'Unix timestamp (seconds) when the conversation started'</span><span class="p">;</span>
+</span><span id="line-30"><span class="linenos">30</span><span class="k">COMMENT</span><span class="w"> </span><span class="k">ON</span><span class="w"> </span><span class="k">COLUMN</span><span class="w"> </span><span class="n">conversation_states</span><span class="p">.</span><span class="n">model</span><span class="w"> </span><span class="k">IS</span><span class="w"> </span><span class="s1">'Model name used for this conversation'</span><span class="p">;</span>
+</span><span id="line-31"><span class="linenos">31</span><span class="k">COMMENT</span><span class="w"> </span><span class="k">ON</span><span class="w"> </span><span class="k">COLUMN</span><span class="w"> </span><span class="n">conversation_states</span><span class="p">.</span><span class="n">provider</span><span class="w"> </span><span class="k">IS</span><span class="w"> </span><span class="s1">'LLM provider (e.g., openai, anthropic, bedrock)'</span><span class="p">;</span>
+</span></code></pre></div>
+</div>
+<p><strong>Using psql:</strong></p>
+<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">psql<span class="w"> </span><span class="nv">$DATABASE_URL</span><span class="w"> </span>-f<span class="w"> </span>docs/db_setup/conversation_states.sql
+</span></code></pre></div>
+</div>
+<p><strong>Using Supabase Dashboard:</strong></p>
+<ol class="arabic simple">
+<li><p>Log in to your Supabase project</p></li>
+<li><p>Navigate to the SQL Editor</p></li>
+<li><p>Copy and paste the SQL from <code class="docutils literal notranslate"><span class="pre">docs/db_setup/conversation_states.sql</span></code></p></li>
+<li><p>Run the query</p></li>
+</ol>
+</section>
+<section id="configuration">
+<h3>Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configuration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Once the database table is created, configure Plano to use PostgreSQL storage:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">state_storage</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">postgres</span>
+</span><span id="line-3"><span class="w">  </span><span class="nt">connection_string</span><span class="p">:</span><span class="w"> </span><span class="s">"postgresql://user:password@host:5432/database"</span>
+</span></code></pre></div>
+</div>
+<p><strong>Using Environment Variables</strong></p>
+<p>You should <strong>never</strong> hardcode credentials. Use environment variables instead:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">state_storage</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">postgres</span>
+</span><span id="line-3"><span class="w">  </span><span class="nt">connection_string</span><span class="p">:</span><span class="w"> </span><span class="s">"postgresql://myuser:$DB_PASSWORD@db.example.com:5432/postgres"</span>
+</span></code></pre></div>
+</div>
+<p>Then set the environment variable before running Plano:</p>
+<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nb">export</span><span class="w"> </span><span class="nv">DB_PASSWORD</span><span class="o">=</span><span class="s2">"your-secure-password"</span>
+</span><span id="line-2"><span class="c1"># Run Plano or config validation</span>
+</span><span id="line-3">./plano
+</span></code></pre></div>
+</div>
+<div class="admonition warning">
+<p class="admonition-title">Warning</p>
+<p><strong>Special Characters in Passwords</strong>: If your password contains special characters like <code class="docutils literal notranslate"><span class="pre">#</span></code>, <code class="docutils literal notranslate"><span class="pre">@</span></code>, or <code class="docutils literal notranslate"><span class="pre">&amp;</span></code>, you must URL-encode them in the connection string. For example, <code class="docutils literal notranslate"><span class="pre">MyPass#123</span></code> becomes <code class="docutils literal notranslate"><span class="pre">MyPass%23123</span></code>.</p>
+</div>
+</section>
+<section id="supabase-connection-strings">
+<h3>Supabase Connection Strings<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#supabase-connection-strings" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#supabase-connection-strings'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
+<p>Supabase requires different connection strings depending on your network setup. Most users should use the <strong>Session Pooler</strong> connection string.</p>
+<p><strong>IPv4 Networks (Most Common)</strong></p>
+<p>Use the Session Pooler connection string (port 5432):</p>
+<div class="highlight-text notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">postgresql://postgres.[PROJECT-REF]:[PASSWORD]@aws-0-[REGION].pooler.supabase.com:5432/postgres
+</span></code></pre></div>
+</div>
+<p><strong>IPv6 Networks</strong></p>
+<p>Use the direct connection (port 5432):</p>
+<div class="highlight-text notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">postgresql://postgres:[PASSWORD]@db.[PROJECT-REF].supabase.co:5432/postgres
+</span></code></pre></div>
+</div>
+<p><strong>Finding Your Connection String</strong></p>
+<ol class="arabic simple">
+<li><p>Go to your Supabase project dashboard</p></li>
+<li><p>Navigate to <strong>Settings → Database → Connection Pooling</strong></p></li>
+<li><p>Copy the <strong>Session mode</strong> connection string</p></li>
+<li><p>Replace <code class="docutils literal notranslate"><span class="pre">[YOUR-PASSWORD]</span></code> with your actual database password</p></li>
+<li><p>URL-encode special characters in the password</p></li>
+</ol>
+<p><strong>Example Configuration</strong></p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">state_storage</span><span class="p">:</span>
+</span><span id="line-2"><span class="w">  </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">postgres</span>
+</span><span id="line-3"><span class="w">  </span><span class="nt">connection_string</span><span class="p">:</span><span class="w"> </span><span class="s">"postgresql://postgres.myproject:$DB_PASSWORD@aws-0-us-west-2.pooler.supabase.com:5432/postgres"</span>
+</span></code></pre></div>
+</div>
+<p>Then set the environment variable:</p>
+<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># If your password is "MyPass#123", encode it as "MyPass%23123"</span>
+</span><span id="line-2"><span class="nb">export</span><span class="w"> </span><span class="nv">DB_PASSWORD</span><span class="o">=</span><span class="s2">"MyPass%23123"</span>
+</span></code></pre></div>
+</div>
+</section>
+</section>
+<section id="troubleshooting">
+<h2>Troubleshooting<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#troubleshooting" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#troubleshooting'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p><strong>“Table ‘conversation_states’ does not exist”</strong></p>
+<p>Run the SQL schema from <code class="docutils literal notranslate"><span class="pre">docs/db_setup/conversation_states.sql</span></code> against your database.</p>
+<p><strong>Connection errors with Supabase</strong></p>
+<ul class="simple">
+<li><p>Verify you’re using the correct connection string format (Session Pooler for IPv4)</p></li>
+<li><p>Check that your password is URL-encoded if it contains special characters</p></li>
+<li><p>Ensure your Supabase project hasn’t paused due to inactivity (free tier)</p></li>
+</ul>
+<p><strong>Permission errors</strong></p>
+<p>Ensure your database user has the following permissions:</p>
+<div class="highlight-sql notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="k">GRANT</span><span class="w"> </span><span class="k">SELECT</span><span class="p">,</span><span class="w"> </span><span class="k">INSERT</span><span class="p">,</span><span class="w"> </span><span class="k">UPDATE</span><span class="p">,</span><span class="w"> </span><span class="k">DELETE</span><span class="w"> </span><span class="k">ON</span><span class="w"> </span><span class="n">conversation_states</span><span class="w"> </span><span class="k">TO</span><span class="w"> </span><span class="n">your_user</span><span class="p">;</span>
+</span></code></pre></div>
+</div>
+<p><strong>State not persisting across requests</strong></p>
+<ul class="simple">
+<li><p>Verify <code class="docutils literal notranslate"><span class="pre">state_storage</span></code> is configured in your <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code></p></li>
+<li><p>Check Plano logs for state storage initialization messages</p></li>
+<li><p>Ensure the client is sending the <code class="docutils literal notranslate"><span class="pre">prev_response_id={$response_id}</span></code> from previous responses</p></li>
+</ul>
+</section>
+<section id="best-practices">
+<h2>Best Practices<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#best-practices" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#best-practices'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<ol class="arabic simple">
+<li><p><strong>Use environment variables for credentials</strong>: Never hardcode database passwords in configuration files.</p></li>
+<li><p><strong>Start with memory storage for development</strong>: Switch to PostgreSQL when moving to production.</p></li>
+<li><p><strong>Implement cleanup policies</strong>: Prevent unbounded growth by regularly archiving or deleting old conversations.</p></li>
+<li><p><strong>Monitor storage usage</strong>: Track conversation state table size and query performance in production.</p></li>
+<li><p><strong>Test failover scenarios</strong>: Ensure your application handles storage backend failures gracefully.</p></li>
+</ol>
+</section>
+<section id="next-steps">
+<h2>Next Steps<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#next-steps" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#next-steps'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<ul class="simple">
+<li><p>Learn more about building <a class="reference internal" href="../concepts/agents.html#agents"><span class="std std-ref">agents</span></a> that leverage conversational state</p></li>
+<li><p>Explore <a class="reference internal" href="../concepts/filter_chain.html#filter-chain"><span class="std std-ref">filter chains</span></a> for enriching conversation context</p></li>
+<li><p>See the <a class="reference internal" href="../concepts/llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">LLM Providers</span></a> guide for configuring model routing</p></li>
+</ul>
+</section>
+</section>
+</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
+<div class="mr-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="prompt_guard.html">
+<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="15 18 9 12 15 6"></polyline>
+</svg>
+        Guardrails
+      </a>
+</div>
+<div class="ml-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../resources/tech_overview/tech_overview.html">
+        Tech Overview
+        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="9 18 15 12 9 6"></polyline>
+</svg>
+</a>
+</div>
+</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
+<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
+<ul>
+<li><a :data-current="activeSection === '#how-it-works'" class="reference internal" href="#how-it-works">How It Works</a></li>
+<li><a :data-current="activeSection === '#configuration-overview'" class="reference internal" href="#configuration-overview">Configuration Overview</a></li>
+<li><a :data-current="activeSection === '#memory-storage-development'" class="reference internal" href="#memory-storage-development">Memory Storage (Development)</a></li>
+<li><a :data-current="activeSection === '#postgresql-storage-production'" class="reference internal" href="#postgresql-storage-production">PostgreSQL Storage (Production)</a><ul>
+<li><a :data-current="activeSection === '#prerequisites'" class="reference internal" href="#prerequisites">Prerequisites</a></li>
+<li><a :data-current="activeSection === '#configuration'" class="reference internal" href="#configuration">Configuration</a></li>
+<li><a :data-current="activeSection === '#supabase-connection-strings'" class="reference internal" href="#supabase-connection-strings">Supabase Connection Strings</a></li>
+</ul>
+</li>
+<li><a :data-current="activeSection === '#troubleshooting'" class="reference internal" href="#troubleshooting">Troubleshooting</a></li>
+<li><a :data-current="activeSection === '#best-practices'" class="reference internal" href="#best-practices">Best Practices</a></li>
+<li><a :data-current="activeSection === '#next-steps'" class="reference internal" href="#next-steps">Next Steps</a></li>
+</ul>
+</div>
+</aside>
+</main>
+</div>
+</div><footer class="py-6 border-t border-border md:py-0">
+<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
+<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
+</div>
+</div>
+</footer>
+</div>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
+<script src="../_static/doctools.js?v=9bcbadda"></script>
+<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
+<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
+<script src="../_static/design-tabs.js?v=f930bc37"></script>
+</body>
+</html>
\ No newline at end of file
diff --git a/includes/llms.txt b/includes/llms.txt
new file mode 100755
index 00000000..2512a90f
--- /dev/null
+++ b/includes/llms.txt
@@ -0,0 +1,5327 @@
+Plano Docs v0.4
+llms.txt (auto-generated)
+Generated (UTC): 2025-12-24T01:15:17.250485+00:00
+
+Table of contents
+- Agents (concepts/agents)
+- Filter Chains (concepts/filter_chain)
+- Listeners (concepts/listeners)
+- Client Libraries (concepts/llm_providers/client_libraries)
+- Model (LLM) Providers (concepts/llm_providers/llm_providers)
+- Model Aliases (concepts/llm_providers/model_aliases)
+- Supported Providers & Configuration (concepts/llm_providers/supported_providers)
+- Prompt Target (concepts/prompt_target)
+- Intro to Plano (get_started/intro_to_plano)
+- Overview (get_started/overview)
+- Quickstart (get_started/quickstart)
+- Function Calling (guides/function_calling)
+- LLM Routing (guides/llm_router)
+- Access Logging (guides/observability/access_logging)
+- Monitoring (guides/observability/monitoring)
+- Observability (guides/observability/observability)
+- Tracing (guides/observability/tracing)
+- Orchestration (guides/orchestration)
+- Guardrails (guides/prompt_guard)
+- Conversational State (guides/state)
+- Welcome to Plano! (index)
+- Configuration Reference (resources/configuration_reference)
+- Deployment (resources/deployment)
+- llms.txt (resources/llms_txt)
+- Bright Staff (resources/tech_overview/model_serving)
+- Request Lifecycle (resources/tech_overview/request_lifecycle)
+- Tech Overview (resources/tech_overview/tech_overview)
+- Threading Model (resources/tech_overview/threading_model)
+
+Agents
+------
+Doc: concepts/agents
+
+Agents
+
+Agents are autonomous systems that handle wide-ranging, open-ended tasks by calling models in a loop until the work is complete. Unlike deterministic prompt targets, agents have access to tools, reason about which actions to take, and adapt their behavior based on intermediate results—making them ideal for complex workflows that require multi-step reasoning, external API calls, and dynamic decision-making.
+
+Plano helps developers build and scale multi-agent systems by managing the orchestration layer—deciding which agent(s) or LLM(s) should handle each request, and in what sequence—while developers focus on implementing agent logic in any language or framework they choose.
+
+Agent Orchestration
+
+Plano-Orchestrator is a family of state-of-the-art routing and orchestration models that decide which agent(s) should handle each request, and in what sequence. Built for real-world multi-agent deployments, it analyzes user intent and conversation context to make precise routing and orchestration decisions while remaining efficient enough for low-latency production use across general chat, coding, and long-context multi-turn conversations.
+
+This allows development teams to:
+
+Scale multi-agent systems: Route requests across multiple specialized agents without hardcoding routing logic in application code.
+
+Improve performance: Direct requests to the most appropriate agent based on intent, reducing unnecessary handoffs and improving response quality.
+
+Enhance debuggability: Centralized routing decisions are observable through Plano’s tracing and logging, making it easier to understand why a particular agent was selected.
+
+Inner Loop vs. Outer Loop
+
+Plano distinguishes between the inner loop (agent implementation logic) and the outer loop (orchestration and routing):
+
+Inner Loop (Agent Logic)
+
+The inner loop is where your agent lives—the business logic that decides which tools to call, how to interpret results, and when the task is complete. You implement this in any language or framework:
+
+Python agents: Using frameworks like LangChain, LlamaIndex, CrewAI, or custom Python code.
+
+JavaScript/TypeScript agents: Using frameworks like LangChain.js or custom Node.js implementations.
+
+Any other AI famreowkr: Agents are just HTTP services that Plano can route to.
+
+Your agent controls:
+
+Which tools or APIs to call in response to a prompt.
+
+How to interpret tool results and decide next steps.
+
+When to call the LLM for reasoning or summarization.
+
+When the task is complete and what response to return.
+
+Making LLM Calls from Agents
+
+When your agent needs to call an LLM for reasoning, summarization, or completion, you should route those calls through Plano’s Model Proxy rather than calling LLM providers directly. This gives you:
+
+Consistent responses: Normalized response formats across all LLM providers, whether you’re using OpenAI, Anthropic, Azure OpenAI, or any OpenAI-compatible provider.
+
+Rich agentic signals: Automatic capture of function calls, tool usage, reasoning steps, and model behavior—surfaced through traces and metrics without instrumenting your agent code.
+
+Smart model routing: Leverage model-based, alias-based, or preference-aligned routing to dynamically select the best model for each task based on cost, performance, or custom policies.
+
+By routing LLM calls through the Model Proxy, your agents remain decoupled from specific providers and can benefit from centralized policy enforcement, observability, and intelligent routing—all managed in the outer loop. For a step-by-step guide, see llm_router in the LLM Router guide.
+
+Outer Loop (Orchestration)
+
+The outer loop is Plano’s orchestration layer—it manages the lifecycle of requests across agents and LLMs:
+
+Intent analysis: Plano-Orchestrator analyzes incoming prompts to determine user intent and conversation context.
+
+Routing decisions: Routes requests to the appropriate agent(s) or LLM(s) based on capabilities, context, and availability.
+
+Sequencing: Determines whether multiple agents need to collaborate and in what order.
+
+Lifecycle management: Handles retries, failover, circuit breaking, and load balancing across agent instances.
+
+By managing the outer loop, Plano allows you to:
+
+Add new agents without changing routing logic in existing agents.
+
+Run multiple versions or variants of agents for A/B testing or canary deployments.
+
+Apply consistent filter chains (guardrails, context enrichment) before requests reach agents.
+
+Monitor and debug multi-agent workflows through centralized observability.
+
+Key Benefits
+
+Language and framework agnostic: Write agents in any language; Plano orchestrates them via HTTP.
+
+Reduced complexity: Agents focus on task logic; Plano handles routing, retries, and cross-cutting concerns.
+
+Better observability: Centralized tracing shows which agents were called, in what sequence, and why.
+
+Easier scaling: Add more agent instances or new agent types without refactoring existing code.
+
+---
+
+Filter Chains
+-------------
+Doc: concepts/filter_chain
+
+Filter Chains
+
+Filter chains are Plano’s way of capturing reusable workflow steps in the dataplane, without duplication and coupling logic into application code. A filter chain is an ordered list of mutations that a request flows through before reaching its final destination —such as an agent, an LLM, or a tool backend. Each filter is a network-addressable service/path that can:
+
+Inspect the incoming prompt, metadata, and conversation state.
+
+Mutate or enrich the request (for example, rewrite queries or build context).
+
+Short-circuit the flow and return a response early (for example, block a request on a compliance failure).
+
+Emit structured logs and traces so you can debug and continuously improve your agents.
+
+In other words, filter chains provide a lightweight programming model over HTTP for building reusable steps
+in your agent architectures.
+
+Typical Use Cases
+
+Without a dataplane programming model, teams tend to spread logic like query rewriting, compliance checks,
+context building, and routing decisions across many agents and frameworks. This quickly becomes hard to reason
+about and even harder to evolve.
+
+Filter chains show up most often in patterns like:
+
+Guardrails and Compliance: Enforcing content policies, stripping or masking sensitive data, and blocking obviously unsafe or off-topic requests before they reach an agent.
+
+Query rewriting, RAG, and Memory: Rewriting user queries for retrieval, normalizing entities, and assembling RAG context envelopes while pulling in relevant memory (for example, conversation history, user profiles, or prior tool results) before calling a model or tool.
+
+Cross-cutting Observability: Injecting correlation IDs, sampling traces, or logging enriched request metadata at consistent points in the request path.
+
+Because these behaviors live in the dataplane rather than inside individual agents, you define them once, attach them to many agents and prompt targets, and can add, remove, or reorder them without changing application code.
+
+Configuration example
+
+The example below shows a configuration where an agent uses a filter chain with two filters: a query rewriter,
+and a context builder that prepares retrieval context before the agent runs.
+
+Example Configuration
+
+version: v0.3.0
+
+agents:
+  - id: rag_agent
+    url: http://host.docker.internal:10505
+
+filters:
+  - id: query_rewriter
+    url: http://host.docker.internal:10501
+    # type: mcp # default is mcp
+    # transport: streamable-http # default is streamable-http
+    # tool: query_rewriter # default name is the filter id
+  - id: context_builder
+    url: http://host.docker.internal:10502
+
+model_providers:
+  - model: openai/gpt-4o-mini
+    access_key: $OPENAI_API_KEY
+    default: true
+  - model: openai/gpt-4o
+    access_key: $OPENAI_API_KEY
+
+model_aliases:
+  fast-llm:
+    target: gpt-4o-mini
+  smart-llm:
+    target: gpt-4o
+
+listeners:
+  - type: agent
+    name: agent_1
+    port: 8001
+    router: arch_agent_router
+    agents:
+      - id: rag_agent
+        description: virtual assistant for retrieval augmented generation tasks
+        filter_chain:
+          - query_rewriter
+          - context_builder
+tracing:
+  random_sampling: 100
+
+
+In this setup:
+
+The filters section defines the reusable filters, each running as its own HTTP/MCP service.
+
+The listeners section wires the rag_agent behind an agent listener and attaches a filter_chain with query_rewriter followed by context_builder.
+
+When a request arrives at agent_1, Plano executes the filters in order before handing control to rag_agent.
+
+Filter Chain Programming Model (HTTP and MCP)
+
+Filters are implemented as simple RESTful endpoints reachable via HTTP. If you want to use the Model Context Protocol (MCP), you can configure that as well, which makes it easy to write filters in any language. However, you can also write a filter as a plain HTTP service.
+
+When defining a filter in Plano configuration, the following fields are optional:
+
+type: Controls the filter runtime. Use mcp for Model Context Protocol filters, or http for plain HTTP filters. Defaults to mcp.
+
+transport: Controls how Plano talks to the filter (defaults to streamable-http for efficient streaming interactions over HTTP). You can omit this for standard HTTP transport.
+
+tool: Names the MCP tool Plano will invoke (by default, the filter id). You can omit this if the tool name matches your filter id.
+
+In practice, you typically only need to specify id and url to get started. Plano’s sensible defaults mean a filter can be as simple as an HTTP endpoint. If you want to customize the runtime or protocol, those fields are there, but they’re optional.
+
+Filters communicate the outcome of their work via HTTP status codes:
+
+HTTP 200 (Success): The filter successfully processed the request. If the filter mutated the request (e.g., rewrote a query or enriched context), those mutations are passed downstream.
+
+HTTP 4xx (User Error): The request violates a filter’s rules or constraints—for example, content moderation policies or compliance checks. The request is terminated, and the error is returned to the caller. This is not a fatal error; it represents expected user-facing policy enforcement.
+
+HTTP 5xx (Fatal Error): An unexpected failure in the filter itself (for example, a crash or misconfiguration). Plano will surface the error back to the caller and record it in logs and traces.
+
+This semantics allows filters to enforce guardrails and policies (4xx) without blocking the entire system, while still surfacing critical failures (5xx) for investigation.
+
+If any filter fails or decides to terminate the request early (for example, after a policy violation), Plano will
+surface that outcome back to the caller and record it in logs and traces. This makes filter chains a safe and
+powerful abstraction for evolving your agent workflows over time.
+
+---
+
+Listeners
+---------
+Doc: concepts/listeners
+
+Listeners
+
+Listeners are a top-level primitive in Plano that bind network traffic to the dataplane. They simplify the
+configuration required to accept incoming connections from downstream clients (edge) and to expose a unified egress
+endpoint for calls from your applications to upstream LLMs.
+
+Plano builds on Envoy’s Listener subsystem to streamline connection management for developers. It hides most of
+Envoy’s complexity behind sensible defaults and a focused configuration surface, so you can bind listeners without
+deep knowledge of Envoy’s configuration model while still getting secure, reliable, and performant connections.
+
+Listeners are modular building blocks: you can configure only inbound listeners (for edge proxying and guardrails),
+only outbound/model-proxy listeners (for LLM routing from your services), or both together. This lets you fit Plano
+cleanly into existing architectures, whether you need it at the edge, behind the firewall, or across the full
+request path.
+
+Network Topology
+
+The diagram below shows how inbound and outbound traffic flow through Plano and how listeners relate to agents,
+prompt targets, and upstream LLMs:
+
+
+
+Inbound (Agent & Prompt Target)
+
+Developers configure inbound listeners to accept connections from clients such as web frontends, backend
+services, or other gateways. An inbound listener acts as the primary entry point for prompt traffic, handling
+initial connection setup, TLS termination, guardrails, and forwarding incoming traffic to the appropriate prompt
+targets or agents.
+
+There are two primary types of inbound connections exposed via listeners:
+
+Agent Inbound (Edge): Clients (web/mobile apps or other services) connect to Plano, send prompts, and receive
+responses. This is typically your public/edge listener where Plano applies guardrails, routing, and orchestration
+before returning results to the caller.
+
+Prompt Target Inbound (Edge): Your application server calls Plano’s internal listener targeting
+prompt targets that can invoke tools and LLMs directly on its behalf.
+
+Inbound listeners are where you attach Filter Chains so that safety and context-building happen
+consistently at the edge.
+
+Outbound (Model Proxy & Egress)
+
+Plano also exposes an egress listener that your applications call when sending requests to upstream LLM providers
+or self-hosted models. From your application’s perspective this looks like a single OpenAI-compatible HTTP endpoint
+(for example, http://127.0.0.1:12000/v1), while Plano handles provider selection, retries, and failover behind
+the scenes.
+
+Under the hood, Plano opens outbound HTTP(S) connections to upstream LLM providers using its unified API surface and
+smart model routing. For more details on how Plano talks to models and how providers are configured, see
+LLM providers.
+
+Configure Listeners
+
+Listeners are configured via the listeners block in your Plano configuration. You can define one or more inbound
+listeners (for example, type:edge) or one or more outbound/model listeners (for example, type:model), or both
+in the same deployment.
+
+To configure an inbound (edge) listener, add a listeners block to your configuration file and define at least one
+listener with address, port, and protocol details:
+
+Example Configuration
+
+version: v0.2.0
+
+listeners:
+  ingress_traffic:
+    address: 0.0.0.0
+    port: 10000
+
+# Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way
+model_providers:
+  - access_key: $OPENAI_API_KEY
+    model: openai/gpt-4o
+    default: true
+
+
+
+When you start Plano, you specify a listener address/port that you want to bind downstream. Plano also exposes a
+predefined internal listener (127.0.0.1:12000) that you can use to proxy egress calls originating from your
+application to LLMs (API-based or hosted) via prompt targets.
+
+---
+
+Client Libraries
+----------------
+Doc: concepts/llm_providers/client_libraries
+
+Client Libraries
+
+Plano provides a unified interface that works seamlessly with multiple client libraries and tools. You can use your preferred client library without changing your existing code - just point it to Plano’s gateway endpoints.
+
+Supported Clients
+
+OpenAI SDK - Full compatibility with OpenAI’s official client
+
+Anthropic SDK - Native support for Anthropic’s client library
+
+cURL - Direct HTTP requests for any programming language
+
+Custom HTTP Clients - Any HTTP client that supports REST APIs
+
+Gateway Endpoints
+
+Plano exposes three main endpoints:
+
+
+
+
+
+Endpoint
+
+Purpose
+
+http://127.0.0.1:12000/v1/chat/completions
+
+OpenAI-compatible chat completions (LLM Gateway)
+
+http://127.0.0.1:12000/v1/responses
+
+OpenAI Responses API with conversational state management (LLM Gateway)
+
+http://127.0.0.1:12000/v1/messages
+
+Anthropic-compatible messages (LLM Gateway)
+
+OpenAI (Python) SDK
+
+The OpenAI SDK works with any provider through Plano’s OpenAI-compatible endpoint.
+
+Installation:
+
+pip install openai
+
+Basic Usage:
+
+from openai import OpenAI
+
+# Point to Plano's LLM Gateway
+client = OpenAI(
+    api_key="test-key",  # Can be any value for local testing
+    base_url="http://127.0.0.1:12000/v1"
+)
+
+# Use any model configured in your arch_config.yaml
+completion = client.chat.completions.create(
+    model="gpt-4o-mini",  # Or use :ref:`model aliases <model_aliases>` like "fast-model"
+    max_tokens=50,
+    messages=[
+        {
+            "role": "user",
+            "content": "Hello, how are you?"
+        }
+    ]
+)
+
+print(completion.choices[0].message.content)
+
+Streaming Responses:
+
+from openai import OpenAI
+
+client = OpenAI(
+    api_key="test-key",
+    base_url="http://127.0.0.1:12000/v1"
+)
+
+stream = client.chat.completions.create(
+    model="gpt-4o-mini",
+    max_tokens=50,
+    messages=[
+        {
+            "role": "user",
+            "content": "Tell me a short story"
+        }
+    ],
+    stream=True
+)
+
+# Collect streaming chunks
+for chunk in stream:
+    if chunk.choices[0].delta.content:
+        print(chunk.choices[0].delta.content, end="")
+
+Using with Non-OpenAI Models:
+
+The OpenAI SDK can be used with any provider configured in Plano:
+
+# Using Claude model through OpenAI SDK
+completion = client.chat.completions.create(
+    model="claude-3-5-sonnet-20241022",
+    max_tokens=50,
+    messages=[
+        {
+            "role": "user",
+            "content": "Explain quantum computing briefly"
+        }
+    ]
+)
+
+# Using Ollama model through OpenAI SDK
+completion = client.chat.completions.create(
+    model="llama3.1",
+    max_tokens=50,
+    messages=[
+        {
+            "role": "user",
+            "content": "What's the capital of France?"
+        }
+    ]
+)
+
+OpenAI Responses API (Conversational State)
+
+The OpenAI Responses API (v1/responses) enables multi-turn conversations with automatic state management. Plano handles conversation history for you, so you don’t need to manually include previous messages in each request.
+
+See managing_conversational_state for detailed configuration and storage backend options.
+
+Installation:
+
+pip install openai
+
+Basic Multi-Turn Conversation:
+
+from openai import OpenAI
+
+# Point to Plano's LLM Gateway
+client = OpenAI(
+    api_key="test-key",
+    base_url="http://127.0.0.1:12000/v1"
+)
+
+# First turn - creates a new conversation
+response = client.chat.completions.create(
+    model="gpt-4o-mini",
+    messages=[
+        {"role": "user", "content": "My name is Alice"}
+    ]
+)
+
+# Extract response_id for conversation continuity
+response_id = response.id
+print(f"Assistant: {response.choices[0].message.content}")
+
+# Second turn - continues the conversation
+# Plano automatically retrieves and merges previous context
+response = client.chat.completions.create(
+    model="gpt-4o-mini",
+    messages=[
+        {"role": "user", "content": "What's my name?"}
+    ],
+    metadata={"response_id": response_id}  # Reference previous conversation
+)
+
+print(f"Assistant: {response.choices[0].message.content}")
+# Output: "Your name is Alice"
+
+Using with Any Provider:
+
+The Responses API works with any LLM provider configured in Plano:
+
+# Multi-turn conversation with Claude
+response = client.chat.completions.create(
+    model="claude-3-5-sonnet-20241022",
+    messages=[
+        {"role": "user", "content": "Let's discuss quantum physics"}
+    ]
+)
+
+response_id = response.id
+
+# Continue conversation - Plano manages state regardless of provider
+response = client.chat.completions.create(
+    model="claude-3-5-sonnet-20241022",
+    messages=[
+        {"role": "user", "content": "Tell me more about entanglement"}
+    ],
+    metadata={"response_id": response_id}
+)
+
+Key Benefits:
+
+Reduced payload size: No need to send full conversation history in each request
+
+Provider flexibility: Use any configured LLM provider with state management
+
+Automatic context merging: Plano handles conversation continuity behind the scenes
+
+Production-ready storage: Configure PostgreSQL or memory storage based on your needs
+
+Anthropic (Python) SDK
+
+The Anthropic SDK works with any provider through Plano’s Anthropic-compatible endpoint.
+
+Installation:
+
+pip install anthropic
+
+Basic Usage:
+
+import anthropic
+
+# Point to Plano's LLM Gateway
+client = anthropic.Anthropic(
+    api_key="test-key",  # Can be any value for local testing
+    base_url="http://127.0.0.1:12000"
+)
+
+# Use any model configured in your arch_config.yaml
+message = client.messages.create(
+    model="claude-3-5-sonnet-20241022",
+    max_tokens=50,
+    messages=[
+        {
+            "role": "user",
+            "content": "Hello, please respond briefly!"
+        }
+    ]
+)
+
+print(message.content[0].text)
+
+Streaming Responses:
+
+import anthropic
+
+client = anthropic.Anthropic(
+    api_key="test-key",
+    base_url="http://127.0.0.1:12000"
+)
+
+with client.messages.stream(
+    model="claude-3-5-sonnet-20241022",
+    max_tokens=50,
+    messages=[
+        {
+            "role": "user",
+            "content": "Tell me about artificial intelligence"
+        }
+    ]
+) as stream:
+    # Collect text deltas
+    for text in stream.text_stream:
+        print(text, end="")
+
+    # Get final assembled message
+    final_message = stream.get_final_message()
+    final_text = "".join(block.text for block in final_message.content if block.type == "text")
+
+Using with Non-Anthropic Models:
+
+The Anthropic SDK can be used with any provider configured in Plano:
+
+# Using OpenAI model through Anthropic SDK
+message = client.messages.create(
+    model="gpt-4o-mini",
+    max_tokens=50,
+    messages=[
+        {
+            "role": "user",
+            "content": "Explain machine learning in simple terms"
+        }
+    ]
+)
+
+# Using Ollama model through Anthropic SDK
+message = client.messages.create(
+    model="llama3.1",
+    max_tokens=50,
+    messages=[
+        {
+            "role": "user",
+            "content": "What is Python programming?"
+        }
+    ]
+)
+
+cURL Examples
+
+For direct HTTP requests or integration with any programming language:
+
+OpenAI-Compatible Endpoint:
+
+# Basic request
+curl -X POST http://127.0.0.1:12000/v1/chat/completions \
+  -H "Content-Type: application/json" \
+  -H "Authorization: Bearer test-key" \
+  -d '{
+    "model": "gpt-4o-mini",
+    "messages": [
+      {"role": "user", "content": "Hello!"}
+    ],
+    "max_tokens": 50
+  }'
+
+# Using :ref:`model aliases <model_aliases>`
+curl -X POST http://127.0.0.1:12000/v1/chat/completions \
+  -H "Content-Type: application/json" \
+  -d '{
+    "model": "fast-model",
+    "messages": [
+      {"role": "user", "content": "Summarize this text..."}
+    ],
+    "max_tokens": 100
+  }'
+
+# Streaming request
+curl -X POST http://127.0.0.1:12000/v1/chat/completions \
+  -H "Content-Type: application/json" \
+  -d '{
+    "model": "gpt-4o-mini",
+    "messages": [
+      {"role": "user", "content": "Tell me a story"}
+    ],
+    "stream": true,
+    "max_tokens": 200
+  }'
+
+Anthropic-Compatible Endpoint:
+
+# Basic request
+curl -X POST http://127.0.0.1:12000/v1/messages \
+  -H "Content-Type: application/json" \
+  -H "x-api-key: test-key" \
+  -H "anthropic-version: 2023-06-01" \
+  -d '{
+    "model": "claude-3-5-sonnet-20241022",
+    "max_tokens": 50,
+    "messages": [
+      {"role": "user", "content": "Hello Claude!"}
+    ]
+  }'
+
+Cross-Client Compatibility
+
+One of Plano’s key features is cross-client compatibility. You can:
+
+Use OpenAI SDK with Claude Models:
+
+# OpenAI client calling Claude model
+from openai import OpenAI
+
+client = OpenAI(base_url="http://127.0.0.1:12000/v1", api_key="test")
+
+response = client.chat.completions.create(
+    model="claude-3-5-sonnet-20241022",  # Claude model
+    messages=[{"role": "user", "content": "Hello"}]
+)
+
+Use Anthropic SDK with OpenAI Models:
+
+# Anthropic client calling OpenAI model
+import anthropic
+
+client = anthropic.Anthropic(base_url="http://127.0.0.1:12000", api_key="test")
+
+response = client.messages.create(
+    model="gpt-4o-mini",  # OpenAI model
+    max_tokens=50,
+    messages=[{"role": "user", "content": "Hello"}]
+)
+
+Mix and Match with Model Aliases:
+
+# Same code works with different underlying models
+def ask_question(client, question):
+    return client.chat.completions.create(
+        model="reasoning-model",  # Alias could point to any provider
+        messages=[{"role": "user", "content": question}]
+    )
+
+# Works regardless of what "reasoning-model" actually points to
+openai_client = OpenAI(base_url="http://127.0.0.1:12000/v1", api_key="test")
+response = ask_question(openai_client, "Solve this math problem...")
+
+Error Handling
+
+OpenAI SDK Error Handling:
+
+from openai import OpenAI
+import openai
+
+client = OpenAI(base_url="http://127.0.0.1:12000/v1", api_key="test")
+
+try:
+    completion = client.chat.completions.create(
+        model="nonexistent-model",
+        messages=[{"role": "user", "content": "Hello"}]
+    )
+except openai.NotFoundError as e:
+    print(f"Model not found: {e}")
+except openai.APIError as e:
+    print(f"API error: {e}")
+
+Anthropic SDK Error Handling:
+
+import anthropic
+
+client = anthropic.Anthropic(base_url="http://127.0.0.1:12000", api_key="test")
+
+try:
+    message = client.messages.create(
+        model="nonexistent-model",
+        max_tokens=50,
+        messages=[{"role": "user", "content": "Hello"}]
+    )
+except anthropic.NotFoundError as e:
+    print(f"Model not found: {e}")
+except anthropic.APIError as e:
+    print(f"API error: {e}")
+
+Best Practices
+
+Use Model Aliases:
+Instead of hardcoding provider-specific model names, use semantic aliases:
+
+# Good - uses semantic alias
+model = "fast-model"
+
+# Less ideal - hardcoded provider model
+model = "openai/gpt-4o-mini"
+
+Environment-Based Configuration:
+Use different model aliases for different environments:
+
+import os
+
+# Development uses cheaper/faster models
+model = os.getenv("MODEL_ALIAS", "dev.chat.v1")
+
+response = client.chat.completions.create(
+    model=model,
+    messages=[{"role": "user", "content": "Hello"}]
+)
+
+Graceful Fallbacks:
+Implement fallback logic for better reliability:
+
+def chat_with_fallback(client, messages, primary_model="smart-model", fallback_model="fast-model"):
+    try:
+        return client.chat.completions.create(model=primary_model, messages=messages)
+    except Exception as e:
+        print(f"Primary model failed, trying fallback: {e}")
+        return client.chat.completions.create(model=fallback_model, messages=messages)
+
+See Also
+
+supported_providers - Configure your providers and see available models
+
+model_aliases - Create semantic model names
+
+llm_router - Intelligent routing capabilities
+
+---
+
+Model (LLM) Providers
+---------------------
+Doc: concepts/llm_providers/llm_providers
+
+Model (LLM) Providers
+
+Model Providers are a top-level primitive in Plano, helping developers centrally define, secure, observe,
+and manage the usage of their models. Plano builds on Envoy’s reliable cluster subsystem to manage egress traffic to models, which includes intelligent routing, retry and fail-over mechanisms,
+ensuring high availability and fault tolerance. This abstraction also enables developers to seamlessly switch between model providers or upgrade model versions, simplifying the integration and scaling of models across applications.
+
+Today, we are enable you to connect to 15+ different AI providers through a unified interface with advanced routing and management capabilities.
+Whether you’re using OpenAI, Anthropic, Azure OpenAI, local Ollama models, or any OpenAI-compatible provider, Plano provides seamless integration with enterprise-grade features.
+
+Please refer to the quickstart guide here to configure and use LLM providers via common client libraries like OpenAI and Anthropic Python SDKs, or via direct HTTP/cURL requests.
+
+Core Capabilities
+
+Multi-Provider Support
+Connect to any combination of providers simultaneously (see supported_providers for full details):
+
+First-Class Providers: Native integrations with OpenAI, Anthropic, DeepSeek, Mistral, Groq, Google Gemini, Together AI, xAI, Azure OpenAI, and Ollama
+
+OpenAI-Compatible Providers: Any provider implementing the OpenAI Chat Completions API standard
+
+Intelligent Routing
+Three powerful routing approaches to optimize model selection:
+
+Model-based Routing: Direct routing to specific models using provider/model names (see supported_providers)
+
+Alias-based Routing: Semantic routing using custom aliases (see model_aliases)
+
+Preference-aligned Routing: Intelligent routing using the Plano-Router model (see preference_aligned_routing)
+
+Unified Client Interface
+Use your preferred client library without changing existing code (see client_libraries for details):
+
+OpenAI Python SDK: Full compatibility with all providers
+
+Anthropic Python SDK: Native support with cross-provider capabilities
+
+cURL & HTTP Clients: Direct REST API access for any programming language
+
+Custom Integrations: Standard HTTP interfaces for seamless integration
+
+Key Benefits
+
+Provider Flexibility: Switch between providers without changing client code
+
+Three Routing Methods: Choose from model-based, alias-based, or preference-aligned routing (using Plano-Router-1.5B) strategies
+
+Cost Optimization: Route requests to cost-effective models based on complexity
+
+Performance Optimization: Use fast models for simple tasks, powerful models for complex reasoning
+
+Environment Management: Configure different models for different environments
+
+Future-Proof: Easy to add new providers and upgrade models
+
+Common Use Cases
+
+Development Teams
+- Use aliases like dev.chat.v1 and prod.chat.v1 for environment-specific models
+- Route simple queries to fast/cheap models, complex tasks to powerful models
+- Test new models safely using canary deployments (coming soon)
+
+Production Applications
+- Implement fallback strategies across multiple providers for reliability
+- Use intelligent routing to optimize cost and performance automatically
+- Monitor usage patterns and model performance across providers
+
+Enterprise Deployments
+- Connect to both cloud providers and on-premises models (Ollama, custom deployments)
+- Apply consistent security and governance policies across all providers
+- Scale across regions using different provider endpoints
+
+Advanced Features
+
+preference_aligned_routing - Learn about preference-aligned dynamic routing and intelligent model selection
+
+Getting Started
+
+Dive into specific areas based on your needs:
+
+---
+
+Model Aliases
+-------------
+Doc: concepts/llm_providers/model_aliases
+
+Model Aliases
+
+Model aliases provide semantic, version-controlled names for your models, enabling cleaner client code, easier model management, and advanced routing capabilities. Instead of using provider-specific model names like gpt-4o-mini or claude-3-5-sonnet-20241022, you can create meaningful aliases like fast-model or arch.summarize.v1.
+
+Benefits of Model Aliases:
+
+Semantic Naming: Use descriptive names that reflect the model’s purpose
+
+Version Control: Implement versioning schemes (e.g., v1, v2) for model upgrades
+
+Environment Management: Different aliases can point to different models across environments
+
+Client Simplification: Clients use consistent, meaningful names regardless of underlying provider
+
+Advanced Routing (Coming Soon): Enable guardrails, fallbacks, and traffic splitting at the alias level
+
+Basic Configuration
+
+Simple Alias Mapping
+
+Basic Model Aliases
+
+llm_providers:
+  - model: openai/gpt-4o-mini
+    access_key: $OPENAI_API_KEY
+
+  - model: openai/gpt-4o
+    access_key: $OPENAI_API_KEY
+
+  - model: anthropic/claude-3-5-sonnet-20241022
+    access_key: $ANTHROPIC_API_KEY
+
+  - model: ollama/llama3.1
+    base_url: http://host.docker.internal:11434
+
+# Define aliases that map to the models above
+model_aliases:
+  # Semantic versioning approach
+  arch.summarize.v1:
+    target: gpt-4o-mini
+
+  arch.reasoning.v1:
+    target: gpt-4o
+
+  arch.creative.v1:
+    target: claude-3-5-sonnet-20241022
+
+  # Functional aliases
+  fast-model:
+    target: gpt-4o-mini
+
+  smart-model:
+    target: gpt-4o
+
+  creative-model:
+    target: claude-3-5-sonnet-20241022
+
+  # Local model alias
+  local-chat:
+    target: llama3.1
+
+Using Aliases
+
+Client Code Examples
+
+Once aliases are configured, clients can use semantic names instead of provider-specific model names:
+
+Python Client Usage
+
+from openai import OpenAI
+
+client = OpenAI(base_url="http://127.0.0.1:12000/")
+
+# Use semantic alias instead of provider model name
+response = client.chat.completions.create(
+    model="arch.summarize.v1",  # Points to gpt-4o-mini
+    messages=[{"role": "user", "content": "Summarize this document..."}]
+)
+
+# Switch to a different capability
+response = client.chat.completions.create(
+    model="arch.reasoning.v1",  # Points to gpt-4o
+    messages=[{"role": "user", "content": "Solve this complex problem..."}]
+)
+
+cURL Example
+
+curl -X POST http://127.0.0.1:12000/v1/chat/completions \
+  -H "Content-Type: application/json" \
+  -d '{
+    "model": "fast-model",
+    "messages": [{"role": "user", "content": "Hello!"}]
+  }'
+
+Naming Best Practices
+
+Semantic Versioning
+
+Use version numbers for backward compatibility and gradual model upgrades:
+
+model_aliases:
+  # Current production version
+  arch.summarize.v1:
+    target: gpt-4o-mini
+
+  # Beta version for testing
+  arch.summarize.v2:
+    target: gpt-4o
+
+  # Stable alias that always points to latest
+  arch.summarize.latest:
+    target: gpt-4o-mini
+
+Purpose-Based Naming
+
+Create aliases that reflect the intended use case:
+
+model_aliases:
+  # Task-specific
+  code-reviewer:
+    target: gpt-4o
+
+  document-summarizer:
+    target: gpt-4o-mini
+
+  creative-writer:
+    target: claude-3-5-sonnet-20241022
+
+  data-analyst:
+    target: gpt-4o
+
+Environment-Specific Aliases
+
+Different environments can use different underlying models:
+
+model_aliases:
+  # Development environment - use faster/cheaper models
+  dev.chat.v1:
+    target: gpt-4o-mini
+
+  # Production environment - use more capable models
+  prod.chat.v1:
+    target: gpt-4o
+
+  # Staging environment - test new models
+  staging.chat.v1:
+    target: claude-3-5-sonnet-20241022
+
+Advanced Features (Coming Soon)
+
+The following features are planned for future releases of model aliases:
+
+Guardrails Integration
+
+Apply safety, cost, or latency rules at the alias level:
+
+Future Feature - Guardrails
+
+model_aliases:
+  arch.reasoning.v1:
+    target: gpt-oss-120b
+    guardrails:
+      max_latency: 5s
+      max_cost_per_request: 0.10
+      block_categories: ["jailbreak", "PII"]
+      content_filters:
+        - type: "profanity"
+        - type: "sensitive_data"
+
+Fallback Chains
+
+Provide a chain of models if the primary target fails or hits quota limits:
+
+Future Feature - Fallbacks
+
+model_aliases:
+  arch.summarize.v1:
+    target: gpt-4o-mini
+    fallbacks:
+      - target: llama3.1
+        conditions: ["quota_exceeded", "timeout"]
+      - target: claude-3-haiku-20240307
+        conditions: ["primary_and_first_fallback_failed"]
+
+Traffic Splitting & Canary Deployments
+
+Distribute traffic across multiple models for A/B testing or gradual rollouts:
+
+Future Feature - Traffic Splitting
+
+model_aliases:
+  arch.v1:
+    targets:
+      - model: llama3.1
+        weight: 80
+      - model: gpt-4o-mini
+        weight: 20
+
+  # Canary deployment
+  arch.experimental.v1:
+    targets:
+      - model: gpt-4o      # Current stable
+        weight: 95
+      - model: o1-preview  # New model being tested
+        weight: 5
+
+Load Balancing
+
+Distribute requests across multiple instances of the same model:
+
+Future Feature - Load Balancing
+
+model_aliases:
+  high-throughput-chat:
+    load_balance:
+      algorithm: "round_robin"  # or "least_connections", "weighted"
+    targets:
+      - model: gpt-4o-mini
+        endpoint: "https://api-1.example.com"
+      - model: gpt-4o-mini
+        endpoint: "https://api-2.example.com"
+      - model: gpt-4o-mini
+        endpoint: "https://api-3.example.com"
+
+Validation Rules
+
+Alias names must be valid identifiers (alphanumeric, dots, hyphens, underscores)
+
+Target models must be defined in the llm_providers section
+
+Circular references between aliases are not allowed
+
+Weights in traffic splitting must sum to 100
+
+See Also
+
+llm_providers - Learn about configuring LLM providers
+
+llm_router - Understand how aliases work with intelligent routing
+
+---
+
+Supported Providers & Configuration
+-----------------------------------
+Doc: concepts/llm_providers/supported_providers
+
+Supported Providers & Configuration
+
+Plano provides first-class support for multiple LLM providers through native integrations and OpenAI-compatible interfaces. This comprehensive guide covers all supported providers, their available chat models, and detailed configuration instructions.
+
+Model Support: Plano supports all chat models from each provider, not just the examples shown in this guide. The configurations below demonstrate common models for reference, but you can use any chat model available from your chosen provider.
+
+Please refer to the quuickstart guide here to configure and use LLM providers via common client libraries like OpenAI and Anthropic Python SDKs, or via direct HTTP/cURL requests.
+
+Configuration Structure
+
+All providers are configured in the llm_providers section of your plano_config.yaml file:
+
+llm_providers:
+  # Provider configurations go here
+  - model: provider/model-name
+    access_key: $API_KEY
+    # Additional provider-specific options
+
+Common Configuration Fields:
+
+model: Provider prefix and model name (format: provider/model-name)
+
+access_key: API key for authentication (supports environment variables)
+
+default: Mark a model as the default (optional, boolean)
+
+name: Custom name for the provider instance (optional)
+
+base_url: Custom endpoint URL (required for some providers, optional for others - see base_url_details)
+
+Provider Categories
+
+First-Class Providers
+Native integrations with built-in support for provider-specific features and authentication.
+
+OpenAI-Compatible Providers
+Any provider that implements the OpenAI API interface can be configured using custom endpoints.
+
+Supported API Endpoints
+
+Plano supports the following standardized endpoints across providers:
+
+
+
+
+
+
+
+Endpoint
+
+Purpose
+
+Supported Clients
+
+/v1/chat/completions
+
+OpenAI-style chat completions
+
+OpenAI SDK, cURL, custom clients
+
+/v1/messages
+
+Anthropic-style messages
+
+Anthropic SDK, cURL, custom clients
+
+/v1/responses
+
+Unified response endpoint for agentic apps
+
+All SDKs, cURL, custom clients
+
+First-Class Providers
+
+OpenAI
+
+Provider Prefix: openai/
+
+API Endpoint: /v1/chat/completions
+
+Authentication: API Key - Get your OpenAI API key from OpenAI Platform.
+
+Supported Chat Models: All OpenAI chat models including GPT-5.2, GPT-5, GPT-4o, and all future releases.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+GPT-5.2
+
+openai/gpt-5.2
+
+Next-generation model (use any model name from OpenAI’s API)
+
+GPT-5
+
+openai/gpt-5
+
+Latest multimodal model
+
+GPT-4o mini
+
+openai/gpt-4o-mini
+
+Fast, cost-effective model
+
+GPT-4o
+
+openai/gpt-4o
+
+High-capability reasoning model
+
+o3-mini
+
+openai/o3-mini
+
+Reasoning-focused model (preview)
+
+o3
+
+openai/o3
+
+Advanced reasoning model (preview)
+
+Configuration Examples:
+
+llm_providers:
+  # Latest models (examples - use any OpenAI chat model)
+  - model: openai/gpt-5.2
+    access_key: $OPENAI_API_KEY
+    default: true
+
+  - model: openai/gpt-5
+    access_key: $OPENAI_API_KEY
+
+  # Use any model name from OpenAI's API
+  - model: openai/gpt-4o
+    access_key: $OPENAI_API_KEY
+
+Anthropic
+
+Provider Prefix: anthropic/
+
+API Endpoint: /v1/messages
+
+Authentication: API Key - Get your Anthropic API key from Anthropic Console.
+
+Supported Chat Models: All Anthropic Claude models including Claude Sonnet 4.5, Claude Opus 4.5, Claude Haiku 4.5, and all future releases.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+Claude Opus 4.5
+
+anthropic/claude-opus-4-5
+
+Most capable model for complex tasks
+
+Claude Sonnet 4.5
+
+anthropic/claude-sonnet-4-5
+
+Balanced performance model
+
+Claude Haiku 4.5
+
+anthropic/claude-haiku-4-5
+
+Fast and efficient model
+
+Claude Sonnet 3.5
+
+anthropic/claude-sonnet-3-5
+
+Complex agents and coding
+
+Configuration Examples:
+
+llm_providers:
+  # Latest models (examples - use any Anthropic chat model)
+  - model: anthropic/claude-opus-4-5
+    access_key: $ANTHROPIC_API_KEY
+
+  - model: anthropic/claude-sonnet-4-5
+    access_key: $ANTHROPIC_API_KEY
+
+  # Use any model name from Anthropic's API
+  - model: anthropic/claude-haiku-4-5
+    access_key: $ANTHROPIC_API_KEY
+
+DeepSeek
+
+Provider Prefix: deepseek/
+
+API Endpoint: /v1/chat/completions
+
+Authentication: API Key - Get your DeepSeek API key from DeepSeek Platform.
+
+Supported Chat Models: All DeepSeek chat models including DeepSeek-Chat, DeepSeek-Coder, and all future releases.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+DeepSeek Chat
+
+deepseek/deepseek-chat
+
+General purpose chat model
+
+DeepSeek Coder
+
+deepseek/deepseek-coder
+
+Code-specialized model
+
+Configuration Examples:
+
+llm_providers:
+  - model: deepseek/deepseek-chat
+    access_key: $DEEPSEEK_API_KEY
+
+  - model: deepseek/deepseek-coder
+    access_key: $DEEPSEEK_API_KEY
+
+Mistral AI
+
+Provider Prefix: mistral/
+
+API Endpoint: /v1/chat/completions
+
+Authentication: API Key - Get your Mistral API key from Mistral AI Console.
+
+Supported Chat Models: All Mistral chat models including Mistral Large, Mistral Small, Ministral, and all future releases.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+Mistral Large
+
+mistral/mistral-large-latest
+
+Most capable model
+
+Mistral Medium
+
+mistral/mistral-medium-latest
+
+Balanced performance
+
+Mistral Small
+
+mistral/mistral-small-latest
+
+Fast and efficient
+
+Ministral 3B
+
+mistral/ministral-3b-latest
+
+Compact model
+
+Configuration Examples:
+Configuration Examples:
+
+llm_providers:
+  - model: mistral/mistral-large-latest
+    access_key: $MISTRAL_API_KEY
+
+  - model: mistral/mistral-small-latest
+    access_key: $MISTRAL_API_KEY
+
+Groq
+
+Provider Prefix: groq/
+
+API Endpoint: /openai/v1/chat/completions (transformed internally)
+
+Authentication: API Key - Get your Groq API key from Groq Console.
+
+Supported Chat Models: All Groq chat models including Llama 4, GPT OSS, Mixtral, Gemma, and all future releases.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+Llama 4 Maverick 17B
+
+groq/llama-4-maverick-17b-128e-instruct
+
+Fast inference Llama model
+
+Llama 4 Scout 8B
+
+groq/llama-4-scout-8b-128e-instruct
+
+Smaller Llama model
+
+GPT OSS 20B
+
+groq/gpt-oss-20b
+
+Open source GPT model
+
+Configuration Examples:
+
+llm_providers:
+  - model: groq/llama-4-maverick-17b-128e-instruct
+    access_key: $GROQ_API_KEY
+
+  - model: groq/llama-4-scout-8b-128e-instruct
+    access_key: $GROQ_API_KEY
+
+  - model: groq/gpt-oss-20b
+    access_key: $GROQ_API_KEY
+
+Google Gemini
+
+Provider Prefix: gemini/
+
+API Endpoint: /v1beta/openai/chat/completions (transformed internally)
+
+Authentication: API Key - Get your Google AI API key from Google AI Studio.
+
+Supported Chat Models: All Google Gemini chat models including Gemini 3 Pro, Gemini 3 Flash, and all future releases.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+Gemini 3 Pro
+
+gemini/gemini-3-pro
+
+Advanced reasoning and creativity
+
+Gemini 3 Flash
+
+gemini/gemini-3-flash
+
+Fast and efficient model
+
+Configuration Examples:
+
+llm_providers:
+  - model: gemini/gemini-3-pro
+    access_key: $GOOGLE_API_KEY
+
+  - model: gemini/gemini-3-flash
+    access_key: $GOOGLE_API_KEY
+
+Together AI
+
+Provider Prefix: together_ai/
+
+API Endpoint: /v1/chat/completions
+
+Authentication: API Key - Get your Together AI API key from Together AI Settings.
+
+Supported Chat Models: All Together AI chat models including Llama, CodeLlama, Mixtral, Qwen, and hundreds of other open-source models.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+Meta Llama 2 7B
+
+together_ai/meta-llama/Llama-2-7b-chat-hf
+
+Open source chat model
+
+Meta Llama 2 13B
+
+together_ai/meta-llama/Llama-2-13b-chat-hf
+
+Larger open source model
+
+Code Llama 34B
+
+together_ai/codellama/CodeLlama-34b-Instruct-hf
+
+Code-specialized model
+
+Configuration Examples:
+
+llm_providers:
+  - model: together_ai/meta-llama/Llama-2-7b-chat-hf
+    access_key: $TOGETHER_API_KEY
+
+  - model: together_ai/codellama/CodeLlama-34b-Instruct-hf
+    access_key: $TOGETHER_API_KEY
+
+xAI
+
+Provider Prefix: xai/
+
+API Endpoint: /v1/chat/completions
+
+Authentication: API Key - Get your xAI API key from xAI Console.
+
+Supported Chat Models: All xAI chat models including Grok Beta and all future releases.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+Grok Beta
+
+xai/grok-beta
+
+Conversational AI model
+
+Configuration Examples:
+
+llm_providers:
+  - model: xai/grok-beta
+    access_key: $XAI_API_KEY
+
+Moonshot AI
+
+Provider Prefix: moonshotai/
+
+API Endpoint: /v1/chat/completions
+
+Authentication: API Key - Get your Moonshot AI API key from Moonshot AI Platform.
+
+Supported Chat Models: All Moonshot AI chat models including Kimi K2, Moonshot v1, and all future releases.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+Kimi K2 Preview
+
+moonshotai/kimi-k2-0905-preview
+
+Foundation model optimized for agentic tasks with 32B activated parameters
+
+Moonshot v1 32K
+
+moonshotai/moonshot-v1-32k
+
+Extended context model with 32K tokens
+
+Moonshot v1 128K
+
+moonshotai/moonshot-v1-128k
+
+Long context model with 128K tokens
+
+Configuration Examples:
+
+llm_providers:
+  # Latest K2 models for agentic tasks
+  - model: moonshotai/kimi-k2-0905-preview
+    access_key: $MOONSHOTAI_API_KEY
+
+  # V1 models with different context lengths
+  - model: moonshotai/moonshot-v1-32k
+    access_key: $MOONSHOTAI_API_KEY
+
+  - model: moonshotai/moonshot-v1-128k
+    access_key: $MOONSHOTAI_API_KEY
+
+Zhipu AI
+
+Provider Prefix: zhipu/
+
+API Endpoint: /api/paas/v4/chat/completions
+
+Authentication: API Key - Get your Zhipu AI API key from Zhipu AI Platform.
+
+Supported Chat Models: All Zhipu AI GLM models including GLM-4, GLM-4 Flash, and all future releases.
+
+
+
+
+
+
+
+Model Name
+
+Model ID for Config
+
+Description
+
+GLM-4.6
+
+zhipu/glm-4.6
+
+Latest and most capable GLM model with enhanced reasoning abilities
+
+GLM-4.5
+
+zhipu/glm-4.5
+
+High-performance model with multimodal capabilities
+
+GLM-4.5 Air
+
+zhipu/glm-4.5-air
+
+Lightweight and fast model optimized for efficiency
+
+Configuration Examples:
+
+llm_providers:
+  # Latest GLM models
+  - model: zhipu/glm-4.6
+    access_key: $ZHIPU_API_KEY
+
+  - model: zhipu/glm-4.5
+    access_key: $ZHIPU_API_KEY
+
+  - model: zhipu/glm-4.5-air
+    access_key: $ZHIPU_API_KEY
+
+Providers Requiring Base URL
+
+The following providers require a base_url parameter to be configured. For detailed information on base URL configuration including path prefix behavior and examples, see base_url_details.
+
+Azure OpenAI
+
+Provider Prefix: azure_openai/
+
+API Endpoint: /openai/deployments/{deployment-name}/chat/completions (constructed automatically)
+
+Authentication: API Key + Base URL - Get your Azure OpenAI API key from Azure Portal → Your OpenAI Resource → Keys and Endpoint.
+
+Supported Chat Models: All Azure OpenAI chat models including GPT-4o, GPT-4, GPT-3.5-turbo deployed in your Azure subscription.
+
+llm_providers:
+  # Single deployment
+  - model: azure_openai/gpt-4o
+    access_key: $AZURE_OPENAI_API_KEY
+    base_url: https://your-resource.openai.azure.com
+
+  # Multiple deployments
+  - model: azure_openai/gpt-4o-mini
+    access_key: $AZURE_OPENAI_API_KEY
+    base_url: https://your-resource.openai.azure.com
+
+Amazon Bedrock
+
+Provider Prefix: amazon_bedrock/
+
+API Endpoint: Plano automatically constructs the endpoint as:
+
+Non-streaming: /model/{model-id}/converse
+
+Streaming: /model/{model-id}/converse-stream
+
+Authentication: AWS Bearer Token + Base URL - Get your API Keys from AWS Bedrock Console → Discover → API Keys.
+
+Supported Chat Models: All Amazon Bedrock foundation models including Claude (Anthropic), Nova (Amazon), Llama (Meta), Mistral AI, and Cohere Command models.
+
+llm_providers:
+  # Amazon Nova models
+  - model: amazon_bedrock/us.amazon.nova-premier-v1:0
+    access_key: $AWS_BEARER_TOKEN_BEDROCK
+    base_url: https://bedrock-runtime.us-west-2.amazonaws.com
+    default: true
+
+  - model: amazon_bedrock/us.amazon.nova-pro-v1:0
+    access_key: $AWS_BEARER_TOKEN_BEDROCK
+    base_url: https://bedrock-runtime.us-west-2.amazonaws.com
+
+  # Claude on Bedrock
+  - model: amazon_bedrock/us.anthropic.claude-3-5-sonnet-20241022-v2:0
+    access_key: $AWS_BEARER_TOKEN_BEDROCK
+    base_url: https://bedrock-runtime.us-west-2.amazonaws.com
+
+Qwen (Alibaba)
+
+Provider Prefix: qwen/
+
+API Endpoint: /v1/chat/completions
+
+Authentication: API Key + Base URL - Get your Qwen API key from Qwen Portal → Your Qwen Resource → Keys and Endpoint.
+
+Supported Chat Models: All Qwen chat models including Qwen3, Qwen3-Coder and all future releases.
+
+llm_providers:
+  # Single deployment
+  - model: qwen/qwen3
+    access_key: $DASHSCOPE_API_KEY
+    base_url: https://dashscope.aliyuncs.com
+
+  # Multiple deployments
+  - model: qwen/qwen3-coder
+    access_key: $DASHSCOPE_API_KEY
+    base_url: "https://dashscope-intl.aliyuncs.com"
+
+Ollama
+
+Provider Prefix: ollama/
+
+API Endpoint: /v1/chat/completions (Ollama’s OpenAI-compatible endpoint)
+
+Authentication: None (Base URL only) - Install Ollama from Ollama.com and pull your desired models.
+
+Supported Chat Models: All chat models available in your local Ollama installation. Use ollama list to see installed models.
+
+llm_providers:
+  # Local Ollama installation
+  - model: ollama/llama3.1
+    base_url: http://localhost:11434
+
+  # Ollama in Docker (from host)
+  - model: ollama/codellama
+    base_url: http://host.docker.internal:11434
+
+OpenAI-Compatible Providers
+
+Supported Models: Any chat models from providers that implement the OpenAI Chat Completions API standard.
+
+For providers that implement the OpenAI API but aren’t natively supported:
+
+llm_providers:
+  # Generic OpenAI-compatible provider
+  - model: custom-provider/custom-model
+    base_url: https://api.customprovider.com
+    provider_interface: openai
+    access_key: $CUSTOM_API_KEY
+
+  # Local deployment
+  - model: local/llama2-7b
+    base_url: http://localhost:8000
+    provider_interface: openai
+
+
+
+Base URL Configuration
+
+The base_url parameter allows you to specify custom endpoints for model providers. It supports both hostname and path components, enabling flexible routing to different API endpoints.
+
+Format: <scheme>://<hostname>[:<port>][/<path>]
+
+Components:
+
+scheme: http or https
+
+hostname: API server hostname or IP address
+
+port: Optional, defaults to 80 for http, 443 for https
+
+path: Optional path prefix that replaces the provider’s default API path
+
+How Path Prefixes Work:
+
+When you include a path in base_url, it replaces the provider’s default path prefix while preserving the endpoint suffix:
+
+Without path prefix: Uses the provider’s default path structure
+
+With path prefix: Your custom path replaces the provider’s default prefix, then the endpoint suffix is appended
+
+Configuration Examples:
+
+llm_providers:
+  # Simple hostname only - uses provider's default path
+  - model: zhipu/glm-4.6
+    access_key: $ZHIPU_API_KEY
+    base_url: https://api.z.ai
+    # Results in: https://api.z.ai/api/paas/v4/chat/completions
+
+  # With custom path prefix - replaces provider's default path
+  - model: zhipu/glm-4.6
+    access_key: $ZHIPU_API_KEY
+    base_url: https://api.z.ai/api/coding/paas/v4
+    # Results in: https://api.z.ai/api/coding/paas/v4/chat/completions
+
+  # Azure with custom path
+  - model: azure_openai/gpt-4
+    access_key: $AZURE_API_KEY
+    base_url: https://mycompany.openai.azure.com/custom/deployment/path
+    # Results in: https://mycompany.openai.azure.com/custom/deployment/path/chat/completions
+
+  # Behind a proxy or API gateway
+  - model: openai/gpt-4o
+    access_key: $OPENAI_API_KEY
+    base_url: https://proxy.company.com/ai-gateway/openai
+    # Results in: https://proxy.company.com/ai-gateway/openai/chat/completions
+
+  # Local endpoint with custom port
+  - model: ollama/llama3.1
+    base_url: http://localhost:8080
+    # Results in: http://localhost:8080/v1/chat/completions
+
+  # Custom provider with path prefix
+  - model: vllm/custom-model
+    access_key: $VLLM_API_KEY
+    base_url: https://vllm.example.com/models/v2
+    provider_interface: openai
+    # Results in: https://vllm.example.com/models/v2/chat/completions
+
+Advanced Configuration
+
+Multiple Provider Instances
+
+Configure multiple instances of the same provider:
+
+llm_providers:
+  # Production OpenAI
+  - model: openai/gpt-4o
+    access_key: $OPENAI_PROD_KEY
+    name: openai-prod
+
+  # Development OpenAI (different key/quota)
+  - model: openai/gpt-4o-mini
+    access_key: $OPENAI_DEV_KEY
+    name: openai-dev
+
+Default Model Configuration
+
+Mark one model as the default for fallback scenarios:
+
+llm_providers:
+  - model: openai/gpt-4o-mini
+    access_key: $OPENAI_API_KEY
+    default: true  # Used when no specific model is requested
+
+Routing Preferences
+
+Configure routing preferences for dynamic model selection:
+
+llm_providers:
+  - model: openai/gpt-5.2
+    access_key: $OPENAI_API_KEY
+    routing_preferences:
+      - name: complex_reasoning
+        description: deep analysis, mathematical problem solving, and logical reasoning
+      - name: code_review
+        description: reviewing and analyzing existing code for bugs and improvements
+
+  - model: anthropic/claude-sonnet-4-5
+    access_key: $ANTHROPIC_API_KEY
+    routing_preferences:
+      - name: creative_writing
+        description: creative content generation, storytelling, and writing assistance
+
+Model Selection Guidelines
+
+For Production Applications:
+- High Performance: OpenAI GPT-5.2, Anthropic Claude Sonnet 4.5
+- Cost-Effective: OpenAI GPT-5, Anthropic Claude Haiku 4.5
+- Code Tasks: DeepSeek Coder, Together AI Code Llama
+- Local Deployment: Ollama with Llama 3.1 or Code Llama
+
+For Development/Testing:
+- Fast Iteration: Groq models (optimized inference)
+- Local Testing: Ollama models
+- Cost Control: Smaller models like GPT-4o or Mistral Small
+
+See Also
+
+client_libraries - Using different client libraries with providers
+
+model_aliases - Creating semantic model names
+
+llm_router - Setting up intelligent routing
+
+client_libraries - Using different client libraries
+
+model_aliases - Creating semantic model names
+
+---
+
+Prompt Target
+-------------
+Doc: concepts/prompt_target
+
+Prompt Target
+
+A Prompt Target is a deterministic, task-specific backend function or API endpoint that your application calls via Plano.
+Unlike agents (which handle wide-ranging, open-ended tasks), prompt targets are designed for focused, specific workloads where Plano can add value through input clarification and validation.
+
+Plano helps by:
+
+Clarifying and validating input: Plano enriches incoming prompts with metadata (e.g., detecting follow-ups or clarifying requests) and can extract structured parameters from natural language before passing them to your backend.
+
+Enabling high determinism: Since the task is specific and well-defined, Plano can reliably extract the information your backend needs without ambiguity.
+
+Reducing backend work: Your backend receives clean, validated, structured inputs—so you can focus on business logic instead of parsing and validation.
+
+For example, a prompt target might be “schedule a meeting” (specific task, deterministic inputs like date, time, attendees) or “retrieve documents” (well-defined RAG query with clear intent). Prompt targets are typically called from your application code via Plano’s internal listener.
+
+
+
+
+
+Capability
+
+Description
+
+Intent Recognition
+
+Identify the purpose of a user prompt.
+
+Parameter Extraction
+
+Extract necessary data from the prompt.
+
+Invocation
+
+Call relevant backend agents or tools (APIs).
+
+Response Handling
+
+Process and return responses to the user.
+
+Key Features
+
+Below are the key features of prompt targets that empower developers to build efficient, scalable, and personalized GenAI solutions:
+
+Design Scenarios: Define prompt targets to effectively handle specific agentic scenarios.
+
+Input Management: Specify required and optional parameters for each target.
+
+Tools Integration: Seamlessly connect prompts to backend APIs or functions.
+
+Error Handling: Direct errors to designated handlers for streamlined troubleshooting.
+
+Multi-Turn Support: Manage follow-up prompts and clarifications in conversational flows.
+
+Basic Configuration
+
+Configuring prompt targets involves defining them in Plano’s configuration file. Each Prompt target specifies how a particular type of prompt should be handled, including the endpoint to invoke and any parameters required. A prompt target configuration includes the following elements:
+
+vale Vale.Spelling = NO
+
+name: A unique identifier for the prompt target.
+
+description: A brief explanation of what the prompt target does.
+
+endpoint: Required if you want to call a tool or specific API. name and path http_method are the three attributes of the endpoint.
+
+parameters (Optional): A list of parameters to extract from the prompt.
+
+
+
+Defining Parameters
+
+Parameters are the pieces of information that Plano needs to extract from the user’s prompt to perform the desired action.
+Each parameter can be marked as required or optional. Here is a full list of parameter attributes that Plano can support:
+
+
+
+
+
+Attribute
+
+Description
+
+name (req.)
+
+Specifies name of the parameter.
+
+description (req.)
+
+Provides a human-readable explanation of the parameter’s purpose.
+
+type (req.)
+
+Specifies the data type. Supported types include: int, str, float, bool, list, set, dict, tuple
+
+in_path
+
+Indicates whether the parameter is part of the path in the endpoint url. Valid values: true or false
+
+default
+
+Specifies a default value for the parameter if not provided by the user.
+
+format
+
+Specifies a format for the parameter value. For example: 2019-12-31 for a date value.
+
+enum
+
+Lists of allowable values for the parameter with data type matching the type attribute. Usage Example: enum: ["celsius`", "fahrenheit"]
+
+items
+
+Specifies the attribute of the elements when type equals list, set, dict, tuple. Usage Example: items: {"type": "str"}
+
+required
+
+Indicates whether the parameter is mandatory or optional. Valid values: true or false
+
+Example Configuration For Tools
+
+Tools and Function Calling Configuration Example
+
+prompt_targets:
+  - name: get_weather
+    description: Get the current weather for a location
+    parameters:
+      - name: location
+        description: The city and state, e.g. San Francisco, New York
+        type: str
+        required: true
+      - name: unit
+        description: The unit of temperature
+        type: str
+        default: fahrenheit
+        enum: [celsius, fahrenheit]
+    endpoint:
+      name: api_server
+      path: /weather
+
+
+
+Multi-Turn
+
+Developers often struggle to efficiently handle
+follow-up or clarification questions. Specifically, when users ask for changes or additions to previous responses, it requires developers to
+re-write prompts using LLMs with precise prompt engineering techniques. This process is slow, manual, error prone and adds latency and token cost for
+common scenarios that can be managed more efficiently.
+
+Plano is highly capable of accurately detecting and processing prompts in multi-turn scenarios so that you can buil fast and accurate agents in minutes.
+Below are some cnversational examples that you can build via Plano. Each example is enriched with annotations (via ** [Plano] ** ) that illustrates how Plano
+processess conversational messages on your behalf.
+
+Example 1: Adjusting Retrieval
+
+User: What are the benefits of renewable energy?
+**[Plano]**: Check if there is an available <prompt_target> that can handle this user query.
+**[Plano]**: Found "get_info_for_energy_source" prompt_target in arch_config.yaml. Forward prompt to the endpoint configured in "get_info_for_energy_source"
+...
+Assistant: Renewable energy reduces greenhouse gas emissions, lowers air pollution, and provides sustainable power sources like solar and wind.
+
+User: Include cost considerations in the response.
+**[Plano]**: Follow-up detected. Forward prompt history to the "get_info_for_energy_source" prompt_target and post the following parameters consideration="cost"
+...
+Assistant: Renewable energy reduces greenhouse gas emissions, lowers air pollution, and provides sustainable power sources like solar and wind. While the initial setup costs can be high, long-term savings from reduced fuel expenses and government incentives make it cost-effective.
+
+Example 2: Switching Intent
+
+User: What are the symptoms of diabetes?
+**[Plano]**: Check if there is an available <prompt_target> that can handle this user query.
+**[Plano]**: Found "diseases_symptoms" prompt_target in arch_config.yaml. Forward disease=diabeteres to "diseases_symptoms" prompt target
+...
+Assistant: Common symptoms include frequent urination, excessive thirst, fatigue, and blurry vision.
+
+User: How is it diagnosed?
+**[Plano]**: New intent detected.
+**[Plano]**: Found "disease_diagnoses" prompt_target in arch_config.yaml. Forward disease=diabeteres to "disease_diagnoses" prompt target
+...
+Assistant: Diabetes is diagnosed through blood tests like fasting blood sugar, A1C, or an oral glucose tolerance test.
+
+Build Multi-Turn RAG Apps
+
+The following section describes how you can easilly add support for multi-turn scenarios via Plano. You process and manage multi-turn prompts
+just like you manage single-turn ones. Plano handles the conpleixity of detecting the correct intent based on the last user prompt and
+the covnersational history, extracts relevant parameters needed by downstream APIs, and dipatches calls to any upstream LLMs to summarize the
+response from your APIs.
+
+
+
+Step 1: Define Plano Config
+
+Plano Config
+
+version: v0.1
+listener:
+  address: 127.0.0.1
+  port: 8080 #If you configure port 443, you'll need to update the listener with tls_certificates
+  message_format: huggingface
+
+# Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way
+llm_providers:
+  - name: OpenAI
+    provider: openai
+    access_key: $OPENAI_API_KEY
+    model: gpt-3.5-turbo
+    default: true
+
+# default system prompt used by all prompt targets
+system_prompt: |
+   You are a helpful assistant and can offer information about energy sources. You will get a JSON object with energy_source and consideration fields. Focus on answering using those fields
+
+prompt_targets:
+  - name: get_info_for_energy_source
+    description: get information about an energy source
+    parameters:
+      - name: energy_source
+        type: str
+        description: a source of energy
+        required: true
+        enum: [renewable, fossil]
+      - name: consideration
+        type: str
+        description: a specific type of consideration for an energy source
+        enum: [cost, economic, technology]
+    endpoint:
+      name: rag_energy_source_agent
+      path: /agent/energy_source_info
+      http_method: POST
+
+
+Step 2: Process Request in Flask
+
+Once the prompt targets are configured as above, handle parameters across multi-turn as if its a single-turn request
+
+Parameter handling with Flask
+
+import os
+import gradio as gr
+
+from fastapi import FastAPI, HTTPException
+from pydantic import BaseModel
+from typing import Optional
+from openai import OpenAI
+from common import create_gradio_app
+
+app = FastAPI()
+
+
+# Define the request model
+class EnergySourceRequest(BaseModel):
+    energy_source: str
+    consideration: Optional[str] = None
+
+
+class EnergySourceResponse(BaseModel):
+    energy_source: str
+    consideration: Optional[str] = None
+
+
+# Post method for device summary
+@app.post("/agent/energy_source_info")
+def get_workforce(request: EnergySourceRequest):
+    """
+    Endpoint to get details about energy source
+    """
+    considertion = "You don't have any specific consideration. Feel free to talk in a more open ended fashion"
+
+    if request.consideration is not None:
+        considertion = f"Add specific focus on the following consideration when you summarize the content for the energy source: {request.consideration}"
+
+    response = {
+        "energy_source": request.energy_source,
+        "consideration": considertion,
+    }
+    return response
+
+
+Demo App
+
+For your convenience, we’ve built a demo app
+that you can test and modify locally for multi-turn RAG scenarios.
+
+
+
+Example multi-turn user conversation showing adjusting retrieval
+
+Summary
+
+By carefully designing prompt targets as deterministic, task-specific entry points, you ensure that prompts are routed to the right workload, necessary parameters are cleanly extracted and validated, and backend services are invoked with structured inputs. This clear separation between prompt handling and business logic simplifies your architecture, makes behavior more predictable and testable, and improves the scalability and maintainability of your agentic applications.
+
+---
+
+Intro to Plano
+--------------
+Doc: get_started/intro_to_plano
+
+Intro to Plano
+
+Building agentic demos is easy. Delivering agentic applications safely, reliably, and repeatably to production is hard. After a quick hack, you end up building the “hidden AI middleware” to reach production: routing logic to reach the right agent, guardrail hooks for safety and moderation, evaluation and observability glue for continuous learning, and model/provider quirks — scattered across frameworks and application code.
+
+Plano solves this by moving core delivery concerns into a unified, out-of-process dataplane. Core capabilities:
+
+🚦 Orchestration: Low-latency orchestration between agents, and add new agents without changing app code. When routing lives inside app code, it becomes hard to evolve and easy to duplicate. Moving orchestration into a centrally managed dataplane lets you change strategies without touching your agents, improving performance and reducing maintenance burden while avoiding tight coupling.
+
+🛡️ Guardrails & Memory Hooks: Apply jailbreak protection, content policies, and context workflows (e.g., rewriting, retrieval, redaction) once via Filter Chains at the dataplane. Instead of re-implementing these in every agentic service, you get centralized governance, reduced code duplication, and consistent behavior across your stack.
+
+🔗 Model Agility: Route by model, alias (semantic names), or automatically via preferences so agents stay decoupled from specific providers. Swap or add models without refactoring prompts, tool-calling, or streaming handlers throughout your codebase by using Plano’s smart routing and unified API.
+
+🕵 Agentic Signals™: Zero-code capture of behavior signals, traces, and metrics consistently across every agent. Rather than stitching together logging and metrics per framework, Plano surfaces traces, token usage, and learning signals in one place so you can iterate safely.
+
+Built by core contributors to the widely adopted Envoy Proxy <https://www.envoyproxy.io/>_, Plano gives you a production‑grade foundation for agentic applications. It helps developers stay focused on the core logic of their agents, helps product teams shorten feedback loops for learning, and helps engineering teams  standardize policy and safety across agents and LLMs. Plano is grounded in open protocols (de facto: OpenAI‑style v1/responses, de jure: MCP) and proven patterns like sidecar deployments, so it plugs in cleanly while remaining robust, scalable, and flexible.
+
+In practice, achieving the above goal is incredibly difficult. Plano attempts to do so by providing the following high level features:
+
+
+
+High-level network flow of where Plano sits in your agentic stack. Designed for both ingress and egress prompt traffic.
+
+Engineered with Task-Specific LLMs (TLMs): Plano is engineered with specialized LLMs that are designed for fast, cost-effective and accurate handling of prompts.
+These LLMs are designed to be best-in-class for critical tasks like:
+
+Agent Orchestration: Plano-Orchestrator is a family of state-of-the-art routing and orchestration models that decide which agent(s) or LLM(s) should handle each request, and in what sequence. Built for real-world multi-agent deployments, it analyzes user intent and conversation context to make precise routing and orchestration decisions while remaining efficient enough for low-latency production use across general chat, coding, and long-context multi-turn conversations.
+
+Function Calling: Plano lets you expose application-specific (API) operations as tools so that your agents can update records, fetch data, or trigger determininistic workflows via prompts. Under the hood this is backed by Arch-Function-Chat; for more details, read Function Calling.
+
+Guardrails: Plano helps you improve the safety of your application by applying prompt guardrails in a centralized way for better governance hygiene.
+With prompt guardrails you can prevent jailbreak attempts present in user’s prompts without having to write a single line of code.
+To learn more about how to configure guardrails available in Plano, read Prompt Guard.
+
+Model Proxy: Plano offers several capabilities for LLM calls originating from your applications, including smart retries on errors from upstream LLMs and automatic cut-over to other LLMs configured in Plano for continuous availability and disaster recovery scenarios. From your application’s perspective you keep using an OpenAI-compatible API, while Plano owns resiliency and failover policies in one place.
+Plano extends Envoy’s cluster subsystem to manage upstream connections to LLMs so that you can build resilient, provider-agnostic AI applications.
+
+Edge Proxy: There is substantial benefit in using the same software at the edge (observability, traffic shaping algorithms, applying guardrails, etc.) as for outbound LLM inference use cases. Plano has the feature set that makes it exceptionally well suited as an edge gateway for AI applications.
+This includes TLS termination, applying guardrails early in the request flow, and intelligently deciding which agent(s) or LLM(s) should handle each request and in what sequence. In practice, you configure listeners and policies once, and every inbound and outbound call flows through the same hardened gateway.
+
+Zero-Code Agent Signals™ & Tracing: Zero-code capture of behavior signals, traces, and metrics consistently across every agent. Plano propagates trace context using the W3C Trace Context standard, specifically through the traceparent header. This allows each component in the system to record its part of the request flow, enabling end-to-end tracing across the entire application. By using OpenTelemetry, Plano ensures that developers can capture this trace data consistently and in a format compatible with various observability tools.
+
+Best-In Class Monitoring: Plano offers several monitoring metrics that help you understand three critical aspects of your application: latency, token usage, and error rates by an upstream LLM provider. Latency measures the speed at which your application is responding to users, which includes metrics like time to first token (TFT), time per output token (TOT) metrics, and the total latency as perceived by users.
+
+Out-of-process architecture, built on Envoy:
+Plano takes a dependency on Envoy and is a self-contained process that is designed to run alongside your application servers. Plano uses Envoy’s HTTP connection management subsystem, HTTP L7 filtering and telemetry capabilities to extend the functionality exclusively for prompts and LLMs.
+This gives Plano several advantages:
+
+Plano builds on Envoy’s proven success. Envoy is used at massive scale by the leading technology companies of our time including AirBnB, Dropbox, Google, Reddit, Stripe, etc. Its battle tested and scales linearly with usage and enables developers to focus on what really matters: application features and business logic.
+
+Plano works with any application language. A single Plano deployment can act as gateway for AI applications written in Python, Java, C++, Go, Php, etc.
+
+Plano can be deployed and upgraded quickly across your infrastructure transparently without the horrid pain of deploying library upgrades in your applications.
+
+---
+
+Overview
+--------
+Doc: get_started/overview
+
+Overview
+
+Plano is delivery infrastructure for agentic apps. A models-native proxy server and data plane designed to help you build agents faster, and deliver them reliably to production.
+
+Plano pulls out the rote plumbing work (the “hidden AI middleware”) and decouples you from brittle, ever‑changing framework abstractions. It centralizes what shouldn’t be bespoke in every codebase like agent routing and orchestration, rich agentic signals and traces for continuous improvement, guardrail filters for safety and moderation, and smart LLM routing APIs for UX and DX agility. Use any language or AI framework, and ship agents to production faster with Plano.
+
+Built by core contributors to the widely adopted Envoy Proxy, Plano gives you a production‑grade foundation for agentic applications. It helps developers stay focused on the core logic of their agents, helps product teams shorten feedback loops for learning, and helps engineering teams  standardize policy and safety across agents and LLMs. Plano is grounded in open protocols (de facto: OpenAI‑style v1/responses, de jure: MCP) and proven patterns like sidecar deployments, so it plugs in cleanly while remaining robust, scalable, and flexible.
+
+In this documentation, you’ll learn how to set up Plano quickly, trigger API calls via prompts, apply guardrails without tight coupling with application code, simplify model and provider integration, and improve observability — so that you can focus on what matters most: the core product logic of your agents.
+
+
+
+High-level network flow of where Plano sits in your agentic stack. Designed for both ingress and egress traffic.
+
+Get Started
+
+This section introduces you to Plano and helps you get set up quickly:
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-apps" viewBox="0 0 16 16" aria-hidden="true"><path d="M1.5 3.25c0-.966.784-1.75 1.75-1.75h2.5c.966 0 1.75.784 1.75 1.75v2.5A1.75 1.75 0 0 1 5.75 7.5h-2.5A1.75 1.75 0 0 1 1.5 5.75Zm7 0c0-.966.784-1.75 1.75-1.75h2.5c.966 0 1.75.784 1.75 1.75v2.5a1.75 1.75 0 0 1-1.75 1.75h-2.5A1.75 1.75 0 0 1 8.5 5.75Zm-7 7c0-.966.784-1.75 1.75-1.75h2.5c.966 0 1.75.784 1.75 1.75v2.5a1.75 1.75 0 0 1-1.75 1.75h-2.5a1.75 1.75 0 0 1-1.75-1.75Zm7 0c0-.966.784-1.75 1.75-1.75h2.5c.966 0 1.75.784 1.75 1.75v2.5a1.75 1.75 0 0 1-1.75 1.75h-2.5a1.75 1.75 0 0 1-1.75-1.75ZM3.25 3a.25.25 0 0 0-.25.25v2.5c0 .138.112.25.25.25h2.5A.25.25 0 0 0 6 5.75v-2.5A.25.25 0 0 0 5.75 3Zm7 0a.25.25 0 0 0-.25.25v2.5c0 .138.112.25.25.25h2.5a.25.25 0 0 0 .25-.25v-2.5a.25.25 0 0 0-.25-.25Zm-7 7a.25.25 0 0 0-.25.25v2.5c0 .138.112.25.25.25h2.5a.25.25 0 0 0 .25-.25v-2.5a.25.25 0 0 0-.25-.25Zm7 0a.25.25 0 0 0-.25.25v2.5c0 .138.112.25.25.25h2.5a.25.25 0 0 0 .25-.25v-2.5a.25.25 0 0 0-.25-.25Z"></path></svg> Overview
+
+Overview of Plano and Doc navigation
+
+overview.html
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-book" viewBox="0 0 16 16" aria-hidden="true"><path d="M0 1.75A.75.75 0 0 1 .75 1h4.253c1.227 0 2.317.59 3 1.501A3.743 3.743 0 0 1 11.006 1h4.245a.75.75 0 0 1 .75.75v10.5a.75.75 0 0 1-.75.75h-4.507a2.25 2.25 0 0 0-1.591.659l-.622.621a.75.75 0 0 1-1.06 0l-.622-.621A2.25 2.25 0 0 0 5.258 13H.75a.75.75 0 0 1-.75-.75Zm7.251 10.324.004-5.073-.002-2.253A2.25 2.25 0 0 0 5.003 2.5H1.5v9h3.757a3.75 3.75 0 0 1 1.994.574ZM8.755 4.75l-.004 7.322a3.752 3.752 0 0 1 1.992-.572H14.5v-9h-3.495a2.25 2.25 0 0 0-2.25 2.25Z"></path></svg> Intro to Plano
+
+Explore Plano’s features and developer workflow
+
+intro_to_plano.html
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-rocket" viewBox="0 0 16 16" aria-hidden="true"><path d="M14.064 0h.186C15.216 0 16 .784 16 1.75v.186a8.752 8.752 0 0 1-2.564 6.186l-.458.459c-.314.314-.641.616-.979.904v3.207c0 .608-.315 1.172-.833 1.49l-2.774 1.707a.749.749 0 0 1-1.11-.418l-.954-3.102a1.214 1.214 0 0 1-.145-.125L3.754 9.816a1.218 1.218 0 0 1-.124-.145L.528 8.717a.749.749 0 0 1-.418-1.11l1.71-2.774A1.748 1.748 0 0 1 3.31 4h3.204c.288-.338.59-.665.904-.979l.459-.458A8.749 8.749 0 0 1 14.064 0ZM8.938 3.623h-.002l-.458.458c-.76.76-1.437 1.598-2.02 2.5l-1.5 2.317 2.143 2.143 2.317-1.5c.902-.583 1.74-1.26 2.499-2.02l.459-.458a7.25 7.25 0 0 0 2.123-5.127V1.75a.25.25 0 0 0-.25-.25h-.186a7.249 7.249 0 0 0-5.125 2.123ZM3.56 14.56c-.732.732-2.334 1.045-3.005 1.148a.234.234 0 0 1-.201-.064.234.234 0 0 1-.064-.201c.103-.671.416-2.273 1.15-3.003a1.502 1.502 0 1 1 2.12 2.12Zm6.94-3.935c-.088.06-.177.118-.266.175l-2.35 1.521.548 1.783 1.949-1.2a.25.25 0 0 0 .119-.213ZM3.678 8.116 5.2 5.766c.058-.09.117-.178.176-.266H3.309a.25.25 0 0 0-.213.119l-1.2 1.95ZM12 5a1 1 0 1 1-2 0 1 1 0 0 1 2 0Z"></path></svg> Quickstart
+
+Learn how to quickly set up and integrate
+
+quickstart.html
+
+Concepts
+
+Deep dive into essential ideas and mechanisms behind Plano:
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-package" viewBox="0 0 16 16" aria-hidden="true"><path d="m8.878.392 5.25 3.045c.54.314.872.89.872 1.514v6.098a1.75 1.75 0 0 1-.872 1.514l-5.25 3.045a1.75 1.75 0 0 1-1.756 0l-5.25-3.045A1.75 1.75 0 0 1 1 11.049V4.951c0-.624.332-1.201.872-1.514L7.122.392a1.75 1.75 0 0 1 1.756 0ZM7.875 1.69l-4.63 2.685L8 7.133l4.755-2.758-4.63-2.685a.248.248 0 0 0-.25 0ZM2.5 5.677v5.372c0 .09.047.171.125.216l4.625 2.683V8.432Zm6.25 8.271 4.625-2.683a.25.25 0 0 0 .125-.216V5.677L8.75 8.432Z"></path></svg> Agents
+
+Learn about how to build and scale agents with Plano
+
+../concepts/agents.html
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-webhook" viewBox="0 0 16 16" aria-hidden="true"><path d="M5.5 4.25a2.25 2.25 0 0 1 4.5 0 .75.75 0 0 0 1.5 0 3.75 3.75 0 1 0-6.14 2.889l-2.272 4.258a.75.75 0 0 0 1.324.706L7 7.25a.75.75 0 0 0-.309-1.015A2.25 2.25 0 0 1 5.5 4.25Z"></path><path d="M7.364 3.607a.75.75 0 0 1 1.03.257l2.608 4.349a3.75 3.75 0 1 1-.628 6.785.75.75 0 0 1 .752-1.299 2.25 2.25 0 1 0-.033-3.88.75.75 0 0 1-1.03-.256L7.107 4.636a.75.75 0 0 1 .257-1.03Z"></path><path d="M2.9 8.776A.75.75 0 0 1 2.625 9.8 2.25 2.25 0 1 0 6 11.75a.75.75 0 0 1 .75-.751h5.5a.75.75 0 0 1 0 1.5H7.425a3.751 3.751 0 1 1-5.55-3.998.75.75 0 0 1 1.024.274Z"></path></svg> Model Providers
+
+Explore Plano’s LLM integration options
+
+../concepts/llm_providers/llm_providers.html
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-workflow" viewBox="0 0 16 16" aria-hidden="true"><path d="M0 1.75C0 .784.784 0 1.75 0h3.5C6.216 0 7 .784 7 1.75v3.5A1.75 1.75 0 0 1 5.25 7H4v4a1 1 0 0 0 1 1h4v-1.25C9 9.784 9.784 9 10.75 9h3.5c.966 0 1.75.784 1.75 1.75v3.5A1.75 1.75 0 0 1 14.25 16h-3.5A1.75 1.75 0 0 1 9 14.25v-.75H5A2.5 2.5 0 0 1 2.5 11V7h-.75A1.75 1.75 0 0 1 0 5.25Zm1.75-.25a.25.25 0 0 0-.25.25v3.5c0 .138.112.25.25.25h3.5a.25.25 0 0 0 .25-.25v-3.5a.25.25 0 0 0-.25-.25Zm9 9a.25.25 0 0 0-.25.25v3.5c0 .138.112.25.25.25h3.5a.25.25 0 0 0 .25-.25v-3.5a.25.25 0 0 0-.25-.25Z"></path></svg> Prompt Target
+
+Understand how Plano handles prompts
+
+../concepts/prompt_target.html
+
+Guides
+
+Step-by-step tutorials for practical Plano use cases and scenarios:
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-shield-check" viewBox="0 0 16 16" aria-hidden="true"><path d="m8.533.133 5.25 1.68A1.75 1.75 0 0 1 15 3.48V7c0 1.566-.32 3.182-1.303 4.682-.983 1.498-2.585 2.813-5.032 3.855a1.697 1.697 0 0 1-1.33 0c-2.447-1.042-4.049-2.357-5.032-3.855C1.32 10.182 1 8.566 1 7V3.48a1.75 1.75 0 0 1 1.217-1.667l5.25-1.68a1.748 1.748 0 0 1 1.066 0Zm-.61 1.429.001.001-5.25 1.68a.251.251 0 0 0-.174.237V7c0 1.36.275 2.666 1.057 3.859.784 1.194 2.121 2.342 4.366 3.298a.196.196 0 0 0 .154 0c2.245-.957 3.582-2.103 4.366-3.297C13.225 9.666 13.5 8.358 13.5 7V3.48a.25.25 0 0 0-.174-.238l-5.25-1.68a.25.25 0 0 0-.153 0ZM11.28 6.28l-3.5 3.5a.75.75 0 0 1-1.06 0l-1.5-1.5a.749.749 0 0 1 .326-1.275.749.749 0 0 1 .734.215l.97.97 2.97-2.97a.751.751 0 0 1 1.042.018.751.751 0 0 1 .018 1.042Z"></path></svg> Guardrails
+
+Instructions on securing and validating prompts
+
+../guides/prompt_guard.html
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-code-square" viewBox="0 0 16 16" aria-hidden="true"><path d="M0 1.75C0 .784.784 0 1.75 0h12.5C15.216 0 16 .784 16 1.75v12.5A1.75 1.75 0 0 1 14.25 16H1.75A1.75 1.75 0 0 1 0 14.25Zm1.75-.25a.25.25 0 0 0-.25.25v12.5c0 .138.112.25.25.25h12.5a.25.25 0 0 0 .25-.25V1.75a.25.25 0 0 0-.25-.25Zm7.47 3.97a.75.75 0 0 1 1.06 0l2 2a.75.75 0 0 1 0 1.06l-2 2a.749.749 0 0 1-1.275-.326.749.749 0 0 1 .215-.734L10.69 8 9.22 6.53a.75.75 0 0 1 0-1.06ZM6.78 6.53 5.31 8l1.47 1.47a.749.749 0 0 1-.326 1.275.749.749 0 0 1-.734-.215l-2-2a.75.75 0 0 1 0-1.06l2-2a.751.751 0 0 1 1.042.018.751.751 0 0 1 .018 1.042Z"></path></svg> LLM Routing
+
+A guide to effective model selection strategies
+
+../guides/llm_router.html
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-issue-opened" viewBox="0 0 16 16" aria-hidden="true"><path d="M8 9.5a1.5 1.5 0 1 0 0-3 1.5 1.5 0 0 0 0 3Z"></path><path d="M8 0a8 8 0 1 1 0 16A8 8 0 0 1 8 0ZM1.5 8a6.5 6.5 0 1 0 13 0 6.5 6.5 0 0 0-13 0Z"></path></svg> State Management
+
+Learn to manage conversation and application state
+
+../guides/state.html
+
+Build with Plano
+
+End to end examples demonstrating how to build agentic applications using Plano:
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-dependabot" viewBox="0 0 16 16" aria-hidden="true"><path d="M5.75 7.5a.75.75 0 0 1 .75.75v1.5a.75.75 0 0 1-1.5 0v-1.5a.75.75 0 0 1 .75-.75Zm5.25.75a.75.75 0 0 0-1.5 0v1.5a.75.75 0 0 0 1.5 0v-1.5Z"></path><path d="M6.25 0h2A.75.75 0 0 1 9 .75V3.5h3.25a2.25 2.25 0 0 1 2.25 2.25V8h.75a.75.75 0 0 1 0 1.5h-.75v2.75a2.25 2.25 0 0 1-2.25 2.25h-8.5a2.25 2.25 0 0 1-2.25-2.25V9.5H.75a.75.75 0 0 1 0-1.5h.75V5.75A2.25 2.25 0 0 1 3.75 3.5H7.5v-2H6.25a.75.75 0 0 1 0-1.5ZM3 5.75v6.5c0 .414.336.75.75.75h8.5a.75.75 0 0 0 .75-.75v-6.5a.75.75 0 0 0-.75-.75h-8.5a.75.75 0 0 0-.75.75Z"></path></svg> Build Agentic Apps
+
+Discover how to create and manage custom agents within Plano
+
+../get_started/quickstart.html#build-agentic-apps-with-plano
+
+<svg version="1.1" width="1.0em" height="1.0em" class="sd-octicon sd-octicon-stack" viewBox="0 0 16 16" aria-hidden="true"><path d="M7.122.392a1.75 1.75 0 0 1 1.756 0l5.003 2.902c.83.481.83 1.68 0 2.162L8.878 8.358a1.75 1.75 0 0 1-1.756 0L2.119 5.456a1.251 1.251 0 0 1 0-2.162ZM8.125 1.69a.248.248 0 0 0-.25 0l-4.63 2.685 4.63 2.685a.248.248 0 0 0 .25 0l4.63-2.685ZM1.601 7.789a.75.75 0 0 1 1.025-.273l5.249 3.044a.248.248 0 0 0 .25 0l5.249-3.044a.75.75 0 0 1 .752 1.298l-5.248 3.044a1.75 1.75 0 0 1-1.756 0L1.874 8.814A.75.75 0 0 1 1.6 7.789Zm0 3.5a.75.75 0 0 1 1.025-.273l5.249 3.044a.248.248 0 0 0 .25 0l5.249-3.044a.75.75 0 0 1 .752 1.298l-5.248 3.044a1.75 1.75 0 0 1-1.756 0l-5.248-3.044a.75.75 0 0 1-.273-1.025Z"></path></svg> Build Multi-LLM Apps
+
+Learn how to route LLM calls through Plano for enhanced control and observability
+
+../get_started/quickstart.html#use-plano-as-a-model-proxy-gateway
+
+---
+
+Quickstart
+----------
+Doc: get_started/quickstart
+
+Quickstart
+
+Follow this guide to learn how to quickly set up Plano and integrate it into your generative AI applications. You can:
+
+Build agents for multi-step workflows (e.g., travel assistants with flights and hotels).
+
+Call deterministic APIs via prompt targets to turn instructions directly into function calls.
+
+Use Plano as a model proxy (Gateway) to standardize access to multiple LLM providers.
+
+This quickstart assumes basic familiarity with agents and prompt targets from the Concepts section. For background, see Agents and Prompt Target.
+
+The full agent and backend API implementations used here are available in the plano-quickstart repository. This guide focuses on wiring and configuring Plano (orchestration, prompt targets, and the model proxy), not application code.
+
+Prerequisites
+
+Before you begin, ensure you have the following:
+
+Docker System (v24)
+
+Docker Compose (v2.29)
+
+Python (v3.10+)
+
+Plano’s CLI allows you to manage and interact with the Plano efficiently. To install the CLI, simply run the following command:
+
+We recommend that developers create a new Python virtual environment to isolate dependencies before installing Plano. This ensures that plano and its dependencies do not interfere with other packages on your system.
+
+$ python -m venv venv
+$ source venv/bin/activate   # On Windows, use: venv\Scripts\activate
+$ pip install plano==0.4.0
+
+Build Agentic Apps with Plano
+
+Plano helps you build agentic applications in two complementary ways:
+
+Orchestrate agents: Let Plano decide which agent or LLM should handle each request and in what sequence.
+
+Call deterministic backends: Use prompt targets to turn natural-language prompts into structured, validated API calls.
+
+
+
+Building agents with Plano orchestration
+
+Agents are where your business logic lives (the “inner loop”). Plano takes care of the “outer loop”—routing, sequencing, and managing calls across agents and LLMs.
+
+At a high level, building agents with Plano looks like this:
+
+Implement your agent in your framework of choice (Python, JS/TS, etc.), exposing it as an HTTP service.
+
+Route LLM calls through Plano’s Model Proxy, so all models share a consistent interface and observability.
+
+Configure Plano to orchestrate: define which agent(s) can handle which kinds of prompts, and let Plano decide when to call an agent vs. an LLM.
+
+This quickstart uses a simplified version of the Travel Booking Assistant; for the full multi-agent walkthrough, see Orchestration.
+
+Step 1. Minimal orchestration config
+
+Here is a minimal configuration that wires Plano-Orchestrator to two HTTP services: one for flights and one for hotels.
+
+version: v0.1.0
+
+agents:
+  - id: flight_agent
+    url: http://host.docker.internal:10520  # your flights service
+  - id: hotel_agent
+    url: http://host.docker.internal:10530  # your hotels service
+
+model_providers:
+  - model: openai/gpt-4o
+    access_key: $OPENAI_API_KEY
+
+listeners:
+  - type: agent
+    name: travel_assistant
+    port: 8001
+    router: plano_orchestrator_v1
+    agents:
+      - id: flight_agent
+        description: Search for flights and provide flight status.
+      - id: hotel_agent
+        description: Find hotels and check availability.
+
+tracing:
+  random_sampling: 100
+
+Step 2. Start your agents and Plano
+
+Run your flight_agent and hotel_agent services (see Orchestration for a full Travel Booking example), then start Plano with the config above:
+
+$ plano up plano_config.yaml
+
+Plano will start the orchestrator and expose an agent listener on port 8001.
+
+Step 3. Send a prompt and let Plano route
+
+Now send a request to Plano using the OpenAI-compatible chat completions API—the orchestrator will analyze the prompt and route it to the right agent based on intent:
+
+$ curl --header 'Content-Type: application/json' \
+  --data '{"messages": [{"role": "user","content": "Find me flights from SFO to JFK tomorrow"}], "model": "openai/gpt-4o"}' \
+  http://localhost:8001/v1/chat/completions
+
+You can then ask a follow-up like “Also book me a hotel near JFK” and Plano-Orchestrator will route to hotel_agent—your agents stay focused on business logic while Plano handles routing.
+
+
+
+Deterministic API calls with prompt targets
+
+Next, we’ll show Plano’s deterministic API calling using a single prompt target. We’ll build a currency exchange backend powered by https://api.frankfurter.dev/, assuming USD as the base currency.
+
+Step 1. Create plano config file
+
+Create plano_config.yaml file with the following content:
+
+version: v0.1.0
+
+listeners:
+  ingress_traffic:
+    address: 0.0.0.0
+    port: 10000
+    message_format: openai
+    timeout: 30s
+
+ model_providers:
+   - access_key: $OPENAI_API_KEY
+     model: openai/gpt-4o
+
+ system_prompt: |
+   You are a helpful assistant.
+
+ prompt_targets:
+   - name: currency_exchange
+     description: Get currency exchange rate from USD to other currencies
+     parameters:
+       - name: currency_symbol
+         description: the currency that needs conversion
+         required: true
+         type: str
+         in_path: true
+     endpoint:
+       name: frankfurther_api
+       path: /v1/latest?base=USD&symbols={currency_symbol}
+     system_prompt: |
+       You are a helpful assistant. Show me the currency symbol you want to convert from USD.
+
+   - name: get_supported_currencies
+     description: Get list of supported currencies for conversion
+     endpoint:
+       name: frankfurther_api
+       path: /v1/currencies
+
+ endpoints:
+   frankfurther_api:
+     endpoint: api.frankfurter.dev:443
+     protocol: https
+
+Step 2. Start plano with currency conversion config
+
+$ plano up plano_config.yaml
+2024-12-05 16:56:27,979 - cli.main - INFO - Starting plano cli version: 0.1.5
+...
+2024-12-05 16:56:28,485 - cli.utils - INFO - Schema validation successful!
+2024-12-05 16:56:28,485 - cli.main - INFO - Starting plano model server and plano gateway
+...
+2024-12-05 16:56:51,647 - cli.core - INFO - Container is healthy!
+
+Once the gateway is up, you can start interacting with it at port 10000 using the OpenAI chat completion API.
+
+Some sample queries you can ask include: what is currency rate for gbp? or show me list of currencies for conversion.
+
+Step 3. Interacting with gateway using curl command
+
+Here is a sample curl command you can use to interact:
+
+$ curl --header 'Content-Type: application/json' \
+  --data '{"messages": [{"role": "user","content": "what is exchange rate for gbp"}], "model": "none"}' \
+  http://localhost:10000/v1/chat/completions | jq ".choices[0].message.content"
+
+"As of the date provided in your context, December 5, 2024, the exchange rate for GBP (British Pound) from USD (United States Dollar) is 0.78558. This means that 1 USD is equivalent to 0.78558 GBP."
+
+And to get the list of supported currencies:
+
+$ curl --header 'Content-Type: application/json' \
+  --data '{"messages": [{"role": "user","content": "show me list of currencies that are supported for conversion"}], "model": "none"}' \
+  http://localhost:10000/v1/chat/completions | jq ".choices[0].message.content"
+
+"Here is a list of the currencies that are supported for conversion from USD, along with their symbols:\n\n1. AUD - Australian Dollar\n2. BGN - Bulgarian Lev\n3. BRL - Brazilian Real\n4. CAD - Canadian Dollar\n5. CHF - Swiss Franc\n6. CNY - Chinese Renminbi Yuan\n7. CZK - Czech Koruna\n8. DKK - Danish Krone\n9. EUR - Euro\n10. GBP - British Pound\n11. HKD - Hong Kong Dollar\n12. HUF - Hungarian Forint\n13. IDR - Indonesian Rupiah\n14. ILS - Israeli New Sheqel\n15. INR - Indian Rupee\n16. ISK - Icelandic Króna\n17. JPY - Japanese Yen\n18. KRW - South Korean Won\n19. MXN - Mexican Peso\n20. MYR - Malaysian Ringgit\n21. NOK - Norwegian Krone\n22. NZD - New Zealand Dollar\n23. PHP - Philippine Peso\n24. PLN - Polish Złoty\n25. RON - Romanian Leu\n26. SEK - Swedish Krona\n27. SGD - Singapore Dollar\n28. THB - Thai Baht\n29. TRY - Turkish Lira\n30. USD - United States Dollar\n31. ZAR - South African Rand\n\nIf you want to convert USD to any of these currencies, you can select the one you are interested in."
+
+
+
+Use Plano as a Model Proxy (Gateway)
+
+Step 1. Create plano config file
+
+Plano operates based on a configuration file where you can define LLM providers, prompt targets, guardrails, etc. Below is an example configuration that defines OpenAI and Mistral LLM providers.
+
+Create plano_config.yaml file with the following content:
+
+ version: v0.1.0
+
+listeners:
+  egress_traffic:
+    address: 0.0.0.0
+    port: 12000
+    message_format: openai
+    timeout: 30s
+
+ model_providers:
+   - access_key: $OPENAI_API_KEY
+     model: openai/gpt-4o
+     default: true
+
+   - access_key: $MISTRAL_API_KEY
+     model: mistralministral-3b-latest
+
+Step 2. Start plano
+
+Once the config file is created, ensure that you have environment variables set up for MISTRAL_API_KEY and OPENAI_API_KEY (or these are defined in a .env file).
+
+Start Plano:
+
+$ plano up plano_config.yaml
+2024-12-05 11:24:51,288 - cli.main - INFO - Starting plano cli version: 0.1.5
+2024-12-05 11:24:51,825 - cli.utils - INFO - Schema validation successful!
+2024-12-05 11:24:51,825 - cli.main - INFO - Starting plano
+...
+2024-12-05 11:25:16,131 - cli.core - INFO - Container is healthy!
+
+Step 3: Interact with LLM
+
+Step 3.1: Using OpenAI Python client
+
+Make outbound calls via the Plano gateway:
+
+from openai import OpenAI
+
+# Use the OpenAI client as usual
+client = OpenAI(
+  # No need to set a specific openai.api_key since it's configured in Plano's gateway
+  api_key='--',
+  # Set the OpenAI API base URL to the Plano gateway endpoint
+  base_url="http://127.0.0.1:12000/v1"
+)
+
+response = client.chat.completions.create(
+    # we select model from plano_config file
+    model="--",
+    messages=[{"role": "user", "content": "What is the capital of France?"}],
+)
+
+print("OpenAI Response:", response.choices[0].message.content)
+
+Step 3.2: Using curl command
+
+$ curl --header 'Content-Type: application/json' \
+  --data '{"messages": [{"role": "user","content": "What is the capital of France?"}], "model": "none"}' \
+  http://localhost:12000/v1/chat/completions
+
+{
+  ...
+  "model": "gpt-4o-2024-08-06",
+  "choices": [
+    {
+      ...
+      "messages": {
+        "role": "assistant",
+        "content": "The capital of France is Paris.",
+      },
+    }
+  ],
+}
+
+Next Steps
+
+Congratulations! You’ve successfully set up Plano and made your first prompt-based request. To further enhance your GenAI applications, explore the following resources:
+
+Full Documentation: Comprehensive guides and references.
+
+GitHub Repository: Access the source code, contribute, and track updates.
+
+Support: Get help and connect with the Plano community .
+
+With Plano, building scalable, fast, and personalized GenAI applications has never been easier. Dive deeper into Plano’s capabilities and start creating innovative AI-driven experiences today!
+
+---
+
+Function Calling
+----------------
+Doc: guides/function_calling
+
+Function Calling
+
+Function Calling is a powerful feature in Plano that allows your application to dynamically execute backend functions or services based on user prompts.
+This enables seamless integration between natural language interactions and backend operations, turning user inputs into actionable results.
+
+What is Function Calling?
+
+Function Calling refers to the mechanism where the user’s prompt is parsed, relevant parameters are extracted, and a designated backend function (or API) is triggered to execute a particular task.
+This feature bridges the gap between generative AI systems and functional business logic, allowing users to interact with the system through natural language while the backend performs the necessary operations.
+
+Function Calling Workflow
+
+Prompt Parsing
+
+When a user submits a prompt, Plano analyzes it to determine the intent. Based on this intent, the system identifies whether a function needs to be invoked and which parameters should be extracted.
+
+Parameter Extraction
+
+Plano’s advanced natural language processing capabilities automatically extract parameters from the prompt that are necessary for executing the function. These parameters can include text, numbers, dates, locations, or other relevant data points.
+
+Function Invocation
+
+Once the necessary parameters have been extracted, Plano invokes the relevant backend function. This function could be an API, a database query, or any other form of backend logic. The function is executed with the extracted parameters to produce the desired output.
+
+Response Handling
+
+After the function has been called and executed, the result is processed and a response is generated. This response is typically delivered in a user-friendly format, which can include text explanations, data summaries, or even a confirmation message for critical actions.
+
+Arch-Function
+
+The Arch-Function collection of large language models (LLMs) is a collection state-of-the-art (SOTA) LLMs specifically designed for function calling tasks.
+The models are designed to understand complex function signatures, identify required parameters, and produce accurate function call outputs based on natural language prompts.
+Achieving performance on par with GPT-4, these models set a new benchmark in the domain of function-oriented tasks, making them suitable for scenarios where automated API interaction and function execution is crucial.
+
+In summary, the Arch-Function collection demonstrates:
+
+State-of-the-art performance in function calling
+
+Accurate parameter identification and suggestion, even in ambiguous or incomplete inputs
+
+High generalization across multiple function calling use cases, from API interactions to automated backend tasks.
+
+Optimized low-latency, high-throughput performance, making it suitable for real-time, production environments.
+
+Key Features
+
+
+
+
+
+Functionality
+
+Definition
+
+Single Function Calling
+
+Call only one function per user prompt
+
+Parallel Function Calling
+
+Call the same function multiple times but with parameter values
+
+Multiple Function Calling
+
+Call different functions per user prompt
+
+Parallel & Multiple
+
+Perform both parallel and multiple function calling
+
+Implementing Function Calling
+
+Here’s a step-by-step guide to configuring function calling within your Plano setup:
+
+Step 1: Define the Function
+
+First, create or identify the backend function you want Plano to call. This could be an API endpoint, a script, or any other executable backend logic.
+
+import requests
+
+def get_weather(location: str, unit: str = "fahrenheit"):
+    if unit not in ["celsius", "fahrenheit"]:
+        raise ValueError("Invalid unit. Choose either 'celsius' or 'fahrenheit'.")
+
+    api_server = "https://api.yourweatherapp.com"
+    endpoint = f"{api_server}/weather"
+
+    params = {
+        "location": location,
+        "unit": unit
+    }
+
+    response = requests.get(endpoint, params=params)
+    return response.json()
+
+# Example usage
+weather_info = get_weather("Seattle, WA", "celsius")
+print(weather_info)
+
+Step 2: Configure Prompt Targets
+
+Next, map the function to a prompt target, defining the intent and parameters that Plano will extract from the user’s prompt.
+Specify the parameters your function needs and how Plano should interpret these.
+
+Prompt Target Example Configuration
+
+prompt_targets:
+  - name: get_weather
+    description: Get the current weather for a location
+    parameters:
+      - name: location
+        description: The city and state, e.g. San Francisco, New York
+        type: str
+        required: true
+      - name: unit
+        description: The unit of temperature to return
+        type: str
+        enum: ["celsius", "fahrenheit"]
+    endpoint:
+      name: api_server
+      path: /weather
+
+For a complete refernce of attributes that you can configure in a prompt target, see here.
+
+Step 3: Plano Takes Over
+
+Once you have defined the functions and configured the prompt targets, Plano takes care of the remaining work.
+It will automatically validate parameters, and ensure that the required parameters (e.g., location) are present in the prompt, and add validation rules if necessary.
+
+
+
+High-level network flow of where Plano sits in your agentic stack. Managing incoming and outgoing prompt traffic
+
+Once a downstream function (API) is called, Plano  takes the response and sends it an upstream LLM to complete the request (for summarization, Q/A, text generation tasks).
+For more details on how Plano  enables you to centralize usage of LLMs, please read LLM providers.
+
+By completing these steps, you enable Plano to manage the process from validation to response, ensuring users receive consistent, reliable results - and that you are focused
+on the stuff that matters most.
+
+Example Use Cases
+
+Here are some common use cases where Function Calling can be highly beneficial:
+
+Data Retrieval: Extracting information from databases or APIs based on user inputs (e.g., checking account balances, retrieving order status).
+
+Transactional Operations: Executing business logic such as placing an order, processing payments, or updating user profiles.
+
+Information Aggregation: Fetching and combining data from multiple sources (e.g., displaying travel itineraries or combining analytics from various dashboards).
+
+Task Automation: Automating routine tasks like setting reminders, scheduling meetings, or sending emails.
+
+User Personalization: Tailoring responses based on user history, preferences, or ongoing interactions.
+
+Best Practices and Tips
+
+When integrating function calling into your generative AI applications, keep these tips in mind to get the most out of our Plano-Function models:
+
+Keep it clear and simple: Your function names and parameters should be straightforward and easy to understand. Think of it like explaining a task to a smart colleague - the clearer you are, the better the results.
+
+Context is king: Don’t skimp on the descriptions for your functions and parameters. The more context you provide, the better the LLM can understand when and how to use each function.
+
+Be specific with your parameters: Instead of using generic types, get specific. If you’re asking for a date, say it’s a date. If you need a number between 1 and 10, spell that out. The more precise you are, the more accurate the LLM’s responses will be.
+
+Expect the unexpected: Test your functions thoroughly, including edge cases. LLMs can be creative in their interpretations, so it’s crucial to ensure your setup is robust and can handle unexpected inputs.
+
+Watch and learn: Pay attention to how the LLM uses your functions. Which ones does it call often? In what contexts? This information can help you optimize your setup over time.
+
+Remember, working with LLMs is part science, part art. Don’t be afraid to experiment and iterate to find what works best for your specific use case.
+
+---
+
+LLM Routing
+-----------
+Doc: guides/llm_router
+
+LLM Routing
+
+With the rapid proliferation of large language models (LLMs) — each optimized for different strengths, style, or latency/cost profile — routing has become an essential technique to operationalize the use of different models. Plano provides three distinct routing approaches to meet different use cases: Model-based routing, Alias-based routing, and Preference-aligned routing. This enables optimal performance, cost efficiency, and response quality by matching requests with the most suitable model from your available LLM fleet.
+
+For details on supported model providers, configuration options, and client libraries, see LLM Providers.
+
+Routing Methods
+
+
+
+Model-based routing
+
+Direct routing allows you to specify exact provider and model combinations using the format provider/model-name:
+
+Use provider-specific names like openai/gpt-5.2 or anthropic/claude-sonnet-4-5
+
+Provides full control and transparency over which model handles each request
+
+Ideal for production workloads where you want predictable routing behavior
+
+Configuration
+
+Configure your LLM providers with specific provider/model names:
+
+Model-based Routing Configuration
+
+listeners:
+  egress_traffic:
+    address: 0.0.0.0
+    port: 12000
+    message_format: openai
+    timeout: 30s
+
+llm_providers:
+  - model: openai/gpt-5.2
+    access_key: $OPENAI_API_KEY
+    default: true
+
+  - model: openai/gpt-5
+    access_key: $OPENAI_API_KEY
+
+  - model: anthropic/claude-sonnet-4-5
+    access_key: $ANTHROPIC_API_KEY
+
+Client usage
+
+Clients specify exact models:
+
+# Direct provider/model specification
+response = client.chat.completions.create(
+    model="openai/gpt-5.2",
+    messages=[{"role": "user", "content": "Hello!"}]
+)
+
+response = client.chat.completions.create(
+    model="anthropic/claude-sonnet-4-5",
+    messages=[{"role": "user", "content": "Write a story"}]
+)
+
+
+
+Alias-based routing
+
+Alias-based routing lets you create semantic model names that decouple your application from specific providers:
+
+Use meaningful names like fast-model, reasoning-model, or plano.summarize.v1 (see model_aliases)
+
+Maps semantic names to underlying provider models for easier experimentation and provider switching
+
+Ideal for applications that want abstraction from specific model names while maintaining control
+
+Configuration
+
+Configure semantic aliases that map to underlying models:
+
+Alias-based Routing Configuration
+
+listeners:
+  egress_traffic:
+    address: 0.0.0.0
+    port: 12000
+    message_format: openai
+    timeout: 30s
+
+llm_providers:
+  - model: openai/gpt-5.2
+    access_key: $OPENAI_API_KEY
+
+  - model: openai/gpt-5
+    access_key: $OPENAI_API_KEY
+
+  - model: anthropic/claude-sonnet-4-5
+    access_key: $ANTHROPIC_API_KEY
+
+model_aliases:
+  # Model aliases - friendly names that map to actual provider names
+  fast-model:
+    target: gpt-5.2
+
+  reasoning-model:
+    target: gpt-5
+
+  creative-model:
+    target: claude-sonnet-4-5
+
+Client usage
+
+Clients use semantic names:
+
+# Using semantic aliases
+response = client.chat.completions.create(
+    model="fast-model",  # Routes to best available fast model
+    messages=[{"role": "user", "content": "Quick summary please"}]
+)
+
+response = client.chat.completions.create(
+    model="reasoning-model",  # Routes to best reasoning model
+    messages=[{"role": "user", "content": "Solve this complex problem"}]
+)
+
+
+
+Preference-aligned routing (Arch-Router)
+
+Preference-aligned routing uses the Arch-Router model to pick the best LLM based on domain, action, and your configured preferences instead of hard-coding a model.
+
+Domain: High-level topic of the request (e.g., legal, healthcare, programming).
+
+Action: What the user wants to do (e.g., summarize, generate code, translate).
+
+Routing preferences: Your mapping from (domain, action) to preferred models.
+
+Arch-Router analyzes each prompt to infer domain and action, then applies your preferences to select a model. This decouples routing policy (how to choose) from model assignment (what to run), making routing transparent, controllable, and easy to extend as you add or swap models.
+
+Configuration
+
+To configure preference-aligned dynamic routing, define routing preferences that map domains and actions to specific models:
+
+Preference-Aligned Dynamic Routing Configuration
+
+listeners:
+  egress_traffic:
+    address: 0.0.0.0
+    port: 12000
+    message_format: openai
+    timeout: 30s
+
+llm_providers:
+  - model: openai/gpt-5.2
+    access_key: $OPENAI_API_KEY
+    default: true
+
+  - model: openai/gpt-5
+    access_key: $OPENAI_API_KEY
+    routing_preferences:
+      - name: code understanding
+        description: understand and explain existing code snippets, functions, or libraries
+      - name: complex reasoning
+        description: deep analysis, mathematical problem solving, and logical reasoning
+
+  - model: anthropic/claude-sonnet-4-5
+    access_key: $ANTHROPIC_API_KEY
+    routing_preferences:
+      - name: creative writing
+        description: creative content generation, storytelling, and writing assistance
+      - name: code generation
+        description: generating new code snippets, functions, or boilerplate based on user prompts
+
+Client usage
+
+Clients can let the router decide or still specify aliases:
+
+# Let Arch-Router choose based on content
+response = client.chat.completions.create(
+    messages=[{"role": "user", "content": "Write a creative story about space exploration"}]
+    # No model specified - router will analyze and choose claude-sonnet-4-5
+)
+
+Arch-Router
+
+The Arch-Router is a state-of-the-art preference-based routing model specifically designed to address the limitations of traditional LLM routing. This compact 1.5B model delivers production-ready performance with low latency and high accuracy while solving key routing challenges.
+
+Addressing Traditional Routing Limitations:
+
+Human Preference Alignment
+Unlike benchmark-driven approaches, Arch-Router learns to match queries with human preferences by using domain-action mappings that capture subjective evaluation criteria, ensuring routing decisions align with real-world user needs.
+
+Flexible Model Integration
+The system supports seamlessly adding new models for routing without requiring retraining or architectural modifications, enabling dynamic adaptation to evolving model landscapes.
+
+Preference-Encoded Routing
+Provides a practical mechanism to encode user preferences through domain-action mappings, offering transparent and controllable routing decisions that can be customized for specific use cases.
+
+To support effective routing, Arch-Router introduces two key concepts:
+
+Domain – the high-level thematic category or subject matter of a request (e.g., legal, healthcare, programming).
+
+Action – the specific type of operation the user wants performed (e.g., summarization, code generation, booking appointment, translation).
+
+Both domain and action configs are associated with preferred models or model variants. At inference time, Arch-Router analyzes the incoming prompt to infer its domain and action using semantic similarity, task indicators, and contextual cues. It then applies the user-defined routing preferences to select the model best suited to handle the request.
+
+In summary, Arch-Router demonstrates:
+
+Structured Preference Routing: Aligns prompt request with model strengths using explicit domain–action mappings.
+
+Transparent and Controllable: Makes routing decisions transparent and configurable, empowering users to customize system behavior.
+
+Flexible and Adaptive: Supports evolving user needs, model updates, and new domains/actions without retraining the router.
+
+Production-Ready Performance: Optimized for low-latency, high-throughput applications in multi-model environments.
+
+Combining Routing Methods
+
+You can combine static model selection with dynamic routing preferences for maximum flexibility:
+
+Hybrid Routing Configuration
+
+llm_providers:
+  - model: openai/gpt-5.2
+    access_key: $OPENAI_API_KEY
+    default: true
+
+  - model: openai/gpt-5
+    access_key: $OPENAI_API_KEY
+    routing_preferences:
+      - name: complex_reasoning
+        description: deep analysis and complex problem solving
+
+  - model: anthropic/claude-sonnet-4-5
+    access_key: $ANTHROPIC_API_KEY
+    routing_preferences:
+      - name: creative_tasks
+        description: creative writing and content generation
+
+model_aliases:
+  # Model aliases - friendly names that map to actual provider names
+  fast-model:
+    target: gpt-5.2
+
+  reasoning-model:
+    target: gpt-5
+
+  # Aliases that can also participate in dynamic routing
+  creative-model:
+    target: claude-sonnet-4-5
+
+This configuration allows clients to:
+
+Use direct model selection: model="fast-model"
+
+Let the router decide: No model specified, router analyzes content
+
+Example Use Cases
+
+Here are common scenarios where Arch-Router excels:
+
+Coding Tasks: Distinguish between code generation requests (“write a Python function”), debugging needs (“fix this error”), and code optimization (“make this faster”), routing each to appropriately specialized models.
+
+Content Processing Workflows: Classify requests as summarization (“summarize this document”), translation (“translate to Spanish”), or analysis (“what are the key themes”), enabling targeted model selection.
+
+Multi-Domain Applications: Accurately identify whether requests fall into legal, healthcare, technical, or general domains, even when the subject matter isn’t explicitly stated in the prompt.
+
+Conversational Routing: Track conversation context to identify when topics shift between domains or when the type of assistance needed changes mid-conversation.
+
+Best practices
+
+💡Consistent Naming:  Route names should align with their descriptions.
+
+❌ Bad:
+`
+{"name": "math", "description": "handle solving quadratic equations"}
+`
+
+✅ Good:
+`
+{"name": "quadratic_equation", "description": "solving quadratic equations"}
+`
+
+💡 Clear Usage Description:  Make your route names and descriptions specific, unambiguous, and minimizing overlap between routes. The Router performs better when it can clearly distinguish between different types of requests.
+
+❌ Bad:
+`
+{"name": "math", "description": "anything closely related to mathematics"}
+`
+
+✅ Good:
+`
+{"name": "math", "description": "solving, explaining math problems, concepts"}
+`
+
+💡Nouns Descriptor: Preference-based routers perform better with noun-centric descriptors, as they offer more stable and semantically rich signals for matching.
+
+💡Domain Inclusion: for best user experience, you should always include a domain route. This helps the router fall back to domain when action is not confidently inferred.
+
+Unsupported Features
+
+The following features are not supported by the Arch-Router model:
+
+Multi-modality: The model is not trained to process raw image or audio inputs. It can handle textual queries about these modalities (e.g., “generate an image of a cat”), but cannot interpret encoded multimedia data directly.
+
+Function calling: Arch-Router is designed for semantic preference matching, not exact intent classification or tool execution. For structured function invocation, use models in the Plano Function Calling collection instead.
+
+System prompt dependency: Arch-Router routes based solely on the user’s conversation history. It does not use or rely on system prompts for routing decisions.
+
+---
+
+Access Logging
+--------------
+Doc: guides/observability/access_logging
+
+Access Logging
+
+Access logging in Plano refers to the logging of detailed information about each request and response that flows through Plano.
+It provides visibility into the traffic passing through Plano, which is crucial for monitoring, debugging, and analyzing the
+behavior of AI applications and their interactions.
+
+Key Features
+
+Per-Request Logging:
+Each request that passes through Plano is logged. This includes important metadata such as HTTP method,
+path, response status code, request duration, upstream host, and more.
+
+Integration with Monitoring Tools:
+Access logs can be exported to centralized logging systems (e.g., ELK stack or Fluentd) or used to feed monitoring and alerting systems.
+
+Structured Logging: where each request is logged as a object, making it easier to parse and analyze using tools like Elasticsearch and Kibana.
+
+How It Works
+
+Plano exposes access logs for every call it manages on your behalf. By default these access logs can be found under ~/plano_logs. For example:
+
+$ tail -F ~/plano_logs/access_*.log
+
+==> /Users/username/plano_logs/access_llm.log <==
+[2024-10-10T03:55:49.537Z] "POST /v1/chat/completions HTTP/1.1" 0 DC 0 0 770 - "-" "OpenAI/Python 1.51.0" "469793af-b25f-9b57-b265-f376e8d8c586" "api.openai.com" "162.159.140.245:443"
+
+==> /Users/username/plano_logs/access_internal.log <==
+[2024-10-10T03:56:03.906Z] "POST /embeddings HTTP/1.1" 200 - 52 21797 54 53 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "model_server" "192.168.65.254:51000"
+[2024-10-10T03:56:03.961Z] "POST /zeroshot HTTP/1.1" 200 - 106 218 87 87 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "model_server" "192.168.65.254:51000"
+[2024-10-10T03:56:04.050Z] "POST /v1/chat/completions HTTP/1.1" 200 - 1301 614 441 441 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "model_server" "192.168.65.254:51000"
+[2024-10-10T03:56:04.492Z] "POST /hallucination HTTP/1.1" 200 - 556 127 104 104 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "model_server" "192.168.65.254:51000"
+[2024-10-10T03:56:04.598Z] "POST /insurance_claim_details HTTP/1.1" 200 - 447 125 17 17 "-" "-" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "api_server" "192.168.65.254:18083"
+
+==> /Users/username/plano_logs/access_ingress.log <==
+[2024-10-10T03:56:03.905Z] "POST /v1/chat/completions HTTP/1.1" 200 - 463 1022 1695 984 "-" "OpenAI/Python 1.51.0" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "plano_llm_listener" "0.0.0.0:12000"
+
+Log Format
+
+What do these logs mean? Let’s break down the log format:
+
+START_TIME METHOD ORIGINAL-PATH PROTOCOL RESPONSE_CODE RESPONSE_FLAGS
+BYTES_RECEIVED BYTES_SENT DURATION UPSTREAM-SERVICE-TIME X-FORWARDED-FOR
+USER-AGENT X-REQUEST-ID  AUTHORITY UPSTREAM_HOST
+
+Most of these fields are self-explanatory, but here are a few key fields to note:
+
+UPSTREAM-SERVICE-TIME: The time taken by the upstream service to process the request.
+
+DURATION: The total time taken to process the request.
+
+For example for following request:
+
+[2024-10-10T03:56:03.905Z] "POST /v1/chat/completions HTTP/1.1" 200 - 463 1022 1695 984 "-" "OpenAI/Python 1.51.0" "604197fe-2a5b-95a2-9367-1d6b30cfc845" "plano_llm_listener" "0.0.0.0:12000"
+
+Total duration was 1695ms, and the upstream service took 984ms to process the request. Bytes received and sent were 463 and 1022 respectively.
+
+---
+
+Monitoring
+----------
+Doc: guides/observability/monitoring
+
+Monitoring
+
+OpenTelemetry is an open-source observability framework providing APIs
+and instrumentation for generating, collecting, processing, and exporting telemetry data, such as traces,
+metrics, and logs. Its flexible design supports a wide range of backends and seamlessly integrates with
+modern application tools.
+
+Plano acts a source for several monitoring metrics related to agents and LLMs natively integrated
+via OpenTelemetry to help you understand three critical aspects of your application:
+latency, token usage, and error rates by an upstream LLM provider. Latency measures the speed at which your application
+is responding to users, which includes metrics like time to first token (TFT), time per output token (TOT) metrics, and
+the total latency as perceived by users. Below are some screenshots how Plano integrates natively with tools like
+Grafana via Promethus
+
+Metrics Dashboard (via Grafana)
+
+
+
+
+
+
+
+Configure Monitoring
+
+Plano publishes stats endpoint at http://localhost:19901/stats. As noted above, Plano is a source for metrics. To view and manipulate dashbaords, you will
+need to configiure Promethus (as a metrics store) and Grafana for dashboards. Below
+are some sample configuration files for both, respectively.
+
+Sample prometheus.yaml config file
+
+global:
+scrape_interval: 15s
+scrape_timeout: 10s
+evaluation_interval: 15s
+alerting:
+alertmanagers:
+    - static_configs:
+        - targets: []
+    scheme: http
+    timeout: 10s
+    api_version: v2
+scrape_configs:
+- job_name: plano
+    honor_timestamps: true
+    scrape_interval: 15s
+    scrape_timeout: 10s
+    metrics_path: /stats
+    scheme: http
+    static_configs:
+    - targets:
+        - host.docker.internal:19901
+    params:
+    format: ["prometheus"]
+
+Sample grafana datasource.yaml config file
+
+apiVersion: 1
+datasources:
+- name: Prometheus
+    type: prometheus
+    url: http://prometheus:9090
+    isDefault: true
+    access: proxy
+    editable: true
+
+---
+
+Observability
+-------------
+Doc: guides/observability/observability
+
+Observability
+
+---
+
+Tracing
+-------
+Doc: guides/observability/tracing
+
+Tracing
+
+Overview
+
+OpenTelemetry is an open-source observability framework providing APIs
+and instrumentation for generating, collecting, processing, and exporting telemetry data, such as traces,
+metrics, and logs. Its flexible design supports a wide range of backends and seamlessly integrates with
+modern application tools. A key feature of OpenTelemetry is its commitment to standards like the
+W3C Trace Context
+
+Tracing is a critical tool that allows developers to visualize and understand the flow of
+requests in an AI application. With tracing, you can capture a detailed view of how requests propagate
+through various services and components, which is crucial for debugging, performance optimization,
+and understanding complex AI agent architectures like Co-pilots.
+
+Plano propagates trace context using the W3C Trace Context standard, specifically through the
+traceparent header. This allows each component in the system to record its part of the request
+flow, enabling end-to-end tracing across the entire application. By using OpenTelemetry, Plano ensures
+that developers can capture this trace data consistently and in a format compatible with various observability
+tools.
+
+
+
+Benefits of Using Traceparent Headers
+
+Standardization: The W3C Trace Context standard ensures compatibility across ecosystem tools, allowing
+traces to be propagated uniformly through different layers of the system.
+
+Ease of Integration: OpenTelemetry’s design allows developers to easily integrate tracing with minimal
+changes to their codebase, enabling quick adoption of end-to-end observability.
+
+Interoperability: Works seamlessly with popular tracing tools like AWS X-Ray, Datadog, Jaeger, and many others,
+making it easy to visualize traces in the tools you’re already usi
+
+How to Initiate A Trace
+
+Enable Tracing Configuration: Simply add the random_sampling in tracing section to 100`` flag to in the listener config
+
+Trace Context Propagation: Plano automatically propagates the traceparent header. When a request is received, Plano will:
+
+Generate a new traceparent header if one is not present.
+
+Extract the trace context from the traceparent header if it exists.
+
+Start a new span representing its processing of the request.
+
+Forward the traceparent header to downstream services.
+
+Sampling Policy: The 100 in random_sampling: 100 means that all the requests as sampled for tracing.
+You can adjust this value from 0-100.
+
+Trace Propagation
+
+Plano uses the W3C Trace Context standard for trace propagation, which relies on the traceparent header.
+This header carries tracing information in a standardized format, enabling interoperability between different
+tracing systems.
+
+Header Format
+
+The traceparent header has the following format:
+
+traceparent: {version}-{trace-id}-{parent-id}-{trace-flags}
+
+{version}: The version of the Trace Context specification (e.g., 00).
+
+{trace-id}: A 16-byte (32-character hexadecimal) unique identifier for the trace.
+
+{parent-id}: An 8-byte (16-character hexadecimal) identifier for the parent span.
+
+{trace-flags}: Flags indicating trace options (e.g., sampling).
+
+Instrumentation
+
+To integrate AI tracing, your application needs to follow a few simple steps. The steps
+below are very common practice, and not unique to Plano, when you reading tracing headers and export
+spans for distributed tracing.
+
+Read the traceparent header from incoming requests.
+
+Start new spans as children of the extracted context.
+
+Include the traceparent header in outbound requests to propagate trace context.
+
+Send tracing data to a collector or tracing backend to export spans
+
+Example with OpenTelemetry in Python
+
+Install OpenTelemetry packages:
+
+$ pip install opentelemetry-api opentelemetry-sdk opentelemetry-exporter-otlp
+$ pip install opentelemetry-instrumentation-requests
+
+Set up the tracer and exporter:
+
+from opentelemetry import trace
+from opentelemetry.exporter.otlp.proto.grpc.trace_exporter import OTLPSpanExporter
+from opentelemetry.instrumentation.requests import RequestsInstrumentor
+from opentelemetry.sdk.resources import Resource
+from opentelemetry.sdk.trace import TracerProvider
+from opentelemetry.sdk.trace.export import BatchSpanProcessor
+
+# Define the service name
+resource = Resource(attributes={
+    "service.name": "customer-support-agent"
+})
+
+# Set up the tracer provider and exporter
+tracer_provider = TracerProvider(resource=resource)
+otlp_exporter = OTLPSpanExporter(endpoint="otel-collector:4317", insecure=True)
+span_processor = BatchSpanProcessor(otlp_exporter)
+tracer_provider.add_span_processor(span_processor)
+trace.set_tracer_provider(tracer_provider)
+
+# Instrument HTTP requests
+RequestsInstrumentor().instrument()
+
+Handle incoming requests:
+
+from opentelemetry import trace
+from opentelemetry.propagate import extract, inject
+import requests
+
+def handle_request(request):
+    # Extract the trace context
+    context = extract(request.headers)
+    tracer = trace.get_tracer(__name__)
+
+    with tracer.start_as_current_span("process_customer_request", context=context):
+        # Example of processing a customer request
+        print("Processing customer request...")
+
+        # Prepare headers for outgoing request to payment service
+        headers = {}
+        inject(headers)
+
+        # Make outgoing request to external service (e.g., payment gateway)
+        response = requests.get("http://payment-service/api", headers=headers)
+
+        print(f"Payment service response: {response.content}")
+
+Integrating with Tracing Tools
+
+AWS X-Ray
+
+To send tracing data to AWS X-Ray :
+
+Configure OpenTelemetry Collector: Set up the collector to export traces to AWS X-Ray.
+
+Collector configuration (otel-collector-config.yaml):
+
+receivers:
+  otlp:
+    protocols:
+      grpc:
+
+processors:
+  batch:
+
+exporters:
+  awsxray:
+    region: <Your-Aws-Region>
+
+service:
+  pipelines:
+    traces:
+      receivers: [otlp]
+      processors: [batch]
+      exporters: [awsxray]
+
+Deploy the Collector: Run the collector as a Docker container, Kubernetes pod, or standalone service.
+
+Ensure AWS Credentials: Provide AWS credentials to the collector, preferably via IAM roles.
+
+Verify Traces: Access the AWS X-Ray console to view your traces.
+
+Datadog
+
+Datadog
+
+To send tracing data to Datadog:
+
+Configure OpenTelemetry Collector: Set up the collector to export traces to Datadog.
+
+Collector configuration (otel-collector-config.yaml):
+
+receivers:
+  otlp:
+    protocols:
+      grpc:
+
+processors:
+  batch:
+
+exporters:
+  datadog:
+    api:
+      key: "${<Your-Datadog-Api-Key>}"
+    site: "${DD_SITE}"
+
+service:
+  pipelines:
+    traces:
+      receivers: [otlp]
+      processors: [batch]
+      exporters: [datadog]
+
+Set Environment Variables: Provide your Datadog API key and site.
+
+$ export <Your-Datadog-Api-Key>=<Your-Datadog-Api-Key>
+$ export DD_SITE=datadoghq.com  # Or datadoghq.eu
+
+Deploy the Collector: Run the collector in your environment.
+
+Verify Traces: Access the Datadog APM dashboard to view your traces.
+
+Langtrace
+
+Langtrace is an observability tool designed specifically for large language models (LLMs). It helps you capture, analyze, and understand how LLMs are used in your applications including those built using Plano.
+
+To send tracing data to Langtrace:
+
+Configure Plano: Make sure Plano is installed and setup correctly. For more information, refer to the installation guide.
+
+Install Langtrace: Install the Langtrace SDK.:
+
+$ pip install langtrace-python-sdk
+
+Set Environment Variables: Provide your Langtrace API key.
+
+$ export LANGTRACE_API_KEY=<Your-Langtrace-Api-Key>
+
+Trace Requests: Once you have Langtrace set up, you can start tracing requests.
+
+Here’s an example of how to trace a request using the Langtrace Python SDK:
+
+import os
+from langtrace_python_sdk import langtrace  # Must precede any llm module imports
+from openai import OpenAI
+
+langtrace.init(api_key=os.environ['LANGTRACE_API_KEY'])
+
+client = OpenAI(api_key=os.environ['OPENAI_API_KEY'], base_url="http://localhost:12000/v1")
+
+response = client.chat.completions.create(
+    model="gpt-4o-mini",
+    messages=[
+        {"role": "system", "content": "You are a helpful assistant"},
+        {"role": "user", "content": "Hello"},
+    ]
+)
+
+print(chat_completion.choices[0].message.content)
+
+Verify Traces: Access the Langtrace dashboard to view your traces.
+
+Best Practices
+
+Consistent Instrumentation: Ensure all services propagate the traceparent header.
+
+Secure Configuration: Protect sensitive data and secure communication between services.
+
+Performance Monitoring: Be mindful of the performance impact and adjust sampling rates accordingly.
+
+Error Handling: Implement proper error handling to prevent tracing issues from affecting your application.
+
+Summary
+
+By leveraging the traceparent header for trace context propagation, Plano enables developers to implement
+tracing efficiently. This approach simplifies the process of collecting and analyzing tracing data in common
+tools like AWS X-Ray and Datadog, enhancing observability and facilitating faster debugging and optimization.
+
+Additional Resources
+
+OpenTelemetry Documentation
+
+W3C Trace Context Specification
+
+AWS X-Ray Exporter
+
+Datadog Exporter
+
+Langtrace Documentation
+
+Replace placeholders such as <Your-Aws-Region> and <Your-Datadog-Api-Key> with your actual configurations.
+
+---
+
+Orchestration
+-------------
+Doc: guides/orchestration
+
+Orchestration
+
+Building multi-agent systems allow you to route requests across multiple specialized agents, each designed to handle specific types of tasks.
+Plano makes it easy to build and scale these systems by managing the orchestration layer—deciding which agent(s) should handle each request—while you focus on implementing individual agent logic.
+
+This guide shows you how to configure and implement multi-agent orchestration in Plano using a real-world example: a Travel Booking Assistant that routes queries to specialized agents for weather and flights.
+
+How It Works
+
+Plano’s orchestration layer analyzes incoming prompts and routes them to the most appropriate agent based on user intent and conversation context. The workflow is:
+
+User submits a prompt: The request arrives at Plano’s agent listener.
+
+Agent selection: Plano uses an LLM to analyze the prompt and determine user intent and complexity. By default, this uses Plano-Orchestrator-30B-A3B, which offers performance of foundation models at 1/10th the cost. The LLM routes the request to the most suitable agent configured in your system—such as a weather agent or flight agent.
+
+Agent handles request: Once the selected agent receives the request object from Plano, it manages its own inner loop until the task is complete. This means the agent autonomously calls models, invokes tools, processes data, and reasons about next steps—all within its specialized domain—before returning the final response.
+
+Seamless handoffs: For multi-turn conversations, Plano repeats the intent analysis for each follow-up query, enabling smooth handoffs between agents as the conversation evolves.
+
+Example: Travel Booking Assistant
+
+Let’s walk through a complete multi-agent system: a Travel Booking Assistant that helps users plan trips by providing weather forecasts and flight information. This system uses two specialized agents:
+
+Weather Agent: Provides real-time weather conditions and multi-day forecasts
+
+Flight Agent: Searches for flights between airports with real-time tracking
+
+Configuration
+
+Configure your agents in the listeners section of your plano_config.yaml:
+
+Travel Booking Multi-Agent Configuration
+
+version: v0.3.0
+
+agents:
+  - id: weather_agent
+    url: http://host.docker.internal:10510
+  - id: flight_agent
+    url: http://host.docker.internal:10520
+
+model_providers:
+  - model: openai/gpt-4o
+    access_key: $OPENAI_API_KEY
+    default: true
+  - model: openai/gpt-4o-mini
+    access_key: $OPENAI_API_KEY # smaller, faster, cheaper model for extracting entities like location
+
+listeners:
+  - type: agent
+    name: travel_booking_service
+    port: 8001
+    router: plano_orchestrator_v1
+    agents:
+      - id: weather_agent
+        description: |
+
+          WeatherAgent is a specialized AI assistant for real-time weather information and forecasts. It provides accurate weather data for any city worldwide using the Open-Meteo API, helping travelers plan their trips with up-to-date weather conditions.
+
+          Capabilities:
+            * Get real-time weather conditions and multi-day forecasts for any city worldwide using Open-Meteo API (free, no API key needed)
+            * Provides current temperature
+            * Provides multi-day forecasts
+            * Provides weather conditions
+            * Provides sunrise/sunset times
+            * Provides detailed weather information
+            * Understands conversation context to resolve location references from previous messages
+            * Handles weather-related questions including "What's the weather in [city]?", "What's the forecast for [city]?", "How's the weather in [city]?"
+            * When queries include both weather and other travel questions (e.g., flights, currency), this agent answers ONLY the weather part
+
+      - id: flight_agent
+        description: |
+
+          FlightAgent is an AI-powered tool specialized in providing live flight information between airports. It leverages the FlightAware AeroAPI to deliver real-time flight status, gate information, and delay updates.
+
+          Capabilities:
+            * Get live flight information between airports using FlightAware AeroAPI
+            * Shows real-time flight status
+            * Shows scheduled/estimated/actual departure and arrival times
+            * Shows gate and terminal information
+            * Shows delays
+            * Shows aircraft type
+            * Shows flight status
+            * Automatically resolves city names to airport codes (IATA/ICAO)
+            * Understands conversation context to infer origin/destination from follow-up questions
+            * Handles flight-related questions including "What flights go from [city] to [city]?", "Do flights go to [city]?", "Are there direct flights from [city]?"
+            * When queries include both flight and other travel questions (e.g., weather, currency), this agent answers ONLY the flight part
+
+tracing:
+  random_sampling: 100
+
+
+Key Configuration Elements:
+
+agent listener: A listener of type: agent tells Plano to perform intent analysis and routing for incoming requests.
+
+agents list: Define each agent with an id, description (used for routing decisions)
+
+router: The plano_orchestrator_v1 router uses Plano-Orchestrator to analyze user intent and select the appropriate agent.
+
+filter_chain: Optionally attach filter chains to agents for guardrails, query rewriting, or context enrichment.
+
+Writing Effective Agent Descriptions
+
+Agent descriptions are critical—they’re used by Plano-Orchestrator to make routing decisions. Effective descriptions should include:
+
+Clear introduction: A concise statement explaining what the agent is and its primary purpose
+
+Capabilities section: A bulleted list of specific capabilities, including:
+
+What APIs or data sources it uses (e.g., “Open-Meteo API”, “FlightAware AeroAPI”)
+
+What information it provides (e.g., “current temperature”, “multi-day forecasts”, “gate information”)
+
+How it handles context (e.g., “Understands conversation context to resolve location references”)
+
+What question patterns it handles (e.g., “What’s the weather in [city]?”)
+
+How it handles multi-part queries (e.g., “When queries include both weather and flights, this agent answers ONLY the weather part”)
+
+Here’s an example of a well-structured agent description:
+
+- id: weather_agent
+  description: |
+
+    WeatherAgent is a specialized AI assistant for real-time weather information
+    and forecasts. It provides accurate weather data for any city worldwide using
+    the Open-Meteo API, helping travelers plan their trips with up-to-date weather
+    conditions.
+
+    Capabilities:
+      * Get real-time weather conditions and multi-day forecasts for any city worldwide
+      * Provides current temperature, weather conditions, sunrise/sunset times
+      * Provides detailed weather information including multi-day forecasts
+      * Understands conversation context to resolve location references from previous messages
+      * Handles weather-related questions including "What's the weather in [city]?"
+      * When queries include both weather and other travel questions (e.g., flights),
+        this agent answers ONLY the weather part
+
+We will soon support “Agents as Tools” via Model Context Protocol (MCP), enabling agents to dynamically discover and invoke other agents as tools. Track progress on GitHub Issue #646.
+
+Implementation
+
+Agents are HTTP services that receive routed requests from Plano. Each agent implements the OpenAI Chat Completions API format, making them compatible with standard LLM clients.
+
+Agent Structure
+
+Let’s examine the Weather Agent implementation:
+
+Weather Agent - Core Structure
+
+@app.post("/v1/chat/completions")
+async def handle_request(request: Request):
+    """HTTP endpoint for chat completions with streaming support."""
+
+    request_body = await request.json()
+    messages = request_body.get("messages", [])
+    logger.info(
+        "messages detail json dumps: %s",
+        json.dumps(messages, indent=2),
+    )
+
+    traceparent_header = request.headers.get("traceparent")
+    return StreamingResponse(
+        invoke_weather_agent(request, request_body, traceparent_header),
+        media_type="text/plain",
+        headers={
+            "content-type": "text/event-stream",
+        },
+    )
+
+
+async def invoke_weather_agent(
+
+
+Key Points:
+
+Agents expose a /v1/chat/completions endpoint that matches OpenAI’s API format
+
+They use Plano’s LLM gateway (via LLM_GATEWAY_ENDPOINT) for all LLM calls
+
+They receive the full conversation history in request_body.messages
+
+Information Extraction with LLMs
+
+Agents use LLMs to extract structured information from natural language queries. This enables them to understand user intent and extract parameters needed for API calls.
+
+The Weather Agent extracts location information:
+
+Weather Agent - Location Extraction
+
+
+    instructions = """Extract the location for WEATHER queries. Return just the city name.
+
+            Rules:
+            1. For multi-part queries, extract ONLY the location mentioned with weather keywords ("weather in [location]")
+            2. If user says "there" or "that city", it typically refers to the DESTINATION city in travel contexts (not the origin)
+            3. For flight queries with weather, "there" means the destination city where they're traveling TO
+            4. Return plain text (e.g., "London", "New York", "Paris, France")
+            5. If no weather location found, return "NOT_FOUND"
+
+            Examples:
+            - "What's the weather in London?" -> "London"
+            - "Flights from Seattle to Atlanta, and show me the weather there" -> "Atlanta"
+            - "Can you get me flights from Seattle to Atlanta tomorrow, and also please show me the weather there" -> "Atlanta"
+            - "What's the weather in Seattle, and what is one flight that goes direct to Atlanta?" -> "Seattle"
+            - User asked about flights to Atlanta, then "what's the weather like there?" -> "Atlanta"
+            - "I'm going to Seattle" -> "Seattle"
+            - "What's happening?" -> "NOT_FOUND"
+
+            Extract location:"""
+
+    try:
+        user_messages = [
+            msg.get("content") for msg in messages if msg.get("role") == "user"
+        ]
+
+        if not user_messages:
+            location = "New York"
+        else:
+            ctx = extract(request.headers)
+            extra_headers = {}
+            inject(extra_headers, context=ctx)
+
+            # For location extraction, pass full conversation for context (e.g., "there" = previous destination)
+            response = await openai_client_via_plano.chat.completions.create(
+                model=LOCATION_MODEL,
+                messages=[
+                    {"role": "system", "content": instructions},
+                    *[
+                        {"role": msg.get("role"), "content": msg.get("content")}
+                        for msg in messages
+                    ],
+                ],
+                temperature=0.1,
+                max_tokens=50,
+                extra_headers=extra_headers if extra_headers else None,
+            )
+
+
+The Flight Agent extracts more complex information—origin, destination, and dates:
+
+Flight Agent - Flight Information Extraction
+
+async def extract_flight_route(messages: list, request: Request) -> dict:
+    """Extract origin, destination, and date from conversation using LLM."""
+
+    extraction_prompt = """Extract flight origin, destination cities, and travel date from the conversation.
+
+    Rules:
+    1. Look for patterns: "flight from X to Y", "flights to Y", "fly from X"
+    2. Extract dates like "tomorrow", "next week", "December 25", "12/25", "on Monday"
+    3. Use conversation context to fill in missing details
+    4. Return JSON: {"origin": "City" or null, "destination": "City" or null, "date": "YYYY-MM-DD" or null}
+
+    Examples:
+    - "Flight from Seattle to Atlanta tomorrow" -> {"origin": "Seattle", "destination": "Atlanta", "date": "2025-12-24"}
+    - "What flights go to New York?" -> {"origin": null, "destination": "New York", "date": null}
+    - "Flights to Miami on Christmas" -> {"origin": null, "destination": "Miami", "date": "2025-12-25"}
+    - "Show me flights from LA to NYC next Monday" -> {"origin": "LA", "destination": "NYC", "date": "2025-12-30"}
+
+    Today is December 23, 2025. Extract flight route and date:"""
+
+    try:
+        ctx = extract(request.headers)
+        extra_headers = {}
+        inject(extra_headers, context=ctx)
+
+        response = await openai_client_via_plano.chat.completions.create(
+            model=EXTRACTION_MODEL,
+            messages=[
+                {"role": "system", "content": extraction_prompt},
+                *[
+                    {"role": msg.get("role"), "content": msg.get("content")}
+                    for msg in messages[-5:]
+                ],
+            ],
+            temperature=0.1,
+            max_tokens=100,
+            extra_headers=extra_headers if extra_headers else None,
+        )
+
+        result = response.choices[0].message.content.strip()
+        if "```json" in result:
+            result = result.split("```json")[1].split("```")[0].strip()
+        elif "```" in result:
+            result = result.split("```")[1].split("```")[0].strip()
+
+        route = json.loads(result)
+        return {
+            "origin": route.get("origin"),
+            "destination": route.get("destination"),
+            "date": route.get("date"),
+        }
+    except Exception as e:
+        logger.error(f"Error extracting flight route: {e}")
+
+
+Key Points:
+
+Use smaller, faster models (like gpt-4o-mini) for extraction tasks
+
+Include conversation context to handle follow-up questions and pronouns
+
+Use structured prompts with clear output formats (JSON)
+
+Handle edge cases with fallback values
+
+Calling External APIs
+
+After extracting information, agents call external APIs to fetch real-time data:
+
+Weather Agent - External API Call
+
+        # Geocode city to get coordinates
+        geocode_url = f"https://geocoding-api.open-meteo.com/v1/search?name={quote(location)}&count=1&language=en&format=json"
+        geocode_response = await http_client.get(geocode_url)
+
+        if geocode_response.status_code != 200 or not geocode_response.json().get(
+            "results"
+        ):
+            logger.warning(f"Could not geocode {location}, using New York")
+            location = "New York"
+            geocode_url = f"https://geocoding-api.open-meteo.com/v1/search?name={quote(location)}&count=1&language=en&format=json"
+            geocode_response = await http_client.get(geocode_url)
+
+        geocode_data = geocode_response.json()
+        if not geocode_data.get("results"):
+            return {
+                "location": location,
+                "weather": {
+                    "date": datetime.now().strftime("%Y-%m-%d"),
+                    "day_name": datetime.now().strftime("%A"),
+                    "temperature_c": None,
+                    "temperature_f": None,
+                    "weather_code": None,
+                    "error": "Could not retrieve weather data",
+                },
+            }
+
+        result = geocode_data["results"][0]
+        location_name = result.get("name", location)
+        latitude = result["latitude"]
+        longitude = result["longitude"]
+
+        logger.info(
+            f"Geocoded '{location}' to {location_name} ({latitude}, {longitude})"
+        )
+
+        # Get weather forecast
+        weather_url = (
+            f"https://api.open-meteo.com/v1/forecast?"
+            f"latitude={latitude}&longitude={longitude}&"
+            f"current=temperature_2m&"
+            f"daily=sunrise,sunset,temperature_2m_max,temperature_2m_min,weather_code&"
+            f"forecast_days={days}&timezone=auto"
+        )
+
+        weather_response = await http_client.get(weather_url)
+        if weather_response.status_code != 200:
+            return {
+                "location": location_name,
+                "weather": {
+                    "date": datetime.now().strftime("%Y-%m-%d"),
+                    "day_name": datetime.now().strftime("%A"),
+                    "temperature_c": None,
+                    "temperature_f": None,
+                    "weather_code": None,
+                    "error": "Could not retrieve weather data",
+                },
+            }
+
+        weather_data = weather_response.json()
+        current_temp = weather_data.get("current", {}).get("temperature_2m")
+        daily = weather_data.get("daily", {})
+
+
+
+The Flight Agent calls FlightAware’s AeroAPI:
+
+Flight Agent - External API Call
+
+async def get_flights(
+    origin_code: str, dest_code: str, travel_date: Optional[str] = None
+) -> Optional[dict]:
+    """Get flights between two airports using FlightAware API.
+
+    Args:
+        origin_code: Origin airport IATA code
+        dest_code: Destination airport IATA code
+        travel_date: Travel date in YYYY-MM-DD format, defaults to today
+
+    Note: FlightAware API limits searches to 2 days in the future.
+    """
+    try:
+        # Use provided date or default to today
+        if travel_date:
+            search_date = travel_date
+        else:
+            search_date = datetime.now().strftime("%Y-%m-%d")
+
+        # Validate date is not too far in the future (FlightAware limit: 2 days)
+        search_date_obj = datetime.strptime(search_date, "%Y-%m-%d")
+        today = datetime.now().replace(hour=0, minute=0, second=0, microsecond=0)
+        days_ahead = (search_date_obj - today).days
+
+        if days_ahead > 2:
+            logger.warning(
+                f"Requested date {search_date} is {days_ahead} days ahead, exceeds FlightAware 2-day limit"
+            )
+            return {
+                "origin_code": origin_code,
+                "destination_code": dest_code,
+                "flights": [],
+                "count": 0,
+                "error": f"FlightAware API only provides flight data up to 2 days in the future. The requested date ({search_date}) is {days_ahead} days ahead. Please search for today, tomorrow, or the day after.",
+            }
+
+        url = f"{AEROAPI_BASE_URL}/airports/{origin_code}/flights/to/{dest_code}"
+        headers = {"x-apikey": AEROAPI_KEY}
+        params = {
+            "start": f"{search_date}T00:00:00Z",
+            "end": f"{search_date}T23:59:59Z",
+            "connection": "nonstop",
+            "max_pages": 1,
+        }
+
+        response = await http_client.get(url, headers=headers, params=params)
+
+        if response.status_code != 200:
+            logger.error(
+                f"FlightAware API error {response.status_code}: {response.text}"
+            )
+            return None
+
+        data = response.json()
+        flights = []
+
+        # Log raw API response for debugging
+        logger.info(f"FlightAware API returned {len(data.get('flights', []))} flights")
+
+        for idx, flight_group in enumerate(
+            data.get("flights", [])[:5]
+        ):  # Limit to 5 flights
+            # FlightAware API nests data in segments array
+            segments = flight_group.get("segments", [])
+            if not segments:
+                continue
+
+            flight = segments[0]  # Get first segment (direct flights only have one)
+
+            # Extract airport codes from nested objects
+            flight_origin = None
+            flight_dest = None
+
+            if isinstance(flight.get("origin"), dict):
+                flight_origin = flight["origin"].get("code_iata")
+
+            if isinstance(flight.get("destination"), dict):
+                flight_dest = flight["destination"].get("code_iata")
+
+            # Build flight object
+            flights.append(
+                {
+                    "airline": flight.get("operator"),
+                    "flight_number": flight.get("ident_iata") or flight.get("ident"),
+                    "departure_time": flight.get("scheduled_out"),
+                    "arrival_time": flight.get("scheduled_in"),
+                    "origin": flight_origin,
+                    "destination": flight_dest,
+                    "aircraft_type": flight.get("aircraft_type"),
+                    "status": flight.get("status"),
+                    "terminal_origin": flight.get("terminal_origin"),
+                    "gate_origin": flight.get("gate_origin"),
+                }
+            )
+
+        return {
+            "origin_code": origin_code,
+            "destination_code": dest_code,
+            "flights": flights,
+            "count": len(flights),
+        }
+    except Exception as e:
+        logger.error(f"Error fetching flights: {e}")
+        return None
+
+
+
+Key Points:
+
+Use async HTTP clients (like httpx.AsyncClient) for non-blocking API calls
+
+Transform external API responses into consistent, structured formats
+
+Handle errors gracefully with fallback values
+
+Cache or validate data when appropriate (e.g., airport code validation)
+
+Preparing Context and Generating Responses
+
+Agents combine extracted information, API data, and conversation history to generate responses:
+
+Weather Agent - Context Preparation and Response Generation
+
+    last_user_msg = get_last_user_content(messages)
+    days = 1
+
+    if "forecast" in last_user_msg or "week" in last_user_msg:
+        days = 7
+    elif "tomorrow" in last_user_msg:
+        days = 2
+
+    # Extract specific number of days if mentioned (e.g., "5 day forecast")
+    import re
+
+    day_match = re.search(r"(\d{1,2})\s+day", last_user_msg)
+    if day_match:
+        requested_days = int(day_match.group(1))
+        days = min(requested_days, 16)  # API supports max 16 days
+
+    # Get live weather data (location extraction happens inside this function)
+    weather_data = await get_weather_data(request, messages, days)
+
+    # Create weather context to append to user message
+    forecast_type = "forecast" if days > 1 else "current weather"
+    weather_context = f"""
+
+Weather data for {weather_data['location']} ({forecast_type}):
+{json.dumps(weather_data, indent=2)}"""
+
+    # System prompt for weather agent
+    instructions = """You are a weather assistant in a multi-agent system. You will receive weather data in JSON format with these fields:
+
+    - "location": City name
+    - "forecast": Array of weather objects, each with date, day_name, temperature_c, temperature_f, temperature_max_c, temperature_min_c, weather_code, sunrise, sunset
+    - weather_code: WMO code (0=clear, 1-3=partly cloudy, 45-48=fog, 51-67=rain, 71-86=snow, 95-99=thunderstorm)
+
+    Your task:
+    1. Present the weather/forecast clearly for the location
+    2. For single day: show current conditions
+    3. For multi-day: show each day with date and conditions
+    4. Include temperature in both Celsius and Fahrenheit
+    5. Describe conditions naturally based on weather_code
+    6. Use conversational language
+
+    Important: If the conversation includes information from other agents (like flight details), acknowledge and build upon that context naturally. Your primary focus is weather, but maintain awareness of the full conversation.
+
+    Remember: Only use the provided data. If fields are null, mention data is unavailable."""
+
+    # Build message history with weather data appended to the last user message
+    response_messages = [{"role": "system", "content": instructions}]
+
+    for i, msg in enumerate(messages):
+        # Append weather data to the last user message
+        if i == len(messages) - 1 and msg.get("role") == "user":
+            response_messages.append(
+                {"role": "user", "content": msg.get("content") + weather_context}
+            )
+        else:
+            response_messages.append(
+                {"role": msg.get("role"), "content": msg.get("content")}
+            )
+
+    try:
+        ctx = extract(request.headers)
+        extra_headers = {"x-envoy-max-retries": "3"}
+        inject(extra_headers, context=ctx)
+
+        stream = await openai_client_via_plano.chat.completions.create(
+            model=WEATHER_MODEL,
+            messages=response_messages,
+            temperature=request_body.get("temperature", 0.7),
+            max_tokens=request_body.get("max_tokens", 1000),
+            stream=True,
+            extra_headers=extra_headers,
+        )
+
+        async for chunk in stream:
+            if chunk.choices:
+                yield f"data: {chunk.model_dump_json()}\n\n"
+
+        yield "data: [DONE]\n\n"
+
+    except Exception as e:
+        logger.error(f"Error generating weather response: {e}")
+
+
+Key Points:
+
+Use system messages to provide structured data to the LLM
+
+Include full conversation history for context-aware responses
+
+Stream responses for better user experience
+
+Route all LLM calls through Plano’s gateway for consistent behavior and observability
+
+Best Practices
+
+Write Clear Agent Descriptions
+
+Agent descriptions are used by Plano-Orchestrator to make routing decisions. Be specific about what each agent handles:
+
+# Good - specific and actionable
+- id: flight_agent
+  description: Get live flight information between airports using FlightAware AeroAPI. Shows real-time flight status, scheduled/estimated/actual departure and arrival times, gate and terminal information, delays, aircraft type, and flight status. Automatically resolves city names to airport codes (IATA/ICAO). Understands conversation context to infer origin/destination from follow-up questions.
+
+# Less ideal - too vague
+- id: flight_agent
+  description: Handles flight queries
+
+Use Conversation Context Effectively
+
+Include conversation history in your extraction and response generation:
+
+# Include conversation context for extraction
+conversation_context = []
+for msg in messages:
+    conversation_context.append({"role": msg.role, "content": msg.content})
+
+# Use recent context (last 10 messages)
+context_messages = conversation_context[-10:] if len(conversation_context) > 10 else conversation_context
+
+Route LLM Calls Through Plano’s Model Proxy
+
+Always route LLM calls through Plano’s Model Proxy for consistent responses, smart routing, and rich observability:
+
+openai_client_via_plano = AsyncOpenAI(
+    base_url=LLM_GATEWAY_ENDPOINT,  # Plano's LLM gateway
+    api_key="EMPTY",
+)
+
+response = await openai_client_via_plano.chat.completions.create(
+    model="openai/gpt-4o",
+    messages=messages,
+    stream=True,
+)
+
+Handle Errors Gracefully
+
+Provide fallback values and clear error messages:
+
+async def get_weather_data(request: Request, messages: list, days: int = 1):
+    try:
+        # ... extraction and API logic ...
+        location = response.choices[0].message.content.strip().strip("\"'`.,!?")
+        if not location or location.upper() == "NOT_FOUND":
+            location = "New York"  # Fallback to default
+        return weather_data
+    except Exception as e:
+        logger.error(f"Error getting weather data: {e}")
+        return {"location": "New York", "weather": {"error": "Could not retrieve weather data"}}
+
+Use Appropriate Models for Tasks
+
+Use smaller, faster models for extraction tasks and larger models for final responses:
+
+# Extraction: Use smaller, faster model
+LOCATION_MODEL = "openai/gpt-4o-mini"
+
+# Final response: Use larger, more capable model
+WEATHER_MODEL = "openai/gpt-4o"
+
+Stream Responses
+
+Stream responses for better user experience:
+
+async def invoke_weather_agent(request: Request, request_body: dict, traceparent_header: str = None):
+    # ... prepare messages with weather data ...
+
+    stream = await openai_client_via_plano.chat.completions.create(
+        model=WEATHER_MODEL,
+        messages=response_messages,
+        temperature=request_body.get("temperature", 0.7),
+        max_tokens=request_body.get("max_tokens", 1000),
+        stream=True,
+        extra_headers=extra_headers,
+    )
+
+    async for chunk in stream:
+        if chunk.choices:
+            yield f"data: {chunk.model_dump_json()}\n\n"
+
+    yield "data: [DONE]\n\n"
+
+Common Use Cases
+
+Multi-agent orchestration is particularly powerful for:
+
+Travel and Booking Systems
+
+Route queries to specialized agents for weather and flights:
+
+agents:
+  - id: weather_agent
+    description: Get real-time weather conditions and forecasts
+  - id: flight_agent
+    description: Search for flights and provide flight status
+
+Customer Support
+
+Route common queries to automated support agents while escalating complex issues:
+
+agents:
+  - id: tier1_support
+    description: Handles common FAQs, password resets, and basic troubleshooting
+  - id: tier2_support
+    description: Handles complex technical issues requiring deep product knowledge
+  - id: human_escalation
+    description: Escalates sensitive issues or unresolved problems to human agents
+
+Sales and Marketing
+
+Direct leads and inquiries to specialized sales agents:
+
+agents:
+  - id: product_recommendation
+    description: Recommends products based on user needs and preferences
+  - id: pricing_agent
+    description: Provides pricing information and quotes
+  - id: sales_closer
+    description: Handles final negotiations and closes deals
+
+Technical Documentation and Support
+
+Combine RAG agents for documentation lookup with specialized troubleshooting agents:
+
+agents:
+  - id: docs_agent
+    description: Retrieves relevant documentation and guides
+    filter_chain:
+      - query_rewriter
+      - context_builder
+  - id: troubleshoot_agent
+    description: Diagnoses and resolves technical issues step by step
+
+Next Steps
+
+Learn more about agents and the inner vs. outer loop model
+
+Explore filter chains for adding guardrails and context enrichment
+
+See observability for monitoring multi-agent workflows
+
+Review the LLM Providers guide for model routing within agents
+
+Check out the complete Travel Booking demo on GitHub
+
+To observe traffic to and from agents, please read more about observability in Plano.
+
+By carefully configuring and managing your Agent routing and hand off, you can significantly improve your application’s responsiveness, performance, and overall user satisfaction.
+
+---
+
+Guardrails
+----------
+Doc: guides/prompt_guard
+
+Guardrails
+
+Guardrails are Plano’s way of applying safety and validation checks to prompts before they reach your application logic. They are typically implemented as
+filters in a Filter Chain attached to an agent, so every request passes through a consistent processing layer.
+
+Why Guardrails
+
+Guardrails are essential for maintaining control over AI-driven applications. They help enforce organizational policies, ensure compliance with regulations
+(like GDPR or HIPAA), and protect users from harmful or inappropriate content. In applications where prompts generate responses or trigger actions, guardrails
+minimize risks like malicious inputs, off-topic queries, or misaligned outputs—adding a consistent layer of input scrutiny that makes interactions safer,
+more reliable, and easier to reason about.
+
+vale Vale.Spelling = NO
+
+Jailbreak Prevention: Detect and filter inputs that attempt to change LLM behavior, expose system prompts, or bypass safety policies.
+
+Domain and Topicality Enforcement: Ensure that agents only respond to prompts within an approved domain (for example, finance-only or healthcare-only use cases) and reject unrelated queries.
+
+Dynamic Error Handling: Provide clear error messages when requests violate policy, helping users correct their inputs.
+
+How Guardrails Work
+
+Guardrails can be implemented as either in-process MCP filters or as HTTP-based filters. HTTP filters are external services that receive the request over HTTP, validate it, and return a response to allow or reject the request. This makes it easy to use filters written in any language or run them as independent services.
+
+Each filter receives the chat messages, evaluates them against policy, and either lets the request continue or raises a ToolError (or returns an error response) to reject it with a helpful error message.
+
+The example below shows an input guard for TechCorp’s customer support system that validates queries are within the company’s domain:
+
+Example domain validation guard using FastMCP
+
+from typing import List
+from fastmcp.exceptions import ToolError
+from . import mcp
+
+@mcp.tool
+async def input_guards(messages: List[ChatMessage]) -> List[ChatMessage]:
+    """Validates queries are within TechCorp's domain."""
+
+    # Get the user's query
+    user_query = next(
+        (msg.content for msg in reversed(messages) if msg.role == "user"),
+        ""
+    )
+
+    # Use an LLM to validate the query scope (simplified)
+    is_valid = await validate_with_llm(user_query)
+
+    if not is_valid:
+        raise ToolError(
+            "I can only assist with questions related to TechCorp and its services. "
+            "Please ask about TechCorp's products, pricing, SLAs, or technical support."
+        )
+
+    return messages
+
+To wire this guardrail into Plano, define the filter and add it to your agent’s filter chain:
+
+Plano configuration with input guard filter
+
+filters:
+  - id: input_guards
+    url: http://localhost:10500
+
+listeners:
+  - type: agent
+    name: agent_1
+    port: 8001
+    router: plano_orchestrator_v1
+    agents:
+      - id: rag_agent
+        description: virtual assistant for retrieval augmented generation tasks
+        filter_chain:
+          - input_guards
+
+When a request arrives at agent_1, Plano invokes the input_guards filter first. If validation passes, the request continues to
+the agent. If validation fails (ToolError raised), Plano returns an error response to the caller.
+
+Testing the Guardrail
+
+Here’s an example of the guardrail in action, rejecting a query about Apple Corporation (outside TechCorp’s domain):
+
+Request that violates the guardrail policy
+
+curl -X POST http://localhost:8001/v1/chat/completions \
+  -H "Content-Type: application/json" \
+  -d '{
+    "model": "gpt-4",
+    "messages": [
+      {
+        "role": "user",
+        "content": "what is sla for apple corporation?"
+      }
+    ],
+    "stream": false
+  }'
+
+Error response from the guardrail
+
+{
+  "error": "ClientError",
+  "agent": "input_guards",
+  "status": 400,
+  "agent_response": "I apologize, but I can only assist with questions related to TechCorp and its services. Your query appears to be outside this scope. The query is about SLA for Apple Corporation, which is unrelated to TechCorp.\n\nPlease ask me about TechCorp's products, services, pricing, SLAs, or technical support."
+}
+
+This prevents out-of-scope queries from reaching your agent while providing clear feedback to users about why their request was rejected.
+
+---
+
+Conversational State
+--------------------
+Doc: guides/state
+
+Conversational State
+
+The OpenAI Responses API (v1/responses) is designed for multi-turn conversations where context needs to persist across requests. Plano provides a unified v1/responses API that works with any LLM provider—OpenAI, Anthropic, Azure OpenAI, DeepSeek, or any OpenAI-compatible provider—while automatically managing conversational state for you.
+
+Unlike the traditional Chat Completions API where you manually manage conversation history by including all previous messages in each request, Plano handles state management behind the scenes. This means you can use the Responses API with any model provider, and Plano will persist conversation context across requests—making it ideal for building conversational agents that remember context without bloating every request with full message history.
+
+How It Works
+
+When a client calls the Responses API:
+
+First request: Plano generates a unique resp_id and stores the conversation state (messages, model, provider, timestamp).
+
+Subsequent requests: The client includes the previous_resp_id from the previous response. Plano retrieves the stored conversation state, merges it with the new input, and sends the combined context to the LLM.
+
+Response: The LLM sees the full conversation history without the client needing to resend all previous messages.
+
+This pattern dramatically reduces bandwidth and makes it easier to build multi-turn agents—Plano handles the state plumbing so you can focus on agent logic.
+
+Example Using OpenAI Python SDK:
+
+from openai import OpenAI
+
+# Point to Plano's Model Proxy endpoint
+client = OpenAI(
+    api_key="test-key",
+    base_url="http://127.0.0.1:12000/v1"
+)
+
+# First turn - Plano creates a new conversation state
+response = client.responses.create(
+    model="claude-sonnet-4-5",  # Works with any configured provider
+    input="My name is Alice and I like Python"
+)
+
+# Save the response_id for conversation continuity
+resp_id = response.id
+print(f"Assistant: {response.output_text}")
+
+# Second turn - Plano automatically retrieves previous context
+resp2 = client.responses.create(
+    model="claude-sonnet-4-5", # Make sure its configured in plano_config.yaml
+    input="Please list all the messages you have received in our conversation, numbering each one.",
+    previous_response_id=resp_id,
+)
+
+print(f"Assistant: {resp2.output_text}")
+# Output: "Your name is Alice and your favorite language is Python"
+
+Notice how the second request only includes the new user message—Plano automatically merges it with the stored conversation history before sending to the LLM.
+
+Configuration Overview
+
+State storage is configured in the state_storage section of your plano_config.yaml:
+
+state_storage:
+  # Type: memory | postgres
+  type: postgres
+
+  # Connection string for postgres type
+  # Environment variables are supported using $VAR_NAME or ${VAR_NAME} syntax
+  # Replace [USER] and [HOST] with your actual database credentials
+  # Variables like $DB_PASSWORD MUST be set before running config validation/rendering
+  # Example: Replace [USER] with 'myuser' and [HOST] with 'db.example.com:5432'
+  connection_string: "postgresql://[USER]:$DB_PASSWORD@[HOST]:5432/postgres"
+
+
+Plano supports two storage backends:
+
+Memory: Fast, ephemeral storage for development and testing. State is lost when Plano restarts.
+
+PostgreSQL: Durable, production-ready storage with support for Supabase and self-hosted PostgreSQL instances.
+
+If you don’t configure state_storage, conversation state management is disabled. The Responses API will still work, but clients must manually include full conversation history in each request (similar to the Chat Completions API behavior).
+
+Memory Storage (Development)
+
+Memory storage keeps conversation state in-memory using a thread-safe HashMap. It’s perfect for local development, demos, and testing, but all state is lost when Plano restarts.
+
+Configuration
+
+Add this to your plano_config.yaml:
+
+state_storage:
+  type: memory
+
+That’s it. No additional setup required.
+
+When to Use Memory Storage
+
+Local development and debugging
+
+Demos and proof-of-concepts
+
+Automated testing environments
+
+Single-instance deployments where persistence isn’t critical
+
+Limitations
+
+State is lost on restart
+
+Not suitable for production workloads
+
+Cannot scale across multiple Plano instances
+
+PostgreSQL Storage (Production)
+
+PostgreSQL storage provides durable, production-grade conversation state management. It works with both self-hosted PostgreSQL and Supabase (PostgreSQL-as-a-service), making it ideal for scaling multi-agent systems in production.
+
+Prerequisites
+
+Before configuring PostgreSQL storage, you need:
+
+A PostgreSQL database (version 12 or later)
+
+Database credentials (host, user, password)
+
+The conversation_states table created in your database
+
+Setting Up the Database
+
+Run the SQL schema to create the required table:
+
+-- Conversation State Storage Table
+-- This table stores conversational context for the OpenAI Responses API
+-- Run this SQL against your PostgreSQL/Supabase database before enabling conversation state storage
+
+CREATE TABLE IF NOT EXISTS conversation_states (
+    response_id TEXT PRIMARY KEY,
+    input_items JSONB NOT NULL,
+    created_at BIGINT NOT NULL,
+    model TEXT NOT NULL,
+    provider TEXT NOT NULL,
+    updated_at TIMESTAMP DEFAULT CURRENT_TIMESTAMP
+);
+
+-- Indexes for common query patterns
+CREATE INDEX IF NOT EXISTS idx_conversation_states_created_at
+    ON conversation_states(created_at);
+
+CREATE INDEX IF NOT EXISTS idx_conversation_states_provider
+    ON conversation_states(provider);
+
+-- Optional: Add a policy for automatic cleanup of old conversations
+-- Uncomment and adjust the retention period as needed
+-- CREATE INDEX IF NOT EXISTS idx_conversation_states_updated_at
+--     ON conversation_states(updated_at);
+
+COMMENT ON TABLE conversation_states IS 'Stores conversation history for OpenAI Responses API continuity';
+COMMENT ON COLUMN conversation_states.response_id IS 'Unique identifier for the conversation state';
+COMMENT ON COLUMN conversation_states.input_items IS 'JSONB array of conversation messages and context';
+COMMENT ON COLUMN conversation_states.created_at IS 'Unix timestamp (seconds) when the conversation started';
+COMMENT ON COLUMN conversation_states.model IS 'Model name used for this conversation';
+COMMENT ON COLUMN conversation_states.provider IS 'LLM provider (e.g., openai, anthropic, bedrock)';
+
+
+Using psql:
+
+psql $DATABASE_URL -f docs/db_setup/conversation_states.sql
+
+Using Supabase Dashboard:
+
+Log in to your Supabase project
+
+Navigate to the SQL Editor
+
+Copy and paste the SQL from docs/db_setup/conversation_states.sql
+
+Run the query
+
+Configuration
+
+Once the database table is created, configure Plano to use PostgreSQL storage:
+
+state_storage:
+  type: postgres
+  connection_string: "postgresql://user:password@host:5432/database"
+
+Using Environment Variables
+
+You should never hardcode credentials. Use environment variables instead:
+
+state_storage:
+  type: postgres
+  connection_string: "postgresql://myuser:$DB_PASSWORD@db.example.com:5432/postgres"
+
+Then set the environment variable before running Plano:
+
+export DB_PASSWORD="your-secure-password"
+# Run Plano or config validation
+./plano
+
+Special Characters in Passwords: If your password contains special characters like #, @, or &, you must URL-encode them in the connection string. For example, MyPass#123 becomes MyPass%23123.
+
+Supabase Connection Strings
+
+Supabase requires different connection strings depending on your network setup. Most users should use the Session Pooler connection string.
+
+IPv4 Networks (Most Common)
+
+Use the Session Pooler connection string (port 5432):
+
+postgresql://postgres.[PROJECT-REF]:[PASSWORD]@aws-0-[REGION].pooler.supabase.com:5432/postgres
+
+IPv6 Networks
+
+Use the direct connection (port 5432):
+
+postgresql://postgres:[PASSWORD]@db.[PROJECT-REF].supabase.co:5432/postgres
+
+Finding Your Connection String
+
+Go to your Supabase project dashboard
+
+Navigate to Settings → Database → Connection Pooling
+
+Copy the Session mode connection string
+
+Replace [YOUR-PASSWORD] with your actual database password
+
+URL-encode special characters in the password
+
+Example Configuration
+
+state_storage:
+  type: postgres
+  connection_string: "postgresql://postgres.myproject:$DB_PASSWORD@aws-0-us-west-2.pooler.supabase.com:5432/postgres"
+
+Then set the environment variable:
+
+# If your password is "MyPass#123", encode it as "MyPass%23123"
+export DB_PASSWORD="MyPass%23123"
+
+Troubleshooting
+
+“Table ‘conversation_states’ does not exist”
+
+Run the SQL schema from docs/db_setup/conversation_states.sql against your database.
+
+Connection errors with Supabase
+
+Verify you’re using the correct connection string format (Session Pooler for IPv4)
+
+Check that your password is URL-encoded if it contains special characters
+
+Ensure your Supabase project hasn’t paused due to inactivity (free tier)
+
+Permission errors
+
+Ensure your database user has the following permissions:
+
+GRANT SELECT, INSERT, UPDATE, DELETE ON conversation_states TO your_user;
+
+State not persisting across requests
+
+Verify state_storage is configured in your plano_config.yaml
+
+Check Plano logs for state storage initialization messages
+
+Ensure the client is sending the prev_response_id={$response_id} from previous responses
+
+Best Practices
+
+Use environment variables for credentials: Never hardcode database passwords in configuration files.
+
+Start with memory storage for development: Switch to PostgreSQL when moving to production.
+
+Implement cleanup policies: Prevent unbounded growth by regularly archiving or deleting old conversations.
+
+Monitor storage usage: Track conversation state table size and query performance in production.
+
+Test failover scenarios: Ensure your application handles storage backend failures gracefully.
+
+Next Steps
+
+Learn more about building agents that leverage conversational state
+
+Explore filter chains for enriching conversation context
+
+See the LLM Providers guide for configuring model routing
+
+---
+
+Welcome to Plano!
+-----------------
+Doc: index
+
+Welcome to Plano!
+
+
+
+Plano is delivery infrastructure for agentic apps. A models-native proxy server and data plane designed to help you build agents faster, and deliver them reliably to production.
+
+Plano pulls out the rote plumbing work (aka “hidden AI middleware”) and decouples you from brittle, ever‑changing framework abstractions. It centralizes what shouldn’t be bespoke in every codebase like agent routing and orchestration, rich agentic signals and traces for continuous improvement, guardrail filters for safety and moderation, and smart LLM routing APIs for UX and DX agility. Use any language or AI framework, and ship agents to production faster with Plano.
+
+Built by contributors to the widely adopted Envoy Proxy, Plano helps developers focus more on the core product logic of agents, product teams accelerate feedback loops for reinforcement learning, and engineering teams standardize policies and access controls across every agent and LLM for safer, more reliable scaling.
+
+Get Started
+
+
+
+Concepts
+
+
+
+Guides
+
+
+
+Resources
+
+---
+
+Configuration Reference
+-----------------------
+Doc: resources/configuration_reference
+
+Configuration Reference
+
+The following is a complete reference of the plano_config.yml that controls the behavior of a single instance of
+the Arch gateway. This where you enable capabilities like routing to upstream LLm providers, defining prompt_targets
+where prompts get routed to, apply guardrails, and enable critical agent observability features.
+
+Plano Configuration - Full Reference
+
+
+# Arch Gateway configuration version
+version: v0.3.0
+
+
+# External HTTP agents - API type is controlled by request path (/v1/responses, /v1/messages, /v1/chat/completions)
+agents:
+  - id: weather_agent  # Example agent for weather
+    url: http://host.docker.internal:10510
+
+  - id: flight_agent   # Example agent for flights
+    url: http://host.docker.internal:10520
+
+
+# MCP filters applied to requests/responses (e.g., input validation, query rewriting)
+filters:
+  - id: input_guards  # Example filter for input validation
+    url: http://host.docker.internal:10500
+    # type: mcp (default)
+    # transport: streamable-http (default)
+    # tool: input_guards (default - same as filter id)
+
+
+# LLM provider configurations with API keys and model routing
+model_providers:
+  - model: openai/gpt-4o
+    access_key: $OPENAI_API_KEY
+    default: true
+
+  - model: openai/gpt-4o-mini
+    access_key: $OPENAI_API_KEY
+
+  - model: anthropic/claude-sonnet-4-0
+    access_key: $ANTHROPIC_API_KEY
+
+  - model: mistral/ministral-3b-latest
+    access_key: $MISTRAL_API_KEY
+
+
+# Model aliases - use friendly names instead of full provider model names
+model_aliases:
+  fast-llm:
+    target: gpt-4o-mini
+
+  smart-llm:
+    target: gpt-4o
+
+
+# HTTP listeners - entry points for agent routing, prompt targets, and direct LLM access
+listeners:
+  # Agent listener for routing requests to multiple agents
+  - type: agent
+    name: travel_booking_service
+    port: 8001
+    router: plano_orchestrator_v1
+    address: 0.0.0.0
+    agents:
+      - id: rag_agent
+        description: virtual assistant for retrieval augmented generation tasks
+        filter_chain:
+          - input_guards
+
+  # Model listener for direct LLM access
+  - type: model
+    name: model_1
+    address: 0.0.0.0
+    port: 12000
+
+  # Prompt listener for function calling (for prompt_targets)
+  - type: prompt
+    name: prompt_function_listener
+    address: 0.0.0.0
+    port: 10000
+    # This listener is used for prompt_targets and function calling
+
+
+# Reusable service endpoints
+endpoints:
+  app_server:
+    endpoint: 127.0.0.1:80
+    connect_timeout: 0.005s
+
+  mistral_local:
+    endpoint: 127.0.0.1:8001
+
+
+# Prompt targets for function calling and API orchestration
+prompt_targets:
+  - name: get_current_weather
+    description: Get current weather at a location.
+    parameters:
+      - name: location
+        description: The location to get the weather for
+        required: true
+        type: string
+        format: City, State
+      - name: days
+        description: the number of days for the request
+        required: true
+        type: int
+    endpoint:
+      name: app_server
+      path: /weather
+      http_method: POST
+
+
+# OpenTelemetry tracing configuration
+tracing:
+  # Random sampling percentage (1-100)
+  random_sampling: 100
+
+---
+
+Deployment
+----------
+Doc: resources/deployment
+
+Deployment
+
+This guide shows how to deploy Plano directly using Docker without the plano CLI, including basic runtime checks for routing and health monitoring.
+
+Docker Deployment
+
+Below is a minimal, production-ready example showing how to deploy the Plano  Docker image directly and run basic runtime checks. Adjust image names, tags, and the plano_config.yaml path to match your environment.
+
+You will need to pass all required environment variables that are referenced in your plano_config.yaml file.
+
+For plano_config.yaml, you can use any sample configuration defined earlier in the documentation. For example, you can try the LLM Routing sample config.
+
+Docker Compose Setup
+
+Create a docker-compose.yml file with the following configuration:
+
+# docker-compose.yml
+services:
+  plano:
+    image: katanemo/plano:0.4.0
+    container_name: plano
+    ports:
+      - "10000:10000" # ingress (client -> plano)
+      - "12000:12000" # egress (plano -> upstream/llm proxy)
+    volumes:
+      - ./plano_config.yaml:/app/plano_config.yaml:ro
+    environment:
+      - OPENAI_API_KEY=${OPENAI_API_KEY:?error}
+      - ANTHROPIC_API_KEY=${ANTHROPIC_API_KEY:?error}
+
+Starting the Stack
+
+Start the services from the directory containing docker-compose.yml and plano_config.yaml:
+
+# Set required environment variables and start services
+OPENAI_API_KEY=xxx ANTHROPIC_API_KEY=yyy docker compose up -d
+
+Check container health and logs:
+
+docker compose ps
+docker compose logs -f plano
+
+Runtime Tests
+
+Perform basic runtime tests to verify routing and functionality.
+
+Gateway Smoke Test
+
+Test the chat completion endpoint with automatic routing:
+
+# Request handled by the gateway. 'model: "none"' lets Plano  decide routing
+curl --header 'Content-Type: application/json' \
+  --data '{"messages":[{"role":"user","content":"tell me a joke"}], "model":"none"}' \
+  http://localhost:12000/v1/chat/completions | jq .model
+
+Expected output:
+
+"gpt-5.2"
+
+Model-Based Routing
+
+Test explicit provider and model routing:
+
+curl -s -H "Content-Type: application/json" \
+  -d '{"messages":[{"role":"user","content":"Explain quantum computing"}], "model":"anthropic/claude-sonnet-4-5"}' \
+  http://localhost:12000/v1/chat/completions | jq .model
+
+Expected output:
+
+"claude-sonnet-4-5"
+
+Troubleshooting
+
+Common Issues and Solutions
+
+Environment Variables
+
+Ensure all environment variables (OPENAI_API_KEY, ANTHROPIC_API_KEY, etc.) used by plano_config.yaml are set before starting services.
+
+TLS/Connection Errors
+
+If you encounter TLS or connection errors to upstream providers:
+
+Check DNS resolution
+
+Verify proxy settings
+
+Confirm correct protocol and port in your plano_config endpoints
+
+Verbose Logging
+
+To enable more detailed logs for debugging:
+
+Run plano with a higher component log level
+
+See the Observability guide for logging and monitoring details
+
+Rebuild the image if required with updated log configuration
+
+CI/Automated Checks
+
+For continuous integration or automated testing, you can use the curl commands above as health checks in your deployment pipeline.
+
+---
+
+llms.txt
+--------
+Doc: resources/llms_txt
+
+llms.txt
+
+This project generates a single plaintext file containing the compiled text of all documentation pages, useful for large context models to reference Plano documentation.
+
+Open it here: llms.txt
+
+---
+
+Bright Staff
+------------
+Doc: resources/tech_overview/model_serving
+
+Bright Staff
+
+Bright Staff is Plano’s memory-efficient, lightweight controller for agentic traffic. It sits inside the Plano
+data plane and makes real-time decisions about how prompts are handled, forwarded, and processed.
+
+Rather than running a separate “model server” subsystem, Plano relies on Envoy’s HTTP connection management
+and cluster subsystem to talk to different models and backends over HTTP(S). Bright Staff uses these primitives to:
+* Inspect prompts, conversation state, and metadata.
+* Decide which upstream model(s), tool backends, or APIs to call, and in what order.
+* Coordinate retries, fallbacks, and traffic splitting across providers and models.
+
+Plano is designed to run alongside your application servers in your cloud VPC, on-premises, or in local
+development. It does not require a GPU itself; GPUs live where your models are hosted (third-party APIs or your
+own deployments), and Plano reaches them via HTTP.
+
+---
+
+Request Lifecycle
+-----------------
+Doc: resources/tech_overview/request_lifecycle
+
+Request Lifecycle
+
+Below we describe the events in the lifecycle of a request passing through a Plano instance. We first
+describe how Plano fits into the request path and then the internal events that take place following
+the arrival of a request at Plano from downstream clients. We follow the request until the corresponding
+dispatch upstream and the response path.
+
+
+
+Network topology
+
+How a request flows through the components in a network (including Plano) depends on the network’s topology.
+Plano can be used in a wide variety of networking topologies. We focus on the inner operations of Plano below,
+but briefly we address how Plano relates to the rest of the network in this section.
+
+Downstream(Ingress) listeners take requests from upstream clients like a web UI or clients that forward
+prompts to you local application responses from the application flow back through Plano to the downstream.
+
+Upstream(Egress) listeners take requests from the application and forward them to LLMs.
+
+High level architecture
+
+Plano is a set of two self-contained processes that are designed to run alongside your application servers
+(or on a separate server connected to your application servers via a network).
+
+The first process is designated to manage HTTP-level networking and connection management concerns (protocol management, request id generation, header sanitization, etc.), and the other process is a controller, which helps Plano make intelligent decisions about the incoming prompts. The controller hosts the purpose-built LLMs to manage several critical, but undifferentiated, prompt related tasks on behalf of developers.
+
+The request processing path in Plano has three main parts:
+
+Listener subsystem which handles downstream and upstream request
+processing. It is responsible for managing the inbound(edge) and outbound(egress) request lifecycle. The downstream and upstream HTTP/2 codec lives here. This also includes the lifecycle of any upstream connection to an LLM provider or tool backend. The listenser subsystmem manages connection pools, load balancing, retries, and failover.
+
+Bright Staff controller subsystem is Plano’s memory-efficient, lightweight controller for agentic traffic. It sits inside the Plano data plane and makes real-time decisions about how prompts are handled, forwarded, and processed.
+
+These two subsystems are bridged with either the HTTP router filter, and the cluster manager subsystems of Envoy.
+
+Also, Plano utilizes Envoy event-based thread model. A main thread is responsible for the server lifecycle, configuration processing, stats, etc. and some number of worker threads process requests. All threads operate around an event loop (libevent) and any given downstream TCP connection will be handled by exactly one worker thread for its lifetime. Each worker thread maintains its own pool of TCP connections to upstream endpoints.
+
+Worker threads rarely share state and operate in a trivially parallel fashion. This threading model
+enables scaling to very high core count CPUs.
+
+Request Flow (Ingress)
+
+A brief outline of the lifecycle of a request and response using the example configuration above:
+
+TCP Connection Establishment:
+A TCP connection from downstream is accepted by an Plano listener running on a worker thread.
+The listener filter chain provides SNI and other pre-TLS information. The transport socket, typically TLS,
+decrypts incoming data for processing.
+
+Routing Decision (Agent vs Prompt Target):
+The decrypted data stream is de-framed by the HTTP/2 codec in Plano’s HTTP connection manager. Plano performs
+intent matching (via the Bright Staff controller and prompt-handling logic) using the configured agents and
+prompt targets, determining whether this request should be handled by an agent workflow
+(with optional Filter Chains) or by a deterministic prompt target.
+
+4a. Agent Path: Orchestration and Filter Chains
+
+If the request is routed to an agent, Plano executes any attached Filter Chains first. These filters can apply guardrails, rewrite prompts, or enrich context (for example, RAG retrieval) before the agent runs. Once filters complete, the Bright Staff controller orchestrates which downstream tools, APIs, or LLMs the agent should call and in what sequence.
+
+Plano may call one or more backend APIs or tools on behalf of the agent.
+
+If an endpoint cluster is identified, load balancing is performed, circuit breakers are checked, and the request is proxied to the appropriate upstream endpoint.
+
+If no specific endpoint is required, the prompt is sent to an upstream LLM using Plano’s model proxy for
+completion or summarization.
+
+For more on agent workflows and orchestration, see Prompt Targets and Agents and
+Agent Filter Chains.
+
+4b. Prompt Target Path: Deterministic Tool/API Calls
+
+If the request is routed to a prompt target, Plano treats it as a deterministic, task-specific call.
+Plano engages its function-calling and parameter-gathering capabilities to extract the necessary details
+from the incoming prompt(s) and produce the structured inputs your backend expects.
+
+Parameter Gathering: Plano extracts and validates parameters defined on the prompt target (for example,
+currency symbols, dates, or entity identifiers) so your backend does not need to parse natural language.
+
+API Call Execution: Plano then routes the call to the configured backend endpoint. If an endpoint cluster is identified, load balancing and circuit-breaker checks are applied before proxying the request upstream.
+
+For more on how to design and configure prompt targets, see Prompt Target.
+
+Error Handling and Forwarding:
+Errors encountered during processing, such as failed function calls or guardrail detections, are forwarded to
+designated error targets. Error details are communicated through specific headers to the application:
+
+X-Function-Error-Code: Code indicating the type of function call error.
+
+X-Prompt-Guard-Error-Code: Code specifying violations detected by prompt guardrails.
+
+Additional headers carry messages and timestamps to aid in debugging and logging.
+
+Response Handling:
+The upstream endpoint’s TLS transport socket encrypts the response, which is then proxied back downstream.
+Responses pass through HTTP filters in reverse order, ensuring any necessary processing or modification before final delivery.
+
+Request Flow (Egress)
+
+A brief outline of the lifecycle of a request and response in the context of egress traffic from an application to Large Language Models (LLMs) via Plano:
+
+HTTP Connection Establishment to LLM:
+Plano initiates an HTTP connection to the upstream LLM service. This connection is handled by Plano’s egress listener running on a worker thread. The connection typically uses a secure transport protocol such as HTTPS, ensuring the prompt data is encrypted before being sent to the LLM service.
+
+Rate Limiting:
+Before sending the request to the LLM, Plano applies rate-limiting policies to ensure that the upstream LLM service is not overwhelmed by excessive traffic. Rate limits are enforced per client or service, ensuring fair usage and preventing accidental or malicious overload. If the rate limit is exceeded, Plano may return an appropriate HTTP error (e.g., 429 Too Many Requests) without sending the prompt to the LLM.
+
+Seamless Request Transformation and Smart Routing:
+After rate limiting, Plano normalizes the outgoing request into a provider-agnostic shape and applies smart routing decisions using the configured LLM Providers. This includes translating client-specific conventions into a unified OpenAI-style contract, enriching or overriding parameters (for example, temperature or max tokens) based on policy, and choosing the best target model or provider using model-based, alias-based, or preference-aligned routing.
+
+Load Balancing to (hosted) LLM Endpoints:
+After smart routing selects the target provider/model, Plano routes the prompt to the appropriate LLM endpoint.
+If multiple LLM provider instances are available, load balancing is performed to distribute traffic evenly
+across the instances. Plano checks the health of the LLM endpoints using circuit breakers and health checks,
+ensuring that the prompt is only routed to a healthy, responsive instance.
+
+Response Reception and Forwarding:
+Once the LLM processes the prompt, Plano receives the response from the LLM service. The response is typically a generated text, completion, or summarization. Upon reception, Plano decrypts (if necessary) and handles the response, passing it through any egress processing pipeline defined by the application, such as logging or additional response filtering.
+
+Post-request processing
+
+Once a request completes, the stream is destroyed. The following also takes places:
+
+The post-request monitoring are updated (e.g. timing, active requests, upgrades, health checks).
+Some statistics are updated earlier however, during request processing. Stats are batched and written by the main
+thread periodically.
+
+Access logs are written to the access log
+
+Trace spans are finalized. If our example request was traced, a
+trace span, describing the duration and details of the request would be created by the HCM when
+processing request headers and then finalized by the HCM during post-request processing.
+
+Configuration
+
+Today, only support a static bootstrap configuration file for simplicity today:
+
+version: v0.2.0
+
+listeners:
+  ingress_traffic:
+    address: 0.0.0.0
+    port: 10000
+
+# Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way
+model_providers:
+  - access_key: $OPENAI_API_KEY
+    model: openai/gpt-4o
+    default: true
+
+prompt_targets:
+  - name: information_extraction
+    default: true
+    description: handel all scenarios that are question and answer in nature. Like summarization, information extraction, etc.
+    endpoint:
+      name: app_server
+      path: /agent/summary
+    # Arch uses the default LLM and treats the response from the endpoint as the prompt to send to the LLM
+    auto_llm_dispatch_on_response: true
+    # override system prompt for this prompt target
+    system_prompt: You are a helpful information extraction assistant. Use the information that is provided to you.
+
+  - name: reboot_network_device
+    description: Reboot a specific network device
+    endpoint:
+      name: app_server
+      path: /agent/action
+    parameters:
+      - name: device_id
+        type: str
+        description: Identifier of the network device to reboot.
+        required: true
+      - name: confirmation
+        type: bool
+        description: Confirmation flag to proceed with reboot.
+        default: false
+        enum: [true, false]
+
+# Arch creates a round-robin load balancing between different endpoints, managed via the cluster subsystem.
+endpoints:
+  app_server:
+    # value could be ip address or a hostname with port
+    # this could also be a list of endpoints for load balancing
+    # for example endpoint: [ ip1:port, ip2:port ]
+    endpoint: 127.0.0.1:80
+    # max time to wait for a connection to be established
+    connect_timeout: 0.005s
+
+---
+
+Tech Overview
+-------------
+Doc: resources/tech_overview/tech_overview
+
+Tech Overview
+
+---
+
+Threading Model
+---------------
+Doc: resources/tech_overview/threading_model
+
+Threading Model
+
+Plano builds on top of Envoy’s single process with multiple threads architecture.
+
+A single primary thread controls various sporadic coordination tasks while some number of worker
+threads perform filtering, and forwarding.
+
+Once a connection is accepted, the connection spends the rest of its lifetime bound to a single worker
+thread. All the functionality around prompt handling from a downstream client is handled in a separate worker thread.
+This allows the majority of Plano to be largely single threaded (embarrassingly parallel) with a small amount
+of more complex code handling coordination between the worker threads.
+
+Generally, Plano is written to be 100% non-blocking.
+
+For most workloads we recommend configuring the number of worker threads to be equal to the number of
+hardware threads on the machine.
+
+---
diff --git a/index.html b/index.html
index 65b31b58..ec248535 100755
--- a/index.html
+++ b/index.html
@@ -7,13 +7,13 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Welcome to Arch! | Arch Docs v0.3.22</title>
-<meta content="Welcome to Arch! | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Welcome to Arch! | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Welcome to Plano! | Plano Docs v0.4</title>
+<meta content="Welcome to Plano! | Plano Docs v0.4" property="og:title"/>
+<meta content="Welcome to Plano! | Plano Docs v0.4" name="twitter:title"/>
 <link href="_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/index.html" rel="canonical"/>
 <link href="_static/favicon.ico" rel="icon"/>
@@ -38,7 +38,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="#">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -55,7 +55,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -74,40 +74,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="#">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -117,27 +110,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="resources/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -151,12 +146,13 @@
 <main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
 <div class="w-full min-w-0 mx-auto">
 <div id="content" role="main">
-<section id="welcome-to-arch">
-<h1>Welcome to Arch!<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#welcome-to-arch"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<a class="reference internal image-reference" href="_images/arch-logo.png"><img alt="_images/arch-logo.png" class="align-center" src="_images/arch-logo.png" style="width: 100%;"/>
+<section id="welcome-to-plano">
+<h1>Welcome to Plano!<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#welcome-to-plano"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<a class="reference internal image-reference" href="_images/PlanoTagline.svg"><img alt="_images/PlanoTagline.svg" class="align-center" src="_images/PlanoTagline.svg" style="width: 100%;"/>
 </a>
-<a href="https://www.producthunt.com/posts/arch-3?embed=true&amp;utm_source=badge-top-post-badge&amp;utm_medium=badge&amp;utm_souce=badge-arch-3" target="_blank"><img alt="Arch - Build fast, hyper-personalized agents with intelligent infra | Product Hunt" height="54" src="https://api.producthunt.com/widgets/embed-image/v1/top-post-badge.svg?post_id=565761&amp;theme=dark&amp;period=daily&amp;t=1742433071161" style="width: 250px; height: 54px;" width="250"/></a><p><a class="reference external" href="https://github.com/katanemo/arch" rel="nofollow noopener">Arch<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> is a models-native edge and LLM proxy/gateway for AI agents - one that is natively designed to handle and process prompts, not just network traffic.</p>
-<p>Built by contributors to the widely adopted <a class="reference external" href="https://www.envoyproxy.io/" rel="nofollow noopener">Envoy Proxy<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, Arch handles the <em>pesky low-level work</em> in building agentic apps — like applying guardrails, clarifying vague user input, routing prompts to the right agent, and unifying access to any LLM. It’s a language and framework friendly infrastructure layer designed to help you build and ship agentic apps faster.</p>
+<p><a class="reference external" href="https://github.com/katanemo/plano" rel="nofollow noopener">Plano<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a> is delivery infrastructure for agentic apps. A models-native proxy server and data plane designed to help you build agents faster, and deliver them reliably to production.</p>
+<p>Plano pulls out the rote plumbing work (aka “hidden AI middleware”) and decouples you from brittle, ever‑changing framework abstractions. It centralizes what shouldn’t be bespoke in every codebase like agent routing and orchestration, rich agentic signals and traces for continuous improvement, guardrail filters for safety and moderation, and smart LLM routing APIs for UX and DX agility. Use any language or AI framework, and ship agents to production faster with Plano.</p>
+<p>Built by contributors to the widely adopted <a class="reference external" href="https://www.envoyproxy.io/" rel="nofollow noopener">Envoy Proxy<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>, Plano <strong>helps developers</strong> focus more on the core product logic of agents, <strong>product teams</strong> accelerate feedback loops for reinforcement learning, and <strong>engineering teams</strong> standardize policies and access controls across every agent and LLM for safer, more reliable scaling.</p>
 <div class="sd-tab-set docutils">
 <input checked="checked" id="sd-tab-item-0" name="sd-tab-set-0" type="radio"/>
 <label class="sd-tab-label" for="sd-tab-item-0">
@@ -165,7 +161,7 @@ Get Started</label><div class="sd-tab-content docutils">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
@@ -177,17 +173,10 @@ Concepts</label><div class="sd-tab-content docutils">
 <div class="toctree-wrapper compound">
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="concepts/tech_overview/tech_overview.html">Tech Overview</a><ul>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="concepts/llm_providers/llm_providers.html">LLM Providers</a><ul>
+<li class="toctree-l1"><a class="reference internal" href="concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1"><a class="reference internal" href="concepts/llm_providers/llm_providers.html">Model (LLM) Providers</a><ul>
 <li class="toctree-l2"><a class="reference internal" href="concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -203,39 +192,35 @@ Guides</label><div class="sd-tab-content docutils">
 <div class="toctree-wrapper compound">
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1"><a class="reference internal" href="guides/observability/observability.html">Observability</a><ul>
 <li class="toctree-l2"><a class="reference internal" href="guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
+<li class="toctree-l1"><a class="reference internal" href="guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/state.html">Conversational State</a></li>
 </ul>
 </div>
 </div>
 <input id="sd-tab-item-3" name="sd-tab-set-0" type="radio"/>
 <label class="sd-tab-label" for="sd-tab-item-3">
-Build with Arch</label><div class="sd-tab-content docutils">
-<div class="toctree-wrapper compound">
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="build_with_arch/multi_turn.html">Multi-Turn</a></li>
-</ul>
-</div>
-</div>
-<input id="sd-tab-item-4" name="sd-tab-set-0" type="radio"/>
-<label class="sd-tab-label" for="sd-tab-item-4">
 Resources</label><div class="sd-tab-content docutils">
 <div class="toctree-wrapper compound">
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1"><a class="reference internal" href="resources/tech_overview/tech_overview.html">Tech Overview</a><ul>
+<li class="toctree-l2"><a class="reference internal" href="resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="resources/llms_txt.html">llms.txt</a></li>
 </ul>
 </div>
 </div>
@@ -256,12 +241,12 @@ Resources</label><div class="sd-tab-content docutils">
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="_static/documentation_options.js?v=3bad885e"></script>
+<script src="_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="_static/doctools.js?v=9bcbadda"></script>
 <script src="_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="_static/theme.js?v=073f68d9"></script>
diff --git a/objects.inv b/objects.inv
index 30a5105a..3d1c6d72 100755
Binary files a/objects.inv and b/objects.inv differ
diff --git a/resources/configuration_reference.html b/resources/configuration_reference.html
index 1df96724..198a3963 100755
--- a/resources/configuration_reference.html
+++ b/resources/configuration_reference.html
@@ -7,17 +7,18 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Configuration Reference | Arch Docs v0.3.22</title>
-<meta content="Configuration Reference | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Configuration Reference | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Configuration Reference | Plano Docs v0.4</title>
+<meta content="Configuration Reference | Plano Docs v0.4" property="og:title"/>
+<meta content="Configuration Reference | Plano Docs v0.4" name="twitter:title"/>
 <link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/resources/configuration_reference.html" rel="canonical"/>
 <link href="../_static/favicon.ico" rel="icon"/>
 <link href="../search.html" rel="search" title="Search"/>
+<link href="llms_txt.html" rel="next" title="llms.txt"/>
 <link href="deployment.html" rel="prev" title="Deployment"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
@@ -38,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -55,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -74,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -117,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul class="current">
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="deployment.html">Deployment</a></li>
 <li class="toctree-l1 current"><a class="current reference internal" href="#">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -152,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -162,111 +158,120 @@
 <div id="content" role="main">
 <section id="configuration-reference">
 <span id="id1"></span><h1>Configuration Reference<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuration-reference"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>The following is a complete reference of the <code class="docutils literal notranslate"><span class="pre">arch_config.yml</span></code> that controls the behavior of a single instance of
+<p>The following is a complete reference of the <code class="docutils literal notranslate"><span class="pre">plano_config.yml</span></code> that controls the behavior of a single instance of
 the Arch gateway. This where you enable capabilities like routing to upstream LLm providers, defining prompt_targets
 where prompts get routed to, apply guardrails, and enable critical agent observability features.</p>
 <div class="literal-block-wrapper docutils container" id="id2">
-<div class="code-block-caption"><span class="caption-text"><a class="reference download internal" download="" href="../_downloads/ca9d3b7116524473d8adbde7cf15d167/arch_config_full_reference.yaml"><code class="xref download docutils literal notranslate"><span class="pre">Arch</span> <span class="pre">Configuration</span> <span class="pre">-</span> <span class="pre">Full</span> <span class="pre">Reference</span></code></a></span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos">  1</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1</span>
-</span><span id="line-2"><span class="linenos">  2</span>
-</span><span id="line-3"><span class="linenos">  3</span><span class="nt">listeners</span><span class="p">:</span>
-</span><span id="line-4"><span class="linenos">  4</span><span class="w">  </span><span class="nt">ingress_traffic</span><span class="p">:</span>
-</span><span id="line-5"><span class="linenos">  5</span><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
-</span><span id="line-6"><span class="linenos">  6</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">10000</span>
-</span><span id="line-7"><span class="linenos">  7</span><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-8"><span class="linenos">  8</span><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">5s</span>
-</span><span id="line-9"><span class="linenos">  9</span><span class="w">  </span><span class="nt">egress_traffic</span><span class="p">:</span>
-</span><span id="line-10"><span class="linenos"> 10</span><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
-</span><span id="line-11"><span class="linenos"> 11</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">12000</span>
-</span><span id="line-12"><span class="linenos"> 12</span><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-13"><span class="linenos"> 13</span><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">5s</span>
-</span><span id="line-14"><span class="linenos"> 14</span>
-</span><span id="line-15"><span class="linenos"> 15</span><span class="c1"># Arch creates a round-robin load balancing between different endpoints, managed via the cluster subsystem.</span>
-</span><span id="line-16"><span class="linenos"> 16</span><span class="nt">endpoints</span><span class="p">:</span>
-</span><span id="line-17"><span class="linenos"> 17</span><span class="w">  </span><span class="nt">app_server</span><span class="p">:</span>
-</span><span id="line-18"><span class="linenos"> 18</span><span class="w">    </span><span class="c1"># value could be ip address or a hostname with port</span>
-</span><span id="line-19"><span class="linenos"> 19</span><span class="w">    </span><span class="c1"># this could also be a list of endpoints for load balancing</span>
-</span><span id="line-20"><span class="linenos"> 20</span><span class="w">    </span><span class="c1"># for example endpoint: [ ip1:port, ip2:port ]</span>
-</span><span id="line-21"><span class="linenos"> 21</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1:80</span>
-</span><span id="line-22"><span class="linenos"> 22</span><span class="w">    </span><span class="c1"># max time to wait for a connection to be established</span>
-</span><span id="line-23"><span class="linenos"> 23</span><span class="w">    </span><span class="nt">connect_timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.005s</span>
-</span><span id="line-24"><span class="linenos"> 24</span>
-</span><span id="line-25"><span class="linenos"> 25</span><span class="w">  </span><span class="nt">mistral_local</span><span class="p">:</span>
-</span><span id="line-26"><span class="linenos"> 26</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1:8001</span>
-</span><span id="line-27"><span class="linenos"> 27</span>
-</span><span id="line-28"><span class="linenos"> 28</span><span class="w">  </span><span class="nt">error_target</span><span class="p">:</span>
-</span><span id="line-29"><span class="linenos"> 29</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">error_target_1</span>
-</span><span id="line-30"><span class="linenos"> 30</span>
-</span><span id="line-31"><span class="linenos"> 31</span><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
-</span><span id="line-32"><span class="linenos"> 32</span><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-33"><span class="linenos"> 33</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
-</span><span id="line-34"><span class="linenos"> 34</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-35"><span class="linenos"> 35</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
-</span><span id="line-36"><span class="linenos"> 36</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+<div class="code-block-caption"><span class="caption-text"><a class="reference download internal" download="" href="../_downloads/ca9d3b7116524473d8adbde7cf15d167/arch_config_full_reference.yaml"><code class="xref download docutils literal notranslate"><span class="pre">Plano</span> <span class="pre">Configuration</span> <span class="pre">-</span> <span class="pre">Full</span> <span class="pre">Reference</span></code></a></span><a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#id2"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></div>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="linenos">  1</span><span class="c1"># Arch Gateway configuration version</span>
+</span><span id="line-2"><span class="linenos">  2</span><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.3.0</span>
+</span><span id="line-3"><span class="linenos">  3</span>
+</span><span id="line-4"><span class="linenos">  4</span>
+</span><span id="line-5"><span class="linenos">  5</span><span class="c1"># External HTTP agents - API type is controlled by request path (/v1/responses, /v1/messages, /v1/chat/completions)</span>
+</span><span id="line-6"><span class="linenos">  6</span><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-7"><span class="linenos">  7</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">weather_agent</span><span class="w">  </span><span class="c1"># Example agent for weather</span>
+</span><span id="line-8"><span class="linenos">  8</span><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10510</span>
+</span><span id="line-9"><span class="linenos">  9</span>
+</span><span id="line-10"><span class="linenos"> 10</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">flight_agent</span><span class="w">   </span><span class="c1"># Example agent for flights</span>
+</span><span id="line-11"><span class="linenos"> 11</span><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10520</span>
+</span><span id="line-12"><span class="linenos"> 12</span>
+</span><span id="line-13"><span class="linenos"> 13</span>
+</span><span id="line-14"><span class="linenos"> 14</span><span class="c1"># MCP filters applied to requests/responses (e.g., input validation, query rewriting)</span>
+</span><span id="line-15"><span class="linenos"> 15</span><span class="nt">filters</span><span class="p">:</span>
+</span><span id="line-16"><span class="linenos"> 16</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">input_guards</span><span class="w">  </span><span class="c1"># Example filter for input validation</span>
+</span><span id="line-17"><span class="linenos"> 17</span><span class="w">    </span><span class="nt">url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://host.docker.internal:10500</span>
+</span><span id="line-18"><span class="linenos"> 18</span><span class="w">    </span><span class="c1"># type: mcp (default)</span>
+</span><span id="line-19"><span class="linenos"> 19</span><span class="w">    </span><span class="c1"># transport: streamable-http (default)</span>
+</span><span id="line-20"><span class="linenos"> 20</span><span class="w">    </span><span class="c1"># tool: input_guards (default - same as filter id)</span>
+</span><span id="line-21"><span class="linenos"> 21</span>
+</span><span id="line-22"><span class="linenos"> 22</span>
+</span><span id="line-23"><span class="linenos"> 23</span><span class="c1"># LLM provider configurations with API keys and model routing</span>
+</span><span id="line-24"><span class="linenos"> 24</span><span class="nt">model_providers</span><span class="p">:</span>
+</span><span id="line-25"><span class="linenos"> 25</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
+</span><span id="line-26"><span class="linenos"> 26</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-27"><span class="linenos"> 27</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-28"><span class="linenos"> 28</span>
+</span><span id="line-29"><span class="linenos"> 29</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o-mini</span>
+</span><span id="line-30"><span class="linenos"> 30</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-31"><span class="linenos"> 31</span>
+</span><span id="line-32"><span class="linenos"> 32</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">anthropic/claude-sonnet-4-0</span>
+</span><span id="line-33"><span class="linenos"> 33</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$ANTHROPIC_API_KEY</span>
+</span><span id="line-34"><span class="linenos"> 34</span>
+</span><span id="line-35"><span class="linenos"> 35</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">mistral/ministral-3b-latest</span>
+</span><span id="line-36"><span class="linenos"> 36</span><span class="w">    </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$MISTRAL_API_KEY</span>
 </span><span id="line-37"><span class="linenos"> 37</span>
-</span><span id="line-38"><span class="linenos"> 38</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$MISTRAL_API_KEY</span>
-</span><span id="line-39"><span class="linenos"> 39</span><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">mistral/mistral-8x7b</span>
-</span><span id="line-40"><span class="linenos"> 40</span>
-</span><span id="line-41"><span class="linenos"> 41</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">mistral/mistral-7b-instruct</span>
-</span><span id="line-42"><span class="linenos"> 42</span><span class="w">    </span><span class="nt">base_url</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">http://mistral_local</span>
+</span><span id="line-38"><span class="linenos"> 38</span>
+</span><span id="line-39"><span class="linenos"> 39</span><span class="c1"># Model aliases - use friendly names instead of full provider model names</span>
+</span><span id="line-40"><span class="linenos"> 40</span><span class="nt">model_aliases</span><span class="p">:</span>
+</span><span id="line-41"><span class="linenos"> 41</span><span class="w">  </span><span class="nt">fast-llm</span><span class="p">:</span>
+</span><span id="line-42"><span class="linenos"> 42</span><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-4o-mini</span>
 </span><span id="line-43"><span class="linenos"> 43</span>
-</span><span id="line-44"><span class="linenos"> 44</span><span class="c1"># Model aliases - friendly names that map to actual provider names</span>
-</span><span id="line-45"><span class="linenos"> 45</span><span class="nt">model_aliases</span><span class="p">:</span>
-</span><span id="line-46"><span class="linenos"> 46</span><span class="w">  </span><span class="c1"># Alias for summarization tasks -&gt; fast/cheap model</span>
-</span><span id="line-47"><span class="linenos"> 47</span><span class="w">  </span><span class="nt">arch.summarize.v1</span><span class="p">:</span>
-</span><span id="line-48"><span class="linenos"> 48</span><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-4o</span>
-</span><span id="line-49"><span class="linenos"> 49</span>
-</span><span id="line-50"><span class="linenos"> 50</span><span class="w">  </span><span class="c1"># Alias for general purpose tasks -&gt; latest model</span>
-</span><span id="line-51"><span class="linenos"> 51</span><span class="w">  </span><span class="nt">arch.v1</span><span class="p">:</span>
-</span><span id="line-52"><span class="linenos"> 52</span><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">mistral-8x7b</span>
-</span><span id="line-53"><span class="linenos"> 53</span>
-</span><span id="line-54"><span class="linenos"> 54</span><span class="c1"># provides a way to override default settings for the arch system</span>
-</span><span id="line-55"><span class="linenos"> 55</span><span class="nt">overrides</span><span class="p">:</span>
-</span><span id="line-56"><span class="linenos"> 56</span><span class="w">  </span><span class="c1"># By default Arch uses an NLI + embedding approach to match an incoming prompt to a prompt target.</span>
-</span><span id="line-57"><span class="linenos"> 57</span><span class="w">  </span><span class="c1"># The intent matching threshold is kept at 0.80, you can override this behavior if you would like</span>
-</span><span id="line-58"><span class="linenos"> 58</span><span class="w">  </span><span class="nt">prompt_target_intent_matching_threshold</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.60</span>
-</span><span id="line-59"><span class="linenos"> 59</span>
-</span><span id="line-60"><span class="linenos"> 60</span><span class="c1"># default system prompt used by all prompt targets</span>
-</span><span id="line-61"><span class="linenos"> 61</span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">You are a network assistant that just offers facts; not advice on manufacturers or purchasing decisions.</span>
-</span><span id="line-62"><span class="linenos"> 62</span>
-</span><span id="line-63"><span class="linenos"> 63</span><span class="nt">prompt_guards</span><span class="p">:</span>
-</span><span id="line-64"><span class="linenos"> 64</span><span class="w">  </span><span class="nt">input_guards</span><span class="p">:</span>
-</span><span id="line-65"><span class="linenos"> 65</span><span class="w">    </span><span class="nt">jailbreak</span><span class="p">:</span>
-</span><span id="line-66"><span class="linenos"> 66</span><span class="w">      </span><span class="nt">on_exception</span><span class="p">:</span>
-</span><span id="line-67"><span class="linenos"> 67</span><span class="w">        </span><span class="nt">message</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Looks like you're curious about my abilities, but I can only provide assistance within my programmed parameters.</span>
-</span><span id="line-68"><span class="linenos"> 68</span>
-</span><span id="line-69"><span class="linenos"> 69</span><span class="nt">prompt_targets</span><span class="p">:</span>
-</span><span id="line-70"><span class="linenos"> 70</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">information_extraction</span>
-</span><span id="line-71"><span class="linenos"> 71</span><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-72"><span class="linenos"> 72</span><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">handel all scenarios that are question and answer in nature. Like summarization, information extraction, etc.</span>
-</span><span id="line-73"><span class="linenos"> 73</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
-</span><span id="line-74"><span class="linenos"> 74</span><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</span><span id="line-75"><span class="linenos"> 75</span><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/summary</span>
-</span><span id="line-76"><span class="linenos"> 76</span><span class="w">      </span><span class="nt">http_method</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">POST</span>
-</span><span id="line-77"><span class="linenos"> 77</span><span class="w">    </span><span class="c1"># Arch uses the default LLM and treats the response from the endpoint as the prompt to send to the LLM</span>
-</span><span id="line-78"><span class="linenos"> 78</span><span class="w">    </span><span class="nt">auto_llm_dispatch_on_response</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-79"><span class="linenos"> 79</span><span class="w">    </span><span class="c1"># override system prompt for this prompt target</span>
-</span><span id="line-80"><span class="linenos"> 80</span><span class="w">    </span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">You are a helpful information extraction assistant. Use the information that is provided to you.</span>
+</span><span id="line-44"><span class="linenos"> 44</span><span class="w">  </span><span class="nt">smart-llm</span><span class="p">:</span>
+</span><span id="line-45"><span class="linenos"> 45</span><span class="w">    </span><span class="nt">target</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">gpt-4o</span>
+</span><span id="line-46"><span class="linenos"> 46</span>
+</span><span id="line-47"><span class="linenos"> 47</span>
+</span><span id="line-48"><span class="linenos"> 48</span><span class="c1"># HTTP listeners - entry points for agent routing, prompt targets, and direct LLM access</span>
+</span><span id="line-49"><span class="linenos"> 49</span><span class="nt">listeners</span><span class="p">:</span>
+</span><span id="line-50"><span class="linenos"> 50</span><span class="w">  </span><span class="c1"># Agent listener for routing requests to multiple agents</span>
+</span><span id="line-51"><span class="linenos"> 51</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">agent</span>
+</span><span id="line-52"><span class="linenos"> 52</span><span class="w">    </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">travel_booking_service</span>
+</span><span id="line-53"><span class="linenos"> 53</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">8001</span>
+</span><span id="line-54"><span class="linenos"> 54</span><span class="w">    </span><span class="nt">router</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">plano_orchestrator_v1</span>
+</span><span id="line-55"><span class="linenos"> 55</span><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
+</span><span id="line-56"><span class="linenos"> 56</span><span class="w">    </span><span class="nt">agents</span><span class="p">:</span>
+</span><span id="line-57"><span class="linenos"> 57</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">id</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">rag_agent</span>
+</span><span id="line-58"><span class="linenos"> 58</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">virtual assistant for retrieval augmented generation tasks</span>
+</span><span id="line-59"><span class="linenos"> 59</span><span class="w">        </span><span class="nt">filter_chain</span><span class="p">:</span>
+</span><span id="line-60"><span class="linenos"> 60</span><span class="w">          </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">input_guards</span>
+</span><span id="line-61"><span class="linenos"> 61</span>
+</span><span id="line-62"><span class="linenos"> 62</span><span class="w">  </span><span class="c1"># Model listener for direct LLM access</span>
+</span><span id="line-63"><span class="linenos"> 63</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">model</span>
+</span><span id="line-64"><span class="linenos"> 64</span><span class="w">    </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">model_1</span>
+</span><span id="line-65"><span class="linenos"> 65</span><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
+</span><span id="line-66"><span class="linenos"> 66</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">12000</span>
+</span><span id="line-67"><span class="linenos"> 67</span>
+</span><span id="line-68"><span class="linenos"> 68</span><span class="w">  </span><span class="c1"># Prompt listener for function calling (for prompt_targets)</span>
+</span><span id="line-69"><span class="linenos"> 69</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">prompt</span>
+</span><span id="line-70"><span class="linenos"> 70</span><span class="w">    </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">prompt_function_listener</span>
+</span><span id="line-71"><span class="linenos"> 71</span><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
+</span><span id="line-72"><span class="linenos"> 72</span><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">10000</span>
+</span><span id="line-73"><span class="linenos"> 73</span><span class="w">    </span><span class="c1"># This listener is used for prompt_targets and function calling</span>
+</span><span id="line-74"><span class="linenos"> 74</span>
+</span><span id="line-75"><span class="linenos"> 75</span>
+</span><span id="line-76"><span class="linenos"> 76</span><span class="c1"># Reusable service endpoints</span>
+</span><span id="line-77"><span class="linenos"> 77</span><span class="nt">endpoints</span><span class="p">:</span>
+</span><span id="line-78"><span class="linenos"> 78</span><span class="w">  </span><span class="nt">app_server</span><span class="p">:</span>
+</span><span id="line-79"><span class="linenos"> 79</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1:80</span>
+</span><span id="line-80"><span class="linenos"> 80</span><span class="w">    </span><span class="nt">connect_timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.005s</span>
 </span><span id="line-81"><span class="linenos"> 81</span>
-</span><span id="line-82"><span class="linenos"> 82</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">reboot_network_device</span>
-</span><span id="line-83"><span class="linenos"> 83</span><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Reboot a specific network device</span>
-</span><span id="line-84"><span class="linenos"> 84</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
-</span><span id="line-85"><span class="linenos"> 85</span><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</span><span id="line-86"><span class="linenos"> 86</span><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/action</span>
-</span><span id="line-87"><span class="linenos"> 87</span><span class="w">    </span><span class="nt">parameters</span><span class="p">:</span>
-</span><span id="line-88"><span class="linenos"> 88</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_id</span>
-</span><span id="line-89"><span class="linenos"> 89</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
-</span><span id="line-90"><span class="linenos"> 90</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Identifier of the network device to reboot.</span>
-</span><span id="line-91"><span class="linenos"> 91</span><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-92"><span class="linenos"> 92</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">confirmation</span>
-</span><span id="line-93"><span class="linenos"> 93</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">bool</span>
-</span><span id="line-94"><span class="linenos"> 94</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Confirmation flag to proceed with reboot.</span>
-</span><span id="line-95"><span class="linenos"> 95</span><span class="w">        </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">false</span>
-</span><span id="line-96"><span class="linenos"> 96</span><span class="w">        </span><span class="nt">enum</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">[</span><span class="nv">true</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">false</span><span class="p p-Indicator">]</span>
-</span><span id="line-97"><span class="linenos"> 97</span>
-</span><span id="line-98"><span class="linenos"> 98</span><span class="nt">tracing</span><span class="p">:</span>
-</span><span id="line-99"><span class="linenos"> 99</span><span class="w">  </span><span class="c1"># sampling rate. Note by default Arch works on OpenTelemetry compatible tracing.</span>
-</span><span id="line-100"><span class="linenos">100</span><span class="w">  </span><span class="nt">sampling_rate</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.1</span>
+</span><span id="line-82"><span class="linenos"> 82</span><span class="w">  </span><span class="nt">mistral_local</span><span class="p">:</span>
+</span><span id="line-83"><span class="linenos"> 83</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1:8001</span>
+</span><span id="line-84"><span class="linenos"> 84</span>
+</span><span id="line-85"><span class="linenos"> 85</span>
+</span><span id="line-86"><span class="linenos"> 86</span><span class="c1"># Prompt targets for function calling and API orchestration</span>
+</span><span id="line-87"><span class="linenos"> 87</span><span class="nt">prompt_targets</span><span class="p">:</span>
+</span><span id="line-88"><span class="linenos"> 88</span><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">get_current_weather</span>
+</span><span id="line-89"><span class="linenos"> 89</span><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Get current weather at a location.</span>
+</span><span id="line-90"><span class="linenos"> 90</span><span class="w">    </span><span class="nt">parameters</span><span class="p">:</span>
+</span><span id="line-91"><span class="linenos"> 91</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">location</span>
+</span><span id="line-92"><span class="linenos"> 92</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">The location to get the weather for</span>
+</span><span id="line-93"><span class="linenos"> 93</span><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-94"><span class="linenos"> 94</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">string</span>
+</span><span id="line-95"><span class="linenos"> 95</span><span class="w">        </span><span class="nt">format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">City, State</span>
+</span><span id="line-96"><span class="linenos"> 96</span><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">days</span>
+</span><span id="line-97"><span class="linenos"> 97</span><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">the number of days for the request</span>
+</span><span id="line-98"><span class="linenos"> 98</span><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-99"><span class="linenos"> 99</span><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">int</span>
+</span><span id="line-100"><span class="linenos">100</span><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
+</span><span id="line-101"><span class="linenos">101</span><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
+</span><span id="line-102"><span class="linenos">102</span><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/weather</span>
+</span><span id="line-103"><span class="linenos">103</span><span class="w">      </span><span class="nt">http_method</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">POST</span>
+</span><span id="line-104"><span class="linenos">104</span>
+</span><span id="line-105"><span class="linenos">105</span>
+</span><span id="line-106"><span class="linenos">106</span><span class="c1"># OpenTelemetry tracing configuration</span>
+</span><span id="line-107"><span class="linenos">107</span><span class="nt">tracing</span><span class="p">:</span>
+</span><span id="line-108"><span class="linenos">108</span><span class="w">  </span><span class="c1"># Random sampling percentage (1-100)</span>
+</span><span id="line-109"><span class="linenos">109</span><span class="w">  </span><span class="nt">random_sampling</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">100</span>
 </span></code></pre></div>
 </div>
 </div>
@@ -280,18 +285,26 @@ where prompts get routed to, apply guardrails, and enable critical agent observa
         Deployment
       </a>
 </div>
+<div class="ml-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="llms_txt.html">
+        llms.txt
+        <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="9 18 15 12 9 6"></polyline>
+</svg>
+</a>
+</div>
 </div></div>
 </main>
 </div>
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../_static/doctools.js?v=9bcbadda"></script>
 <script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
diff --git a/resources/deployment.html b/resources/deployment.html
index 4ef9326d..9a0bc2b5 100755
--- a/resources/deployment.html
+++ b/resources/deployment.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Deployment | Arch Docs v0.3.22</title>
-<meta content="Deployment | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Deployment | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Deployment | Plano Docs v0.4</title>
+<meta content="Deployment | Plano Docs v0.4" property="og:title"/>
+<meta content="Deployment | Plano Docs v0.4" name="twitter:title"/>
 <link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
 <link href="./docs/resources/deployment.html" rel="canonical"/>
 <link href="../_static/favicon.ico" rel="icon"/>
 <link href="../search.html" rel="search" title="Search"/>
 <link href="configuration_reference.html" rel="next" title="Configuration Reference"/>
-<link href="../build_with_arch/multi_turn.html" rel="prev" title="Multi-Turn"/>
+<link href="tech_overview/threading_model.html" rel="prev" title="Threading Model"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,40 +75,33 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -118,27 +111,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul class="current">
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1 current"><a class="current reference internal" href="#">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -163,28 +158,28 @@
 <div id="content" role="main">
 <section id="deployment">
 <span id="id1"></span><h1>Deployment<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#deployment"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>This guide shows how to deploy Arch directly using Docker without the archgw CLI, including basic runtime checks for routing and health monitoring.</p>
+<p>This guide shows how to deploy Plano directly using Docker without the <code class="docutils literal notranslate"><span class="pre">plano</span></code> CLI, including basic runtime checks for routing and health monitoring.</p>
 <section id="docker-deployment">
 <h2>Docker Deployment<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#docker-deployment" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#docker-deployment'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Below is a minimal, production-ready example showing how to deploy the Arch Docker image directly and run basic runtime checks. Adjust image names, tags, and the <code class="docutils literal notranslate"><span class="pre">arch_config.yaml</span></code> path to match your environment.</p>
+<p>Below is a minimal, production-ready example showing how to deploy the Plano  Docker image directly and run basic runtime checks. Adjust image names, tags, and the <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code> path to match your environment.</p>
 <div class="admonition note">
 <p class="admonition-title">Note</p>
-<p>You will need to pass all required environment variables that are referenced in your <code class="docutils literal notranslate"><span class="pre">arch_config.yaml</span></code> file.</p>
+<p>You will need to pass all required environment variables that are referenced in your <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code> file.</p>
 </div>
-<p>For <code class="docutils literal notranslate"><span class="pre">arch_config.yaml</span></code>, you can use any sample configuration defined earlier in the documentation. For example, you can try the <a class="reference internal" href="../guides/llm_router.html#llm-router"><span class="std std-ref">LLM Routing</span></a> sample config.</p>
+<p>For <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code>, you can use any sample configuration defined earlier in the documentation. For example, you can try the <a class="reference internal" href="../guides/llm_router.html#llm-router"><span class="std std-ref">LLM Routing</span></a> sample config.</p>
 <section id="docker-compose-setup">
 <h3>Docker Compose Setup<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#docker-compose-setup" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#docker-compose-setup'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <p>Create a <code class="docutils literal notranslate"><span class="pre">docker-compose.yml</span></code> file with the following configuration:</p>
 <div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># docker-compose.yml</span>
 </span><span id="line-2"><span class="nt">services</span><span class="p">:</span>
-</span><span id="line-3"><span class="w">  </span><span class="nt">archgw</span><span class="p">:</span>
-</span><span id="line-4"><span class="w">    </span><span class="nt">image</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">katanemo/archgw:0.3.22</span>
-</span><span id="line-5"><span class="w">    </span><span class="nt">container_name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">archgw</span>
+</span><span id="line-3"><span class="w">  </span><span class="nt">plano</span><span class="p">:</span>
+</span><span id="line-4"><span class="w">    </span><span class="nt">image</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">katanemo/plano:0.4.0</span>
+</span><span id="line-5"><span class="w">    </span><span class="nt">container_name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">plano</span>
 </span><span id="line-6"><span class="w">    </span><span class="nt">ports</span><span class="p">:</span>
-</span><span id="line-7"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="s">"10000:10000"</span><span class="w"> </span><span class="c1"># ingress (client -&gt; arch)</span>
-</span><span id="line-8"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="s">"12000:12000"</span><span class="w"> </span><span class="c1"># egress (arch -&gt; upstream/llm proxy)</span>
+</span><span id="line-7"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="s">"10000:10000"</span><span class="w"> </span><span class="c1"># ingress (client -&gt; plano)</span>
+</span><span id="line-8"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="s">"12000:12000"</span><span class="w"> </span><span class="c1"># egress (plano -&gt; upstream/llm proxy)</span>
 </span><span id="line-9"><span class="w">    </span><span class="nt">volumes</span><span class="p">:</span>
-</span><span id="line-10"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">./arch_config.yaml:/app/arch_config.yaml:ro</span>
+</span><span id="line-10"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">./plano_config.yaml:/app/plano_config.yaml:ro</span>
 </span><span id="line-11"><span class="w">    </span><span class="nt">environment</span><span class="p">:</span>
 </span><span id="line-12"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">OPENAI_API_KEY=${OPENAI_API_KEY:?error}</span>
 </span><span id="line-13"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">ANTHROPIC_API_KEY=${ANTHROPIC_API_KEY:?error}</span>
@@ -193,14 +188,14 @@
 </section>
 <section id="starting-the-stack">
 <h3>Starting the Stack<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#starting-the-stack" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#starting-the-stack'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
-<p>Start the services from the directory containing <code class="docutils literal notranslate"><span class="pre">docker-compose.yml</span></code> and <code class="docutils literal notranslate"><span class="pre">arch_config.yaml</span></code>:</p>
+<p>Start the services from the directory containing <code class="docutils literal notranslate"><span class="pre">docker-compose.yml</span></code> and <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code>:</p>
 <div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Set required environment variables and start services</span>
 </span><span id="line-2"><span class="nv">OPENAI_API_KEY</span><span class="o">=</span>xxx<span class="w"> </span><span class="nv">ANTHROPIC_API_KEY</span><span class="o">=</span>yyy<span class="w"> </span>docker<span class="w"> </span>compose<span class="w"> </span>up<span class="w"> </span>-d
 </span></code></pre></div>
 </div>
 <p>Check container health and logs:</p>
 <div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">docker<span class="w"> </span>compose<span class="w"> </span>ps
-</span><span id="line-2">docker<span class="w"> </span>compose<span class="w"> </span>logs<span class="w"> </span>-f<span class="w"> </span>archgw
+</span><span id="line-2">docker<span class="w"> </span>compose<span class="w"> </span>logs<span class="w"> </span>-f<span class="w"> </span>plano
 </span></code></pre></div>
 </div>
 </section>
@@ -211,14 +206,14 @@
 <section id="gateway-smoke-test">
 <h3>Gateway Smoke Test<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#gateway-smoke-test" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#gateway-smoke-test'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <p>Test the chat completion endpoint with automatic routing:</p>
-<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Request handled by the gateway. 'model: "none"' lets Arch decide routing</span>
+<div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="c1"># Request handled by the gateway. 'model: "none"' lets Plano  decide routing</span>
 </span><span id="line-2">curl<span class="w"> </span>--header<span class="w"> </span><span class="s1">'Content-Type: application/json'</span><span class="w"> </span><span class="se">\</span>
 </span><span id="line-3"><span class="w">  </span>--data<span class="w"> </span><span class="s1">'{"messages":[{"role":"user","content":"tell me a joke"}], "model":"none"}'</span><span class="w"> </span><span class="se">\</span>
 </span><span id="line-4"><span class="w">  </span>http://localhost:12000/v1/chat/completions<span class="w"> </span><span class="p">|</span><span class="w"> </span>jq<span class="w"> </span>.model
 </span></code></pre></div>
 </div>
 <p>Expected output:</p>
-<div class="highlight-json notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="s2">"gpt-4o-2024-08-06"</span>
+<div class="highlight-json notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="s2">"gpt-5.2"</span>
 </span></code></pre></div>
 </div>
 </section>
@@ -226,12 +221,12 @@
 <h3>Model-Based Routing<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#model-based-routing" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#model-based-routing'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <p>Test explicit provider and model routing:</p>
 <div class="highlight-bash notranslate"><div class="highlight"><pre><span></span><code><span id="line-1">curl<span class="w"> </span>-s<span class="w"> </span>-H<span class="w"> </span><span class="s2">"Content-Type: application/json"</span><span class="w"> </span><span class="se">\</span>
-</span><span id="line-2"><span class="w">  </span>-d<span class="w"> </span><span class="s1">'{"messages":[{"role":"user","content":"Explain quantum computing"}], "model":"anthropic/claude-3-5-sonnet-20241022"}'</span><span class="w"> </span><span class="se">\</span>
+</span><span id="line-2"><span class="w">  </span>-d<span class="w"> </span><span class="s1">'{"messages":[{"role":"user","content":"Explain quantum computing"}], "model":"anthropic/claude-sonnet-4-5"}'</span><span class="w"> </span><span class="se">\</span>
 </span><span id="line-3"><span class="w">  </span>http://localhost:12000/v1/chat/completions<span class="w"> </span><span class="p">|</span><span class="w"> </span>jq<span class="w"> </span>.model
 </span></code></pre></div>
 </div>
 <p>Expected output:</p>
-<div class="highlight-json notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="s2">"claude-3-5-sonnet-20241022"</span>
+<div class="highlight-json notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="s2">"claude-sonnet-4-5"</span>
 </span></code></pre></div>
 </div>
 </section>
@@ -241,18 +236,18 @@
 <section id="common-issues-and-solutions">
 <h3>Common Issues and Solutions<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#common-issues-and-solutions" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#common-issues-and-solutions'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
 <dl class="simple">
-<dt><strong>Environment Variables</strong></dt><dd><p>Ensure all environment variables (<code class="docutils literal notranslate"><span class="pre">OPENAI_API_KEY</span></code>, <code class="docutils literal notranslate"><span class="pre">ANTHROPIC_API_KEY</span></code>, etc.) used by <code class="docutils literal notranslate"><span class="pre">arch_config.yaml</span></code> are set before starting services.</p>
+<dt><strong>Environment Variables</strong></dt><dd><p>Ensure all environment variables (<code class="docutils literal notranslate"><span class="pre">OPENAI_API_KEY</span></code>, <code class="docutils literal notranslate"><span class="pre">ANTHROPIC_API_KEY</span></code>, etc.) used by <code class="docutils literal notranslate"><span class="pre">plano_config.yaml</span></code> are set before starting services.</p>
 </dd>
 <dt><strong>TLS/Connection Errors</strong></dt><dd><p>If you encounter TLS or connection errors to upstream providers:</p>
 <ul class="simple">
 <li><p>Check DNS resolution</p></li>
 <li><p>Verify proxy settings</p></li>
-<li><p>Confirm correct protocol and port in your <code class="docutils literal notranslate"><span class="pre">arch_config</span></code> endpoints</p></li>
+<li><p>Confirm correct protocol and port in your <code class="docutils literal notranslate"><span class="pre">plano_config</span></code> endpoints</p></li>
 </ul>
 </dd>
 <dt><strong>Verbose Logging</strong></dt><dd><p>To enable more detailed logs for debugging:</p>
 <ul class="simple">
-<li><p>Run archgw with a higher component log level</p></li>
+<li><p>Run plano with a higher component log level</p></li>
 <li><p>See the <a class="reference internal" href="../guides/observability/observability.html#observability"><span class="std std-ref">Observability</span></a> guide for logging and monitoring details</p></li>
 <li><p>Rebuild the image if required with updated log configuration</p></li>
 </ul>
@@ -265,11 +260,11 @@
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../build_with_arch/multi_turn.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="tech_overview/threading_model.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Multi-Turn
+        Threading Model
       </a>
 </div>
 <div class="ml-auto">
@@ -305,12 +300,12 @@
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../_static/doctools.js?v=9bcbadda"></script>
 <script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
diff --git a/resources/llms_txt.html b/resources/llms_txt.html
new file mode 100755
index 00000000..0782715d
--- /dev/null
+++ b/resources/llms_txt.html
@@ -0,0 +1,189 @@
+<!DOCTYPE html>
+
+<html :class="{'dark': darkMode === 'dark' || (darkMode === 'system' &amp;&amp; window.matchMedia('(prefers-color-scheme: dark)').matches)}" class="scroll-smooth" data-content_root="../" lang="en" x-data="{ darkMode: localStorage.getItem('darkMode') || localStorage.setItem('darkMode', 'system'), activeSection: '' }" x-init="$watch('darkMode', val =&gt; localStorage.setItem('darkMode', val))">
+<head>
+<meta content="width=device-width, initial-scale=1.0" name="viewport"/>
+<meta charset="utf-8"/>
+<meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
+<meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
+<meta content="width=device-width, initial-scale=1" name="viewport"/>
+<title>llms.txt | Plano Docs v0.4</title>
+<meta content="llms.txt | Plano Docs v0.4" property="og:title"/>
+<meta content="llms.txt | Plano Docs v0.4" name="twitter:title"/>
+<link href="../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
+<link href="../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
+<link href="../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
+<link href="../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
+<link href="./docs/resources/llms_txt.html" rel="canonical"/>
+<link href="../_static/favicon.ico" rel="icon"/>
+<link href="../search.html" rel="search" title="Search"/>
+<link href="configuration_reference.html" rel="prev" title="Configuration Reference"/>
+<script>
+    <!-- Prevent Flash of wrong theme -->
+      const userPreference = localStorage.getItem('darkMode');
+      let mode;
+      if (userPreference === 'dark' || window.matchMedia('(prefers-color-scheme: dark)').matches) {
+        mode = 'dark';
+        document.documentElement.classList.add('dark');
+      } else {
+        mode = 'light';
+      }
+      if (!userPreference) {localStorage.setItem('darkMode', mode)}
+    </script>
+</head>
+<body :class="{ 'overflow-hidden': showSidebar }" class="min-h-screen font-sans antialiased bg-background text-foreground" x-data="{ showSidebar: false, showScrollTop: false }">
+<div @click.self="showSidebar = false" class="fixed inset-0 z-50 overflow-hidden bg-background/80 backdrop-blur-sm md:hidden" x-cloak="" x-show="showSidebar"></div><div class="relative flex flex-col min-h-screen" id="page"><a class="absolute top-0 left-0 z-[100] block bg-background p-4 text-xl transition -translate-x-full opacity-0 focus:translate-x-0 focus:opacity-100" href="#content">
+      Skip to content
+    </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
+<div class="hidden mr-4 md:flex">
+<a class="flex items-center mr-6" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
+<svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
+</svg>
+<span class="sr-only">Toggle navigation menu</span>
+</button>
+<div class="flex items-center justify-between flex-1 space-x-2 sm:space-x-4 md:justify-end">
+<div class="flex-1 w-full md:w-auto md:flex-none"><form @keydown.k.window.meta="$refs.search.focus()" action="../search.html" class="relative flex items-center group" id="searchbox" method="get">
+<input aria-label="Search the docs" class="inline-flex items-center font-medium transition-colors bg-transparent focus-visible:outline-none focus-visible:ring-2 focus-visible:ring-ring focus-visible:ring-offset-2 ring-offset-background border border-input hover:bg-accent focus:bg-accent hover:text-accent-foreground focus:text-accent-foreground hover:placeholder-accent-foreground py-2 px-4 relative h-9 w-full justify-start rounded-[0.5rem] text-sm text-muted-foreground sm:pr-12 md:w-40 lg:w-64" id="search-input" name="q" placeholder="Search ..." type="search" x-ref="search"/>
+<kbd class="pointer-events-none absolute right-1.5 top-2 hidden h-5 select-none text-muted-foreground items-center gap-1 rounded border border-border bg-muted px-1.5 font-mono text-[10px] font-medium opacity-100 sm:flex group-hover:bg-accent group-hover:text-accent-foreground">
+<span class="text-xs">⌘</span>
+    K
+  </kbd>
+</form>
+</div>
+<nav class="flex items-center space-x-1">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
+<div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
+<svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
+</div>
+</a>
+<button @click="darkMode = darkMode === 'light' ? 'dark' : 'light'" aria-label="Color theme switcher" class="relative inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md hover:bg-accent hover:text-accent-foreground h-9 w-9" type="button">
+<svg class="absolute transition-all scale-100 rotate-0 dark:-rotate-90 dark:scale-0" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 685q45.456 0 77.228-31.772Q589 621.456 589 576q0-45.456-31.772-77.228Q525.456 467 480 467q-45.456 0-77.228 31.772Q371 530.544 371 576q0 45.456 31.772 77.228Q434.544 685 480 685Zm0 91q-83 0-141.5-58.5T280 576q0-83 58.5-141.5T480 376q83 0 141.5 58.5T680 576q0 83-58.5 141.5T480 776ZM80 621.5q-19.152 0-32.326-13.174T34.5 576q0-19.152 13.174-32.326T80 530.5h80q19.152 0 32.326 13.174T205.5 576q0 19.152-13.174 32.326T160 621.5H80Zm720 0q-19.152 0-32.326-13.174T754.5 576q0-19.152 13.174-32.326T800 530.5h80q19.152 0 32.326 13.174T925.5 576q0 19.152-13.174 32.326T880 621.5h-80Zm-320-320q-19.152 0-32.326-13.174T434.5 256v-80q0-19.152 13.174-32.326T480 130.5q19.152 0 32.326 13.174T525.5 176v80q0 19.152-13.174 32.326T480 301.5Zm0 720q-19.152 0-32.326-13.17Q434.5 995.152 434.5 976v-80q0-19.152 13.174-32.326T480 850.5q19.152 0 32.326 13.174T525.5 896v80q0 19.152-13.174 32.33-13.174 13.17-32.326 13.17ZM222.174 382.065l-43-42Q165.5 327.391 166 308.239t13.174-33.065q13.435-13.674 32.587-13.674t32.065 13.674l42.239 43q12.674 13.435 12.555 31.706-.12 18.272-12.555 31.946-12.674 13.674-31.445 13.413-18.772-.261-32.446-13.174Zm494 494.761-42.239-43q-12.674-13.435-12.674-32.087t12.674-31.565Q686.609 756.5 705.38 757q18.772.5 32.446 13.174l43 41.761Q794.5 824.609 794 843.761t-13.174 33.065Q767.391 890.5 748.239 890.5t-32.065-13.674Zm-42-494.761Q660.5 369.391 661 350.62q.5-18.772 13.174-32.446l41.761-43Q728.609 261.5 747.761 262t33.065 13.174q13.674 13.435 13.674 32.587t-13.674 32.065l-43 42.239q-13.435 12.674-31.706 12.555-18.272-.12-31.946-12.555Zm-495 494.761Q165.5 863.391 165.5 844.239t13.674-32.065l43-42.239q13.435-12.674 32.087-12.674t31.565 12.674Q299.5 782.609 299 801.38q-.5 18.772-13.174 32.446l-41.761 43Q231.391 890.5 212.239 890t-33.065-13.174ZM480 576Z"></path>
+</svg>
+<svg class="absolute transition-all scale-0 rotate-90 dark:rotate-0 dark:scale-100" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 936q-151 0-255.5-104.5T120 576q0-138 90-239.5T440 218q25-3 39 18t-1 44q-17 26-25.5 55t-8.5 61q0 90 63 153t153 63q31 0 61.5-9t54.5-25q21-14 43-1.5t19 39.5q-14 138-117.5 229T480 936Zm0-80q88 0 158-48.5T740 681q-20 5-40 8t-40 3q-123 0-209.5-86.5T364 396q0-20 3-40t8-40q-78 32-126.5 102T200 576q0 116 82 198t198 82Zm-10-270Z"></path>
+</svg>
+</button>
+</nav>
+</div>
+</div>
+</header>
+<div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
+<a class="!justify-start text-sm md:!hidden bg-background" href="../index.html">
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
+</a>
+<div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
+<div class="overflow-y-auto h-full w-full relative pr-6">
+
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
+<script>
+  window.dataLayer = window.dataLayer || [];
+  function gtag(){dataLayer.push(arguments);}
+  gtag('js', new Date());
+
+  gtag('config', 'G-EH2VW19FXE');
+</script>
+<nav class="table w-full min-w-full my-6 lg:my-8">
+<p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/overview.html">Overview</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/intro_to_plano.html">Intro to Plano</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html">Quickstart</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../get_started/quickstart.html#next-steps">Next Steps</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../concepts/prompt_target.html">Prompt Target</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Guides</span></p>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../guides/orchestration.html">Orchestration</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/tracing.html">Tracing</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/monitoring.html">Monitoring</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../guides/observability/access_logging.html">Access Logging</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../guides/state.html">Conversational State</a></li>
+</ul>
+<p class="caption" role="heading"><span class="caption-text">Resources</span></p>
+<ul class="current">
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview/tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1 current"><a class="current reference internal" href="#">llms.txt</a></li>
+</ul>
+</nav>
+</div>
+</div>
+<button @click="showSidebar = false" class="absolute md:hidden right-4 top-4 rounded-sm opacity-70 transition-opacity hover:opacity-100" type="button">
+<svg class="h-4 w-4" fill="currentColor" height="24" stroke="none" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
+<path d="M480 632 284 828q-11 11-28 11t-28-11q-11-11-11-28t11-28l196-196-196-196q-11-11-11-28t11-28q11-11 28-11t28 11l196 196 196-196q11-11 28-11t28 11q11 11 11 28t-11 28L536 576l196 196q11 11 11 28t-11 28q-11 11-28 11t-28-11L480 632Z"></path>
+</svg>
+</button>
+</aside>
+<main class="relative py-6 lg:gap-10 lg:py-8 xl:grid xl:grid-cols-[1fr_300px]">
+<div class="w-full min-w-0 mx-auto">
+<nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
+<a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../index.html">
+<span class="hidden md:inline">Plano Docs v0.4</span>
+<svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
+<path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
+</svg>
+</a>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">llms.txt</span>
+</nav>
+<div id="content" role="main">
+<section id="llms-txt">
+<h1>llms.txt<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#llms-txt"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p>This project generates a single plaintext file containing the compiled text of all documentation pages, useful for large context models to reference Plano documentation.</p>
+<p>Open it here: <a class="reference external" href="../includes/llms.txt" rel="nofollow noopener">llms.txt<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a></p>
+</section>
+</div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
+<div class="mr-auto">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="configuration_reference.html">
+<svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
+<polyline points="15 18 9 12 15 6"></polyline>
+</svg>
+        Configuration Reference
+      </a>
+</div>
+</div></div>
+</main>
+</div>
+</div><footer class="py-6 border-t border-border md:py-0">
+<div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
+<div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
+</div>
+</div>
+</footer>
+</div>
+<script src="../_static/documentation_options.js?v=8cf1ab6b"></script>
+<script src="../_static/doctools.js?v=9bcbadda"></script>
+<script src="../_static/sphinx_highlight.js?v=dc90522c"></script>
+<script defer="defer" src="../_static/theme.js?v=073f68d9"></script>
+<script src="../_static/design-tabs.js?v=f930bc37"></script>
+</body>
+</html>
\ No newline at end of file
diff --git a/concepts/tech_overview/model_serving.html b/resources/tech_overview/model_serving.html
similarity index 68%
rename from concepts/tech_overview/model_serving.html
rename to resources/tech_overview/model_serving.html
index 79cabf16..7e1e10ee 100755
--- a/concepts/tech_overview/model_serving.html
+++ b/resources/tech_overview/model_serving.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Model Serving | Arch Docs v0.3.22</title>
-<meta content="Model Serving | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Model Serving | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Bright Staff | Plano Docs v0.4</title>
+<meta content="Bright Staff | Plano Docs v0.4" property="og:title"/>
+<meta content="Bright Staff | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/concepts/tech_overview/model_serving.html" rel="canonical"/>
+<link href="./docs/resources/tech_overview/model_serving.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
 <link href="../../search.html" rel="search" title="Search"/>
-<link href="request_lifecycle.html" rel="next" title="Request Lifecycle"/>
-<link href="prompt.html" rel="prev" title="Prompts"/>
+<link href="threading_model.html" rel="next" title="Threading Model"/>
+<link href="request_lifecycle.html" rel="prev" title="Request Lifecycle"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,70 +75,65 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul class="current">
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html">Prompts</a></li>
-<li class="toctree-l2 current"><a class="current reference internal" href="#">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html">Error Target</a></li>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
 </ul>
 </li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../prompt_target.html">Prompt Target</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/prompt_target.html">Prompt Target</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<ul class="current">
+<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2 current"><a class="current reference internal" href="#">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,89 +148,59 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
 </a>
 <div class="mr-1">/</div><a class="hover:text-foreground overflow-hidden text-ellipsis whitespace-nowrap" href="tech_overview.html">Tech Overview</a>
-<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Model Serving</span>
+<div class="mr-1">/</div><span aria-current="page" class="font-medium text-foreground overflow-hidden text-ellipsis whitespace-nowrap">Bright Staff</span>
 </nav>
 <div id="content" role="main">
-<section id="model-serving">
-<span id="id1"></span><h1>Model Serving<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#model-serving"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Arch is a set of <cite>two</cite> self-contained processes that are designed to run alongside your application
-servers (or on a separate host connected via a network). The first process is designated to manage low-level
-networking and HTTP related concerns, and the other process is for model serving, which helps Arch make
-intelligent decisions about the incoming prompts. The model server is designed to call the purpose-built
-LLMs in Arch.</p>
-<a class="reference internal image-reference" href="../../_images/arch-system-architecture.jpg"><img alt="../../_images/arch-system-architecture.jpg" class="align-center" src="../../_images/arch-system-architecture.jpg" style="width: 40%;"/>
+<section id="bright-staff">
+<span id="id1"></span><h1>Bright Staff<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#bright-staff"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
+<p>Bright Staff is Plano’s memory-efficient, lightweight controller for agentic traffic. It sits inside the Plano
+data plane and makes real-time decisions about how prompts are handled, forwarded, and processed.</p>
+<p>Rather than running a separate “model server” subsystem, Plano relies on Envoy’s HTTP connection management
+and cluster subsystem to talk to different models and backends over HTTP(S). Bright Staff uses these primitives to:
+* Inspect prompts, conversation state, and metadata.
+* Decide which upstream model(s), tool backends, or APIs to call, and in what order.
+* Coordinate retries, fallbacks, and traffic splitting across providers and models.</p>
+<p>Plano is designed to run alongside your application servers in your cloud VPC, on-premises, or in local
+development. It does not require a GPU itself; GPUs live where your models are hosted (third-party APIs or your
+own deployments), and Plano reaches them via HTTP.</p>
+<a class="reference internal image-reference" href="../../_images/plano-system-architecture.png"><img alt="../../_images/plano-system-architecture.png" class="align-center" src="../../_images/plano-system-architecture.png" style="width: 40%;"/>
 </a>
-<p>Arch’ is designed to be deployed in your cloud VPC, on a on-premises host, and can work on devices that don’t
-have a GPU. Note, GPU devices are need for fast and cost-efficient use, so that Arch (model server, specifically)
-can process prompts quickly and forward control back to the application host. There are three modes in which Arch
-can be configured to run its <strong>model server</strong> subsystem:</p>
-<section id="local-serving-cpu-moderate">
-<h2>Local Serving (CPU - Moderate)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#local-serving-cpu-moderate" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#local-serving-cpu-moderate'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>The following bash commands enable you to configure the model server subsystem in Arch to run local on device
-and only use CPU devices. This will be the slowest option but can be useful in dev/test scenarios where GPUs
-might not be available.</p>
-<div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="gp">$ </span>archgw<span class="w"> </span>up<span class="w"> </span>--local-cpu
-</span></code></pre></div>
-</div>
-</section>
-<section id="cloud-serving-gpu-blazing-fast">
-<h2>Cloud Serving (GPU - Blazing Fast)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#cloud-serving-gpu-blazing-fast" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#cloud-serving-gpu-blazing-fast'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>The command below instructs Arch to intelligently use GPUs locally for fast intent detection, but default to
-cloud serving for function calling and guardrails scenarios to dramatically improve the speed and overall performance
-of your applications.</p>
-<div class="highlight-console notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="gp">$ </span>archgw<span class="w"> </span>up
-</span></code></pre></div>
-</div>
-<div class="admonition note">
-<p class="admonition-title">Note</p>
-<p>Arch’s model serving in the cloud is priced at $0.05M/token (156x cheaper than GPT-4o) with average latency
-of 200ms (10x faster than GPT-4o). Please refer to our <a class="reference internal" href="../../get_started/quickstart.html#quickstart"><span class="std std-ref">Get Started</span></a> to know
-how to generate API keys for model serving</p>
-</div>
-</section>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="prompt.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="request_lifecycle.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Prompts
+        Request Lifecycle
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="request_lifecycle.html">
-        Request Lifecycle
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="threading_model.html">
+        Threading Model
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
 </a>
 </div>
-</div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
-<div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
-<ul>
-<li><a :data-current="activeSection === '#local-serving-cpu-moderate'" class="reference internal" href="#local-serving-cpu-moderate">Local Serving (CPU - Moderate)</a></li>
-<li><a :data-current="activeSection === '#cloud-serving-gpu-blazing-fast'" class="reference internal" href="#cloud-serving-gpu-blazing-fast">Cloud Serving (GPU - Blazing Fast)</a></li>
-</ul>
-</div>
-</aside>
+</div></div>
 </main>
 </div>
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/concepts/tech_overview/request_lifecycle.html b/resources/tech_overview/request_lifecycle.html
similarity index 68%
rename from concepts/tech_overview/request_lifecycle.html
rename to resources/tech_overview/request_lifecycle.html
index e27114b9..da8488b1 100755
--- a/concepts/tech_overview/request_lifecycle.html
+++ b/resources/tech_overview/request_lifecycle.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Request Lifecycle | Arch Docs v0.3.22</title>
-<meta content="Request Lifecycle | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Request Lifecycle | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Request Lifecycle | Plano Docs v0.4</title>
+<meta content="Request Lifecycle | Plano Docs v0.4" property="og:title"/>
+<meta content="Request Lifecycle | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/concepts/tech_overview/request_lifecycle.html" rel="canonical"/>
+<link href="./docs/resources/tech_overview/request_lifecycle.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
 <link href="../../search.html" rel="search" title="Search"/>
-<link href="error_target.html" rel="next" title="Error Target"/>
-<link href="model_serving.html" rel="prev" title="Model Serving"/>
+<link href="model_serving.html" rel="next" title="Bright Staff"/>
+<link href="tech_overview.html" rel="prev" title="Tech Overview"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,70 +75,65 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul class="current">
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Model Serving</a></li>
-<li class="toctree-l2 current"><a class="current reference internal" href="#">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html">Error Target</a></li>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
 </ul>
 </li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../prompt_target.html">Prompt Target</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/prompt_target.html">Prompt Target</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<ul class="current">
+<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
+<li class="toctree-l2 current"><a class="current reference internal" href="#">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -164,155 +159,80 @@
 <div id="content" role="main">
 <section id="request-lifecycle">
 <span id="lifecycle-of-a-request"></span><h1>Request Lifecycle<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#request-lifecycle"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Below we describe the events in the lifecycle of a request passing through an Arch gateway instance. We first
-describe how Arch fits into the request path and then the internal events that take place following
-the arrival of a request at Arch from downstream clients. We follow the request until the corresponding
+<p>Below we describe the events in the lifecycle of a request passing through a Plano instance. We first
+describe how Plano fits into the request path and then the internal events that take place following
+the arrival of a request at Plano from downstream clients. We follow the request until the corresponding
 dispatch upstream and the response path.</p>
-<a class="reference internal image-reference" href="../../_images/network-topology-ingress-egress.jpg"><img alt="../../_images/network-topology-ingress-egress.jpg" class="align-center" src="../../_images/network-topology-ingress-egress.jpg" style="width: 100%;"/>
+<a class="reference internal image-reference" href="../../_images/network-topology-ingress-egress.png"><img alt="../../_images/network-topology-ingress-egress.png" class="align-center" src="../../_images/network-topology-ingress-egress.png" style="width: 100%;"/>
 </a>
-<section id="terminology">
-<h2>Terminology<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#terminology" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#terminology'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>We recommend that you get familiar with some of the <a class="reference internal" href="terminology.html#arch-terminology"><span class="std std-ref">terminology</span></a> used in Arch
-before reading this section.</p>
-</section>
 <section id="network-topology">
 <h2>Network topology<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#network-topology" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#network-topology'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>How a request flows through the components in a network (including Arch) depends on the network’s topology.
-Arch can be used in a wide variety of networking topologies. We focus on the inner operation of Arch below,
-but briefly we address how Arch relates to the rest of the network in this section.</p>
+<p>How a request flows through the components in a network (including Plano) depends on the network’s topology.
+Plano can be used in a wide variety of networking topologies. We focus on the inner operations of Plano below,
+but briefly we address how Plano relates to the rest of the network in this section.</p>
 <ul class="simple">
 <li><p><strong>Downstream(Ingress)</strong> listeners take requests from upstream clients like a web UI or clients that forward
-prompts to you local application responses from the application flow back through Arch to the downstream.</p></li>
+prompts to you local application responses from the application flow back through Plano to the downstream.</p></li>
 <li><p><strong>Upstream(Egress)</strong> listeners take requests from the application and forward them to LLMs.</p></li>
 </ul>
-<a class="reference internal image-reference" href="../../_images/network-topology-ingress-egress.jpg"><img alt="../../_images/network-topology-ingress-egress.jpg" class="align-center" src="../../_images/network-topology-ingress-egress.jpg" style="width: 100%;"/>
-</a>
-<p>In practice, Arch can be deployed on the edge and as an internal load balancer between AI agents. A request path may
-traverse multiple Arch gateways:</p>
-<a class="reference internal image-reference" href="../../_images/network-topology-agent.jpg"><img alt="../../_images/network-topology-agent.jpg" class="align-center" src="../../_images/network-topology-agent.jpg" style="width: 100%;"/>
-</a>
 </section>
 <section id="high-level-architecture">
 <h2>High level architecture<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#high-level-architecture" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#high-level-architecture'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Arch is a set of <strong>two</strong> self-contained processes that are designed to run alongside your application servers
-(or on a separate server connected to your application servers via a network). The first process is designated
-to manage HTTP-level networking and connection management concerns (protocol management, request id generation,
-header sanitization, etc.), and the other process is for <strong>model serving</strong>, which helps Arch make intelligent
-decisions about the incoming prompts. The model server hosts the purpose-built LLMs to
-manage several critical, but undifferentiated, prompt related tasks on behalf of developers.</p>
-<p>The request processing path in Arch has three main parts:</p>
+<p>Plano is a set of <strong>two</strong> self-contained processes that are designed to run alongside your application servers
+(or on a separate server connected to your application servers via a network).</p>
+<p>The first process is designated to manage HTTP-level networking and connection management concerns (protocol management, request id generation, header sanitization, etc.), and the other process is a <strong>controller</strong>, which helps Plano make intelligent decisions about the incoming prompts. The controller hosts the purpose-built LLMs to manage several critical, but undifferentiated, prompt related tasks on behalf of developers.</p>
+<p>The request processing path in Plano has three main parts:</p>
 <ul class="simple">
-<li><p><a class="reference internal" href="listener.html#arch-overview-listeners"><span class="std std-ref">Listener subsystem</span></a> which handles <strong>downstream</strong> and <strong>upstream</strong> request
-processing. It is responsible for managing the downstream (ingress) and the upstream (egress) request
-lifecycle. The downstream and upstream HTTP/2 codec lives here.</p></li>
-<li><p><a class="reference internal" href="prompt.html#arch-overview-prompt-handling"><span class="std std-ref">Prompt handler subsystem</span></a> which is responsible for selecting and
-forwarding prompts <code class="docutils literal notranslate"><span class="pre">prompt_targets</span></code> and establishes the lifecycle of any <strong>upstream</strong> connection to a
-hosted endpoint that implements domain-specific business logic for incoming prompts. This is where knowledge
-of targets and endpoint health, load balancing and connection pooling exists.</p></li>
-<li><p><a class="reference internal" href="model_serving.html#model-serving"><span class="std std-ref">Model serving subsystem</span></a> which helps Arch make intelligent decisions about the
-incoming prompts. The model server is designed to call the purpose-built LLMs in Arch.</p></li>
+<li><p><a class="reference internal" href="../../concepts/listeners.html#plano-overview-listeners"><span class="std std-ref">Listener subsystem</span></a> which handles <strong>downstream</strong> and <strong>upstream</strong> request
+processing. It is responsible for managing the inbound(edge) and outbound(egress) request lifecycle. The downstream and upstream HTTP/2 codec lives here. This also includes the lifecycle of any <strong>upstream</strong> connection to an LLM provider or tool backend. The listenser subsystmem manages connection pools, load balancing, retries, and failover.</p></li>
+<li><p><a class="reference internal" href="model_serving.html#bright-staff"><span class="std std-ref">Bright Staff controller subsystem</span></a> is Plano’s memory-efficient, lightweight controller for agentic traffic. It sits inside the Plano data plane and makes real-time decisions about how prompts are handled, forwarded, and processed.</p></li>
 </ul>
-<p>The three subsystems are bridged with either the HTTP router filter, and the cluster manager subsystems of Envoy.</p>
-<p>Also, Arch utilizes <a class="reference external" href="https://blog.envoyproxy.io/envoy-threading-model-a8d44b922310" rel="nofollow noopener">Envoy event-based thread model<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>.
-A main thread is responsible for the server lifecycle, configuration processing, stats, etc. and some number of
-<a class="reference internal" href="threading_model.html#arch-overview-threading"><span class="std std-ref">worker threads</span></a> process requests. All threads operate around an event loop (<a class="reference external" href="https://libevent.org/" rel="nofollow noopener">libevent<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>)
-and any given downstream TCP connection will be handled by exactly one worker thread for its lifetime. Each worker
-thread maintains its own pool of TCP connections to upstream endpoints.</p>
+<p>These two subsystems are bridged with either the HTTP router filter, and the cluster manager subsystems of Envoy.</p>
+<p>Also, Plano utilizes <a class="reference external" href="https://blog.envoyproxy.io/envoy-threading-model-a8d44b922310" rel="nofollow noopener">Envoy event-based thread model<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>. A main thread is responsible for the server lifecycle, configuration processing, stats, etc. and some number of <a class="reference internal" href="threading_model.html#arch-overview-threading"><span class="std std-ref">worker threads</span></a> process requests. All threads operate around an event loop (<a class="reference external" href="https://libevent.org/" rel="nofollow noopener">libevent<svg fill="currentColor" height="1em" stroke="none" viewbox="0 96 960 960" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M188 868q-11-11-11-28t11-28l436-436H400q-17 0-28.5-11.5T360 336q0-17 11.5-28.5T400 296h320q17 0 28.5 11.5T760 336v320q0 17-11.5 28.5T720 696q-17 0-28.5-11.5T680 656V432L244 868q-11 11-28 11t-28-11Z"></path></svg></a>) and any given downstream TCP connection will be handled by exactly one worker thread for its lifetime. Each worker thread maintains its own pool of TCP connections to upstream endpoints.</p>
 <p>Worker threads rarely share state and operate in a trivially parallel fashion. This threading model
 enables scaling to very high core count CPUs.</p>
 </section>
-<section id="configuration">
-<h2>Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configuration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>Today, only support a static bootstrap configuration file for simplicity today:</p>
-<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.1.0</span>
-</span><span id="line-2">
-</span><span id="line-3"><span class="nt">listeners</span><span class="p">:</span>
-</span><span id="line-4"><span class="w">  </span><span class="nt">ingress_traffic</span><span class="p">:</span>
-</span><span id="line-5"><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
-</span><span id="line-6"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">10000</span>
-</span><span id="line-7"><span class="w">    </span><span class="nt">message_format</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai</span>
-</span><span id="line-8"><span class="w">    </span><span class="nt">timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">30s</span>
-</span><span id="line-9">
-</span><span id="line-10"><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
-</span><span id="line-11"><span class="nt">llm_providers</span><span class="p">:</span>
-</span><span id="line-12"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
-</span><span id="line-13"><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
-</span><span id="line-14"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-15">
-</span><span id="line-16"><span class="c1"># default system prompt used by all prompt targets</span>
-</span><span id="line-17"><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">You are a network assistant that just offers facts; not advice on manufacturers or purchasing decisions.</span>
-</span><span id="line-18">
-</span><span id="line-19"><span class="nt">prompt_guards</span><span class="p">:</span>
-</span><span id="line-20"><span class="w">  </span><span class="nt">input_guards</span><span class="p">:</span>
-</span><span id="line-21"><span class="w">    </span><span class="nt">jailbreak</span><span class="p">:</span>
-</span><span id="line-22"><span class="w">      </span><span class="nt">on_exception</span><span class="p">:</span>
-</span><span id="line-23"><span class="w">        </span><span class="nt">message</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Looks like you're curious about my abilities, but I can only provide assistance within my programmed parameters.</span>
-</span><span id="line-24">
-</span><span id="line-25"><span class="nt">prompt_targets</span><span class="p">:</span>
-</span><span id="line-26"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">information_extraction</span>
-</span><span id="line-27"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-28"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">handel all scenarios that are question and answer in nature. Like summarization, information extraction, etc.</span>
-</span><span id="line-29"><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
-</span><span id="line-30"><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</span><span id="line-31"><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/summary</span>
-</span><span id="line-32"><span class="w">    </span><span class="c1"># Arch uses the default LLM and treats the response from the endpoint as the prompt to send to the LLM</span>
-</span><span id="line-33"><span class="w">    </span><span class="nt">auto_llm_dispatch_on_response</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-34"><span class="w">    </span><span class="c1"># override system prompt for this prompt target</span>
-</span><span id="line-35"><span class="w">    </span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">You are a helpful information extraction assistant. Use the information that is provided to you.</span>
-</span><span id="line-36">
-</span><span id="line-37"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">reboot_network_device</span>
-</span><span id="line-38"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Reboot a specific network device</span>
-</span><span id="line-39"><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
-</span><span id="line-40"><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
-</span><span id="line-41"><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/action</span>
-</span><span id="line-42"><span class="w">    </span><span class="nt">parameters</span><span class="p">:</span>
-</span><span id="line-43"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_id</span>
-</span><span id="line-44"><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
-</span><span id="line-45"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Identifier of the network device to reboot.</span>
-</span><span id="line-46"><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
-</span><span id="line-47"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">confirmation</span>
-</span><span id="line-48"><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">bool</span>
-</span><span id="line-49"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Confirmation flag to proceed with reboot.</span>
-</span><span id="line-50"><span class="w">        </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">false</span>
-</span><span id="line-51"><span class="w">        </span><span class="nt">enum</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">[</span><span class="nv">true</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">false</span><span class="p p-Indicator">]</span>
-</span><span id="line-52">
-</span><span id="line-53"><span class="c1"># Arch creates a round-robin load balancing between different endpoints, managed via the cluster subsystem.</span>
-</span><span id="line-54"><span class="nt">endpoints</span><span class="p">:</span>
-</span><span id="line-55"><span class="w">  </span><span class="nt">app_server</span><span class="p">:</span>
-</span><span id="line-56"><span class="w">    </span><span class="c1"># value could be ip address or a hostname with port</span>
-</span><span id="line-57"><span class="w">    </span><span class="c1"># this could also be a list of endpoints for load balancing</span>
-</span><span id="line-58"><span class="w">    </span><span class="c1"># for example endpoint: [ ip1:port, ip2:port ]</span>
-</span><span id="line-59"><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1:80</span>
-</span><span id="line-60"><span class="w">    </span><span class="c1"># max time to wait for a connection to be established</span>
-</span><span id="line-61"><span class="w">    </span><span class="nt">connect_timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.005s</span>
-</span></code></pre></div>
-</div>
-</section>
 <section id="request-flow-ingress">
 <h2>Request Flow (Ingress)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#request-flow-ingress" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#request-flow-ingress'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
 <p>A brief outline of the lifecycle of a request and response using the example configuration above:</p>
 <ol class="arabic simple">
 <li><p><strong>TCP Connection Establishment</strong>:
-A TCP connection from downstream is accepted by an Arch listener running on a worker thread.
+A TCP connection from downstream is accepted by an Plano listener running on a worker thread.
 The listener filter chain provides SNI and other pre-TLS information. The transport socket, typically TLS,
 decrypts incoming data for processing.</p></li>
-<li><p><strong>Prompt Guardrails Check</strong>:
-Arch first checks the incoming prompts for guardrails such as jailbreak attempts. This ensures
-that harmful or unwanted behaviors are detected early in the request processing pipeline.</p></li>
-<li><p><strong>Intent Matching</strong>:
-The decrypted data stream is de-framed by the HTTP/2 codec in Arch’s HTTP connection manager. Arch performs
-intent matching via is <strong>prompt-handler</strong> subsystem using the name and description of the defined prompt targets,
-determining which endpoint should handle the prompt.</p></li>
-<li><p><strong>Parameter Gathering with Arch-Function</strong>:
-If a prompt target requires specific parameters, Arch engages Arch-FC to extract the necessary details
-from the incoming prompt(s). This process gathers the critical information needed for downstream API calls.</p></li>
-<li><p><strong>API Call Execution</strong>:
-Arch routes the prompt to the appropriate backend API or function call. If an endpoint cluster is identified,
-load balancing is performed, circuit breakers are checked, and the request is proxied to the upstream endpoint.</p></li>
-<li><p><strong>Default Summarization by Upstream LLM</strong>:
-By default, if no specific endpoint processing is needed, the prompt is sent to an upstream LLM for summarization.
-This ensures that responses are concise and relevant, enhancing user experience in RAG (Retrieval Augmented Generation)
-and agentic applications.</p></li>
+</ol>
+<ol class="arabic simple" start="3">
+<li><p><strong>Routing Decision (Agent vs Prompt Target)</strong>:
+The decrypted data stream is de-framed by the HTTP/2 codec in Plano’s HTTP connection manager. Plano performs
+intent matching (via the Bright Staff controller and prompt-handling logic) using the configured agents and
+<a class="reference internal" href="../../concepts/prompt_target.html#prompt-target"><span class="std std-ref">prompt targets</span></a>, determining whether this request should be handled by an agent workflow
+(with optional <a class="reference internal" href="../../concepts/filter_chain.html#filter-chain"><span class="std std-ref">Filter Chains</span></a>) or by a deterministic prompt target.</p></li>
+</ol>
+<p>4a. <strong>Agent Path: Orchestration and Filter Chains</strong></p>
+<blockquote>
+<div><p>If the request is routed to an <strong>agent</strong>, Plano executes any attached <a class="reference internal" href="../../concepts/filter_chain.html#filter-chain"><span class="std std-ref">Filter Chains</span></a> first. These filters can apply guardrails, rewrite prompts, or enrich context (for example, RAG retrieval) before the agent runs. Once filters complete, the Bright Staff controller orchestrates which downstream tools, APIs, or LLMs the agent should call and in what sequence.</p>
+<ul class="simple">
+<li><p>Plano may call one or more backend APIs or tools on behalf of the agent.</p></li>
+<li><p>If an endpoint cluster is identified, load balancing is performed, circuit breakers are checked, and the request is proxied to the appropriate upstream endpoint.</p></li>
+<li><p>If no specific endpoint is required, the prompt is sent to an upstream LLM using Plano’s model proxy for
+completion or summarization.</p></li>
+</ul>
+<p>For more on agent workflows and orchestration, see <a class="reference internal" href="../../concepts/prompt_target.html#prompt-target"><span class="std std-ref">Prompt Targets and Agents</span></a> and
+<a class="reference internal" href="../../concepts/filter_chain.html#filter-chain"><span class="std std-ref">Agent Filter Chains</span></a>.</p>
+</div></blockquote>
+<p>4b. <strong>Prompt Target Path: Deterministic Tool/API Calls</strong></p>
+<blockquote>
+<div><p>If the request is routed to a <strong>prompt target</strong>, Plano treats it as a deterministic, task-specific call.
+Plano engages its function-calling and parameter-gathering capabilities to extract the necessary details
+from the incoming prompt(s) and produce the structured inputs your backend expects.</p>
+<ul class="simple">
+<li><p><strong>Parameter Gathering</strong>: Plano extracts and validates parameters defined on the prompt target (for example,
+currency symbols, dates, or entity identifiers) so your backend does not need to parse natural language.</p></li>
+<li><p><strong>API Call Execution</strong>: Plano then routes the call to the configured backend endpoint. If an endpoint cluster is identified, load balancing and circuit-breaker checks are applied before proxying the request upstream.</p></li>
+</ul>
+<p>For more on how to design and configure prompt targets, see <a class="reference internal" href="../../concepts/prompt_target.html#prompt-target"><span class="std std-ref">Prompt Target</span></a>.</p>
+</div></blockquote>
+<ol class="arabic simple" start="5">
 <li><p><strong>Error Handling and Forwarding</strong>:
 Errors encountered during processing, such as failed function calls or guardrail detections, are forwarded to
 designated error targets. Error details are communicated through specific headers to the application:</p>
@@ -329,26 +249,21 @@ Responses pass through HTTP filters in reverse order, ensuring any necessary pro
 </section>
 <section id="request-flow-egress">
 <h2>Request Flow (Egress)<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#request-flow-egress" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#request-flow-egress'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
-<p>A brief outline of the lifecycle of a request and response in the context of egress traffic from an application to Large Language Models (LLMs) via Arch:</p>
+<p>A brief outline of the lifecycle of a request and response in the context of egress traffic from an application to Large Language Models (LLMs) via Plano:</p>
 <ol class="arabic simple">
 <li><p><strong>HTTP Connection Establishment to LLM</strong>:
-Arch initiates an HTTP connection to the upstream LLM service. This connection is handled by Arch’s egress listener
-running on a worker thread. The connection typically uses a secure transport protocol such as HTTPS, ensuring the
-prompt data is encrypted before being sent to the LLM service.</p></li>
+Plano initiates an HTTP connection to the upstream LLM service. This connection is handled by Plano’s egress listener running on a worker thread. The connection typically uses a secure transport protocol such as HTTPS, ensuring the prompt data is encrypted before being sent to the LLM service.</p></li>
 <li><p><strong>Rate Limiting</strong>:
-Before sending the request to the LLM, Arch applies rate-limiting policies to ensure that the upstream LLM service
-is not overwhelmed by excessive traffic. Rate limits are enforced per client or service, ensuring fair usage and
-preventing accidental or malicious overload. If the rate limit is exceeded, Arch may return an appropriate HTTP
-error (e.g., 429 Too Many Requests) without sending the prompt to the LLM.</p></li>
+Before sending the request to the LLM, Plano applies rate-limiting policies to ensure that the upstream LLM service is not overwhelmed by excessive traffic. Rate limits are enforced per client or service, ensuring fair usage and preventing accidental or malicious overload. If the rate limit is exceeded, Plano may return an appropriate HTTP error (e.g., 429 Too Many Requests) without sending the prompt to the LLM.</p></li>
+<li><p><strong>Seamless Request Transformation and Smart Routing</strong>:
+After rate limiting, Plano normalizes the outgoing request into a provider-agnostic shape and applies smart routing decisions using the configured <a class="reference internal" href="../../concepts/llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">LLM Providers</span></a>. This includes translating client-specific conventions into a unified OpenAI-style contract, enriching or overriding parameters (for example, temperature or max tokens) based on policy, and choosing the best target model or provider using <a class="reference internal" href="../../concepts/llm_providers/llm_providers.html#llm-providers"><span class="std std-ref">model-based, alias-based, or preference-aligned routing</span></a>.</p></li>
 <li><p><strong>Load Balancing to (hosted) LLM Endpoints</strong>:
-After passing the rate-limiting checks, Arch routes the prompt to the appropriate LLM endpoint.
-If multiple LLM providers instances are available, load balancing is performed to distribute traffic evenly
-across the instances. Arch checks the health of the LLM endpoints using circuit breakers and health checks,
+After smart routing selects the target provider/model, Plano routes the prompt to the appropriate LLM endpoint.
+If multiple LLM provider instances are available, load balancing is performed to distribute traffic evenly
+across the instances. Plano checks the health of the LLM endpoints using circuit breakers and health checks,
 ensuring that the prompt is only routed to a healthy, responsive instance.</p></li>
 <li><p><strong>Response Reception and Forwarding</strong>:
-Once the LLM processes the prompt, Arch receives the response from the LLM service. The response is typically a
-generated text, completion, or summarization. Upon reception, Arch decrypts (if necessary) and handles the response,
-passing it through any egress processing pipeline defined by the application, such as logging or additional response filtering.</p></li>
+Once the LLM processes the prompt, Plano receives the response from the LLM service. The response is typically a generated text, completion, or summarization. Upon reception, Plano decrypts (if necessary) and handles the response, passing it through any egress processing pipeline defined by the application, such as logging or additional response filtering.</p></li>
 </ol>
 <section id="post-request-processing">
 <h3>Post-request processing<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#post-request-processing" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#post-request-processing'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h3>
@@ -364,19 +279,75 @@ processing request headers and then finalized by the HCM during post-request pro
 </ul>
 </section>
 </section>
+<section id="configuration">
+<h2>Configuration<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#configuration" x-intersect.margin.0%.0%.-70%.0%="activeSection = '#configuration'"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h2>
+<p>Today, only support a static bootstrap configuration file for simplicity today:</p>
+<div class="highlight-yaml notranslate"><div class="highlight"><pre><span></span><code><span id="line-1"><span class="nt">version</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">v0.2.0</span>
+</span><span id="line-2">
+</span><span id="line-3"><span class="nt">listeners</span><span class="p">:</span>
+</span><span id="line-4"><span class="w">  </span><span class="nt">ingress_traffic</span><span class="p">:</span>
+</span><span id="line-5"><span class="w">    </span><span class="nt">address</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.0.0.0</span>
+</span><span id="line-6"><span class="w">    </span><span class="nt">port</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">10000</span>
+</span><span id="line-7">
+</span><span id="line-8"><span class="c1"># Centralized way to manage LLMs, manage keys, retry logic, failover and limits in a central way</span>
+</span><span id="line-9"><span class="nt">model_providers</span><span class="p">:</span>
+</span><span id="line-10"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">access_key</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">$OPENAI_API_KEY</span>
+</span><span id="line-11"><span class="w">    </span><span class="nt">model</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">openai/gpt-4o</span>
+</span><span id="line-12"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-13">
+</span><span id="line-14"><span class="nt">prompt_targets</span><span class="p">:</span>
+</span><span id="line-15"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">information_extraction</span>
+</span><span id="line-16"><span class="w">    </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-17"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">handel all scenarios that are question and answer in nature. Like summarization, information extraction, etc.</span>
+</span><span id="line-18"><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
+</span><span id="line-19"><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
+</span><span id="line-20"><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/summary</span>
+</span><span id="line-21"><span class="w">    </span><span class="c1"># Arch uses the default LLM and treats the response from the endpoint as the prompt to send to the LLM</span>
+</span><span id="line-22"><span class="w">    </span><span class="nt">auto_llm_dispatch_on_response</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-23"><span class="w">    </span><span class="c1"># override system prompt for this prompt target</span>
+</span><span id="line-24"><span class="w">    </span><span class="nt">system_prompt</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">You are a helpful information extraction assistant. Use the information that is provided to you.</span>
+</span><span id="line-25">
+</span><span id="line-26"><span class="w">  </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">reboot_network_device</span>
+</span><span id="line-27"><span class="w">    </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Reboot a specific network device</span>
+</span><span id="line-28"><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span>
+</span><span id="line-29"><span class="w">      </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">app_server</span>
+</span><span id="line-30"><span class="w">      </span><span class="nt">path</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">/agent/action</span>
+</span><span id="line-31"><span class="w">    </span><span class="nt">parameters</span><span class="p">:</span>
+</span><span id="line-32"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">device_id</span>
+</span><span id="line-33"><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">str</span>
+</span><span id="line-34"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Identifier of the network device to reboot.</span>
+</span><span id="line-35"><span class="w">        </span><span class="nt">required</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">true</span>
+</span><span id="line-36"><span class="w">      </span><span class="p p-Indicator">-</span><span class="w"> </span><span class="nt">name</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">confirmation</span>
+</span><span id="line-37"><span class="w">        </span><span class="nt">type</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">bool</span>
+</span><span id="line-38"><span class="w">        </span><span class="nt">description</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">Confirmation flag to proceed with reboot.</span>
+</span><span id="line-39"><span class="w">        </span><span class="nt">default</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">false</span>
+</span><span id="line-40"><span class="w">        </span><span class="nt">enum</span><span class="p">:</span><span class="w"> </span><span class="p p-Indicator">[</span><span class="nv">true</span><span class="p p-Indicator">,</span><span class="w"> </span><span class="nv">false</span><span class="p p-Indicator">]</span>
+</span><span id="line-41">
+</span><span id="line-42"><span class="c1"># Arch creates a round-robin load balancing between different endpoints, managed via the cluster subsystem.</span>
+</span><span id="line-43"><span class="nt">endpoints</span><span class="p">:</span>
+</span><span id="line-44"><span class="w">  </span><span class="nt">app_server</span><span class="p">:</span>
+</span><span id="line-45"><span class="w">    </span><span class="c1"># value could be ip address or a hostname with port</span>
+</span><span id="line-46"><span class="w">    </span><span class="c1"># this could also be a list of endpoints for load balancing</span>
+</span><span id="line-47"><span class="w">    </span><span class="c1"># for example endpoint: [ ip1:port, ip2:port ]</span>
+</span><span id="line-48"><span class="w">    </span><span class="nt">endpoint</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">127.0.0.1:80</span>
+</span><span id="line-49"><span class="w">    </span><span class="c1"># max time to wait for a connection to be established</span>
+</span><span id="line-50"><span class="w">    </span><span class="nt">connect_timeout</span><span class="p">:</span><span class="w"> </span><span class="l l-Scalar l-Scalar-Plain">0.005s</span>
+</span></code></pre></div>
+</div>
+</section>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="model_serving.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="tech_overview.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Model Serving
+        Tech Overview
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="error_target.html">
-        Error Target
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="model_serving.html">
+        Bright Staff
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -385,15 +356,14 @@ processing request headers and then finalized by the HCM during post-request pro
 </div></div><aside class="hidden text-sm xl:block" id="right-sidebar">
 <div class="sticky top-16 -mt-10 max-h-[calc(100vh-5rem)] overflow-y-auto pt-6 space-y-2"><p class="font-medium">On this page</p>
 <ul>
-<li><a :data-current="activeSection === '#terminology'" class="reference internal" href="#terminology">Terminology</a></li>
 <li><a :data-current="activeSection === '#network-topology'" class="reference internal" href="#network-topology">Network topology</a></li>
 <li><a :data-current="activeSection === '#high-level-architecture'" class="reference internal" href="#high-level-architecture">High level architecture</a></li>
-<li><a :data-current="activeSection === '#configuration'" class="reference internal" href="#configuration">Configuration</a></li>
 <li><a :data-current="activeSection === '#request-flow-ingress'" class="reference internal" href="#request-flow-ingress">Request Flow (Ingress)</a></li>
 <li><a :data-current="activeSection === '#request-flow-egress'" class="reference internal" href="#request-flow-egress">Request Flow (Egress)</a><ul>
 <li><a :data-current="activeSection === '#post-request-processing'" class="reference internal" href="#post-request-processing">Post-request processing</a></li>
 </ul>
 </li>
+<li><a :data-current="activeSection === '#configuration'" class="reference internal" href="#configuration">Configuration</a></li>
 </ul>
 </div>
 </aside>
@@ -402,12 +372,12 @@ processing request headers and then finalized by the HCM during post-request pro
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/concepts/tech_overview/tech_overview.html b/resources/tech_overview/tech_overview.html
similarity index 78%
rename from concepts/tech_overview/tech_overview.html
rename to resources/tech_overview/tech_overview.html
index a0dd2ac8..62515bda 100755
--- a/concepts/tech_overview/tech_overview.html
+++ b/resources/tech_overview/tech_overview.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Tech Overview | Arch Docs v0.3.22</title>
-<meta content="Tech Overview | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Tech Overview | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Tech Overview | Plano Docs v0.4</title>
+<meta content="Tech Overview | Plano Docs v0.4" property="og:title"/>
+<meta content="Tech Overview | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/concepts/tech_overview/tech_overview.html" rel="canonical"/>
+<link href="./docs/resources/tech_overview/tech_overview.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
 <link href="../../search.html" rel="search" title="Search"/>
-<link href="terminology.html" rel="next" title="Terminology"/>
-<link href="../../get_started/quickstart.html" rel="prev" title="Quickstart"/>
+<link href="request_lifecycle.html" rel="next" title="Request Lifecycle"/>
+<link href="../../guides/state.html" rel="prev" title="Conversational State"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,70 +75,65 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul class="current">
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="current reference internal expandable" href="#">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html">Error Target</a></li>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
 </ul>
 </li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../prompt_target.html">Prompt Target</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/prompt_target.html">Prompt Target</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<ul class="current">
+<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="current reference internal expandable" href="#">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -165,56 +160,31 @@
 <span id="id1"></span><h1>Tech Overview<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#tech-overview"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
 <div class="toctree-wrapper compound">
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="terminology.html">Terminology</a></li>
-<li class="toctree-l1"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
-<li class="toctree-l1"><a class="reference internal" href="listener.html">Listener</a><ul>
-<li class="toctree-l2"><a class="reference internal" href="listener.html#downstream-ingress">Downstream (Ingress)</a></li>
-<li class="toctree-l2"><a class="reference internal" href="listener.html#upstream-egress">Upstream (Egress)</a></li>
-<li class="toctree-l2"><a class="reference internal" href="listener.html#configure-listener">Configure Listener</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="prompt.html">Prompts</a><ul>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html#messages">Messages</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html#prompt-guard">Prompt Guard</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html#prompt-targets">Prompt Targets</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html#prompting-llms">Prompting LLMs</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="model_serving.html">Model Serving</a><ul>
-<li class="toctree-l2"><a class="reference internal" href="model_serving.html#local-serving-cpu-moderate">Local Serving (CPU - Moderate)</a></li>
-<li class="toctree-l2"><a class="reference internal" href="model_serving.html#cloud-serving-gpu-blazing-fast">Cloud Serving (GPU - Blazing Fast)</a></li>
-</ul>
-</li>
 <li class="toctree-l1"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a><ul>
-<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html#terminology">Terminology</a></li>
 <li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html#network-topology">Network topology</a></li>
 <li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html#high-level-architecture">High level architecture</a></li>
-<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html#configuration">Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html#request-flow-ingress">Request Flow (Ingress)</a></li>
 <li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html#request-flow-egress">Request Flow (Egress)</a></li>
+<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html#configuration">Configuration</a></li>
 </ul>
 </li>
-<li class="toctree-l1"><a class="reference internal" href="error_target.html">Error Target</a><ul>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html#key-concepts">Key Concepts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html#error-header-example">Error Header Example</a></li>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html#best-practices-and-tips">Best Practices and Tips</a></li>
-</ul>
-</li>
+<li class="toctree-l1"><a class="reference internal" href="model_serving.html">Bright Staff</a></li>
+<li class="toctree-l1"><a class="reference internal" href="threading_model.html">Threading Model</a></li>
 </ul>
 </div>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../../get_started/quickstart.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../../guides/state.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Quickstart
+        Conversational State
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="terminology.html">
-        Terminology
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="request_lifecycle.html">
+        Request Lifecycle
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -226,12 +196,12 @@
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/concepts/tech_overview/threading_model.html b/resources/tech_overview/threading_model.html
similarity index 85%
rename from concepts/tech_overview/threading_model.html
rename to resources/tech_overview/threading_model.html
index 09fb7421..86687374 100755
--- a/concepts/tech_overview/threading_model.html
+++ b/resources/tech_overview/threading_model.html
@@ -7,19 +7,19 @@
 <meta content="white" media="(prefers-color-scheme: light)" name="theme-color"/>
 <meta content="black" media="(prefers-color-scheme: dark)" name="theme-color"/>
 <meta content="width=device-width, initial-scale=1" name="viewport"/>
-<title>Threading Model | Arch Docs v0.3.22</title>
-<meta content="Threading Model | Arch Docs v0.3.22" property="og:title"/>
-<meta content="Threading Model | Arch Docs v0.3.22" name="twitter:title"/>
+<title>Threading Model | Plano Docs v0.4</title>
+<meta content="Threading Model | Plano Docs v0.4" property="og:title"/>
+<meta content="Threading Model | Plano Docs v0.4" name="twitter:title"/>
 <link href="../../_static/pygments.css?v=466e7b45" rel="stylesheet" type="text/css"/>
 <link href="../../_static/theme.css?v=42baaae4" rel="stylesheet" type="text/css"/>
-<link href="../../_static/_static/custom.css" rel="stylesheet" type="text/css"/>
 <link href="../../_static/sphinx-design.min.css?v=95c83b7e" rel="stylesheet" type="text/css"/>
+<link href="../../_static/css/custom.css?v=2929376a" rel="stylesheet" type="text/css"/>
 <link href="../../_static/awesome-sphinx-design.css?v=15e0fffa" rel="stylesheet" type="text/css"/>
-<link href="./docs/concepts/tech_overview/threading_model.html" rel="canonical"/>
+<link href="./docs/resources/tech_overview/threading_model.html" rel="canonical"/>
 <link href="../../_static/favicon.ico" rel="icon"/>
 <link href="../../search.html" rel="search" title="Search"/>
-<link href="listener.html" rel="next" title="Listener"/>
-<link href="terminology.html" rel="prev" title="Terminology"/>
+<link href="../deployment.html" rel="next" title="Deployment"/>
+<link href="model_serving.html" rel="prev" title="Bright Staff"/>
 <script>
     <!-- Prevent Flash of wrong theme -->
       const userPreference = localStorage.getItem('darkMode');
@@ -39,7 +39,7 @@
     </a><header class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
 <div class="hidden mr-4 md:flex">
 <a class="flex items-center mr-6" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="24" src="../../_static/favicon.ico" width="24"/><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a></div><button @click="showSidebar = true" class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden" type="button">
 <svg aria-hidden="true" fill="currentColor" height="24" viewbox="0 96 960 960" width="24" xmlns="http://www.w3.org/2000/svg">
 <path d="M152.587 825.087q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440Zm0-203.587q-19.152 0-32.326-13.174T107.087 576q0-19.152 13.174-32.326t32.326-13.174h320q19.152 0 32.326 13.174T518.087 576q0 19.152-13.174 32.326T472.587 621.5h-320Zm0-203.587q-19.152 0-32.326-13.174t-13.174-32.326q0-19.152 13.174-32.326t32.326-13.174h440q19.152 0 32.326 13.174t13.174 32.326q0 19.152-13.174 32.326t-32.326 13.174h-440ZM708.913 576l112.174 112.174q12.674 12.674 12.674 31.826t-12.674 31.826Q808.413 764.5 789.261 764.5t-31.826-12.674l-144-144Q600 594.391 600 576t13.435-31.826l144-144q12.674-12.674 31.826-12.674t31.826 12.674q12.674 12.674 12.674 31.826t-12.674 31.826L708.913 576Z"></path>
@@ -56,7 +56,7 @@
 </form>
 </div>
 <nav class="flex items-center space-x-1">
-<a href="https://github.com/katanemo/arch" rel="noopener nofollow" title="Visit repository on GitHub">
+<a href="https://github.com/katanemo/plano" rel="noopener nofollow" title="Visit repository on GitHub">
 <div class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
 <svg fill="currentColor" height="26px" style="margin-top:-2px;display:inline" viewbox="0 0 45 44" xmlns="http://www.w3.org/2000/svg"><path clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor" fill-rule="evenodd"></path></svg>
 </div>
@@ -75,70 +75,65 @@
 </header>
 <div class="flex-1"><div class="container flex-1 items-start md:grid md:grid-cols-[220px_minmax(0,1fr)] md:gap-6 lg:grid-cols-[240px_minmax(0,1fr)] lg:gap-10"><aside :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }" class="fixed inset-y-0 left-0 md:top-14 z-50 md:z-30 bg-background md:bg-transparent transition-all duration-100 -translate-x-full md:translate-x-0 ml-0 p-6 md:p-0 md:-ml-2 md:h-[calc(100vh-3.5rem)] w-5/6 md:w-full shrink-0 overflow-y-auto border-r border-border md:sticky" id="left-sidebar">
 <a class="!justify-start text-sm md:!hidden bg-background" href="../../index.html">
-<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+<img alt="Logo" class="mr-2 dark:invert" height="16" src="../../_static/favicon.ico" width="16"/><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
 </a>
 <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
 <div class="overflow-y-auto h-full w-full relative pr-6">
 
-<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async="" src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
 <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
-<ul class="current">
-<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="terminology.html">Terminology</a></li>
-<li class="toctree-l2 current"><a class="current reference internal" href="#">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="error_target.html">Error Target</a></li>
+<ul>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../concepts/llm_providers/llm_providers.html">Model (LLM) Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
+<li class="toctree-l2"><a class="reference internal" href="../../concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
 </ul>
 </li>
-<li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../llm_providers/llm_providers.html">LLM Providers<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/client_libraries.html">Client Libraries</a></li>
-<li class="toctree-l2"><a class="reference internal" href="../llm_providers/model_aliases.html">Model Aliases</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="../prompt_target.html">Prompt Target</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../concepts/prompt_target.html">Prompt Target</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="../../guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="../../guides/observability/observability.html">Observability<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul x-show="expanded">
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="../../guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../../guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/deployment.html">Deployment</a></li>
-<li class="toctree-l1"><a class="reference internal" href="../../resources/configuration_reference.html">Configuration Reference</a></li>
+<ul class="current">
+<li class="toctree-l1 current" x-data="{ expanded: $el.classList.contains('current') ? true : false }"><a :class="{ 'expanded' : expanded }" @click="expanded = !expanded" class="reference internal expandable" href="tech_overview.html">Tech Overview<button @click.prevent.stop="expanded = !expanded" type="button"><span class="sr-only"></span><svg fill="currentColor" height="18px" stroke="none" viewbox="0 0 24 24" width="18px" xmlns="http://www.w3.org/2000/svg"><path d="M10 6L8.59 7.41 13.17 12l-4.58 4.59L10 18l6-6z"></path></svg></button></a><ul class="current" x-show="expanded">
+<li class="toctree-l2"><a class="reference internal" href="request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2 current"><a class="current reference internal" href="#">Threading Model</a></li>
+</ul>
+</li>
+<li class="toctree-l1"><a class="reference internal" href="../deployment.html">Deployment</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="../llms_txt.html">llms.txt</a></li>
 </ul>
 </nav>
 </div>
@@ -153,7 +148,7 @@
 <div class="w-full min-w-0 mx-auto">
 <nav aria-label="breadcrumbs" class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
 <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground" href="../../index.html">
-<span class="hidden md:inline">Arch Docs v0.3.22</span>
+<span class="hidden md:inline">Plano Docs v0.4</span>
 <svg aria-label="Home" class="md:hidden" fill="currentColor" height="18" stroke="none" viewbox="0 96 960 960" width="18" xmlns="http://www.w3.org/2000/svg">
 <path d="M240 856h120V616h240v240h120V496L480 316 240 496v360Zm-80 80V456l320-240 320 240v480H520V696h-80v240H160Zm320-350Z"></path>
 </svg>
@@ -164,14 +159,14 @@
 <div id="content" role="main">
 <section id="threading-model">
 <span id="arch-overview-threading"></span><h1>Threading Model<a @click.prevent="window.navigator.clipboard.writeText($el.href); $el.setAttribute('data-tooltip', 'Copied!'); setTimeout(() =&gt; $el.setAttribute('data-tooltip', 'Copy link to this element'), 2000)" aria-label="Copy link to this element" class="headerlink" data-tooltip="Copy link to this element" href="#threading-model"><svg height="1em" viewbox="0 0 24 24" width="1em" xmlns="http://www.w3.org/2000/svg"><path d="M3.9 12c0-1.71 1.39-3.1 3.1-3.1h4V7H7c-2.76 0-5 2.24-5 5s2.24 5 5 5h4v-1.9H7c-1.71 0-3.1-1.39-3.1-3.1zM8 13h8v-2H8v2zm9-6h-4v1.9h4c1.71 0 3.1 1.39 3.1 3.1s-1.39 3.1-3.1 3.1h-4V17h4c2.76 0 5-2.24 5-5s-2.24-5-5-5z"></path></svg></a></h1>
-<p>Arch builds on top of Envoy’s single process with multiple threads architecture.</p>
+<p>Plano builds on top of Envoy’s single process with multiple threads architecture.</p>
 <p>A single <em>primary</em> thread controls various sporadic coordination tasks while some number of <em>worker</em>
 threads perform filtering, and forwarding.</p>
 <p>Once a connection is accepted, the connection spends the rest of its lifetime bound to a single worker
 thread. All the functionality around prompt handling from a downstream client is handled in a separate worker thread.
-This allows the majority of Arch to be largely single threaded (embarrassingly parallel) with a small amount
+This allows the majority of Plano to be largely single threaded (embarrassingly parallel) with a small amount
 of more complex code handling coordination between the worker threads.</p>
-<p>Generally, Arch is written to be 100% non-blocking.</p>
+<p>Generally, Plano is written to be 100% non-blocking.</p>
 <div class="admonition tip">
 <p class="admonition-title">Tip</p>
 <p>For most workloads we recommend configuring the number of worker threads to be equal to the number of
@@ -180,16 +175,16 @@ hardware threads on the machine.</p>
 </section>
 </div><div class="flex justify-between items-center pt-6 mt-12 border-t border-border gap-4">
 <div class="mr-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="terminology.html">
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="model_serving.html">
 <svg class="mr-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="15 18 9 12 15 6"></polyline>
 </svg>
-        Terminology
+        Bright Staff
       </a>
 </div>
 <div class="ml-auto">
-<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="listener.html">
-        Listener
+<a class="inline-flex items-center justify-center rounded-md text-sm font-medium transition-colors border border-input hover:bg-accent hover:text-accent-foreground py-2 px-4" href="../deployment.html">
+        Deployment
         <svg class="ml-2 h-4 w-4" fill="none" height="24" stroke="currentColor" stroke-linecap="round" stroke-linejoin="round" stroke-width="2" viewbox="0 0 24 24" width="24" xmlns="http://www.w3.org/2000/svg">
 <polyline points="9 18 15 12 9 6"></polyline>
 </svg>
@@ -201,12 +196,12 @@ hardware threads on the machine.</p>
 </div><footer class="py-6 border-t border-border md:py-0">
 <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
 <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 23, 2025. </p>
+<p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc Last updated: Dec 24, 2025. </p>
 </div>
 </div>
 </footer>
 </div>
-<script src="../../_static/documentation_options.js?v=3bad885e"></script>
+<script src="../../_static/documentation_options.js?v=8cf1ab6b"></script>
 <script src="../../_static/doctools.js?v=9bcbadda"></script>
 <script src="../../_static/sphinx_highlight.js?v=dc90522c"></script>
 <script defer="defer" src="../../_static/theme.js?v=073f68d9"></script>
diff --git a/search.html b/search.html
index 8d8a6218..325ad3f0 100755
--- a/search.html
+++ b/search.html
@@ -12,13 +12,13 @@
   <meta name="theme-color" media="(prefers-color-scheme: light)" content="white" />
   <meta name="theme-color" media="(prefers-color-scheme: dark)" content="black" />
   
-    <title>Search | Arch Docs v0.3.22</title>
-    <meta property="og:title" content="Search | Arch Docs v0.3.22" />
-    <meta name="twitter:title" content="Search | Arch Docs v0.3.22" />
+    <title>Search | Plano Docs v0.4</title>
+    <meta property="og:title" content="Search | Plano Docs v0.4" />
+    <meta name="twitter:title" content="Search | Plano Docs v0.4" />
       <link rel="stylesheet" type="text/css" href="_static/pygments.css?v=466e7b45" />
       <link rel="stylesheet" type="text/css" href="_static/theme.css?v=42baaae4" />
-      <link rel="stylesheet" type="text/css" href="_static/_static/custom.css" />
       <link rel="stylesheet" type="text/css" href="_static/sphinx-design.min.css?v=95c83b7e" />
+      <link rel="stylesheet" type="text/css" href="_static/css/custom.css?v=2929376a" />
       <link rel="stylesheet" type="text/css" href="_static/awesome-sphinx-design.css?v=15e0fffa" />
     <link rel="canonical" href="./docs/search.html" />
       <link rel="icon" href="_static/favicon.ico" />
@@ -44,7 +44,7 @@
   class="sticky top-0 z-40 w-full border-b shadow-sm border-border supports-backdrop-blur:bg-background/60 bg-background/95 backdrop-blur"><div class="container flex items-center h-14">
     <div class="hidden mr-4 md:flex">
       <a href="index.html" class="flex items-center mr-6">
-          <img height="24" width="24" class="mr-2 dark:invert" src="_static/favicon.ico" alt="Logo" /><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+          <img height="24" width="24" class="mr-2 dark:invert" src="_static/favicon.ico" alt="Logo" /><span class="hidden font-bold sm:inline-block text-clip whitespace-nowrap">Plano Docs v0.4</span>
       </a></div><button
       class="inline-flex items-center justify-center h-10 px-0 py-2 mr-2 text-base font-medium transition-colors rounded-md hover:text-accent-foreground hover:bg-transparent md:hidden"
       type="button" @click="showSidebar = true">
@@ -75,7 +75,7 @@
 </form>
       </div>
       <nav class="flex items-center space-x-1">
-        <a href="https://github.com/katanemo/arch" title="Visit repository on GitHub" rel="noopener nofollow">
+        <a href="https://github.com/katanemo/plano" title="Visit repository on GitHub" rel="noopener nofollow">
           <div
             class="inline-flex items-center justify-center px-0 text-sm font-medium transition-colors rounded-md disabled:opacity-50 disabled:pointer-events-none hover:bg-accent hover:text-accent-foreground h-9 w-9">
             <svg height="26px" style="margin-top:-2px;display:inline" viewBox="0 0 45 44" fill="currentColor" xmlns="http://www.w3.org/2000/svg"><path fill-rule="evenodd" clip-rule="evenodd" d="M22.477.927C10.485.927.76 10.65.76 22.647c0 9.596 6.223 17.736 14.853 20.608 1.087.2 1.483-.47 1.483-1.047 0-.516-.019-1.881-.03-3.693-6.04 1.312-7.315-2.912-7.315-2.912-.988-2.51-2.412-3.178-2.412-3.178-1.972-1.346.149-1.32.149-1.32 2.18.154 3.327 2.24 3.327 2.24 1.937 3.318 5.084 2.36 6.321 1.803.197-1.403.759-2.36 1.379-2.903-4.823-.548-9.894-2.412-9.894-10.734 0-2.37.847-4.31 2.236-5.828-.224-.55-.969-2.759.214-5.748 0 0 1.822-.584 5.972 2.226 1.732-.482 3.59-.722 5.437-.732 1.845.01 3.703.25 5.437.732 4.147-2.81 5.967-2.226 5.967-2.226 1.185 2.99.44 5.198.217 5.748 1.392 1.517 2.232 3.457 2.232 5.828 0 8.344-5.078 10.18-9.916 10.717.779.67 1.474 1.996 1.474 4.021 0 2.904-.027 5.247-.027 5.96 0 .58.392 1.256 1.493 1.044C37.981 40.375 44.2 32.24 44.2 22.647c0-11.996-9.726-21.72-21.722-21.72" fill="currentColor"/></svg>
@@ -107,42 +107,35 @@
   :aria-hidden="!showSidebar" :class="{ 'translate-x-0': showSidebar }">
 
     <a href="index.html" class="!justify-start text-sm md:!hidden bg-background">
-        <img height="16" width="16" class="mr-2 dark:invert" src="_static/favicon.ico" alt="Logo" /><span class="font-bold text-clip whitespace-nowrap">Arch Docs v0.3.22</span>
+        <img height="16" width="16" class="mr-2 dark:invert" src="_static/favicon.ico" alt="Logo" /><span class="font-bold text-clip whitespace-nowrap">Plano Docs v0.4</span>
     </a>
 
     <div class="relative overflow-hidden md:overflow-auto my-4 md:my-0 h-[calc(100vh-8rem)] md:h-auto">
       <div class="overflow-y-auto h-full w-full relative pr-6"><!-- _templates/analytics.html -->
 
 <!-- Google tag (gtag.js) -->
-<script async src="https://www.googletagmanager.com/gtag/js?id=G-K2LXXSX6HB"></script>
+<script async src="https://www.googletagmanager.com/gtag/js?id=G-EH2VW19FXE"></script>
 <script>
   window.dataLayer = window.dataLayer || [];
   function gtag(){dataLayer.push(arguments);}
   gtag('js', new Date());
 
-  gtag('config', 'G-K2LXXSX6HB');
+  gtag('config', 'G-EH2VW19FXE');
 </script>
 <nav class="table w-full min-w-full my-6 lg:my-8">
   <p class="caption" role="heading"><span class="caption-text">Get Started</span></p>
 <ul>
 <li class="toctree-l1"><a class="reference internal" href="get_started/overview.html">Overview</a></li>
-<li class="toctree-l1"><a class="reference internal" href="get_started/intro_to_arch.html">Intro to Arch</a></li>
+<li class="toctree-l1"><a class="reference internal" href="get_started/intro_to_plano.html">Intro to Plano</a></li>
 <li class="toctree-l1"><a class="reference internal" href="get_started/quickstart.html">Quickstart</a></li>
 <li class="toctree-l1"><a class="reference internal" href="get_started/quickstart.html#next-steps">Next Steps</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Concepts</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="concepts/tech_overview/tech_overview.html">Tech Overview</a><ul>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/terminology.html">Terminology</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/threading_model.html">Threading Model</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/listener.html">Listener</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/prompt.html">Prompts</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/model_serving.html">Model Serving</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
-<li class="toctree-l2"><a class="reference internal" href="concepts/tech_overview/error_target.html">Error Target</a></li>
-</ul>
-</li>
-<li class="toctree-l1"><a class="reference internal" href="concepts/llm_providers/llm_providers.html">LLM Providers</a><ul>
+<li class="toctree-l1"><a class="reference internal" href="concepts/listeners.html">Listeners</a></li>
+<li class="toctree-l1"><a class="reference internal" href="concepts/agents.html">Agents</a></li>
+<li class="toctree-l1"><a class="reference internal" href="concepts/filter_chain.html">Filter Chains</a></li>
+<li class="toctree-l1"><a class="reference internal" href="concepts/llm_providers/llm_providers.html">Model (LLM) Providers</a><ul>
 <li class="toctree-l2"><a class="reference internal" href="concepts/llm_providers/supported_providers.html">Supported Providers &amp; Configuration</a></li>
 <li class="toctree-l2"><a class="reference internal" href="concepts/llm_providers/client_libraries.html">Client Libraries</a></li>
 <li class="toctree-l2"><a class="reference internal" href="concepts/llm_providers/model_aliases.html">Model Aliases</a></li>
@@ -152,27 +145,29 @@
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Guides</span></p>
 <ul>
-<li class="toctree-l1"><a class="reference internal" href="guides/prompt_guard.html">Prompt Guard</a></li>
-<li class="toctree-l1"><a class="reference internal" href="guides/agent_routing.html">Agent Routing and Hand Off</a></li>
-<li class="toctree-l1"><a class="reference internal" href="guides/function_calling.html">Function Calling</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/orchestration.html">Orchestration</a></li>
 <li class="toctree-l1"><a class="reference internal" href="guides/llm_router.html">LLM Routing</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/function_calling.html">Function Calling</a></li>
 <li class="toctree-l1"><a class="reference internal" href="guides/observability/observability.html">Observability</a><ul>
 <li class="toctree-l2"><a class="reference internal" href="guides/observability/tracing.html">Tracing</a></li>
 <li class="toctree-l2"><a class="reference internal" href="guides/observability/monitoring.html">Monitoring</a></li>
 <li class="toctree-l2"><a class="reference internal" href="guides/observability/access_logging.html">Access Logging</a></li>
 </ul>
 </li>
-</ul>
-<p class="caption" role="heading"><span class="caption-text">Build with Arch</span></p>
-<ul>
-<li class="toctree-l1"><a class="reference internal" href="build_with_arch/agent.html">Agentic Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="build_with_arch/rag.html">RAG Apps</a></li>
-<li class="toctree-l1"><a class="reference internal" href="build_with_arch/multi_turn.html">Multi-Turn</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/prompt_guard.html">Guardrails</a></li>
+<li class="toctree-l1"><a class="reference internal" href="guides/state.html">Conversational State</a></li>
 </ul>
 <p class="caption" role="heading"><span class="caption-text">Resources</span></p>
 <ul>
+<li class="toctree-l1"><a class="reference internal" href="resources/tech_overview/tech_overview.html">Tech Overview</a><ul>
+<li class="toctree-l2"><a class="reference internal" href="resources/tech_overview/request_lifecycle.html">Request Lifecycle</a></li>
+<li class="toctree-l2"><a class="reference internal" href="resources/tech_overview/model_serving.html">Bright Staff</a></li>
+<li class="toctree-l2"><a class="reference internal" href="resources/tech_overview/threading_model.html">Threading Model</a></li>
+</ul>
+</li>
 <li class="toctree-l1"><a class="reference internal" href="resources/deployment.html">Deployment</a></li>
 <li class="toctree-l1"><a class="reference internal" href="resources/configuration_reference.html">Configuration Reference</a></li>
+<li class="toctree-l1"><a class="reference internal" href="resources/llms_txt.html">llms.txt</a></li>
 </ul>
 
 </nav>
@@ -193,7 +188,7 @@
      class="flex items-center mb-4 space-x-1 text-sm text-muted-foreground">
   <a class="overflow-hidden text-ellipsis whitespace-nowrap hover:text-foreground"
      href="index.html">
-    <span class="hidden md:inline">Arch Docs v0.3.22</span>
+    <span class="hidden md:inline">Plano Docs v0.4</span>
     <svg xmlns="http://www.w3.org/2000/svg"
          height="18"
          width="18"
@@ -221,14 +216,14 @@
     </div><footer class="py-6 border-t border-border md:py-0">
     <div class="container flex flex-col items-center justify-between gap-4 md:h-24 md:flex-row">
       <div class="flex flex-col items-center gap-4 px-8 md:flex-row md:gap-2 md:px-0">
-        <p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc&nbsp;Last updated: Dec 23, 2025.&nbsp;</p>
+        <p class="text-sm leading-loose text-center text-muted-foreground md:text-left">© 2025, Katanemo Labs, Inc&nbsp;Last updated: Dec 24, 2025.&nbsp;</p>
 </div>
 </div>
 </footer>
   </div>
   
   
-    <script src="_static/documentation_options.js?v=3bad885e"></script>
+    <script src="_static/documentation_options.js?v=8cf1ab6b"></script>
     <script src="_static/doctools.js?v=9bcbadda"></script>
     <script src="_static/sphinx_highlight.js?v=dc90522c"></script>
     <script defer="defer" src="_static/theme.js?v=073f68d9"></script>
diff --git a/searchindex.js b/searchindex.js
index 3679b02d..d41d6505 100755
--- a/searchindex.js
+++ b/searchindex.js
@@ -1 +1 @@
-Search.setIndex({"alltitles":{"AI Agent Tracing Visualization Example":[[25,"ai-agent-tracing-visualization-example"]],"AWS X-Ray":[[25,"aws-x-ray"]],"Access Logging":[[22,null]],"Additional Resources":[[25,"additional-resources"]],"Advanced Configuration":[[6,"advanced-configuration"]],"Advanced Features":[[4,"advanced-features"]],"Advanced Features (Coming Soon)":[[5,"advanced-features-coming-soon"]],"Agent Routing and Hand Off":[[19,null]],"Agentic Apps":[[0,null]],"Agentic Apps via Prompt Targets":[[11,"agentic-apps-via-prompt-targets"]],"Alias-based Routing":[[21,"alias-based-routing"]],"Alias-based Routing Workflow":[[21,"alias-based-routing-workflow"]],"Amazon Bedrock":[[6,"amazon-bedrock"]],"Anthropic":[[6,"anthropic"]],"Anthropic (Python) SDK":[[3,"anthropic-python-sdk"]],"Arch-Function":[[20,"arch-function"]],"Arch-Router":[[21,"id1"]],"Azure OpenAI":[[6,"azure-openai"]],"Base URL Configuration":[[6,"base-url-configuration"]],"Basic Configuration":[[5,"basic-configuration"],[7,"basic-configuration"]],"Benefits of Using Arch Guard":[[26,"benefits-of-using-arch-guard"]],"Benefits of Using Traceparent Headers":[[25,"benefits-of-using-traceparent-headers"]],"Best Practices":[[3,"best-practices"],[25,"best-practices"]],"Best Practices and Tips":[[8,"best-practices-and-tips"],[19,"best-practices-and-tips"],[20,"best-practices-and-tips"]],"Best practicesm":[[21,"best-practicesm"]],"Build AI Agent with Arch Gateway":[[18,"build-ai-agent-with-arch-gateway"]],"Build Multi-Turn RAG Apps":[[1,"build-multi-turn-rag-apps"]],"Build with Arch":[[17,"build-with-arch"],[27,null]],"Client Libraries":[[3,null]],"Cloud Serving (GPU - Blazing Fast)":[[10,"cloud-serving-gpu-blazing-fast"]],"Combining Routing Methods":[[21,"combining-routing-methods"]],"Common Issues and Solutions":[[29,"common-issues-and-solutions"]],"Common Use Cases":[[4,"common-use-cases"]],"Concepts":[[17,"concepts"],[27,null]],"Configuration":[[12,"configuration"]],"Configuration Reference":[[28,null]],"Configuration Structure":[[6,"configuration-structure"]],"Configure Listener":[[9,"configure-listener"]],"Configure Monitoring":[[23,"configure-monitoring"]],"Configuring Prompt Targets":[[7,"configuring-prompt-targets"]],"Core Capabilities":[[4,"core-capabilities"]],"Cross-Client Compatibility":[[3,"cross-client-compatibility"]],"Datadog":[[25,"datadog"]],"DeepSeek":[[6,"deepseek"]],"Default Model Configuration":[[6,"default-model-configuration"]],"Default Targets":[[7,"default-targets"]],"Defining Parameters":[[7,"defining-parameters"]],"Demo App":[[1,"demo-app"]],"Deployment":[[29,null]],"Docker Compose Setup":[[29,"docker-compose-setup"]],"Docker Deployment":[[29,"docker-deployment"]],"Downstream (Ingress)":[[9,"downstream-ingress"]],"Error Handling":[[3,"error-handling"]],"Error Header Example":[[8,"error-header-example"]],"Error Target":[[8,null]],"Example 1: Adjusting Retrieval":[[1,"example-1-adjusting-retrieval"]],"Example 2: Switching Intent":[[1,"example-2-switching-intent"]],"Example Configuration":[[26,"example-configuration"]],"Example Configuration For Agents":[[7,"example-configuration-for-agents"]],"Example Configuration For Tools":[[7,"example-configuration-for-tools"]],"Example Use Cases":[[19,"example-use-cases"],[20,"example-use-cases"],[21,"example-use-cases"]],"Example with OpenTelemetry in Python":[[25,"example-with-opentelemetry-in-python"]],"Example: Using OpenAI Client with Arch as an Egress Gateway":[[11,"example-using-openai-client-with-arch-as-an-egress-gateway"]],"First-Class Providers":[[6,"first-class-providers"]],"Function Calling":[[20,null]],"Function Calling Workflow":[[20,"function-calling-workflow"]],"Gateway Endpoints":[[3,"gateway-endpoints"]],"Gateway Smoke Test":[[29,"gateway-smoke-test"]],"Get Started":[[17,"get-started"],[27,null]],"Getting Started":[[4,"getting-started"]],"Google Gemini":[[6,"google-gemini"]],"Groq":[[6,"groq"]],"Guides":[[17,"guides"],[27,null]],"Header Format":[[25,"header-format"]],"High level architecture":[[12,"high-level-architecture"]],"How Arch-Guard Works":[[26,"how-arch-guard-works"]],"How It Works":[[22,"how-it-works"]],"How to Initiate A Trace":[[25,"how-to-initiate-a-trace"]],"Implementing Function Calling":[[20,"implementing-function-calling"]],"Implementing Routing":[[21,"implementing-routing"]],"Instrumentation":[[25,"instrumentation"]],"Integrating with Tracing Tools":[[25,"integrating-with-tracing-tools"]],"Intent Matching":[[7,"intent-matching"],[11,"intent-matching"]],"Intro to Arch":[[16,null]],"Key Benefits":[[4,"key-benefits"]],"Key Concepts":[[8,"key-concepts"]],"Key Features":[[7,"key-features"],[20,"key-features"],[22,"key-features"]],"LLM Providers":[[4,null]],"LLM Routing":[[21,null]],"Langtrace":[[25,"langtrace"]],"Listener":[[9,null]],"Local Serving (CPU - Moderate)":[[10,"local-serving-cpu-moderate"]],"Log Format":[[22,"log-format"]],"Messages":[[11,"messages"]],"Metrics Dashboard (via Grafana)":[[23,"metrics-dashboard-via-grafana"]],"Mistral AI":[[6,"mistral-ai"]],"Model Aliases":[[5,null]],"Model Selection Guidelines":[[6,"model-selection-guidelines"]],"Model Serving":[[10,null]],"Model-Based Routing":[[29,"model-based-routing"]],"Model-based Routing":[[21,"model-based-routing"]],"Model-based Routing Workflow":[[21,"model-based-routing-workflow"]],"Monitoring":[[23,null]],"Moonshot AI":[[6,"moonshot-ai"]],"Multi-Turn":[[1,null]],"Multi-Turn RAG (Follow-up Questions)":[[2,"multi-turn-rag-follow-up-questions"]],"Multiple Provider Instances":[[6,"multiple-provider-instances"]],"Naming Best Practices":[[5,"naming-best-practices"]],"Network topology":[[12,"network-topology"]],"Next Steps":[[18,"next-steps"]],"Observability":[[24,null]],"Ollama":[[6,"ollama"]],"OpenAI":[[6,"openai"]],"OpenAI (Python) SDK":[[3,"openai-python-sdk"]],"OpenAI-Compatible Providers":[[6,"openai-compatible-providers"]],"Overview":[[17,null],[25,"overview"]],"Parallel & Multiple Function Calling":[[0,"parallel-multiple-function-calling"]],"Parameter Extraction for RAG":[[2,"parameter-extraction-for-rag"]],"Post-request processing":[[12,"post-request-processing"]],"Preference-aligned Routing (Arch-Router)":[[21,"preference-aligned-routing-arch-router"]],"Preference-aligned Routing Workflow (Arch-Router)":[[21,"preference-aligned-routing-workflow-arch-router"]],"Prerequisites":[[18,"prerequisites"]],"Prompt Guard":[[11,"prompt-guard"],[26,null]],"Prompt Target":[[7,null]],"Prompt Targets":[[11,"prompt-targets"]],"Prompting LLMs":[[11,"prompting-llms"]],"Prompts":[[11,null]],"Provider Categories":[[6,"provider-categories"]],"Providers Requiring Base URL":[[6,"providers-requiring-base-url"]],"Quickstart":[[18,null]],"Qwen (Alibaba)":[[6,"qwen-alibaba"]],"RAG Apps":[[2,null]],"Request Flow (Egress)":[[12,"request-flow-egress"]],"Request Flow (Ingress)":[[12,"request-flow-ingress"]],"Request Lifecycle":[[12,null]],"Resources":[[27,null]],"Routing Logic":[[7,"routing-logic"]],"Routing Methods":[[21,"routing-methods"]],"Routing Preferences":[[6,"routing-preferences"]],"Runtime Tests":[[29,"runtime-tests"]],"See Also":[[3,"see-also"],[5,"see-also"],[6,"see-also"]],"Single Function Call":[[0,"single-function-call"]],"Starting the Stack":[[29,"starting-the-stack"]],"Step 1. Create arch config file":[[18,"step-1-create-arch-config-file"],[18,"id2"]],"Step 1: Define Arch Config":[[1,"step-1-define-arch-config"]],"Step 1: Define Prompt Targets":[[0,"step-1-define-prompt-targets"],[0,"id1"],[2,"step-1-define-prompt-targets"]],"Step 1: Define the Function":[[20,"step-1-define-the-function"]],"Step 2. Start arch gateway":[[18,"step-2-start-arch-gateway"]],"Step 2. Start arch gateway with currency conversion config":[[18,"step-2-start-arch-gateway-with-currency-conversion-config"]],"Step 2: Configure Prompt Targets":[[20,"step-2-configure-prompt-targets"]],"Step 2: Process Request Parameters":[[0,"step-2-process-request-parameters"]],"Step 2: Process Request Parameters in Flask":[[2,"step-2-process-request-parameters-in-flask"]],"Step 2: Process Request in Flask":[[1,"step-2-process-request-in-flask"]],"Step 3. Interacting with gateway using curl command":[[18,"step-3-interacting-with-gateway-using-curl-command"]],"Step 3.1: Using OpenAI Python client":[[18,"step-3-1-using-openai-python-client"]],"Step 3.2: Using curl command":[[18,"step-3-2-using-curl-command"]],"Step 3: Arch Takes Over":[[20,"step-3-arch-takes-over"]],"Step 3: Interact with LLM":[[18,"step-3-interact-with-llm"]],"Summary":[[7,"summary"],[25,"summary"],[26,"summary"]],"Supported API Endpoints":[[6,"supported-api-endpoints"]],"Supported Clients":[[3,"supported-clients"]],"Supported Providers & Configuration":[[6,null]],"Tech Overview":[[13,null]],"Terminology":[[12,"terminology"],[14,null]],"Threading Model":[[15,null]],"Together AI":[[6,"together-ai"]],"Trace Breakdown:":[[25,"trace-breakdown"]],"Trace Propagation":[[25,"trace-propagation"]],"Tracing":[[25,null]],"Troubleshooting":[[29,"troubleshooting"]],"Upstream (Egress)":[[9,"upstream-egress"]],"Use Arch Gateway as LLM Router":[[18,"use-arch-gateway-as-llm-router"]],"Using Aliases":[[5,"using-aliases"]],"Validation Rules":[[5,"validation-rules"]],"Welcome to Arch!":[[27,null]],"What Are Prompt Targets?":[[7,"what-are-prompt-targets"]],"What Is Arch-Guard":[[26,"what-is-arch-guard"]],"What is Function Calling?":[[20,"what-is-function-calling"]],"What is Retrieval-Augmented Generation (RAG)?":[[2,"what-is-retrieval-augmented-generation-rag"]],"Why Prompt Guard":[[26,"why-prompt-guard"]],"Zhipu AI":[[6,"zhipu-ai"]],"cURL Examples":[[3,"curl-examples"]],"xAI":[[6,"xai"]]},"docnames":["build_with_arch/agent","build_with_arch/multi_turn","build_with_arch/rag","concepts/llm_providers/client_libraries","concepts/llm_providers/llm_providers","concepts/llm_providers/model_aliases","concepts/llm_providers/supported_providers","concepts/prompt_target","concepts/tech_overview/error_target","concepts/tech_overview/listener","concepts/tech_overview/model_serving","concepts/tech_overview/prompt","concepts/tech_overview/request_lifecycle","concepts/tech_overview/tech_overview","concepts/tech_overview/terminology","concepts/tech_overview/threading_model","get_started/intro_to_arch","get_started/overview","get_started/quickstart","guides/agent_routing","guides/function_calling","guides/llm_router","guides/observability/access_logging","guides/observability/monitoring","guides/observability/observability","guides/observability/tracing","guides/prompt_guard","index","resources/configuration_reference","resources/deployment"],"envversion":{"sphinx":65,"sphinx.domains.c":3,"sphinx.domains.changeset":1,"sphinx.domains.citation":1,"sphinx.domains.cpp":9,"sphinx.domains.index":1,"sphinx.domains.javascript":3,"sphinx.domains.math":2,"sphinx.domains.python":4,"sphinx.domains.rst":2,"sphinx.domains.std":2,"sphinx.ext.intersphinx":1,"sphinx.ext.viewcode":1},"filenames":["build_with_arch/agent.rst","build_with_arch/multi_turn.rst","build_with_arch/rag.rst","concepts/llm_providers/client_libraries.rst","concepts/llm_providers/llm_providers.rst","concepts/llm_providers/model_aliases.rst","concepts/llm_providers/supported_providers.rst","concepts/prompt_target.rst","concepts/tech_overview/error_target.rst","concepts/tech_overview/listener.rst","concepts/tech_overview/model_serving.rst","concepts/tech_overview/prompt.rst","concepts/tech_overview/request_lifecycle.rst","concepts/tech_overview/tech_overview.rst","concepts/tech_overview/terminology.rst","concepts/tech_overview/threading_model.rst","get_started/intro_to_arch.rst","get_started/overview.rst","get_started/quickstart.rst","guides/agent_routing.rst","guides/function_calling.rst","guides/llm_router.rst","guides/observability/access_logging.rst","guides/observability/monitoring.rst","guides/observability/observability.rst","guides/observability/tracing.rst","guides/prompt_guard.rst","index.rst","resources/configuration_reference.rst","resources/deployment.rst"],"indexentries":{},"objects":{},"objnames":{},"objtypes":{},"terms":{"":[0,1,2,3,4,5,6,7,8,9,10,11,12,14,15,16,17,18,19,20,21,22,25,27,29],"0":[0,1,3,5,6,9,10,11,12,14,18,21,22,25,28,29],"00":25,"003":11,"005":[0,11,12,28],"01":3,"03":22,"04":22,"05":18,"050z":22,"05m":10,"06":[3,18,29],"08":[18,29],"0905":6,"1":[3,4,5,6,8,9,11,12,14,19,21,22,23,25,28],"10":[5,18,20,22,23,25],"100":[3,5,15,25],"10000":[9,11,12,18,28,29],"1022":22,"104":22,"106":22,"10m":11,"10t03":22,"10x":[10,11],"11":[4,18],"11434":[5,6],"12":[7,18],"12000":[3,5,6,9,11,14,18,21,22,25,28,29],"120b":5,"1234":8,"125":22,"127":[0,1,3,5,9,11,12,14,18,22,28],"128k":6,"12x":0,"1301":22,"131":18,"13b":6,"140":22,"15":23,"156x":10,"159":22,"16":[18,25],"162":22,"168":22,"1695":22,"1695m":22,"17":22,"18083":[0,22],"192":22,"19901":23,"1d6b30cfc845":22,"2":[5,6,12,19,25],"20":5,"200":[0,2,3,22],"200m":[0,10,11],"2019":7,"2023":3,"2024":[18,22,29],"20240229":6,"20240307":[5,6],"20241022":[3,5,6,21,29],"2025":26,"21797":22,"218":22,"22":[18,29],"24":18,"245":22,"25":18,"254":22,"27":18,"28":18,"288":18,"29":18,"2a5b":22,"3":[0,1,3,5,6,11,19,21,25,29],"30":[6,9,11,12,18,21],"31":7,"32":25,"32768":6,"32b":6,"32k":6,"34b":6,"3b":[6,18],"4":[6,11,20,25],"400":[0,2,8],"429":12,"4317":25,"441":22,"443":[0,1,6,14,18,22],"447":22,"44x":[0,11],"463":22,"469793af":22,"485":18,"49":22,"492z":22,"4o":[0,3,5,6,9,10,11,12,18,21,25,28,29],"5":[0,1,3,5,6,11,18,21,25,28,29],"50":3,"51":[18,22],"51000":22,"52":22,"53":22,"537z":22,"54":22,"55":22,"556":22,"56":[18,22],"598z":22,"5b":[4,21],"6":[6,25],"60":28,"604197fe":22,"614":22,"647":18,"65":22,"7":[0,2,25],"70b":6,"770":22,"78558":18,"7b":[6,28],"8":25,"80":[5,6,11,12,14,28],"8000":6,"8001":28,"8080":[0,1,6],"8192":6,"825":18,"87":22,"8b":6,"8x7b":[6,28],"9":25,"905z":22,"906z":22,"9090":23,"9367":22,"95":5,"95a2":22,"961z":22,"979":18,"984":22,"984m":22,"9b57":22,"A":[0,2,5,7,8,9,12,14,15,16,17,20,24],"And":18,"As":[11,18,23,26],"At":21,"Be":[20,25],"But":[14,16],"By":[2,7,11,12,16,19,20,22,25,26,28],"FOR":22,"For":[1,2,3,6,9,11,14,15,16,17,18,19,20,21,22,25,29],"If":[0,1,8,11,12,20,21,26,29],"In":[0,2,7,11,12,16,17,18,20,21,26],"It":[0,6,12,17,18,20,21,24,25,26,27],"Its":[16,23,25],"No":[11,18,21],"On":18,"One":3,"Or":[3,25],"The":[1,2,3,5,6,7,8,10,11,12,14,18,19,20,21,22,25,28],"There":[10,16],"These":[0,2,8,14,16,20],"To":[2,9,11,14,16,18,19,21,23,25,29],"With":[0,2,6,16,18,21,25,26],"__init__":19,"__main__":[0,2],"__name__":[0,2,25],"a1c":1,"abil":[6,7,11,12,14,18,26,28],"abl":14,"about":[0,1,3,4,5,10,11,12,14,16,17,18,19,21,22,26,28],"abov":[0,1,2,5,12,16,19,23,29],"abstract":[4,21],"accept":[9,11,12,15],"access":[4,12,17,18,23,24,25,27],"access_":22,"access_ingress":22,"access_intern":22,"access_kei":[0,1,5,6,9,11,12,18,21,28],"access_llm":22,"accident":12,"accordingli":25,"account":[20,25],"accur":[1,2,7,11,14,16,20,21,26],"accuraci":[2,11,21,26],"achiev":[11,16,20],"across":[0,1,4,5,6,12,16,20,25],"act":[7,9,16,23],"action":[0,7,8,11,12,14,20,21,26,28],"activ":[6,12,18],"actual":[0,2,3,11,19,21,25,28],"ad":[0,14,16,21,26],"adapt":21,"add":[1,2,4,9,11,20,21,25],"add_span_processor":25,"addit":[1,2,6,7,9,11,12,24,26],"addition":1,"address":[0,1,6,9,11,12,14,18,21,28],"adher":26,"adilhafeez":22,"adjust":[14,19,25,29],"adopt":[17,25,27],"advanc":[11,20],"advantag":[11,16],"advic":[0,9,11,12,28],"affect":25,"afraid":20,"african":18,"after":[12,19,20,21,25],"against":[11,26],"agent":[1,2,6,12,14,16,17,20,22,26,27,28],"agent_nam":19,"aggreg":20,"agnost":17,"agre":19,"ai":[2,4,7,12,14,16,17,19,20,22,27],"aid":12,"air":[1,6],"airbnb":16,"alert":[22,23],"alertmanag":23,"algorithm":[5,16],"alia":[3,4,5,28],"alias":[3,4,6,21,27,28],"align":[4,11],"aliyunc":6,"all":[0,1,4,6,9,11,12,15,16,17,25,28,29],"allow":[0,5,6,7,11,14,15,16,18,20,21,25],"along":18,"alongsid":[10,12,14,16],"alphanumer":5,"alreadi":25,"also":[0,4,7,11,12,14,18,21,28],"altern":[8,26],"alwai":[5,8,19,21],"amazon_bedrock":6,"amazonaw":6,"ambigu":20,"amount":15,"an":[0,1,2,8,12,14,16,18,20,21,23,25,26,28],"analysi":[6,8,19,21],"analyst":5,"analyt":20,"analyz":[0,6,7,11,16,19,20,21,22,25,26],"analyze_earnings_report":21,"ang":11,"ani":[0,1,3,4,6,7,11,12,16,17,18,20,25,26,27,29],"annot":1,"answer":[1,11,12,19,28],"anthrop":[4,5,21,29],"anthropic_api_kei":[5,6,21,29],"anyth":21,"api":[0,1,3,4,5,7,8,9,10,11,12,14,16,17,18,19,20,22,23,25],"api_bas":11,"api_kei":[3,6,11,18,25],"api_serv":[7,20,22],"api_vers":23,"apierror":3,"apivers":23,"apm":25,"app":[14,16,17,19,27,29],"app_serv":[0,11,12,28],"append":[0,2,6],"appli":[4,5,11,12,14,16,17,21,26,27,28],"applic":[0,2,3,4,5,6,7,8,9,10,11,12,14,16,17,18,19,20,21,22,23,25,26,29],"appoint":21,"approach":[2,4,5,7,11,21,25,28],"appropri":[2,7,8,11,12,14,19,21],"ar":[0,1,2,3,4,5,6,8,9,10,11,12,14,16,18,19,20,21,22,23,25,26,28,29],"arch":[0,2,3,4,5,6,7,8,9,10,12,14,15,19,22,23,25,28,29],"arch_config":[1,3,6,11,14,18,28,29],"arch_llm_listen":22,"archgw":[10,18,23,29],"archgw_log":22,"architectur":[7,13,14,15,16,21,25,26],"area":4,"aren":6,"aris":[8,14],"around":[12,15],"arriv":12,"art":[11,18,20,21],"artifici":3,"ask":[1,2,18,20,25],"ask_quest":3,"aspect":[16,23],"assembl":3,"assign":[19,21],"assist":[0,1,6,9,11,12,18,19,21,25,26,28],"associ":21,"assum":[1,11,18],"attach":7,"attack":26,"attempt":[11,12,16,26],"attent":[19,20],"attribut":[7,20,25],"aud":18,"augment":12,"australian":18,"authent":6,"author":[3,22],"auto_llm_dispatch_on_respons":[11,12,28],"autom":[19,20,29],"automat":[4,6,9,11,16,19,20,21,25,26,29],"avail":[1,3,4,6,10,12,16,21,25],"averag":10,"aw":6,"awar":21,"aws_bearer_token_bedrock":6,"awsxrai":25,"azur":4,"azure_api_kei":6,"azure_openai":6,"azure_openai_api_kei":6,"b":5,"b25f":22,"b265":22,"back":[10,12,19,21,25],"backend":[7,11,12,14,16,20,23,25],"backward":5,"bad":[8,21],"baht":18,"balanc":[0,5,6,11,12,20,21,28],"base":[0,1,2,3,4,5,7,9,11,12,14,16,18,19,20,26],"base_url":[3,5,6,18,25,28],"basemodel":1,"bash":10,"basic":[3,4,29],"batch":[12,25],"batchspanprocessor":25,"battl":16,"bearer":[3,6],"becom":21,"been":[18,20],"befor":[11,12,14,18,26,29],"begin":18,"behalf":[1,11,12,22],"behavior":[6,11,12,14,21,22,26,28],"behind":[6,17],"being":[5,11,12,14],"belief":16,"below":[1,6,7,10,11,12,18,23,25,29],"benchmark":[20,21],"benefici":[19,20],"benefit":[1,5,16,24],"best":[4,7,9,13,16,24],"beta":[5,6],"better":[3,7,16,20,21,26],"between":[0,4,5,7,11,12,15,19,20,21,25,28],"bgn":18,"bill":25,"bin":18,"bind":[9,11,14],"black":21,"blaze":13,"block":[3,15],"block_categori":5,"blood":1,"blurri":1,"bodi":11,"bog":16,"boilerpl":21,"book":21,"bool":[7,11,12,28],"boolean":6,"bootstrap":12,"borrow":14,"both":[4,6,11,16,17,20,21,23,26],"bound":[11,15],"box":21,"brazilian":18,"break":22,"breaker":12,"bridg":[7,12,20],"brief":[7,12],"briefli":[3,12],"british":18,"brl":18,"broader":0,"bug":6,"buil":[1,2],"build":[0,2,4,7,9,11,14,15,16,26],"build_with_arch":17,"built":[0,1,6,10,11,12,14,16,17,25,27],"bulgarian":18,"busi":[0,11,12,14,16,20],"bypass":26,"byte":[22,25],"bytes_receiv":22,"bytes_s":22,"c":16,"cad":18,"call":[1,3,7,8,10,11,12,14,16,17,18,19,22,27],"call_openai":19,"campaign":[0,14,16],"can":[0,1,2,3,5,6,7,8,9,10,11,12,14,16,18,19,20,21,22,25,26,28,29],"canadian":18,"canari":[4,5],"capabl":[0,1,2,3,5,6,7,11,14,16,18,19,20,28],"capit":[3,11,18],"captur":[8,16,21,25],"care":20,"carefulli":[7,19],"carri":[12,25],"case":[0,5,16,17],"categor":8,"categori":[4,11,19,21],"cathedr":18,"caus":8,"caveat":19,"celsiu":[7,20],"center":18,"central":[0,1,4,9,11,12,16,20,22,28],"centric":21,"certain":26,"chain":[5,12],"challeng":21,"chanc":26,"chang":[1,2,3,4,11,14,21,25],"charact":25,"chat":[3,4,5,6,18,19,21,22,25,29],"chat_complet":25,"chat_with_fallback":3,"chatbot":25,"chatcompletionsrequest":19,"cheap":[4,28],"cheaper":[3,5,10,11],"cheapest":0,"check":[1,9,11,12,14,20,25,29],"chf":18,"children":25,"chines":18,"choic":[3,11,18,25],"choos":[4,20,21],"chosen":6,"chunk":[2,3],"ci":29,"circuit":12,"circular":5,"citi":[7,18,20],"claim":[0,16],"clarif":[1,2],"clarifi":[14,16,17,27],"class":[1,4,16,19],"classif":[11,21],"classifi":[21,26],"claud":[3,5,6,11,21,29],"cleaner":5,"clear":[8,19,20,21,26],"clearer":[7,20],"clearli":[7,19,21],"cli":[18,29],"client":[4,5,6,8,9,12,14,15,21,25,27,29],"close":21,"closest":11,"cloud":[4,13],"cluster":[0,4,11,12,16,28],"cnversat":1,"cny":18,"co":25,"code":[3,4,5,6,7,8,11,12,14,15,16,18,21,22],"code_review":6,"codebas":25,"codec":12,"codellama":6,"coder":6,"coher":6,"colleagu":20,"collect":[0,3,11,20,23,25],"collector":25,"com":[5,6,20,22,25],"combin":[2,4,7,20],"come":4,"command":[6,10,29],"comment":[11,14],"commit":25,"common":[0,1,6,19,20,21,25],"commun":[8,12,14,18,25],"compact":[6,11,21],"compani":[6,16],"compar":11,"compat":[4,5,16,25,28],"complet":[3,4,5,6,11,12,18,19,20,21,22,25,28,29],"completion_api":19,"complex":[0,4,5,6,7,9,14,15,19,20,21,25,26],"complex_reason":[6,21],"complianc":26,"compon":[6,12,16,25,29],"compos":18,"composit":11,"comprehens":[6,18],"comput":[3,29],"concept":[1,7,13,14,21],"concern":[10,12],"concis":12,"condit":5,"confid":26,"config":[6,14,21,23,25,29],"configiur":23,"configur":[0,1,2,3,4,10,11,13,14,15,16,18,19,21,24,25,27,29],"confirm":[11,12,20,25,28,29],"congratul":18,"connect":[0,4,7,9,10,11,12,14,15,16,18,28,29],"connect_timeout":[0,11,12,28],"conpleix":1,"consid":[11,19],"consider":1,"considert":1,"consist":[4,5,14,16,20,21,25],"consol":[6,25],"constraint":26,"construct":6,"contain":[10,11,12,14,16,18,25,26,29],"container_nam":29,"content":[1,3,5,6,8,11,18,19,21,25,26,29],"content_filt":5,"context":[6,7,8,12,14,16,18,20,21,25],"contextu":[2,21],"continu":[16,29],"contribut":18,"contributor":[16,17,27],"control":[5,6,10,14,15,21,28],"conveni":1,"conver":2,"convers":[0,1,2,6,11,16,21],"convert":18,"coordin":[7,15],"core":[1,7,12,18],"corpu":26,"correct":[1,8,11,19,26,29],"correctli":25,"correspond":12,"cost":[1,4,5,6,10,16,21],"could":[0,3,11,12,14,20,28],"count":12,"cover":[6,7],"covners":1,"cpu":[12,13],"creat":[0,3,5,6,8,9,11,12,14,16,17,20,21,25,28,29],"create_gradio_app":1,"createus":8,"creativ":[5,6,20,21],"creative_task":21,"creative_writ":6,"credenti":25,"crime":11,"criteria":[21,26],"critic":[0,2,8,11,12,14,16,20,23,25,28],"cross":4,"crucial":[20,22,25],"cryptic":21,"cue":21,"cultur":18,"curiou":[11,12,18,26,28],"curl":[4,5,6,29],"currency_exchang":18,"currency_symbol":18,"current":[5,7,20],"custom":[3,4,6,11,16,17,19,21,25],"custom_api_kei":6,"customprovid":6,"cut":16,"cycl":16,"czech":18,"czk":18,"d":[3,5,29],"dai":[0,2],"dame":18,"danish":18,"dashbaord":23,"dashboard":[20,24,25],"dashscop":6,"dashscope_api_kei":6,"data":[0,2,5,7,11,12,14,16,18,20,23,25,26,29],"databas":[0,2,20],"datadoghq":25,"dataset":0,"datasourc":23,"date":[7,18,20],"davinci":11,"dc":22,"dd_site":25,"de":12,"debug":[0,2,12,21,22,25,29],"decemb":18,"decid":[21,29],"decis":[0,9,10,11,12,14,21,28],"decoupl":[7,21],"decrypt":12,"dedic":11,"deep":[6,9,17,21],"deeper":18,"deepseek":4,"deepseek_api_kei":6,"def":[0,1,2,3,19,20,25],"default":[0,1,2,9,10,11,12,18,21,22,28],"defens":26,"defin":[4,5,11,12,18,19,21,25,28,29],"definit":[14,19,20],"degrad":8,"deliv":[20,21],"deliveri":12,"delta":3,"demo":[16,18,19],"demonstr":[6,19,20,21],"depend":[0,12,16,18],"deploi":[6,10,12,16,25,29],"deploy":[4,5,6,16,27],"describ":[1,2,8,12,19],"descript":[0,1,2,5,6,7,11,12,18,19,20,21,28],"descriptor":21,"design":[0,7,8,10,11,12,14,16,17,19,20,21,23,25,27],"desir":[6,7,20],"destroi":12,"detail":[1,2,4,6,7,8,9,11,12,14,16,20,22,25,29],"detect":[1,2,10,11,12,14,26],"determin":[7,12,19,20,21,25],"dev":[3,4,5,6,10,18],"develop":[1,2,3,4,5,6,7,8,9,11,12,16,17,18,25,26],"devic":[0,1,2,7,10,11,12,28],"device_id":[0,2,11,12,28],"device_reboot":0,"device_summari":[0,2],"diabet":1,"diabeter":1,"diagnos":1,"dict":7,"dictionari":11,"differ":[0,3,4,5,6,11,12,20,21,25,28],"difficult":16,"dipatch":1,"direct":[3,4,7,9,19,21],"directli":[11,14,21,29],"directori":29,"disast":16,"discov":[6,17],"diseas":1,"disease_diagnos":1,"diseases_symptom":1,"dispatch":12,"displai":[8,20],"distinct":21,"distinguish":21,"distribut":[5,11,12,25],"dive":[4,14,17,18],"divers":26,"dkk":18,"dn":[9,29],"do":[7,11,16,18,22],"doc":[2,17],"docker":[0,5,6,18,23,25],"document":[5,7,11,14,17,18,21,25,29],"doe":[0,7,20,26],"dollar":18,"domain":[12,20,21],"don":[1,10,20],"dot":5,"down":[16,22],"downstream":[1,11,12,13,14,15,20,25],"dramat":10,"draw":11,"driven":[7,17,18,21],"dropbox":16,"due":8,"durat":[12,22],"dure":[8,12],"dynam":[4,6,11,20,21,26],"e":[3,5,7,12,14,20,21,22,25],"each":[1,6,7,11,12,16,19,20,21,22,25],"earli":[11,12,16],"earlier":[12,29],"eas":25,"easi":[4,16,18,20,25],"easier":[5,7,9,18,21,22],"easili":[16,25],"easilli":1,"econom":1,"ecosystem":25,"edg":[12,16,17,20,27],"edit":23,"effect":[1,4,6,7,11,16,17,19,21],"effici":[0,1,2,6,7,10,11,18,19,21,25],"egress":[4,13,14,16,17,29],"egress_traff":[6,18,21,28],"eiffel":18,"either":[7,12,20],"elasticsearch":22,"element":7,"elk":22,"email":20,"embarrassingli":15,"embed":[11,22,26,28],"emiss":1,"empow":[7,21,26],"enabl":[0,2,4,5,6,7,8,10,12,14,16,19,20,21,25,28,29],"encod":21,"encount":[12,14,29],"encrypt":12,"end":[1,3,7,16,25],"endpoint":[0,1,2,4,5,7,8,11,12,14,18,19,20,21,23,25,28,29],"energi":1,"energy_sourc":1,"energy_source_info":1,"energysourcerequest":1,"energysourcerespons":1,"enforc":12,"engag":[0,12,16,19],"engin":[1,2,11,16,21],"enhanc":[2,6,7,11,12,18,19,25,26],"enough":8,"enrich":[1,7,11,14],"ensur":[0,4,7,8,9,11,12,16,18,19,20,21,25,26,29],"enterpris":4,"entir":[16,25],"entri":9,"enum":[1,7,11,12,20,28],"env":18,"environ":[3,4,5,6,18,20,21,25,29],"envoi":[4,9,11,12,14,15,16,17,27],"equal":[7,15],"equat":21,"equival":18,"error":[0,1,2,4,7,12,13,14,16,21,23,25,26,27,29],"error_target":28,"error_target_1":28,"escal":[7,19],"escalate_to_human":[7,19],"essenti":[7,9,17,21,26],"establish":[0,11,12,28],"etc":[11,12,14,16,18,28,29],"ethnic":26,"eu":25,"eur":18,"euro":18,"evalu":[11,21],"evaluation_interv":23,"even":[0,20,21],"evenli":12,"event":[0,12],"everi":22,"evolv":21,"exact":21,"exactli":12,"exampl":[0,4,5,6,9,12,13,14,16,18,22,28,29],"exceed":12,"excel":[21,26],"except":[3,8],"exception":16,"excess":[1,12],"exchang":18,"exclus":16,"execut":[8,11,12,20,26],"exist":[3,4,6,12,16,21,25],"expect":[8,20,29],"expens":1,"experi":[2,7,11,12,16,18,20,21,26],"experiment":[5,21],"expert":6,"explain":[3,20,21,29],"explan":[7,20],"explanatori":22,"explicit":[21,29],"explicitli":[21,26],"explor":[17,18,21],"export":[22,23,25],"expos":[0,3,9,11,16,22,26],"extend":[6,16,17],"extern":[0,2,25],"extract":[0,1,7,11,12,14,16,20,25,28],"extrem":11,"f":[0,1,2,3,19,20,22,25,29],"f376e8d8c586":22,"face":[11,21],"facilit":25,"fact":[0,9,11,12,28],"fahrenheit":[7,20],"fail":[3,4,5,8,12,14,21],"failov":[0,1,9,11,12,28],"failur":8,"fair":12,"fall":21,"fallback":[3,4,5,6,8],"fallback_model":3,"fals":[2,7,11,12,28],"familiar":12,"familizari":1,"faser":0,"fashion":[1,12,18],"fast":[1,2,3,4,5,6,7,11,13,14,16,18,21,28],"fastapi":[1,19],"faster":[2,3,5,10,11,14,17,21,25,27],"fastest":[0,6],"fatigu":1,"fault":[4,11],"fc":12,"fc1b":0,"featur":[3,6,11,14,16,17,19,24,25,26,28],"feed":22,"feedback":[8,11,26],"feel":1,"fetch":[18,20],"few":[14,22,25],"field":[1,6,22],"file":[1,6,7,9,11,12,23,29],"filter":[9,12,15,16,26],"final":[3,12],"final_messag":3,"final_text":3,"financ":21,"find":20,"first":[1,4,10,11,12,16,18,20,23,26],"firtst":1,"fit":12,"fix":21,"flag":[11,12,25,26,28],"flash":6,"flask":0,"fleet":21,"flexibl":[4,6,16,21,23,25],"float":7,"flow":[8,13,16,17,20,22,25],"fluentd":22,"focu":[0,1,12,16],"focus":[6,11,20],"follow":[1,5,6,7,8,10,11,12,14,16,18,19,22,25,28,29],"forint":18,"form":20,"format":[6,7,11,16,20,21,23,24,26],"forward":[1,8,10,11,12,14,15,21,22,25],"fossil":1,"found":[1,3,22],"foundat":6,"frame":12,"framework":[11,16,17,21,23,25,27],"franc":[3,11,18],"francisco":[7,20],"frankfurt":18,"frankfurther_api":18,"free":1,"frequent":1,"friendli":[17,20,21,27,28],"from":[0,1,2,3,4,5,6,7,8,9,11,12,14,15,16,18,19,20,21,25,28,29],"front":16,"frontier":[0,11],"fuel":1,"full":[3,4,7,11,14,18,19,21,28],"fulli":7,"function":[5,7,8,10,11,12,14,15,16,17,21,27,29],"function_cal":17,"functionvalidationerror":8,"fundament":21,"further":[11,18,19,25],"futur":[4,5,6,7],"g":[5,7,12,14,20,21,22,25],"ga":1,"gap":20,"gastronomi":18,"gatewai":[4,6,12,14,16,17,19,20,21,22,23,25,27,28],"gather":[0,12,16],"gbp":18,"gemini":4,"gemma":6,"genai":[7,18],"gener":[0,6,7,10,12,14,15,18,19,20,21,23,25,26,28],"get":[0,1,2,3,6,7,10,12,16,18,19,20,25,28],"get_device_statist":2,"get_device_summari":[0,2],"get_final_messag":3,"get_info_for_energy_sourc":1,"get_json":[0,2],"get_supported_curr":18,"get_system_prompt":19,"get_trac":25,"get_weath":[7,20],"get_workforc":1,"getenv":3,"github":[7,11,18,19],"give":16,"given":[0,2,12],"glm":6,"global":[18,23],"glucos":1,"go":[0,6,16],"goal":16,"good":[3,21],"googl":[4,16],"google_api_kei":6,"got":8,"govern":[1,4,16],"gpt":[0,1,3,5,6,9,10,11,12,18,20,21,25,28,29],"gpu":13,"gr":1,"grace":[3,8],"gracefulli":[8,14],"grade":4,"gradio":1,"gradual":5,"grafana":24,"greenhous":1,"grok":6,"groq":4,"groq_api_kei":6,"group":0,"grpc":25,"guard":[12,13,16,17,27],"guardrail":[5,8,9,10,11,12,14,16,17,18,26,27,28],"guid":[6,11,18,20,25,26,29],"guidelin":[4,26],"h":[3,5,29],"ha":[11,12,16,18,19,20,21,25],"hack":16,"haiku":[5,6],"hallucin":22,"hand":[7,27],"handel":[11,12,28],"handl":[0,1,2,4,7,8,9,11,12,14,15,16,17,19,20,21,25,26,27,29],"handle_request":25,"handler":[7,11,12],"handoff":16,"hardcod":3,"hardwar":15,"harm":[11,12,26],"hate":11,"have":[1,10,16,18,20,25,26],"haven":1,"hazard":11,"hcm":12,"header":[12,13,14,16,18,24,29],"health":[12,29],"healthcar":21,"healthi":[12,18],"hello":[3,5,21,25],"help":[0,1,2,4,7,8,10,11,12,14,16,17,18,20,21,23,25,26,27,28],"here":[0,2,6,7,9,11,12,14,18,20,21,22,25,26],"hexadecim":25,"hf":6,"high":[0,1,4,5,6,7,11,13,16,17,20,21],"higher":29,"highli":[1,2,20],"hint":18,"histori":[1,20],"hit":5,"hkd":18,"hong":18,"honor_timestamp":23,"horrid":16,"host":[0,5,6,9,10,11,12,14,22,23],"hostnam":[0,6,11,12,28],"how":[0,1,2,3,5,6,7,10,11,12,16,17,18,19,20,23,24,29],"howev":12,"html":17,"http":[3,4,5,6,8,10,11,12,16,18,20,22,23,25,28,29],"http_method":[1,7,28],"httpexcept":1,"huf":18,"hug":11,"huggingfac":[0,1],"human":[7,8,19,21],"hundr":6,"hungarian":18,"hybrid":21,"hygien":16,"hyphen":5,"i":[0,1,3,6,7,9,10,11,12,14,15,16,17,18,19,21,22,23,25,27,28,29],"iam":25,"iceland":18,"icon":18,"id":[0,2,6,8,12,22,25],"idea":17,"ideal":[3,21],"identif":20,"identifi":[0,2,5,7,8,11,12,20,21,25,28],"idr":18,"il":18,"illustr":1,"imag":29,"immedi":26,"impact":25,"implement":[3,4,5,6,11,12,19,25,26],"import":[0,1,2,3,5,11,18,20,22,25],"improp":8,"improv":[2,6,10,16,17,19,26],"in_path":[7,18],"incent":1,"includ":[1,2,4,6,7,9,12,16,18,20,21,22,23,25,26,29],"inclus":21,"incom":[9,10,11,12,14,20,21,25,28],"incomplet":20,"incredibli":16,"independ":21,"indian":18,"indic":[7,12,21,25],"indonesian":18,"industri":11,"infer":[6,16,21],"info":[18,19],"inform":[0,1,2,6,7,8,11,12,14,16,20,22,25,26,28],"information_extract":[11,12,26,28],"infrastructur":[16,17,27],"ingress":[11,13,14,16,17,29],"ingress_traff":[9,11,12,18,28],"init":25,"initi":[1,9,12,14,24],"inject":[25,26],"inner":12,"innov":[16,18],"input":[0,7,8,16,17,19,20,26,27],"input_guard":[11,12,18,28],"inquiri":19,"inr":18,"insecur":25,"instal":[3,6,18,25],"instanc":[5,12,14,28],"instead":[3,5,20],"instruct":[6,10,17,19,28],"instrument":23,"insur":[0,16],"insurance_claim_detail":22,"int":[0,2,7],"integ":[0,2,8],"integr":[2,3,4,5,6,7,16,17,18,20,21,22,23,24,29],"intellig":[3,4,5,6,7,10,12,14,16,19,21],"intend":[5,26],"intent":[0,10,12,16,19,20,21,28],"interact":[0,11,17,19,20,22,25,26],"intercept":11,"interest":18,"interfac":[3,4,6],"interfer":18,"intern":[0,2,5,6,12,14,23,25],"interoper":25,"interpret":[7,8,20],"intl":6,"intro":[17,27],"intro_to_arch":17,"introduc":[14,17,21],"invalid":[8,20],"inventori":25,"investig":25,"invoc":[7,20],"invok":[0,7,8,20],"involv":[0,7,16,21],"ip":[0,6,11,12,28],"ip1":[0,11,12,28],"ip2":[0,11,12,28],"isdefault":23,"isinst":[0,2],"isk":18,"isn":21,"isol":18,"isra":18,"issu":[7,8,14,19,25],"issues_and_repair":[7,19],"item":[7,11],"iter":[6,20],"itinerari":20,"its":[0,1,7,10,11,12,14,15,16,18,21,25],"itself":[9,11],"jaeger":25,"jailbreak":[5,11,12,16,18,26,28],"japanes":18,"java":16,"job_nam":23,"join":3,"joke":29,"jpy":18,"jq":[18,29],"json":[1,3,5,8,18,20,29],"jsonifi":[0,2],"just":[0,1,3,6,9,11,12,17,21,27,28],"k2":6,"katanemo":29,"keep":[14,20],"kei":[0,1,3,6,9,10,11,12,13,19,21,24,25,28],"kept":28,"kibana":22,"kimi":6,"kind":8,"king":20,"know":[10,11,14,26],"knowledg":[1,2,9,12,17],"known":[11,18],"kong":18,"korean":18,"koruna":18,"krona":18,"krone":18,"krw":18,"kr\u00f3na":18,"kubernet":25,"l7":16,"label":21,"landmark":18,"landscap":21,"langtrace_api_kei":25,"langtrace_python_sdk":25,"languag":[2,3,4,12,16,20,21,25,27],"larg":[2,6,11,12,15,20,21,25],"larger":6,"last":[0,1,2],"latenc":[0,1,2,5,10,11,16,20,21,23],"later":8,"latest":[5,6,18,28],"layer":[11,17,25,26,27],"lead":[7,11,16,19,26],"learn":[3,4,5,16,17,18,20,21],"least_connect":5,"left":16,"legal":21,"length":6,"less":[3,19],"let":[21,22,29],"leu":18,"lev":18,"level":[4,5,7,9,10,13,16,17,20,21,27,29],"leverag":[2,25],"libev":12,"librari":[4,6,16,21,27],"lifecycl":[13,27],"lifetim":[12,15],"lightweight":[0,6,16],"like":[1,2,3,4,5,6,7,8,9,11,12,14,16,17,18,20,21,22,23,25,26,27,28],"limit":[0,1,5,9,11,12,21,28],"line":16,"linearli":16,"lira":18,"list":[0,2,6,7,11,12,18,28],"listen":[0,1,6,11,12,13,14,18,21,25,27,28],"live":12,"ll":[0,1],"llama":6,"llama2":6,"llama3":[3,5,6],"llm":[0,1,2,3,5,6,9,10,12,13,14,16,17,20,23,25,26,27,28,29],"llm_provid":[0,1,5,6,9,11,12,17,18,21,28],"load":[0,5,11,12,21,28],"load_bal":5,"local":[1,3,4,5,6,9,11,12,13],"localhost":[6,18,23,25,29],"locat":[7,14,20],"lock":16,"log":[8,12,14,21,23,24,25,27,29],"logger":19,"logic":[0,1,3,6,8,9,11,12,14,16,17,20,21,26,28],"long":[1,6],"look":[11,12,18,26,28],"loop":12,"louvr":18,"low":[0,10,16,17,20,21,27],"lower":1,"m":18,"machin":[3,15],"made":18,"mai":[0,12,26],"main":[3,12,14,18,25],"maintain":[7,12,16,21],"mainten":7,"major":[15,18],"make":[0,1,8,9,10,11,12,14,16,18,20,21,22,25],"malaysian":18,"malici":[12,26],"manag":[0,1,4,5,7,8,9,10,11,12,16,17,18,19,20,22,25,26,28],"mandatori":7,"mani":[12,25],"manipul":[0,16,23],"manual":[1,2],"manufactur":[0,9,11,12,28],"map":[5,20,21,28],"mark":[6,7],"market":19,"massiv":16,"match":[0,3,12,21,28,29],"math":[3,21],"mathemat":[6,21],"matter":[16,20,21],"max":[0,11,12,28],"max_cost_per_request":5,"max_lat":5,"max_token":3,"maximum":21,"me":[3,18,29],"mean":[18,22,25],"meaning":[5,8,21],"measur":[16,23],"mechan":[0,4,17,20,21,26],"medium":6,"meet":[20,21,26],"messag":[1,3,5,6,8,12,13,18,19,20,21,25,26,28,29],"message_format":[0,1,6,9,11,12,18,21,28],"meta":6,"metadata":[7,11,14,19,21,22],"method":[0,1,2,4,22],"metric":[14,16,21,24,25],"metrics_path":23,"mexican":18,"mi":19,"mid":21,"might":[10,26],"mind":[20,25],"mini":[3,5,6,21,25],"minim":[9,21,25,26,29],"ministr":[6,18],"minor":26,"minut":[1,2],"misalign":26,"misformat":26,"miss":[0,8,16],"mistak":26,"mistral":[4,18,28],"mistral_api_kei":[6,18,28],"mistral_loc":28,"mistralministr":18,"misus":26,"mix":3,"mixtral":6,"mixtur":6,"mode":10,"model":[0,1,2,3,4,9,11,12,13,14,18,20,25,26,27,28],"model_alia":3,"model_alias":[3,5,21,28],"model_dump":19,"model_serv":22,"moder":13,"modern":[23,25],"modif":[12,21],"modifi":1,"modul":25,"modular":7,"monitor":[4,11,12,16,17,19,21,22,24,25,27,29],"moonshotai":6,"moonshotai_api_kei":6,"more":[0,1,2,5,7,9,11,14,15,16,19,20,21,22,25,26,29],"most":[0,6,7,11,15,18,19,20,21,22],"multi":[4,21,27],"multimod":6,"multipl":[3,4,5,12,15,20,21],"museum":18,"must":[0,2,5,25],"mxn":18,"my":[11,12,18,26,28],"mycompani":6,"myr":18,"n":[18,19],"n1":18,"n10":18,"n11":18,"n12":18,"n13":18,"n14":18,"n15":18,"n16":18,"n17":18,"n18":18,"n19":18,"n2":18,"n20":18,"n21":18,"n22":18,"n23":18,"n24":18,"n25":18,"n26":18,"n27":18,"n28":18,"n29":18,"n3":18,"n30":18,"n31":18,"n4":18,"n5":18,"n6":18,"n7":18,"n8":18,"n9":18,"name":[0,1,2,3,4,6,7,11,12,14,18,19,20,21,23,25,26,28,29],"nativ":[3,4,6,16,17,23,27],"natur":[8,11,12,19,20,28],"navig":[17,26],"necessari":[7,11,12,19,20],"need":[0,1,4,7,10,11,12,14,17,18,19,20,21,23,25,29],"network":[0,9,10,11,13,14,16,17,20,27,28],"network_qa":0,"network_summari":0,"never":18,"new":[1,4,5,7,16,18,20,21,25],"next":[6,20,27],"nif":18,"nli":[11,28],"nok":18,"non":[3,6,15],"none":[1,6,18,29],"nonexist":3,"norwegian":18,"note":[10,14,22,23,28],"notfounderror":3,"notr":18,"noun":21,"nova":6,"nuanc":16,"number":[2,5,12,15,20],"nzd":18,"o":[1,3,25],"o1":5,"o3":6,"object":[1,22],"observ":[4,16,17,19,23,25,27,28,29],"occur":8,"off":[7,27],"offer":[0,1,7,9,11,12,14,16,19,21,26,28],"offici":3,"often":[1,2,20,21],"ollama":[3,4,5],"on_except":[11,12,18,26,28],"onc":[0,1,2,5,11,12,15,18,20,21,25],"one":[0,6,12,14,17,18,20,21,25,27],"ones":[1,20],"ongo":20,"onli":[6,7,9,10,11,12,18,20,26,28],"opaqu":16,"open":[1,6,7,23,25],"openai":[0,1,4,5,9,12,21,22,25,28],"openai_api_kei":[0,1,5,6,9,11,12,18,21,25,28,29],"openai_cli":3,"openai_dev_kei":6,"openai_prod_kei":6,"opentelemetri":[16,23,28],"oper":[12,14,16,18,20,21],"operation":21,"optim":[4,6,9,19,20,21,25],"option":[0,1,2,6,7,10,17,21,25],"opu":6,"oral":1,"orchestr":7,"order":[12,20,25],"organ":7,"orient":20,"origin":[11,14,16,21,22],"oss":5,"otel":25,"other":[6,8,10,12,14,16,18,20,25],"otlp":25,"otlp_export":25,"otlpspanexport":25,"our":[2,10,11,12,16,19,20,21],"out":[14,16,20],"outbound":[16,18,25],"outgo":[9,20,25],"outlin":12,"outperform":0,"output":[2,16,19,20,23,26,29],"outsid":16,"over":[0,2,4,16,21],"overal":[7,10,19],"overlap":21,"overload":12,"overrid":[7,11,12,18,28],"overview":[1,24,27],"overwhelm":12,"own":12,"p":29,"p50":[0,11],"paa":6,"packag":[18,25],"page":11,"pai":20,"pain":16,"pair":11,"par":20,"parallel":[12,15,20],"param":[20,23],"paramet":[1,6,8,11,12,16,18,20,26,28],"parent":25,"pari":18,"pars":[8,20,22],"part":[7,12,16,20,25],"particip":21,"particular":[7,20],"particularli":19,"pass":[11,12,22,29],"past":[2,16],"path":[0,1,2,6,7,11,12,14,18,20,22,28,29],"pattern":4,"payment":[20,25],"pend":25,"per":[12,16,20,22,23],"perceiv":[16,23],"perform":[0,4,6,7,9,10,11,12,15,16,19,20,21,25,26,29],"period":12,"person":[0,7,16,18,19,20],"peski":[17,27],"peso":18,"phase":11,"philippin":18,"php":[16,18],"piec":[7,11],"pii":5,"pilot":25,"pip":[3,18,25],"pipelin":[12,25,29],"place":[12,20],"placehold":[0,2,19,25],"plan":[5,7,11,14],"platform":6,"pleas":[2,3,7,8,10,11,19,20,21],"pln":18,"plumb":16,"pod":25,"point":[3,5,9,11,20],"polici":[4,12,21,25],"polish":18,"pollut":1,"pool":[12,21],"popul":18,"popular":25,"port":[0,1,6,9,11,12,14,18,21,28,29],"portal":6,"possibl":8,"post":[0,1,2,3,5,19,22,28],"potenti":[11,19],"pound":18,"power":[1,4,11,20,25,26],"practic":[4,7,9,12,13,16,17,21,24],"pre":[12,26],"preced":25,"precis":[1,2,7,20],"predefin":[0,11,14,16],"predict":21,"prefer":[3,4,20,25],"prefix":6,"premier":6,"premis":[4,10],"prepar":25,"presenc":9,"present":[2,16,20,25],"preserv":6,"prevent":[12,16,25,26],"preview":[5,6],"previou":[1,2],"price":[10,18,19],"primari":[3,5,9,11,15],"primary_and_first_fallback_fail":5,"primary_model":3,"primit":[4,9,11,14],"principl":21,"print":[3,11,18,20,25],"priorit":26,"privaci":11,"pro":6,"problem":[3,5,6,16,21],"proce":[11,12,28],"process":[7,8,10,11,14,15,16,17,19,20,21,22,23,25,26,27],"process_customer_request":25,"processess":1,"processor":25,"prod":[4,5,6],"produc":20,"product":[4,5,6,20,21,25,29],"profan":5,"profil":[20,21],"program":[3,4,11,12,21,26,28],"prolifer":21,"prometheu":23,"promethu":23,"prompt":[1,8,9,10,12,13,14,15,16,17,18,19,21,23,27,28],"prompt_guard":[11,12,17,18,28],"prompt_target":[0,1,2,7,11,12,14,17,18,19,20,26,28],"prompt_target_intent_matching_threshold":28,"prone":[1,2],"proof":4,"propag":[11,16,24],"proper":25,"properli":14,"propos":19,"protect":[25,26],"proto":25,"protocol":[12,17,18,22,25,29],"proven":16,"provid":[0,1,2,3,5,7,8,9,11,12,14,16,17,18,20,21,22,23,25,26,27,28,29],"provider_interfac":6,"proxi":[6,11,12,14,16,17,23,27,29],"publish":23,"pull":[2,6],"purchas":[0,7,9,11,12,19,28],"purpos":[0,3,5,6,7,10,11,12,14,21,28],"pydant":1,"python":[4,5,16,21,22],"q":[0,20],"q1":26,"quadrat":21,"quadratic_equ":21,"qualiti":[2,21],"quantum":[3,29],"queri":[0,1,2,4,7,18,19,20,21,25],"question":[1,3,11,12,28],"quick":[16,21,25],"quickli":[10,16,17,18],"quickstart":[17,27],"quot":19,"quota":[5,6],"quota_exceed":5,"qwen3":6,"rag":[11,12,14,17,27],"rag_energy_source_ag":1,"rais":20,"rand":18,"random_sampl":25,"rang":[0,2,14,23,25],"rapid":21,"rare":12,"rate":[11,12,16,18,23,25,28],"rather":[2,8,16,21],"raw":8,"rch":14,"re":[1,2,4,11,12,18,20,25,26,28],"reach":26,"read":[1,9,11,12,16,19,20,25],"readabl":[7,8,21],"readi":[21,29],"real":[0,2,18,20,21],"realli":16,"reason":[3,4,5,6,21],"reboot":[0,2,7,11,12,28],"reboot_devic":[0,7],"reboot_network_devic":[11,12,28],"rebuild":29,"receiv":[8,11,12,14,19,20,22,25,26],"recept":12,"recognit":[7,11],"recommend":[1,7,12,15,18],"record":[16,25],"recoveri":16,"reddit":16,"reduc":[1,19,26],"ref":3,"refer":[0,2,5,6,10,14,18,20,22,25,27],"referenc":29,"refernc":20,"refin":11,"reflect":5,"refund":[7,19],"regardless":[3,5],"region":[4,25],"regularli":19,"reject":11,"relat":[0,7,10,12,14,19,21,23,25],"releas":[5,6],"relev":[0,1,2,7,12,20],"reli":[2,11,25],"reliabl":[2,3,4,9,20,26],"remain":20,"rememb":20,"remind":20,"renew":1,"renminbi":18,"repair":[7,19],"repeat":[19,26],"replac":[6,25],"report":19,"repositori":18,"repres":[11,25],"req":[7,19],"request":[3,4,5,6,8,9,11,13,14,16,18,19,20,21,22,25,26,27,29],"requestsinstrumentor":25,"requir":[0,1,2,4,7,8,9,11,12,16,18,19,20,21,28,29],"resili":[11,16],"resolut":[21,29],"resolv":[21,25],"resourc":[6,18,24],"respect":[22,23],"respond":[3,8,16,19,23,25],"respons":[0,1,2,3,5,7,8,11,12,14,17,18,19,20,21,22,25,26,28],"response_cod":22,"response_flag":22,"rest":[3,4,12,15],"result":[0,6,20],"retrain":21,"retri":[0,1,4,9,11,12,16,28],"retriev":[0,11,12,14,20],"return":[0,1,2,3,7,12,14,19,20,21,25,26],"reveal":19,"revers":12,"review":[5,6,19],"rich":21,"ridicul":19,"right":[17,27],"ringgit":18,"risk":26,"ro":29,"roadmap":11,"robin":[0,11,12,28],"robust":[11,16,20,26],"role":[3,5,8,11,18,19,21,25,29],"rollout":5,"romanian":18,"ron":18,"root":8,"round":[0,11,12,28],"round_robin":5,"rout":[0,2,3,4,5,9,11,12,14,16,17,27,28],"router":[4,7,12],"routin":[19,20],"routing_prefer":[6,21],"rule":[4,20],"run":[0,2,10,11,12,14,16,18,25,29],"runtim":6,"runtimeerror":8,"rupe":18,"rupiah":18,"safe":[4,16],"safer":26,"safeti":[5,9,11,16,26],"sai":20,"sale":[7,19],"sales_ag":[7,19],"same":[3,5,6,16,20],"sampl":[18,23,25,28,29],"sampling_r":28,"san":[7,20],"sanit":[12,26],"satisfact":19,"save":1,"scalabl":[7,18],"scale":[4,12,16],"scenario":[0,1,2,6,7,10,11,12,16,17,19,20,21,28],"schedul":[11,14,20],"schema":18,"scheme":[5,6,23],"scienc":20,"score":[11,21],"scrape_config":23,"scrape_interv":23,"scrape_timeout":23,"screenshot":23,"script":[18,20],"scrutini":26,"sdk":[4,6,25],"seamless":[4,19,20],"seamlessli":[3,4,7,21,23,25],"seattl":20,"section":[1,2,5,6,7,9,11,12,17,19,25],"secur":[4,9,11,12,16,17,25,26],"see":[4,7,9,11,14,19,20,21,29],"sek":18,"select":[4,12,18,19,21],"self":[10,12,14,16,19,22],"semant":[3,4,5,6,11,21],"send":[11,12,14,20,25,28],"sender":11,"sensit":[19,25],"sensitive_data":5,"sent":[12,21,22],"sentenc":19,"separ":[10,12,14,15,21],"sequenti":0,"serv":[12,13,14,27],"server":[6,8,10,12,14,16,18],"servic":[0,2,7,11,12,20,22,25,29],"set":[6,7,9,10,11,12,14,16,17,18,20,25,28,29],"set_tracer_provid":25,"setup":[1,9,11,20,25],"sever":[12,16,23],"sgd":18,"shape":16,"share":[11,12,14],"sheqel":18,"shift":21,"ship":[17,27],"short":3,"should":[7,8,12,20,21],"show":[1,18,29],"shown":6,"side":8,"signal":21,"signatur":20,"signifc":2,"signific":21,"significantli":19,"similar":[11,21],"simpl":[3,4,5,6,11,14,18,19,20,25],"simpli":[9,11,18,25],"simplic":12,"simplif":[5,9],"simplifi":[4,7,9,17,25],"simul":[0,2],"simultan":[0,4],"sinc":[11,18],"singapor":18,"singl":[1,6,11,14,15,16,20,28],"sit":[16,17,20],"site":25,"skimp":20,"slow":[1,2,16],"slowest":10,"small":[6,15],"smaller":6,"smart":[3,5,16,17,20],"smarter":2,"sni":12,"snippet":21,"so":[0,1,2,10,11,14,16,20],"socket":12,"softwar":[11,16],"solar":1,"sole":2,"solut":[7,19],"solv":[3,5,6,16,21],"some":[1,6,12,15,18,20,23],"sonnet":[3,5,6,11,21,29],"soon":4,"sophist":7,"sota":20,"sourc":[1,2,6,18,20,23,25],"south":18,"space":21,"span":[12,25],"span_processor":25,"spanish":21,"special":[6,7,16,17,19,21,25],"specif":[0,1,2,3,4,5,6,7,8,10,11,12,16,18,19,20,21,25,26,28],"specifi":[2,6,7,12,14,20,21],"speed":[2,10,16,23],"spell":20,"spend":15,"split":5,"sporad":15,"sql":2,"stabl":[5,21],"stack":[8,16,17,20,22],"staff":19,"stage":[5,26],"stai":11,"standalon":25,"standard":[4,6,16,25],"start":[9,10,11,14,25],"start_as_current_span":25,"start_tim":22,"stat":[0,2,12,23],"state":[7,11,12,18,20,21],"static":[12,21],"static_config":23,"statist":[0,2,12],"statu":[20,22,25],"step":[17,25,27],"stock":25,"store":[2,23],"stori":[3,21],"storytel":[6,21],"str":[1,7,11,12,18,19,20,28],"straightforward":[20,21],"strategi":4,"stream":[3,6,12,19],"streamlin":[2,7,9],"strength":21,"string":8,"strip":11,"stripe":16,"structur":[2,4,21,22],"struggl":[1,2],"stuck":16,"studio":6,"stuff":20,"style":[6,21],"subject":21,"submit":[19,20,21],"subscript":6,"substanti":16,"subsystem":[0,4,9,10,11,12,16,28],"success":[16,18],"successfulli":18,"suffix":6,"sugar":1,"suggest":[20,26],"suit":[16,21],"suitabl":[7,19,20,21],"sum":5,"summar":[1,3,5,7,11,12,14,20,21,28],"summari":[1,11,12,20,21,24,28],"support":[0,1,4,7,11,12,14,16,18,19,21,23,25,26,27],"sure":[8,25],"sustain":1,"swedish":18,"swiss":18,"switch":[4,5,21],"symbol":18,"symptom":1,"system":[0,1,8,9,11,12,16,18,19,20,21,22,25,26,28],"system_prompt":[0,1,9,11,12,18,19,28],"t":[1,6,10,20,21],"tag":29,"tail":22,"tailor":[0,16,20,26],"take":[11,12,14,16],"taken":22,"talk":1,"target":[1,5,9,12,13,14,17,18,19,21,23,27,28],"task":[0,4,5,6,7,11,12,14,15,16,20,21,28],"tcp":12,"team":4,"tech":[1,17,25,27],"tech_overview":17,"technic":[19,21,25],"techniqu":[1,2,11,21,26],"technologi":[1,16,17],"telemetri":[16,23,25],"tell":[3,29],"temperatur":[7,20],"term":[1,3],"termin":16,"terminologi":[13,27],"test":[1,3,4,5,6,10,16,20],"text":[3,11,12,20],"text_stream":3,"tft":[16,23],"thai":18,"than":[0,2,10,16,21],"thb":18,"thei":[7,19,21,26],"them":[2,7,8,12,14,19,20,26],"themat":21,"theme":21,"thi":[0,1,2,3,4,5,6,7,8,9,10,11,12,14,15,16,17,18,19,20,21,22,25,28,29],"thing":14,"think":20,"thirst":1,"thoroughli":20,"those":[0,1,2,14,25],"thread":[12,13,27],"three":[4,7,10,12,16,21,23],"threshold":28,"thrill":16,"through":[1,3,4,6,11,12,16,20,21,22,25],"throughput":[0,5,20,21],"time":[0,2,11,12,16,19,20,21,22,23,28],"time_rang":[0,2],"timeout":[5,6,9,11,12,18,21,23,28],"timestamp":12,"tip":13,"tl":[12,16,29],"tlm":16,"tls_certif":[0,1],"todai":[4,7,12,18,26],"togeth":4,"together_ai":6,"together_api_kei":6,"token":[1,6,10,11,16,23],"toler":[1,4,11],"too":12,"took":22,"tool":[3,8,16,22,23,24,26],"top":[4,9,15],"topic":21,"topologi":13,"tot":[16,23],"total":[16,22,23],"tower":18,"trace":[8,11,12,14,16,23,24,27,28],"trace_export":25,"tracepar":[16,24],"tracer":25,"tracer_provid":25,"tracerprovid":25,"track":[18,21],"tradit":[16,21],"traffic":[4,5,9,11,12,16,17,19,20,22,27],"train":26,"transact":20,"transform":6,"translat":21,"transpar":[16,21],"transport":12,"travel":[11,14,20],"travers":12,"treat":[11,12,28],"trigger":[0,14,17,20],"trivial":12,"troubleshoot":[7,8,17],"true":[0,1,2,3,6,7,9,11,12,18,20,21,23,25,28],"try":[3,18,29],"tupl":7,"turbo":[0,1,6],"turkish":18,"turn":[20,27],"tutori":17,"two":[0,3,7,10,11,12,14,21],"type":[0,1,2,3,5,7,8,11,12,14,18,20,21,23,28,29],"typic":[12,20],"typo":26,"u":6,"ui":12,"unambigu":21,"unauthor":11,"under":22,"underli":[3,5,21],"underscor":5,"understand":[5,16,17,20,21,23,25,26],"undifferenti":[12,14],"unexpect":20,"unifi":[3,4,16,17,27],"uniformli":25,"uniqu":[7,25],"unit":[7,18,20],"unknown_ag":19,"unlik":21,"until":12,"unwant":12,"up":[1,6,10,14,17,18,19,25,29],"updat":[0,1,12,16,18,20,21,29],"upgrad":[4,5,12,16],"upon":12,"upstream":[1,11,12,13,14,16,17,20,22,23,28,29],"upstream_host":22,"urin":1,"url":[4,7,11,18,23],"us":[0,1,2,3,6,7,9,10,12,14,16,17,22,24,28,29],"usabl":26,"usag":[1,3,4,5,7,12,16,20,21,23],"usd":18,"use_agent_orchestr":7,"user":[0,1,2,3,5,7,8,11,12,14,16,17,18,19,20,21,22,23,25,26,27,29],"user_id":8,"usi":25,"usual":[11,18],"util":[7,11,12,18],"v":25,"v0":[0,1,6,9,11,12,18,28],"v1":[3,4,5,6,9,18,19,21,22,25,28,29],"v1beta":6,"v2":[5,6,18,23],"v24":18,"v3":18,"v4":6,"vagu":[17,27],"valid":[0,2,4,7,16,17,18,20,21,26],"validationerror":8,"valu":[0,3,7,11,12,20,25,26,28],"valueerror":20,"variabl":[6,18,25,29],"variant":21,"varieti":12,"variou":[15,16,20,25],"ve":[1,18],"vector":[2,11],"venv":18,"verbos":29,"veri":[12,25],"verif":25,"verifi":[25,29],"version":[0,1,3,4,5,6,9,11,12,18,25,28],"via":[0,1,7,8,10,12,14,16,17,18,19,24,25,26,28],"view":[23,25],"violat":[8,12,14,26],"violent":11,"virtual":18,"visibl":22,"vision":1,"visit":11,"vllm":6,"vllm_api_kei":6,"volum":29,"vpc":10,"w3c":[16,25],"wa":[12,20,22],"wai":[0,1,2,9,11,12,16,28],"wait":[0,11,12,19,28],"want":[0,7,14,16,18,20,21],"wast":16,"watch":20,"we":[1,4,7,12,14,15,18,26],"weather":[7,20],"weather_info":20,"web":[12,14],"weight":5,"well":16,"were":22,"west":6,"what":[1,3,8,11,16,18,21,22],"when":[0,1,2,6,7,8,9,11,12,14,16,19,20,21,25],"where":[0,7,10,11,12,14,16,17,18,20,21,22,26,28],"whether":[4,7,20,21],"which":[0,2,4,9,10,12,16,20,21,22,23,25],"while":[1,6,11,15,17,19,20,21,26],"wide":[12,14,17,23,25,27],"wind":1,"window":18,"within":[0,7,11,12,17,19,20,26,28],"without":[3,4,6,9,12,16,17,21,29],"won":18,"work":[3,5,6,10,14,16,17,20,24,25,27,28],"worker":[12,15],"workflow":[17,19],"workload":[15,21],"world":21,"would":[0,2,12,28],"write":[0,1,2,6,16,17,21],"writer":5,"written":[12,15,16],"x":[3,5,8,12,14,18,22],"xai":4,"xai_api_kei":6,"xxx":29,"yaml":[0,1,3,6,14,18,23,25,29],"yen":18,"yml":[11,28,29],"york":[7,20],"you":[0,1,2,3,4,5,6,7,8,9,10,11,12,14,16,17,18,19,20,21,23,25,26,27,28,29],"your":[0,1,2,3,4,5,6,7,9,10,11,12,14,16,17,18,19,20,21,22,23,25,26,29],"yourself":1,"yourweatherapp":20,"yuan":18,"yyi":29,"z":6,"zar":18,"zealand":18,"zeroshot":22,"zhipu_api_kei":6,"z\u0142oti":18},"titles":["Agentic Apps","Multi-Turn","RAG Apps","Client Libraries","LLM Providers","Model Aliases","Supported Providers &amp; Configuration","Prompt Target","Error Target","Listener","Model Serving","Prompts","Request Lifecycle","Tech Overview","Terminology","Threading Model","Intro to Arch","Overview","Quickstart","Agent Routing and Hand Off","Function Calling","LLM Routing","Access Logging","Monitoring","Observability","Tracing","Prompt Guard","Welcome to Arch!","Configuration Reference","Deployment"],"titleterms":{"1":[0,1,2,18,20],"2":[0,1,2,18,20],"3":[18,20],"A":25,"For":7,"It":22,"access":22,"addit":25,"adjust":1,"advanc":[4,5,6],"agent":[0,7,11,18,19,25],"ai":[6,18,25],"alia":21,"alias":5,"alibaba":6,"align":21,"also":[3,5,6],"amazon":6,"an":11,"anthrop":[3,6],"api":6,"app":[0,1,2,11],"ar":7,"arch":[1,11,16,17,18,20,21,26,27],"architectur":12,"augment":2,"aw":25,"azur":6,"base":[6,21,29],"basic":[5,7],"bedrock":6,"benefit":[4,25,26],"best":[3,5,8,19,20,21,25],"blaze":10,"breakdown":25,"build":[1,17,18,27],"call":[0,20],"capabl":4,"case":[4,19,20,21],"categori":6,"class":6,"client":[3,11,18],"cloud":10,"combin":21,"come":5,"command":18,"common":[4,29],"compat":[3,6],"compos":29,"concept":[8,17,27],"config":[1,18],"configur":[5,6,7,9,12,20,23,26,28],"convers":18,"core":4,"cpu":10,"creat":18,"cross":3,"curl":[3,18],"currenc":18,"dashboard":23,"datadog":25,"deepseek":6,"default":[6,7],"defin":[0,1,2,7,20],"demo":1,"deploy":29,"docker":29,"downstream":9,"egress":[9,11,12],"endpoint":[3,6],"error":[3,8],"exampl":[1,3,7,8,11,19,20,21,25,26],"extract":2,"fast":10,"featur":[4,5,7,20,22],"file":18,"first":6,"flask":[1,2],"flow":12,"follow":2,"format":[22,25],"function":[0,20],"gatewai":[3,11,18,29],"gemini":6,"gener":2,"get":[4,17,27],"googl":6,"gpu":10,"grafana":23,"groq":6,"guard":[11,26],"guid":[17,27],"guidelin":6,"hand":19,"handl":3,"header":[8,25],"high":12,"how":[22,25,26],"i":[2,20,26],"implement":[20,21],"ingress":[9,12],"initi":25,"instanc":6,"instrument":25,"integr":25,"intent":[1,7,11],"interact":18,"intro":16,"issu":29,"kei":[4,7,8,20,22],"langtrac":25,"level":12,"librari":3,"lifecycl":12,"listen":9,"llm":[4,11,18,21],"local":10,"log":22,"logic":7,"match":[7,11],"messag":11,"method":21,"metric":23,"mistral":6,"model":[5,6,10,15,21,29],"moder":10,"monitor":23,"moonshot":6,"multi":[1,2],"multipl":[0,6],"name":5,"network":12,"next":18,"observ":24,"off":19,"ollama":6,"openai":[3,6,11,18],"opentelemetri":25,"over":20,"overview":[13,17,25],"parallel":0,"paramet":[0,2,7],"post":12,"practic":[3,5,8,19,20,25],"practicesm":21,"prefer":[6,21],"prerequisit":18,"process":[0,1,2,12],"prompt":[0,2,7,11,20,26],"propag":25,"provid":[4,6],"python":[3,18,25],"question":2,"quickstart":18,"qwen":6,"rag":[1,2],"rai":25,"refer":28,"request":[0,1,2,12],"requir":6,"resourc":[25,27],"retriev":[1,2],"rout":[6,7,19,21,29],"router":[18,21],"rule":5,"runtim":29,"sdk":3,"see":[3,5,6],"select":6,"serv":10,"setup":29,"singl":0,"smoke":29,"solut":29,"soon":5,"stack":29,"start":[4,17,18,27,29],"step":[0,1,2,18,20],"structur":6,"summari":[7,25,26],"support":[3,6],"switch":1,"take":20,"target":[0,2,7,8,11,20],"tech":13,"terminologi":[12,14],"test":29,"thread":15,"tip":[8,19,20],"togeth":6,"tool":[7,25],"topologi":12,"trace":25,"tracepar":25,"troubleshoot":29,"turn":[1,2],"up":2,"upstream":9,"url":6,"us":[4,5,11,18,19,20,21,25,26],"valid":5,"via":[11,23],"visual":25,"welcom":27,"what":[2,7,20,26],"why":26,"work":[22,26],"workflow":[20,21],"x":25,"xai":6,"zhipu":6}})
\ No newline at end of file
+Search.setIndex({"alltitles":{"AWS X-Ray":[[16,"aws-x-ray"]],"Access Logging":[[13,null]],"Additional Resources":[[16,"additional-resources"]],"Advanced Configuration":[[6,"advanced-configuration"]],"Advanced Features":[[4,"advanced-features"]],"Advanced Features (Coming Soon)":[[5,"advanced-features-coming-soon"]],"Agent Orchestration":[[0,"agent-orchestration"]],"Agent Structure":[[17,"agent-structure"]],"Agents":[[0,null]],"Alias-based routing":[[12,"alias-based-routing"]],"Amazon Bedrock":[[6,"amazon-bedrock"]],"Anthropic":[[6,"anthropic"]],"Anthropic (Python) SDK":[[3,"anthropic-python-sdk"]],"Arch-Function":[[11,"arch-function"]],"Arch-Router":[[12,"id7"]],"Azure OpenAI":[[6,"azure-openai"]],"Base URL Configuration":[[6,"base-url-configuration"]],"Basic Configuration":[[5,"basic-configuration"],[7,"basic-configuration"]],"Benefits of Using Traceparent Headers":[[16,"benefits-of-using-traceparent-headers"]],"Best Practices":[[3,"best-practices"],[16,"best-practices"],[17,"best-practices"],[19,"best-practices"]],"Best Practices and Tips":[[11,"best-practices-and-tips"]],"Best practices":[[12,"best-practices"]],"Bright Staff":[[24,null]],"Build Agentic Apps with Plano":[[10,"build-agentic-apps-with-plano"]],"Build Multi-Turn RAG Apps":[[7,"build-multi-turn-rag-apps"]],"Build with Plano":[[9,"build-with-plano"]],"Building agents with Plano orchestration":[[10,"building-agents-with-plano-orchestration"]],"Calling External APIs":[[17,"calling-external-apis"]],"Client Libraries":[[3,null]],"Client usage":[[12,"client-usage"],[12,"id4"],[12,"id6"]],"Combining Routing Methods":[[12,"combining-routing-methods"]],"Common Issues and Solutions":[[22,"common-issues-and-solutions"]],"Common Use Cases":[[4,"common-use-cases"],[17,"common-use-cases"]],"Concepts":[[9,"concepts"],[20,null]],"Configuration":[[12,"configuration"],[12,"id3"],[12,"id5"],[17,"configuration"],[19,"configuration"],[25,"configuration"]],"Configuration Overview":[[19,"configuration-overview"]],"Configuration Reference":[[21,null]],"Configuration Structure":[[6,"configuration-structure"]],"Configuration example":[[1,"configuration-example"]],"Configure Listeners":[[2,"configure-listeners"]],"Configure Monitoring":[[14,"configure-monitoring"]],"Conversational State":[[19,null]],"Core Capabilities":[[4,"core-capabilities"]],"Cross-Client Compatibility":[[3,"cross-client-compatibility"]],"Datadog":[[16,"datadog"]],"DeepSeek":[[6,"deepseek"]],"Default Model Configuration":[[6,"default-model-configuration"]],"Defining Parameters":[[7,"defining-parameters"]],"Demo App":[[7,"demo-app"]],"Deployment":[[22,null]],"Deterministic API calls with prompt targets":[[10,"deterministic-api-calls-with-prompt-targets"]],"Docker Compose Setup":[[22,"docker-compose-setup"]],"Docker Deployment":[[22,"docker-deployment"]],"Error Handling":[[3,"error-handling"]],"Example 2: Switching Intent":[[7,"example-2-switching-intent"]],"Example Configuration For Tools":[[7,"example-configuration-for-tools"]],"Example Use Cases":[[11,"example-use-cases"],[12,"example-use-cases"]],"Example with OpenTelemetry in Python":[[16,"example-with-opentelemetry-in-python"]],"Example: Travel Booking Assistant":[[17,"example-travel-booking-assistant"]],"Filter Chain Programming Model (HTTP and MCP)":[[1,"filter-chain-programming-model-http-and-mcp"]],"Filter Chains":[[1,null]],"First-Class Providers":[[6,"first-class-providers"]],"Function Calling":[[11,null]],"Function Calling Workflow":[[11,"function-calling-workflow"]],"Gateway Endpoints":[[3,"gateway-endpoints"]],"Gateway Smoke Test":[[22,"gateway-smoke-test"]],"Get Started":[[9,"get-started"],[20,null]],"Getting Started":[[4,"getting-started"]],"Google Gemini":[[6,"google-gemini"]],"Groq":[[6,"groq"]],"Guardrails":[[18,null]],"Guides":[[9,"guides"],[20,null]],"Header Format":[[16,"header-format"]],"High level architecture":[[25,"high-level-architecture"]],"How Guardrails Work":[[18,"how-guardrails-work"]],"How It Works":[[13,"how-it-works"],[17,"how-it-works"],[19,"how-it-works"]],"How to Initiate A Trace":[[16,"how-to-initiate-a-trace"]],"Implementation":[[17,"implementation"]],"Implementing Function Calling":[[11,"implementing-function-calling"]],"Inbound (Agent & Prompt Target)":[[2,"inbound-agent-prompt-target"]],"Information Extraction with LLMs":[[17,"information-extraction-with-llms"]],"Inner Loop (Agent Logic)":[[0,"inner-loop-agent-logic"]],"Inner Loop vs. Outer Loop":[[0,"inner-loop-vs-outer-loop"]],"Instrumentation":[[16,"instrumentation"]],"Integrating with Tracing Tools":[[16,"integrating-with-tracing-tools"]],"Intro to Plano":[[8,null]],"Key Benefits":[[0,"key-benefits"],[4,"key-benefits"]],"Key Features":[[7,"key-features"],[11,"key-features"],[13,"key-features"]],"LLM Routing":[[12,null]],"Langtrace":[[16,"langtrace"]],"Listeners":[[2,null]],"Log Format":[[13,"log-format"]],"Memory Storage (Development)":[[19,"memory-storage-development"]],"Metrics Dashboard (via Grafana)":[[14,"metrics-dashboard-via-grafana"]],"Mistral AI":[[6,"mistral-ai"]],"Model (LLM) Providers":[[4,null]],"Model Aliases":[[5,null]],"Model Selection Guidelines":[[6,"model-selection-guidelines"]],"Model-Based Routing":[[22,"model-based-routing"]],"Model-based routing":[[12,"model-based-routing"]],"Monitoring":[[14,null]],"Moonshot AI":[[6,"moonshot-ai"]],"Multi-Turn":[[7,"multi-turn"]],"Multiple Provider Instances":[[6,"multiple-provider-instances"]],"Naming Best Practices":[[5,"naming-best-practices"]],"Network Topology":[[2,"network-topology"]],"Network topology":[[25,"network-topology"]],"Next Steps":[[10,"next-steps"],[17,"next-steps"],[19,"next-steps"]],"Observability":[[15,null]],"Ollama":[[6,"ollama"]],"OpenAI":[[6,"openai"]],"OpenAI (Python) SDK":[[3,"openai-python-sdk"]],"OpenAI Responses API (Conversational State)":[[3,"openai-responses-api-conversational-state"]],"OpenAI-Compatible Providers":[[6,"openai-compatible-providers"]],"Orchestration":[[17,null]],"Outbound (Model Proxy & Egress)":[[2,"outbound-model-proxy-egress"]],"Outer Loop (Orchestration)":[[0,"outer-loop-orchestration"]],"Overview":[[9,null],[16,"overview"]],"Post-request processing":[[25,"post-request-processing"]],"PostgreSQL Storage (Production)":[[19,"postgresql-storage-production"]],"Preference-aligned routing (Arch-Router)":[[12,"preference-aligned-routing-arch-router"]],"Preparing Context and Generating Responses":[[17,"preparing-context-and-generating-responses"]],"Prerequisites":[[10,"prerequisites"],[19,"prerequisites"]],"Prompt Target":[[7,null]],"Provider Categories":[[6,"provider-categories"]],"Providers Requiring Base URL":[[6,"providers-requiring-base-url"]],"Quickstart":[[10,null]],"Qwen (Alibaba)":[[6,"qwen-alibaba"]],"Request Flow (Egress)":[[25,"request-flow-egress"]],"Request Flow (Ingress)":[[25,"request-flow-ingress"]],"Request Lifecycle":[[25,null]],"Resources":[[20,null]],"Routing Methods":[[12,"routing-methods"]],"Routing Preferences":[[6,"routing-preferences"]],"Runtime Tests":[[22,"runtime-tests"]],"See Also":[[3,"see-also"],[5,"see-also"],[6,"see-also"]],"Starting the Stack":[[22,"starting-the-stack"]],"Step 1. Create plano config file":[[10,"step-1-create-plano-config-file"],[10,"id2"]],"Step 1. Minimal orchestration config":[[10,"step-1-minimal-orchestration-config"]],"Step 1: Define Plano Config":[[7,"step-1-define-plano-config"]],"Step 1: Define the Function":[[11,"step-1-define-the-function"]],"Step 2. Start plano":[[10,"step-2-start-plano"]],"Step 2. Start plano with currency conversion config":[[10,"step-2-start-plano-with-currency-conversion-config"]],"Step 2. Start your agents and Plano":[[10,"step-2-start-your-agents-and-plano"]],"Step 2: Configure Prompt Targets":[[11,"step-2-configure-prompt-targets"]],"Step 2: Process Request in Flask":[[7,"step-2-process-request-in-flask"]],"Step 3. Interacting with gateway using curl command":[[10,"step-3-interacting-with-gateway-using-curl-command"]],"Step 3. Send a prompt and let Plano route":[[10,"step-3-send-a-prompt-and-let-plano-route"]],"Step 3.1: Using OpenAI Python client":[[10,"step-3-1-using-openai-python-client"]],"Step 3.2: Using curl command":[[10,"step-3-2-using-curl-command"]],"Step 3: Interact with LLM":[[10,"step-3-interact-with-llm"]],"Step 3: Plano Takes Over":[[11,"step-3-plano-takes-over"]],"Summary":[[7,"summary"],[16,"summary"]],"Supabase Connection Strings":[[19,"supabase-connection-strings"]],"Supported API Endpoints":[[6,"supported-api-endpoints"]],"Supported Clients":[[3,"supported-clients"]],"Supported Providers & Configuration":[[6,null]],"Tech Overview":[[26,null]],"Testing the Guardrail":[[18,"testing-the-guardrail"]],"Threading Model":[[27,null]],"Together AI":[[6,"together-ai"]],"Trace Propagation":[[16,"trace-propagation"]],"Tracing":[[16,null]],"Troubleshooting":[[19,"troubleshooting"],[22,"troubleshooting"]],"Typical Use Cases":[[1,"typical-use-cases"]],"Unsupported Features":[[12,"unsupported-features"]],"Use Plano as a Model Proxy (Gateway)":[[10,"use-plano-as-a-model-proxy-gateway"]],"Using Aliases":[[5,"using-aliases"]],"Validation Rules":[[5,"validation-rules"]],"Welcome to Plano!":[[20,null]],"What is Function Calling?":[[11,"what-is-function-calling"]],"Why Guardrails":[[18,"why-guardrails"]],"Zhipu AI":[[6,"zhipu-ai"]],"cURL Examples":[[3,"curl-examples"]],"llms.txt":[[23,null]],"xAI":[[6,"xai"]]},"docnames":["concepts/agents","concepts/filter_chain","concepts/listeners","concepts/llm_providers/client_libraries","concepts/llm_providers/llm_providers","concepts/llm_providers/model_aliases","concepts/llm_providers/supported_providers","concepts/prompt_target","get_started/intro_to_plano","get_started/overview","get_started/quickstart","guides/function_calling","guides/llm_router","guides/observability/access_logging","guides/observability/monitoring","guides/observability/observability","guides/observability/tracing","guides/orchestration","guides/prompt_guard","guides/state","index","resources/configuration_reference","resources/deployment","resources/llms_txt","resources/tech_overview/model_serving","resources/tech_overview/request_lifecycle","resources/tech_overview/tech_overview","resources/tech_overview/threading_model"],"envversion":{"sphinx":65,"sphinx.domains.c":3,"sphinx.domains.changeset":1,"sphinx.domains.citation":1,"sphinx.domains.cpp":9,"sphinx.domains.index":1,"sphinx.domains.javascript":3,"sphinx.domains.math":2,"sphinx.domains.python":4,"sphinx.domains.rst":2,"sphinx.domains.std":2,"sphinx.ext.intersphinx":1,"sphinx.ext.viewcode":1},"filenames":["concepts/agents.rst","concepts/filter_chain.rst","concepts/listeners.rst","concepts/llm_providers/client_libraries.rst","concepts/llm_providers/llm_providers.rst","concepts/llm_providers/model_aliases.rst","concepts/llm_providers/supported_providers.rst","concepts/prompt_target.rst","get_started/intro_to_plano.rst","get_started/overview.rst","get_started/quickstart.rst","guides/function_calling.rst","guides/llm_router.rst","guides/observability/access_logging.rst","guides/observability/monitoring.rst","guides/observability/observability.rst","guides/observability/tracing.rst","guides/orchestration.rst","guides/prompt_guard.rst","guides/state.rst","index.rst","resources/configuration_reference.rst","resources/deployment.rst","resources/llms_txt.rst","resources/tech_overview/model_serving.rst","resources/tech_overview/request_lifecycle.rst","resources/tech_overview/tech_overview.rst","resources/tech_overview/threading_model.rst"],"indexentries":{},"objects":{},"objnames":{},"objtypes":{},"terms":{"":[0,1,2,3,4,5,6,7,8,9,10,11,12,13,16,17,18,19,22,24,25,27],"0":[1,2,3,5,6,7,10,12,13,16,17,19,21,22,25],"00":[16,17],"005":[21,25],"00z":17,"01":3,"03":13,"04":13,"05":10,"050z":13,"06":[3,10],"08":10,"0905":6,"1":[2,3,4,5,6,12,13,14,17,19,21,25],"10":[5,10,11,13,14,17],"100":[1,3,5,10,16,17,21,27],"1000":17,"10000":[2,10,21,22,25],"1022":13,"104":13,"10500":[18,21],"10501":1,"10502":1,"10505":1,"10510":[17,21],"10520":[10,17,21],"10530":10,"106":13,"10t03":13,"10th":17,"11":10,"11434":[5,6],"12":[7,10,17,19],"12000":[2,3,5,10,12,13,16,19,21,22],"120b":5,"123":19,"125":13,"127":[2,3,5,7,10,13,19,21,25],"128e":6,"128k":6,"1301":13,"131":10,"13b":6,"140":13,"15":[4,14],"159":13,"16":[10,16,17],"162":13,"168":13,"1695":13,"1695m":13,"17":13,"17b":6,"18083":13,"192":13,"19901":14,"1d6b30cfc845":13,"2":[2,5,6,12,17,19,22,25],"20":5,"200":[1,3,13,17],"2019":7,"2023":3,"2024":[10,13],"20240307":5,"20241022":[3,5,6],"2025":17,"20b":6,"21797":13,"218":13,"23":17,"23123":19,"24":[10,17],"245":13,"25":[10,17],"254":13,"27":10,"28":10,"288":10,"29":10,"2a5b":13,"3":[1,3,5,6,7,17,21],"30":[10,12,17],"30b":17,"31":7,"32":16,"32b":6,"32k":6,"34b":6,"3b":[6,10,21],"4":[6,10,11,12,17,18,19,21,22],"400":18,"429":25,"4317":16,"441":13,"443":[6,7,10,13],"447":13,"45":17,"463":13,"469793af":13,"48":17,"485":10,"49":13,"492z":13,"4a":25,"4b":25,"4o":[1,2,3,5,6,10,16,17,21,25],"4xx":1,"5":[3,5,6,7,10,12,17,19,22],"50":[3,17],"51":[10,13,17],"51000":13,"52":13,"53":13,"537z":13,"54":13,"5432":19,"55":13,"556":13,"56":[10,13],"59":17,"598z":13,"59z":17,"5b":[4,12],"5xx":1,"6":[6,17],"604197fe":13,"614":13,"646":17,"647":10,"65":13,"67":17,"7":17,"71":17,"770":13,"78558":10,"7b":6,"8":16,"80":[5,6,21,25],"8000":6,"8001":[1,10,17,18,21],"8080":[6,7],"825":10,"86":17,"87":13,"8b":6,"905z":13,"906z":13,"9090":14,"9367":13,"95":[5,17],"95a2":13,"961z":13,"979":10,"984":13,"984m":13,"99":17,"9b57":13,"A":[0,1,5,7,8,9,11,15,17,19,20,25,27],"And":10,"As":[10,14],"At":[10,12],"Be":[11,16,17],"By":[0,7,8,11,13,16,17],"FOR":13,"For":[0,2,3,6,10,11,12,13,16,17,19,22,25,27],"IF":19,"If":[1,7,11,17,18,19,22,25],"In":[1,8,9,11,12,18],"It":[2,6,8,9,11,12,15,16,20,24,25],"Its":[8,14,16],"NOT":19,"No":[3,10,12,19],"Not":19,"ON":19,"On":10,"One":3,"Or":[3,16],"TO":[17,19],"That":19,"The":[0,1,2,3,5,6,7,10,11,12,13,16,17,18,19,21,25],"Then":19,"There":[2,8],"These":[8,11,25],"To":[2,8,10,12,14,16,17,18,22],"With":[6,8,10,12,16],"_":8,"__name__":16,"a1c":7,"a3b":17,"abil":6,"about":[0,1,3,4,5,7,8,9,12,13,17,18,19,24,25],"abov":[5,7,8,10,14,22,25],"abstract":[1,4,9,12,20],"acceler":20,"accept":[2,25,27],"access":[0,4,10,14,15,16,20,21,25],"access_":13,"access_ingress":13,"access_intern":13,"access_kei":[1,2,5,6,7,10,12,17,21,25],"access_llm":13,"accident":25,"accordingli":16,"account":11,"accur":[7,8,11,12,17],"accuraci":12,"achiev":[8,11],"acknowledg":17,"across":[0,1,2,4,5,6,7,8,9,10,11,16,17,19,20,24,25],"act":[2,8,14],"action":[0,7,11,12,17,18,25],"activ":[6,10,25],"actual":[3,12,16,17,19],"ad":[12,17,18],"adapt":[0,12],"add":[0,1,2,4,7,8,11,12,16,18,19],"add_span_processor":16,"addit":[6,7,15,19,25],"address":[1,2,6,7,10,12,21,25],"adjust":[7,16,19,22],"adopt":[8,9,16,20],"advanc":11,"advantag":8,"aeroapi":17,"aeroapi_base_url":17,"aeroapi_kei":17,"affect":16,"afraid":11,"african":10,"after":[1,8,11,17,25],"against":[18,19],"agent":[1,6,7,8,9,11,13,14,16,18,19,20,21,24,25],"agent_1":[1,18],"agent_respons":18,"aggreg":11,"agil":[8,9,20],"agnost":[0,8,25],"ahead":17,"ai":[0,4,8,9,10,11,13,16,17,18,20],"aid":25,"air":[6,7],"airbnb":8,"aircraft":17,"aircraft_typ":17,"airlin":17,"airport":17,"aka":20,"alert":[13,14],"alertmanag":14,"algorithm":[5,8],"alia":[0,3,4,5,8,25],"alias":[3,4,6,12,20,21],"alic":[3,19],"align":[0,4,25],"aliyunc":6,"all":[0,4,6,7,10,16,17,19,22,23,25,27],"allow":[0,1,5,6,7,8,10,11,12,16,17,18,27],"along":10,"alongsid":[8,24,25],"alphanumer":5,"alreadi":16,"also":[1,2,4,10,12,17,25],"alwai":[5,12,17],"amazon_bedrock":6,"amazonaw":6,"ambigu":[7,11],"amount":27,"an":[0,1,2,7,8,10,11,12,14,16,17,18,25],"analysi":[0,6,12,17],"analyst":5,"analyt":11,"analyz":[0,6,8,10,11,12,13,16,17],"ani":[0,1,3,4,6,7,8,9,10,11,16,17,18,19,20,22,25],"annot":7,"answer":[7,17,25],"anthrop":[0,4,5,12,19,21,22],"anthropic_api_kei":[5,6,12,21,22],"anyth":12,"api":[0,2,4,5,7,8,9,11,13,14,16,19,20,21,24,25],"api_kei":[3,6,10,16,17,19],"api_serv":[7,11,13],"api_vers":14,"apierror":3,"apikei":17,"apivers":14,"apm":16,"apolog":18,"app":[2,6,8,9,17,20,22],"app_serv":[21,25],"appear":18,"append":[6,17],"appl":18,"appli":[0,2,4,5,8,9,12,18,21,25],"applic":[0,1,2,3,4,5,6,7,8,9,10,11,12,13,14,16,17,18,19,22,24,25],"appoint":12,"approach":[4,5,12,16],"appropri":[0,2,12,17,25],"approv":18,"ar":[0,1,2,3,4,5,6,7,8,10,11,12,13,14,16,17,18,19,22,24,25],"arch":[4,5,8,21,25],"arch_agent_rout":1,"arch_config":[3,7],"architectur":[1,2,7,8,12,16,26,27],"archiv":19,"area":4,"aren":6,"arg":17,"around":[25,27],"arrai":[17,19],"arriv":[1,17,18,25],"arrival_tim":17,"art":[0,8,11,12],"artifici":3,"ask":[7,10,11,17,18],"ask_quest":3,"aspect":[8,14],"assembl":[1,3],"assign":12,"assist":[1,3,6,7,10,12,16,18,19,21,25],"associ":12,"assum":10,"async":[17,18],"asynccli":17,"asyncopenai":17,"atlanta":17,"attach":[1,2,17,18,25],"attempt":[8,18],"attende":7,"attent":11,"attribut":[7,11,16],"aud":10,"audio":12,"augment":[1,18,21],"australian":10,"authent":6,"author":[3,13],"auto":17,"auto_llm_dispatch_on_respons":25,"autom":[11,17,19,22],"automat":[0,3,4,6,8,11,16,17,19,22],"autonom":[0,17],"avail":[0,3,4,6,7,8,10,12,25],"avoid":8,"aw":[6,19],"await":[17,18],"awar":17,"aws_bearer_token_bedrock":6,"awsxrai":16,"azur":[0,4,19],"azure_api_kei":6,"azure_openai":6,"azure_openai_api_kei":6,"b":[0,5],"b25f":13,"b265":13,"back":[1,8,12,25],"backend":[1,2,3,7,10,11,14,16,19,24,25],"background":10,"backward":5,"bad":12,"baht":10,"balanc":[0,5,6,11,25],"bandwidth":19,"base":[0,2,3,4,5,7,10,11,17,18,25],"base_url":[3,5,6,10,16,17,19],"basemodel":7,"basic":[3,4,10,17,22],"batch":[16,25],"batchspanprocessor":16,"battl":8,"bearer":[3,6],"becaus":1,"becom":[1,8,12,19],"bedrock":19,"been":[10,11],"befor":[0,1,2,7,10,17,18,19,22,25],"begin":10,"behalf":[2,7,13,25],"behavior":[0,1,6,7,8,12,13,17,18,19,21],"behind":[1,2,3,6,9,19],"being":[5,25],"below":[1,2,6,7,10,14,16,18,22,25],"benchmark":[11,12],"benefici":11,"benefit":[3,5,7,8,15],"bespok":[9,20],"best":[0,4,8,15,25],"beta":[5,6],"better":[0,3,8,11,12,17],"between":[0,4,5,7,8,11,12,16,17,25,27],"bgn":10,"bigint":19,"bin":10,"bind":2,"bloat":19,"block":[1,2,3,17,27],"block_categori":5,"blood":7,"blurri":7,"boilerpl":12,"book":[10,12],"bool":[7,25],"boolean":6,"bootstrap":25,"both":[2,4,6,8,9,11,12,14,17,19],"bound":27,"brazilian":10,"break":[0,13],"breaker":25,"bridg":[11,25],"brief":[7,25],"briefli":[3,25],"bright":[20,25,26],"british":10,"brittl":[9,20],"brl":10,"bug":6,"buil":7,"build":[0,1,2,4,8,17,19,20,27],"builder":1,"built":[0,6,7,8,9,16,20,25],"bulgarian":10,"bullet":17,"burden":8,"busi":[0,7,8,10,11],"bypass":18,"byte":[13,16],"bytes_receiv":13,"bytes_s":13,"c":8,"cach":17,"cad":10,"call":[0,1,2,3,7,8,9,12,13,19,20,21,24,25],"caller":[1,2,18],"can":[0,1,2,3,5,6,7,8,9,10,11,12,13,16,17,18,19,22,25],"canadian":10,"canari":[0,4,5],"cannot":[12,19],"capabl":[0,3,5,6,7,8,10,11,17,21,25],"capit":[3,10],"captur":[0,1,8,12,16],"care":[10,11],"carefulli":[7,17],"carri":[16,25],"case":[5,8,9,18],"cat":12,"categori":[4,12],"celsiu":[7,11,17],"central":[0,2,4,7,8,9,11,13,20,25],"centric":12,"chain":[0,2,5,8,17,18,19,20,25],"challeng":12,"chang":[0,1,3,4,7,8,9,12,16,18,20],"charact":[16,19],"chat":[0,3,4,5,6,8,10,12,13,16,17,18,19,21,22],"chat_complet":16,"chat_with_fallback":3,"chatmessag":18,"cheap":4,"cheaper":[3,5,17],"check":[1,7,10,11,17,18,19,22,25],"chf":10,"children":16,"chines":10,"choic":[3,10,16,17],"choos":[0,4,11,12,25],"chosen":6,"christma":17,"chunk":[3,17],"ci":22,"circuit":[0,1,25],"circular":5,"citi":[7,11,17,21],"clarif":7,"clarifi":7,"class":[4,7,8],"classif":12,"classifi":12,"claud":[3,5,6,12,19,21,22],"clean":7,"cleaner":5,"cleanli":[2,7,8,9],"cleanup":19,"clear":[7,11,12,17,18],"clearer":11,"clearli":[12,17],"cli":[10,22],"client":[2,4,5,6,16,17,19,20,22,25,27],"clienterror":18,"close":[12,17],"cloud":[4,24],"cloudi":17,"cluster":[4,8,24,25],"cnversat":7,"cny":10,"co":[16,19],"code":[0,1,3,4,5,6,7,8,9,10,12,13,17,25,27],"code_iata":17,"code_review":6,"codebas":[8,9,16,20],"codec":25,"codellama":6,"coder":6,"coher":6,"collabor":0,"colleagu":11,"collect":[3,11,12,14,16],"collector":16,"column":19,"com":[5,6,11,13,16,17,19],"combin":[4,11,17,19],"come":4,"command":[6,22],"comment":19,"commit":16,"common":[6,7,11,12,16,19],"commun":[1,10,16,25],"compact":[6,12],"compani":[6,8,18],"compat":[0,2,4,5,8,10,16,17,19],"compil":23,"complementari":10,"complet":[0,3,4,5,6,10,11,12,13,16,17,18,19,21,22,25],"complex":[0,2,4,5,6,11,12,16,17,27],"complex_reason":[6,12],"complianc":[1,18],"compon":[6,8,16,22,25],"compos":10,"comprehens":[6,10],"comput":[3,22],"concept":[10,12,19],"concern":[0,8,25],"concis":17,"condit":[5,17],"confid":12,"config":[6,12,14,16,19,22],"configiur":14,"configur":[3,4,8,10,15,16,18,20,22,26,27],"confirm":[11,22,25],"congratul":10,"connect":[2,4,7,8,10,17,22,24,25,27],"connect_timeout":[21,25],"connection_str":19,"conpleix":7,"consider":7,"considert":7,"consist":[0,1,2,4,5,8,10,11,12,16,17,18],"consol":[6,16],"constraint":1,"construct":6,"contain":[8,10,16,19,22,23,25],"container_nam":22,"content":[1,3,5,6,7,8,10,12,16,17,18,22],"content_filt":5,"context":[0,1,2,3,6,8,10,11,12,16,19,23,25],"context_build":[1,17],"context_messag":17,"contextu":12,"continu":[1,3,8,9,17,18,19,20,22],"contract":25,"contribut":10,"contributor":[8,9,20],"control":[0,1,5,6,9,12,18,20,21,24,25,27],"conveni":7,"convent":25,"convers":[0,1,4,6,7,8,9,12,17,20,24],"conversation_context":17,"conversation_st":19,"convert":10,"coordin":[17,24,27],"copi":19,"core":[8,9,10,17,20,25],"corpor":18,"correct":[7,18,19,22],"correctli":16,"correl":1,"correspond":25,"cost":[0,4,5,6,7,8,12,17],"could":[3,11,17,25],"count":[17,25],"coupl":[1,8,9],"cover":6,"covners":7,"cpu":25,"crash":1,"creat":[3,5,6,9,11,12,16,17,19,22,25],"create_gradio_app":7,"created_at":19,"creativ":[5,6,11,12],"creative_task":12,"creative_writ":6,"credenti":[16,19],"crewai":0,"criteria":12,"critic":[1,8,11,14,16,17,19,21,25],"cross":[0,1,4],"crucial":[11,13,16],"ctx":17,"cue":12,"curl":[4,5,6,18,22],"currenc":[17,25],"currency_exchang":10,"currency_symbol":10,"current":[5,7,11,17,21],"current_temp":17,"current_timestamp":19,"custom":[0,1,3,4,6,9,12,16,17,18],"custom_api_kei":6,"customprovid":6,"cut":[0,1,8],"czech":10,"czk":10,"d":[3,5,17,18,22],"dai":[17,21],"daili":17,"danish":10,"dashbaord":14,"dashboard":[11,15,16,19],"dashscop":6,"dashscope_api_kei":6,"data":[1,5,7,8,9,10,11,12,14,16,17,20,22,24,25],"databas":[11,19],"database_url":19,"datadoghq":16,"dataplan":[1,2,8],"datasourc":14,"date":[7,10,11,17,25],"datetim":17,"day_match":17,"day_nam":17,"days_ahead":17,"db":19,"db_password":19,"db_setup":19,"dc":13,"dd":17,"dd_site":16,"de":[8,9,25],"deal":17,"debug":[0,1,12,13,16,17,19,22,25],"debugg":0,"decemb":[10,17],"decid":[0,1,8,10,12,17,22,24],"decis":[0,1,8,12,17,24,25],"decoupl":[0,8,9,12,20],"decrypt":25,"deep":[2,6,9,12,17],"deeper":10,"deepseek":[4,19],"deepseek_api_kei":6,"def":[3,7,11,16,17,18],"default":[1,2,7,10,12,13,17,19,21,25],"defin":[1,2,4,5,10,12,16,17,18,21,22,25],"definit":11,"delai":17,"delet":19,"deliv":[8,9,11,12,17,20],"deliveri":[8,9,20,25],"delta":3,"demo":[8,17,19],"demonstr":[6,9,11,12],"departur":17,"departure_tim":17,"depend":[8,10,12,19,25],"deploi":[6,8,16,22],"deploy":[0,2,4,5,6,8,9,19,20,24],"describ":[7,17,25],"descript":[1,5,6,7,10,11,12,17,18,21,25],"descriptor":12,"design":[7,8,9,11,12,14,16,17,19,20,24,25],"desir":[6,7,11],"dest_cod":17,"destin":[1,17],"destination_cod":17,"destroi":25,"detail":[2,3,4,6,7,8,11,12,13,16,17,22,25],"detect":[7,18,25],"determin":[0,7,11,17,25],"determininist":8,"determinist":[0,7,25],"dev":[3,4,5,6,10],"develop":[0,2,3,4,5,6,7,8,9,10,16,20,24,25],"devic":[7,25],"device_id":25,"diabet":7,"diabeter":7,"diagnos":[7,17],"diagram":2,"dict":[7,17],"differ":[3,4,5,6,11,12,16,19,24,25],"difficult":8,"dipatch":7,"direct":[0,3,4,6,7,12,17,19,21],"directli":[0,2,10,12,22],"directori":22,"disabl":19,"disast":8,"discov":[6,9,17],"discuss":3,"diseas":7,"disease_diagnos":7,"diseases_symptom":7,"dispatch":25,"displai":11,"distinct":12,"distinguish":[0,12],"distribut":[5,16,25],"dive":[4,9,10],"dkk":10,"dn":22,"do":[8,10,12,13,17],"doc":[9,19],"docker":[1,5,6,10,14,16,17,21],"docs_ag":17,"document":[5,7,9,10,12,16,17,22,23],"doe":[7,11,12,19,24,25],"dollar":10,"domain":[11,12,17,18],"don":[3,7,11,19],"done":17,"dot":5,"down":13,"downstream":[1,2,7,11,16,25,27],"dramat":19,"driven":[10,12,18],"dropbox":8,"due":19,"dump":17,"duplic":[1,8],"durabl":19,"durat":[13,25],"dure":25,"dx":[9,20],"dynam":[0,4,6,11,12,17,18],"e":[1,3,5,7,8,10,11,12,13,16,17,19,21,25],"each":[0,1,3,6,7,8,10,11,12,13,16,17,18,19,25],"earli":[1,8],"earlier":[22,25],"eas":16,"easi":[1,4,8,11,12,16,17,18],"easier":[0,5,10,12,13,18,19],"easili":16,"easilli":7,"econom":7,"ecosystem":16,"edg":[2,8,11,17,25],"edit":14,"editor":19,"effect":[4,6,7,8,9,12,17],"effici":[0,1,6,7,8,10,12,16,24,25],"egress":[4,8,9,22,26],"egress_traff":[10,12],"either":[11,18,25],"elasticsearch":13,"element":[7,17],"elif":17,"elk":13,"els":17,"email":11,"embarrassingli":27,"embed":13,"emiss":7,"emit":1,"empow":[7,12],"empti":17,"en":17,"enabl":[3,4,5,6,7,8,11,12,16,17,19,21,22,25],"encod":[12,19],"encount":[22,25],"encrypt":25,"end":[0,3,7,8,9,16,17],"endpoint":[1,2,4,5,7,10,11,14,16,17,19,21,22,25],"energi":7,"energy_sourc":7,"energy_source_info":7,"energysourcerequest":7,"energysourcerespons":7,"enforc":[0,1,18,25],"engag":25,"engin":[7,8,9,20],"enhanc":[0,6,9,10,16],"enough":[0,8],"enrich":[0,1,7,17,19,25],"ensur":[4,7,8,10,11,12,16,18,19,22,25],"entangl":3,"enterpris":4,"entir":[1,8,16],"entiti":[1,17,25],"entri":[2,7,21],"enum":[7,11,25],"enumer":17,"env":10,"envelop":1,"environ":[3,4,5,6,10,11,12,16,19,22],"envoi":[2,4,8,9,17,20,24,25,27],"envoyproxi":8,"ephemer":19,"equal":[7,27],"equat":12,"equival":10,"error":[1,4,7,8,12,14,16,17,18,19,22,25],"escal":17,"essenti":[9,12,18],"establish":25,"estim":17,"etc":[8,10,22,25],"eu":16,"eur":10,"euro":10,"evalu":[8,12,18],"evaluation_interv":14,"even":[1,11,12],"evenli":25,"event":[17,25],"ever":[9,20],"everi":[8,9,13,18,19,20],"evolv":[1,8,12,17],"exact":12,"exactli":25,"examin":17,"exampl":[2,4,5,6,9,10,13,18,19,21,22,25],"exce":17,"exceed":25,"excel":12,"except":[3,17,18],"exception":8,"excess":[7,25],"exchang":10,"exclus":8,"execut":[1,11,12,25],"exist":[0,2,3,4,6,12,16,19],"expect":[1,11,22,25],"expens":7,"experi":[10,11,12,17],"experiment":[5,12],"explain":[3,11,12,17,22],"explan":[7,11],"explanatori":13,"explicit":[12,22],"explicitli":12,"explor":[9,10,12,17,19],"export":[13,14,16,19],"expos":[2,3,8,10,13,17,18],"extend":[6,8,12],"extern":[0,16,18,21],"extra_head":17,"extract":[3,7,11,16,25],"extract_flight_rout":17,"extraction_model":17,"extraction_prompt":17,"f":[3,7,11,13,16,17,19,22],"f376e8d8c586":13,"face":1,"facilit":16,"facto":[8,9],"fahrenheit":[7,11,17],"fail":[1,3,4,5,18,25],"failov":[0,2,7,8,19,25],"failur":[1,19],"fair":25,"fall":12,"fallback":[3,4,5,6,17,24],"fallback_model":3,"fals":[7,18,25],"famili":[0,8],"familiar":10,"famreowkr":0,"faq":17,"far":17,"fashion":[7,25],"fast":[1,3,4,5,6,7,8,10,12,19,21],"fastapi":7,"faster":[3,5,9,12,16,17,20],"fastmcp":18,"fatal":1,"fatigu":7,"fault":4,"favorit":19,"featur":[3,6,8,9,15,16,21],"feed":13,"feedback":[8,9,18,20],"feel":7,"fetch":[8,11,17],"few":[13,16],"field":[1,6,7,13,17],"file":[2,6,7,14,19,22,23,25],"fill":17,"filter":[0,2,8,9,17,18,19,20,21,25,27],"filter_chain":[1,17,18,21],"final":[1,3,17,25],"final_messag":3,"final_text":3,"financ":18,"find":[10,11,19],"firewal":2,"first":[3,4,8,10,11,14,17,18,19,25],"fit":[2,25],"fix":12,"flag":[16,25],"flash":6,"fleet":12,"flexibl":[3,4,6,8,9,12,14,16],"flight":[10,17,21],"flight_ag":[10,17,21],"flight_dest":17,"flight_group":17,"flight_numb":17,"flight_origin":17,"flightag":17,"flightawar":17,"float":7,"flow":[1,2,7,8,9,11,13,16,26],"fluentd":13,"fly":17,"focu":[0,7,8,9,17,19,20,25],"focus":[2,6,7,8,9,10,11],"fog":17,"follow":[1,5,6,7,8,10,12,13,16,17,19,21,22,25],"forecast":17,"forecast_dai":17,"forecast_typ":17,"forint":10,"form":11,"format":[0,6,7,8,11,12,14,15,17,19,21],"forward":[2,7,13,16,24,25,27],"fossil":7,"found":[3,7,13,17],"foundat":[6,8,9,17],"frame":25,"framework":[0,1,8,9,10,14,16,20],"franc":[3,10,17],"francisco":[7,11],"frankfurt":10,"frankfurther_api":10,"free":[7,17,19],"frequent":7,"friendli":[11,12,21],"from":[0,2,3,4,5,6,7,8,9,10,11,12,16,17,18,19,20,22,25,27],"frontend":2,"fuel":7,"full":[2,3,4,7,10,12,17,19,21],"function":[0,5,7,8,10,12,17,20,21,22,25,27],"further":10,"futur":[4,5,6,17],"g":[1,5,7,8,10,11,12,13,16,17,19,21,25],"ga":7,"gap":11,"gate":17,"gate_origin":17,"gatewai":[2,4,6,8,9,16,17,21],"gather":25,"gbp":10,"gdpr":18,"gemini":4,"gemma":6,"genai":[7,10],"gener":[0,1,6,8,10,11,12,14,16,18,19,21,23,25,27],"geocod":17,"geocode_data":17,"geocode_respons":17,"geocode_url":17,"get":[1,2,3,6,7,8,10,11,16,17,18,21],"get_current_weath":21,"get_final_messag":3,"get_flight":17,"get_info_for_energy_sourc":7,"get_last_user_cont":17,"get_start":9,"get_supported_curr":10,"get_trac":16,"get_weath":[7,11],"get_weather_data":17,"get_workforc":7,"getenv":3,"github":[10,17],"give":[0,8,9],"given":25,"glm":6,"global":14,"glucos":7,"glue":8,"go":[6,8,17,19],"goal":8,"goe":17,"good":[3,12,17],"googl":[4,8],"google_api_kei":6,"govern":[4,7,8],"gpt":[1,2,3,5,6,7,10,11,12,16,17,18,21,22,25],"gpu":24,"gr":7,"grace":3,"gracefulli":[17,19],"grade":[4,8,9,19],"gradio":7,"gradual":5,"grafana":15,"grant":19,"greenhous":7,"grok":6,"groq":4,"groq_api_kei":6,"ground":[8,9],"group":17,"growth":19,"grpc":16,"guard":[8,18,25],"guardrail":[0,1,2,5,8,9,10,17,20,21,25],"guid":[0,4,6,10,11,16,17,19,22],"guidelin":4,"h":[3,5,18,22],"ha":[8,10,11,12,16,19,25],"hack":8,"haiku":[5,6],"hallucin":13,"hand":[1,17],"handel":25,"handl":[0,2,4,7,8,9,10,11,12,16,17,18,19,22,24,25,27],"handle_request":[16,17],"handler":[7,8],"handoff":[0,17],"happen":[2,17],"hard":[1,8,12],"hardcod":[0,3,19],"harden":8,"harder":1,"hardwar":27,"harm":18,"hashmap":19,"hasn":19,"have":[0,7,8,10,11,16,17,19],"hcm":25,"header":[8,10,15,17,22,25],"health":[22,25],"healthcar":[12,18],"healthi":[10,25],"hello":[3,5,12,16],"help":[0,4,7,8,9,10,11,12,14,16,17,18,20,25],"here":[4,6,7,10,11,12,13,16,17,18,23,25],"hexadecim":16,"hf":6,"hidden":[8,9,20],"hide":2,"high":[4,5,6,7,8,9,10,11,12,26],"higher":22,"highli":[7,11],"hipaa":18,"histori":[1,3,7,11,12,17,19],"hit":5,"hkd":10,"hong":10,"honor_timestamp":14,"hood":[2,8],"hook":8,"horrid":8,"host":[1,2,5,6,10,13,14,17,19,21,24,25],"hostnam":[6,25],"hotel":10,"hotel_ag":10,"hour":17,"how":[0,1,2,3,5,6,7,8,9,10,11,12,14,15,22,24,25],"howev":[1,25],"html":9,"http":[0,2,3,4,5,6,8,10,11,13,14,16,17,18,19,21,22,24,25],"http_client":17,"http_method":[7,21],"httpexcept":7,"httpx":17,"huf":10,"huggingfac":7,"human":[7,12,17],"human_escal":17,"hundr":6,"hungarian":10,"hybrid":12,"hygien":8,"hyphen":5,"i":[0,1,2,3,6,7,8,9,10,12,13,14,16,17,18,19,20,21,22,24,25,27],"iam":16,"iata":17,"icao":17,"iceland":10,"id":[1,3,6,10,13,16,17,18,19,21,25],"idea":9,"ideal":[0,3,12,17,19],"ident":17,"ident_iata":17,"identif":11,"identifi":[5,7,11,12,16,19,25],"idr":10,"idx":17,"idx_conversation_states_created_at":19,"idx_conversation_states_provid":19,"idx_conversation_states_updated_at":19,"il":10,"illustr":7,"imag":[12,22],"impact":16,"implement":[0,1,3,4,5,6,8,10,16,18,19],"import":[3,5,7,10,11,13,16,17,18,19],"improv":[0,1,6,7,8,9,17,20],"in_path":[7,10],"inact":19,"inappropri":18,"inbound":[8,25],"incent":7,"includ":[3,4,6,7,8,10,11,12,13,14,16,17,19,22,25],"inclus":12,"incom":[0,1,2,7,11,12,16,17,25],"incomplet":11,"incredibli":8,"indent":17,"independ":18,"index":19,"indian":10,"indic":[7,12,16,25],"individu":[1,17],"indonesian":10,"infer":[6,8,12,17],"info":[10,17],"inform":[6,7,11,13,16,25],"information_extract":25,"infrastructur":[8,9,20],"ingress":[8,9,22,26],"ingress_traff":[2,10,25],"init":16,"initi":[2,7,15,19,25],"inject":[1,16,17],"inner":[10,17,25],"innov":10,"input":[7,11,12,18,19,21,25],"input_guard":[18,21],"input_item":19,"inquiri":17,"inr":10,"insecur":16,"insert":19,"insid":[1,8,17,24,25],"inspect":[1,24],"instal":[3,6,10,16],"instanc":[0,5,19,21,25],"instead":[3,5,7,8,11,12,19,21],"instruct":[6,9,10,17],"instrument":[0,14],"insurance_claim_detail":13,"int":[7,17,21],"integr":[3,4,5,6,7,9,10,11,12,13,14,15,22],"intellig":[0,3,4,5,6,8,25],"intend":5,"intent":[0,8,10,11,12,17,25],"interact":[1,11,13,18],"interest":10,"interfac":[3,4,6,10],"interfer":10,"intermedi":0,"intern":[1,2,5,6,7,10,14,17,21,25],"interoper":16,"interpret":[0,11,12],"intl":6,"intro":[9,20],"intro_to_plano":9,"introduc":[9,12],"introduct":17,"invalid":11,"investig":1,"invoc":[7,11,12],"invok":[1,2,7,11,17,18],"invoke_weather_ag":17,"involv":7,"io":8,"ip":[6,25],"ip1":25,"ip2":25,"ipv4":19,"ipv6":19,"is_valid":18,"isdefault":14,"isinst":17,"isk":10,"isn":[12,19],"isol":10,"isra":10,"issu":[16,17],"item":7,"iter":[6,8,11],"itinerari":11,"its":[1,2,7,8,10,12,16,17,18,19,25,27],"itself":[1,24],"j":[0,10],"jaeger":16,"jailbreak":[5,8,18],"japanes":10,"java":8,"javascript":0,"jfk":10,"job_nam":14,"join":3,"joke":22,"jpy":10,"jq":[10,22],"json":[3,5,7,10,11,17,18,22],"jsonb":19,"jure":[8,9],"just":[0,3,6,7,17],"k2":6,"katanemo":22,"keep":[8,11,19],"kei":[2,3,6,12,15,16,17,19,21,25],"keyword":17,"kibana":13,"kimi":6,"kind":10,"king":11,"knowledg":[2,17],"kong":10,"korean":10,"koruna":10,"krona":10,"krone":10,"krw":10,"kr\u00f3na":10,"kubernet":16,"l7":8,"la":17,"landscap":12,"langchain":0,"langtrace_api_kei":16,"langtrace_python_sdk":16,"languag":[0,1,3,4,7,8,9,10,11,12,16,17,18,19,20,25],"larg":[6,11,12,16,23,25,27],"larger":[6,17],"last":[7,17],"last_user_msg":17,"latenc":[0,5,7,8,11,12,14],"later":19,"latest":[5,6,10,21],"latitud":17,"layer":[0,16,17,18],"lead":[8,17],"learn":[3,4,5,8,9,10,11,12,17,19,20],"least":2,"least_connect":5,"legal":12,"len":17,"length":6,"less":[3,17],"let":[2,3,8,12,13,17,18,22],"leu":10,"lev":10,"level":[2,4,5,8,9,10,11,12,22,26],"leverag":[0,16,17,19],"libev":25,"librari":[4,6,8,12,20],"lifecycl":[0,20,26],"lifetim":[25,27],"lightweight":[1,6,24,25],"like":[0,1,2,3,4,5,6,7,8,9,10,11,12,13,14,16,17,18,19,20,21,25],"limit":[2,5,7,12,17,19,25],"line":8,"linearli":8,"lira":10,"list":[1,6,7,10,17,18,19,25],"listen":[1,7,8,10,12,16,17,18,20,21,25],"listens":25,"live":[0,1,8,10,17,24,25],"ll":[7,9,10],"llama":6,"llama2":6,"llama3":[3,5,6],"llamaindex":0,"llm":[0,1,2,3,5,6,7,8,9,11,14,16,18,19,20,21,22,25],"llm_gateway_endpoint":17,"llm_provid":[5,6,7,9,12],"llm_router":9,"load":[0,5,17,25],"load_bal":5,"local":[3,4,5,6,7,19,24,25],"localhost":[6,10,14,16,18,22],"locat":[7,11,17,21],"location_model":17,"location_nam":17,"log":[0,1,8,14,15,16,17,19,20,22,25],"logger":17,"logic":[1,2,3,6,7,8,9,10,11,12,17,18,19,20,25],"london":17,"long":[0,6,7,8],"longitud":17,"look":[2,10,17],"lookup":17,"loop":[8,9,10,17,20,25],"lost":19,"low":[0,8,11,12],"lower":7,"m":[10,17],"machin":[3,27],"made":10,"mai":25,"main":[3,10,25],"maintain":[7,12,17,18,25],"mainten":8,"major":27,"make":[0,1,7,8,10,11,12,13,16,17,18,19,24,25],"malaysian":10,"malici":[18,25],"manag":[0,2,3,4,5,7,8,9,10,11,13,17,19,24,25],"mandatori":7,"mani":[1,16,25],"manipul":14,"manual":[3,7,19],"map":[5,11,12],"mark":[6,7],"market":17,"mask":1,"massiv":8,"match":[1,3,7,12,17,22,25],"math":[3,12],"mathemat":[6,12],"matter":[8,9,11,12],"maverick":6,"max":[17,25],"max_cost_per_request":5,"max_lat":5,"max_pag":17,"max_token":[3,17],"maximum":12,"mcp":[8,9,17,18,21],"me":[3,10,17,18,22],"mean":[1,10,13,16,17,19],"meaning":[5,12],"measur":[8,14],"mechan":[4,9,11,12],"media_typ":17,"medium":6,"meet":[7,11,12],"memori":[1,3,8,24,25],"mention":17,"merg":[3,19],"messag":[3,5,6,7,10,11,12,16,17,18,19,21,22,25],"message_format":[7,10,12],"meta":6,"metadata":[1,3,7,13,24],"meteo":17,"method":[4,7,13],"metric":[0,8,15,16],"metrics_path":14,"mexican":10,"miami":17,"microsecond":17,"mid":12,"middlewar":[8,9,20],"might":7,"min":17,"mind":[11,16],"mini":[1,3,5,6,16,17,21],"minim":[12,16,18,22],"ministr":[6,21],"minut":[7,17],"misalign":18,"misconfigur":1,"miss":17,"mistral":[4,10,21],"mistral_api_kei":[6,10,21],"mistral_loc":21,"mistralministr":10,"mix":3,"mixtral":6,"mm":17,"mobil":2,"modal":12,"mode":19,"model":[0,3,7,8,9,11,16,17,18,19,20,21,23,24,25,26],"model_1":21,"model_alia":3,"model_alias":[1,3,5,12,21],"model_dump_json":17,"model_provid":[1,2,10,17,21,25],"model_serv":13,"moder":[1,8,9,20],"modern":[14,16],"modif":[12,25],"modifi":7,"modul":16,"modular":2,"mondai":17,"monitor":[0,4,8,13,15,16,17,19,20,22,25],"moonshotai":6,"moonshotai_api_kei":6,"more":[0,2,3,5,7,8,11,12,13,16,17,18,19,20,22,25,27],"most":[0,1,2,6,9,11,12,13,17,19,27],"move":[8,19],"msg":[17,18],"multi":[0,3,4,8,9,10,12,17,19],"multimedia":12,"multimod":6,"multipl":[0,3,4,5,10,11,17,19,21,25,27],"must":[5,16,19],"mutat":1,"mxn":10,"my":[3,19],"mycompani":6,"mypass":19,"myproject":19,"myr":10,"myuser":19,"n":[10,17,18],"n1":10,"n10":10,"n11":10,"n12":10,"n13":10,"n14":10,"n15":10,"n16":10,"n17":10,"n18":10,"n19":10,"n2":10,"n20":10,"n21":10,"n22":10,"n23":10,"n24":10,"n25":10,"n26":10,"n27":10,"n28":10,"n29":10,"n3":10,"n30":10,"n31":10,"n4":10,"n5":10,"n6":10,"n7":10,"n8":10,"n9":10,"name":[1,3,4,6,7,8,10,11,12,14,16,17,18,19,21,22,25],"nativ":[3,4,6,9,14,20],"natur":[7,10,11,17,25],"navig":[9,19],"necessari":[7,11,25],"need":[0,1,2,3,4,7,10,11,12,14,16,17,19,22,25],"negoti":17,"nest":17,"network":[1,8,9,11,19,26],"never":[10,19],"new":[0,3,4,5,7,8,10,11,12,16,17,19],"next":[0,6,11,18,20],"nif":10,"node":0,"nok":10,"non":[3,6,17,27],"none":[6,7,10,17,22],"nonexist":3,"nonstop":17,"normal":[0,1,25],"norwegian":10,"not_found":17,"note":[13,14,17],"notfounderror":3,"notic":19,"noun":12,"nova":6,"now":[10,17],"npleas":18,"null":[17,19],"number":[5,11,17,19,21,25,27],"nyc":17,"nzd":10,"o":[3,7,16],"o1":5,"o3":6,"object":[7,13,17],"observ":[0,1,4,8,9,10,14,16,17,20,21,22],"obvious":1,"off":[1,17,18],"offer":[7,8,12,17],"offici":3,"often":[1,7,11],"old":19,"ollama":[3,4,5],"omit":1,"onc":[1,5,7,8,10,11,16,17,19,25,27],"one":[2,6,8,10,11,16,17,19,25],"ones":[7,11],"ongo":11,"onli":[1,2,6,11,17,18,19,25],"open":[0,2,6,7,8,9,14,16,17,23],"openai":[0,1,2,4,5,7,8,9,12,13,16,17,19,21,25],"openai_api_kei":[1,2,5,6,7,10,12,16,17,21,22,25],"openai_cli":3,"openai_client_via_plano":17,"openai_dev_kei":6,"openai_prod_kei":6,"opentelemetri":[8,14,21],"oper":[8,10,11,12,17,25],"operation":12,"optim":[4,6,11,12,16],"option":[1,3,6,7,9,12,16,17,19,25],"opu":6,"oral":7,"orchestr":[2,8,9,20,21,25],"order":[0,1,11,24,25],"organiz":18,"orient":11,"origin":[2,8,13,17],"origin_cod":17,"oss":[5,6],"otel":16,"other":[0,1,2,6,8,10,11,16,17,25],"otlp":16,"otlp_export":16,"otlpspanexport":16,"our":[8,11,19,25],"out":[8,9,11,17,18,20],"outbound":[8,10,16,25],"outcom":1,"outer":[10,17],"outgo":[11,16,25],"outlin":25,"output":[3,8,11,14,17,18,19,22],"output_text":19,"outsid":18,"over":[1,4,8,12,18,24],"overal":17,"overlap":12,"overload":25,"overrid":25,"overview":[15,20],"overwhelm":25,"own":[1,8,17,24,25],"p":22,"paa":6,"packag":[10,16],"page":23,"pai":11,"pain":8,"par":11,"parallel":[11,25,27],"param":[11,14,17],"paramet":[6,10,11,17,21,25],"parent":16,"pari":[10,17],"pars":[7,11,13,25],"part":[7,8,11,16,17,25],"parti":24,"particip":12,"particular":[0,7,11],"particularli":17,"partli":17,"pass":[1,7,13,17,18,22,25],"password":[17,19],"past":19,"path":[1,2,6,7,10,11,13,21,22,25],"pattern":[1,4,8,9,17,19],"paus":19,"payload":3,"payment":[11,16],"per":[8,11,13,14,25],"perceiv":[8,14],"percentag":21,"perfect":19,"perform":[0,2,4,6,7,8,11,12,16,17,19,22,25,27],"period":[19,25],"permiss":19,"persist":19,"person":[7,10,11],"perspect":[2,8],"peso":10,"philippin":10,"php":[8,10],"physic":3,"pick":12,"piec":7,"pii":5,"pilot":16,"pip":[3,10,16],"pipelin":[16,22,25],"place":[8,11,25],"placehold":16,"plain":[1,17],"plaintext":23,"plan":[5,17],"plane":[9,20,24,25],"plano":[0,1,2,3,4,6,12,13,14,16,17,18,19,21,22,23,24,25,27],"plano_config":[6,10,17,19,21,22],"plano_llm_listen":13,"plano_log":13,"plano_orchestrator_v1":[10,17,18,21],"platform":6,"pleas":[3,4,6,11,12,17,18,19],"pln":10,"plug":[8,9],"plumb":[9,19,20],"pod":16,"point":[1,2,3,5,7,11,17,19,21],"polici":[0,1,4,8,9,12,16,18,19,20,25],"polish":10,"pollut":7,"pool":[19,25],"pooler":19,"popular":16,"port":[1,2,6,7,10,12,17,18,19,21,22,25],"portal":6,"post":[3,5,7,13,17,18,21],"postgr":19,"postgresql":3,"pound":10,"power":[1,4,7,10,11,17],"practic":[1,4,8,9,15],"pre":25,"preced":16,"precis":[0,7,8,11],"predefin":2,"predict":[7,12],"prefer":[0,3,4,8,11,16,17,25],"prefix":6,"premier":6,"premis":[4,24],"prepar":[1,16],"present":[8,11,16,17],"preserv":6,"prev_response_id":19,"prevent":[8,16,18,19,25],"preview":[5,6],"previou":[3,7,17,19],"previous_resp_id":19,"previous_response_id":19,"price":[17,18],"pricing_ag":17,"primari":[2,3,5,17,19,27],"primary_and_first_fallback_fail":5,"primary_model":3,"primit":[2,4,24],"print":[3,10,11,16,19],"prior":1,"pro":6,"problem":[3,5,6,12,17],"proce":25,"process":[1,8,11,12,13,14,16,17,18,24,27],"process_customer_request":16,"processess":7,"processor":16,"prod":[4,5,6],"produc":[11,25],"product":[0,3,4,5,6,8,9,11,12,17,18,20,22],"product_recommend":17,"profan":5,"profil":[1,11,12],"program":[3,4,12],"progress":17,"project":[19,23],"prolifer":12,"prometheu":14,"promethu":14,"prompt":[0,1,8,9,12,17,18,20,21,24,25,27],"prompt_function_listen":21,"prompt_guard":9,"prompt_target":[7,9,10,11,21,25],"prone":7,"pronoun":17,"proof":[4,19],"propag":[8,15],"proper":16,"protect":[8,16,18],"proto":16,"protocol":[1,2,8,9,10,13,16,17,22,25],"proven":[8,9],"provid":[0,1,2,3,5,7,8,9,10,11,12,13,14,16,17,18,19,20,21,22,24,25],"provider_interfac":6,"proxi":[0,6,8,9,14,17,19,20,22,25],"psql":19,"public":2,"publish":14,"pull":[1,6,9,20],"purpos":[3,5,6,7,17,25],"pydant":7,"python":[0,4,5,6,8,12,13,19],"q":11,"quadrat":12,"quadratic_equ":12,"qualiti":[0,12],"quantum":[3,22],"queri":[1,4,7,10,11,12,17,18,19,21],"query_rewrit":[1,17],"question":[3,7,17,18,25],"quick":[8,12,16],"quickli":[1,8,9,10],"quickstart":[4,9,20],"quirk":8,"quot":17,"quota":[5,6],"quota_exceed":5,"quuickstart":6,"qwen3":6,"r":17,"rag":[1,17,25],"rag_ag":[1,18,21],"rag_energy_source_ag":7,"rain":17,"rais":[11,18],"rand":10,"random":21,"random_sampl":[1,10,16,17,21],"rang":[0,7,14,16],"rapid":12,"rare":25,"rate":[8,10,14,16,25],"rather":[0,1,8,24],"raw":[12,17],"re":[0,1,4,7,8,11,16,17,19],"reach":[0,1,8,18,24],"reachabl":1,"read":[8,11,16,17],"readabl":7,"readi":[3,12,19,22],"real":[0,8,10,11,12,17,24,25],"realli":8,"reason":[0,1,3,4,5,6,12,17,18],"reboot":25,"reboot_network_devic":25,"rebuild":22,"receiv":[2,7,11,13,16,17,18,19,25],"recent":17,"recept":25,"recognit":7,"recommend":[10,17,27],"record":[1,8,16],"recoveri":8,"redact":8,"reddit":8,"reduc":[0,3,7,8,19],"ref":[3,19],"refactor":[0,8],"refer":[3,4,5,6,10,11,13,16,17,20,23],"referenc":22,"refernc":11,"reflect":5,"regardless":[3,5],"region":[4,16,19],"regul":18,"regularli":19,"reinforc":20,"reject":18,"relat":[2,12,14,17,18,25],"releas":[5,6],"relev":[1,7,11,17],"reli":[12,16,24],"reliabl":[2,3,4,7,8,9,11,18,20],"remain":[0,8,9,11],"rememb":[11,17,19],"remind":11,"remov":1,"render":19,"renew":7,"renminbi":10,"reorder":1,"repeat":[8,17],"replac":[6,16,17,19],"repositori":10,"repres":[1,16],"req":7,"request":[0,1,2,3,4,5,6,8,10,11,12,13,16,17,18,19,20,21,22,26],"request_bodi":17,"requested_dai":17,"requestsinstrumentor":16,"requir":[0,2,4,7,10,11,12,17,19,21,22,24,25],"resend":19,"reset":17,"resili":8,"resolut":22,"resolv":17,"resourc":[6,10,15],"resp2":19,"resp_id":19,"respect":[13,14],"respond":[3,8,14,18],"respons":[0,1,2,4,5,6,7,8,9,10,11,12,13,16,18,19,21,25],"response_cod":13,"response_flag":13,"response_id":[3,19],"response_messag":17,"rest":[1,3,4,25,27],"restart":19,"result":[0,1,2,6,11,17],"retent":19,"retrain":12,"retri":[0,2,4,7,8,17,24,25],"retriev":[1,3,7,8,11,17,18,19,21,25],"return":[0,1,2,3,7,11,17,18,25],"reusabl":[1,21],"revers":[18,25],"review":[5,6,17],"rewrit":[1,8,17,21,25],"rewrot":1,"rich":[0,9,12,17,20],"right":[7,8,10],"ringgit":10,"risk":18,"ro":22,"robin":25,"robust":[8,9,11],"role":[3,5,10,12,16,17,18,22],"rollout":5,"romanian":10,"ron":10,"rote":[9,20],"round":25,"round_robin":5,"rout":[0,1,2,3,4,5,7,8,9,17,19,20,21,25],"router":[0,1,4,10,17,18,21,25],"routin":11,"routing_prefer":[6,12],"rule":[1,4,11,17],"run":[0,1,8,10,12,16,18,19,22,24,25],"runtim":[1,6],"rupe":10,"rupiah":10,"safe":[1,4,8,19],"safer":[18,20],"safeti":[2,5,8,9,18,20],"sai":[11,17],"sale":17,"sales_clos":17,"same":[2,3,5,6,8,11,21],"sampl":[1,10,14,16,21,22],"san":[7,11],"sanit":25,"satisfact":17,"save":[7,19],"scalabl":[7,8,9,10],"scale":[0,4,8,9,17,19,20,25],"scatter":8,"scenario":[6,7,8,9,11,12,19,25],"scene":[2,3,19],"schedul":[7,11,17],"scheduled_in":17,"scheduled_out":17,"schema":[10,19],"scheme":[5,6,14],"scienc":11,"scope":18,"scout":6,"scrape_config":14,"scrape_interv":14,"scrape_timeout":14,"screenshot":14,"script":[10,11],"scrutini":18,"sdk":[4,6,16,19],"seamless":[4,11,17,25],"seamlessli":[3,4,7,12,14,16],"search":[10,17],"search_dat":17,"search_date_obj":17,"seattl":[11,17],"second":[3,17,19],"section":[1,5,6,7,9,10,16,17,19,25],"secur":[2,4,9,16,19,25],"see":[0,2,4,10,11,12,17,19,22,25],"segment":17,"sek":10,"select":[0,2,4,9,10,12,17,19,25],"self":[2,8,13,19,25],"semant":[1,3,4,5,6,8,12],"send":[2,3,11,16,19,25],"sensibl":[1,2],"sensit":[1,16,17],"sensitive_data":5,"sent":[13,25],"separ":[7,24,25,27],"sequenc":[0,8,10,25],"server":[2,6,8,9,10,20,24,25],"servic":[0,1,2,7,8,10,11,13,16,17,18,19,21,22,25],"session":19,"set":[6,7,8,9,10,11,16,19,22,25],"set_tracer_provid":16,"setup":[1,2,7,11,16,19],"sever":[8,14,25],"sfo":10,"sgd":10,"shape":[8,25],"share":[10,25],"sheqel":10,"shift":12,"ship":[9,20],"short":[1,3],"shorten":[8,9],"should":[0,7,8,10,11,12,17,19,25],"shouldn":[9,20],"show":[0,1,2,7,10,17,18,22],"shown":6,"sidecar":[8,9],"signal":[0,8,9,12,20],"signatur":11,"significantli":17,"similar":[12,19],"simpl":[1,3,4,5,6,11,16],"simpli":[10,16],"simplic":25,"simplif":5,"simplifi":[2,4,7,9,10,16,18],"simultan":4,"sinc":[7,10],"singapor":10,"singl":[2,6,7,8,10,11,17,19,21,23,27],"sit":[8,9,11,24,25],"site":16,"size":[3,19],"skimp":11,"sla":18,"slow":7,"small":[6,27],"smaller":[6,17],"smart":[0,1,2,3,5,8,9,11,17,20,21,25],"smooth":17,"sni":25,"snippet":12,"snow":17,"so":[1,2,3,7,8,9,10,11,18,19,25],"socket":25,"softwar":8,"solar":7,"sole":12,"solut":7,"solv":[3,5,6,8,12],"some":[6,7,10,11,14,25,27],"sonnet":[3,5,6,12,19,21,22],"soon":[4,17],"sota":11,"sourc":[6,7,10,11,14,16,17],"south":10,"space":12,"span":[16,25],"span_processor":16,"spanish":12,"special":[0,6,8,12,17,19],"specif":[0,3,4,5,6,7,8,10,11,12,16,17,25],"specifi":[1,2,6,7,11,12,25],"speed":[8,14],"spell":11,"spend":27,"split":[5,17,24],"sporad":27,"spread":1,"sql":19,"stabl":[5,12],"stack":[8,9,11,13],"staff":[20,25,26],"stage":5,"stai":[8,9,10],"standalon":16,"standard":[1,4,6,8,9,10,16,17,20],"start":[1,2,16,17,19],"start_as_current_span":16,"start_tim":13,"stat":[14,25],"state":[0,1,4,7,8,9,10,11,12,20,21,24,25],"state_storag":19,"statement":17,"static":[12,25],"static_config":14,"statist":25,"statu":[1,10,11,13,17,18],"status_cod":17,"step":[0,1,9,16,20],"still":[1,2,12,19],"stitch":8,"storag":3,"store":[14,19],"stori":[3,12],"storytel":[6,12],"str":[7,10,11,17,25],"straightforward":11,"strategi":[4,8,9],"stream":[1,3,6,8,17,18,25],"streamabl":[1,21],"streamingrespons":17,"streamlin":[2,7],"strength":12,"strftime":17,"string":21,"strip":[1,17],"stripe":8,"strptime":17,"structur":[1,4,7,10,12,13,25],"struggl":7,"studio":6,"stuff":11,"style":[6,8,9,12,25],"subject":12,"submit":[11,17],"subscript":6,"subsequ":19,"substanti":8,"subsystem":[2,4,8,24,25],"subsystmem":25,"success":[1,8,10],"successfulli":[1,10],"suffix":6,"sugar":7,"suggest":11,"suit":[8,12],"suitabl":[11,12,17,19],"sum":5,"summar":[0,3,5,7,11,12,25],"summari":[11,12,15,25],"sunris":17,"sunset":17,"support":[4,7,10,12,14,16,17,18,19,20,25],"sure":[16,19],"surfac":[0,1,2,8],"sustain":7,"swap":[8,12],"swedish":10,"swiss":10,"switch":[4,5,12,19],"symbol":[10,25],"symptom":7,"syntax":19,"system":[0,1,7,8,10,11,12,13,16,17,18,19,25],"system_prompt":[7,10,25],"t":[3,6,7,9,10,11,12,19,20],"t00":17,"t23":17,"tabl":19,"tag":22,"tail":13,"tailor":11,"take":[0,8,10,25],"taken":13,"talk":[1,2,7,24],"target":[0,1,5,9,12,14,20,21,25],"task":[0,1,4,5,6,7,8,11,12,17,18,21,25,27],"tcp":25,"team":[0,1,4,8,9,20],"tech":20,"techcorp":18,"technic":[12,17,18],"techniqu":[7,12],"technologi":[7,8],"telemetri":[8,14,16],"tell":[3,17,22],"temperatur":[7,11,17,25],"temperature_2m":17,"temperature_2m_max":17,"temperature_2m_min":17,"temperature_c":17,"temperature_f":17,"temperature_max_c":17,"temperature_min_c":17,"tend":1,"term":[3,7],"termin":[1,2,8,17],"terminal_origin":17,"test":[0,3,4,5,6,7,8,11,19],"testabl":7,"text":[3,11,17,19,23,25],"text_stream":3,"textual":12,"tft":[8,14],"thai":10,"than":[0,1,8,24],"thb":10,"thei":[0,1,2,12,17,18],"them":[0,1,7,9,11,17,18,19,20,24,25],"themat":12,"theme":12,"thi":[0,1,2,3,4,5,6,7,8,9,10,11,12,13,16,17,18,19,21,22,23,25,27],"think":11,"third":24,"thirst":7,"thoroughli":11,"those":[0,1,7,16],"thread":[19,20,25,26],"three":[3,4,7,8,12,14,25],"through":[0,1,2,3,4,6,7,8,9,10,11,12,13,16,17,18,25],"throughout":8,"throughput":[5,11,12],"thunderstorm":17,"tier":19,"tier1_support":17,"tier2_support":17,"tight":[8,9],"time":[1,7,8,11,12,13,14,17,24,25],"timeout":[5,10,12,14],"timestamp":[19,25],"timezon":17,"tl":[2,8,22,25],"tlm":8,"tls_certif":7,"todai":[4,10,17,25],"togeth":[2,4,8],"together_ai":6,"together_api_kei":6,"token":[6,7,8,14,25],"toler":[4,7],"tomorrow":[10,17],"too":[17,25],"took":13,"tool":[0,1,2,3,8,12,13,14,15,17,18,21,24,25],"toolerror":18,"top":[2,4,27],"topic":[1,12,18],"topologi":26,"tot":[8,14],"total":[8,13,14],"touch":8,"trace":[0,1,8,9,10,14,15,17,20,21,25],"trace_export":16,"tracepar":[8,15,17],"traceparent_head":17,"tracer":16,"tracer_provid":16,"tracerprovid":16,"track":[10,12,17,19],"tradit":[12,19],"traffic":[2,4,5,8,9,11,13,17,24,25],"train":12,"transact":11,"transform":[6,17,25],"translat":[12,25],"transpar":[8,12],"transport":[1,21,25],"travel":[10,11],"travel_assist":10,"travel_booking_servic":[17,21],"travel_d":17,"treat":25,"trigger":[8,9,11,18],"trip":17,"trivial":25,"troubleshoot":[7,17],"troubleshoot_ag":17,"true":[1,2,3,6,7,10,11,12,14,16,17,21,25],"try":[3,10,17,22],"tupl":7,"turbo":[6,7],"turkish":10,"turn":[0,3,8,10,11,17,19],"tutori":9,"two":[1,2,10,12,17,19,25],"txt":20,"type":[0,1,2,3,5,7,10,11,12,14,17,18,19,21,22,25],"typescript":0,"typic":[2,7,11,17,18,25],"u":[6,19],"ui":25,"unambigu":12,"unavail":17,"unbound":19,"uncom":19,"under":[2,8,13],"underli":[3,5,12],"underscor":5,"understand":[0,5,8,9,11,12,14,16,17],"undifferenti":25,"unexpect":[1,11],"unifi":[2,3,4,6,8,19,25],"uniformli":16,"uniqu":[7,16,19],"unit":[7,10,11],"unix":19,"unlik":[0,7,12,19],"unnecessari":0,"unrel":18,"unresolv":17,"unsaf":1,"until":[0,17,25],"up":[1,6,7,8,9,10,16,17,19,22],"updat":[7,8,10,11,12,17,19,22,25],"updated_at":19,"upgrad":[4,5,8,25],"upon":[17,25],"upper":17,"upstream":[2,7,8,11,13,14,21,22,24,25],"upstream_host":13,"urin":7,"url":[1,4,7,10,14,17,18,19,21],"us":[0,2,3,6,7,8,9,13,15,18,19,20,21,22,23,24,25],"usag":[0,3,4,5,7,8,11,14,19,25],"usd":10,"user":[0,1,3,5,7,8,10,11,12,13,14,16,17,18,19,22],"user_messag":17,"user_queri":18,"usernam":13,"usi":16,"usual":10,"util":[10,25],"ux":[9,20],"v":[10,17,25],"v0":[1,2,7,10,17,21,25],"v1":[2,3,4,5,6,8,9,10,12,13,16,17,18,19,21,22],"v1beta":6,"v2":[5,6,10,14],"v24":10,"v3":10,"v4":6,"vagu":17,"valid":[4,7,9,10,11,17,18,19,21,25],"validate_with_llm":18,"valu":[3,7,11,16,17,25],"valueerror":11,"var_nam":19,"variabl":[6,10,16,19,22],"variant":[0,12],"varieti":25,"variou":[8,11,16,27],"ve":[7,10],"venv":10,"verbos":22,"veri":[16,25],"verifi":[16,19,22],"version":[0,1,2,3,4,5,7,10,16,17,19,21,25],"via":[0,1,2,4,6,7,8,9,10,15,16,17,24,25],"view":[14,16],"violat":[1,18,25],"virtual":[1,10,18,21],"visibl":13,"vision":7,"visual":16,"vllm":6,"vllm_api_kei":6,"volum":22,"vpc":24,"w3c":[8,16],"wa":[0,11,13,18,25],"wai":[1,2,7,8,10,18,25],"wait":25,"walk":17,"walkthrough":10,"want":[1,2,7,10,11,12],"warn":17,"watch":11,"we":[4,7,10,17,25,27],"weather":[7,11,17,21],"weather_ag":[17,21],"weather_cod":17,"weather_context":17,"weather_data":17,"weather_info":11,"weather_model":17,"weather_respons":17,"weather_url":17,"weatherag":17,"web":[2,25],"week":17,"weight":5,"well":[1,7,8,17],"were":[0,13],"west":[6,19],"what":[0,3,7,8,9,10,12,13,17,18,20,24,25],"when":[0,1,2,6,7,8,10,11,12,16,17,18,19,25],"where":[0,1,2,7,8,9,10,11,12,13,17,18,19,21,24],"whether":[0,2,4,7,11,12,25],"which":[0,1,4,7,8,10,11,12,13,14,16,17,18,24,25],"while":[0,1,2,6,7,8,9,10,11,12,17,18,19,27],"why":0,"wide":[0,7,8,9,14,16,20,25],"wind":7,"window":10,"wire":[1,10,18],"within":[9,11,17,18],"without":[0,1,2,3,4,6,7,8,9,12,19,22,25],"wmo":17,"won":10,"word":1,"work":[0,1,3,5,6,7,8,9,11,15,16,20],"worker":[25,27],"workflow":[0,1,8,9,10,12,17,25],"workload":[7,12,19,27],"world":[0,8,12,17],"worldwid":17,"would":25,"write":[0,1,6,7,8,12,17],"writer":5,"written":[8,18,25,27],"www":8,"x":[3,5,13,17,18,25],"xai":4,"xai_api_kei":6,"xxx":22,"y":17,"yaml":[3,6,7,10,14,16,17,19,22],"yen":10,"yield":17,"yml":[21,22],"york":[7,11,17],"you":[0,1,2,3,4,5,6,7,8,9,10,11,12,14,16,17,19,20,21,22,25],"your":[0,1,2,3,4,5,6,7,8,9,11,12,13,14,16,17,18,19,22,24,25],"your_us":19,"yourweatherapp":11,"yuan":10,"yyi":22,"yyyi":17,"z":6,"zar":10,"zealand":10,"zero":8,"zeroshot":13,"zhipu_api_kei":6,"z\u0142oti":10},"titles":["Agents","Filter Chains","Listeners","Client Libraries","Model (LLM) Providers","Model Aliases","Supported Providers &amp; Configuration","Prompt Target","Intro to Plano","Overview","Quickstart","Function Calling","LLM Routing","Access Logging","Monitoring","Observability","Tracing","Orchestration","Guardrails","Conversational State","Welcome to Plano!","Configuration Reference","Deployment","llms.txt","Bright Staff","Request Lifecycle","Tech Overview","Threading Model"],"titleterms":{"1":[7,10,11],"2":[7,10,11],"3":[10,11],"A":16,"For":7,"It":[13,17,19],"access":13,"addit":16,"advanc":[4,5,6],"agent":[0,2,10,17],"ai":6,"alia":12,"alias":5,"alibaba":6,"align":12,"also":[3,5,6],"amazon":6,"anthrop":[3,6],"api":[3,6,10,17],"app":[7,10],"arch":[11,12],"architectur":25,"assist":17,"aw":16,"azur":6,"base":[6,12,22],"basic":[5,7],"bedrock":6,"benefit":[0,4,16],"best":[3,5,11,12,16,17,19],"book":17,"bright":24,"build":[7,9,10],"call":[10,11,17],"capabl":4,"case":[1,4,11,12,17],"categori":6,"chain":1,"class":6,"client":[3,10,12],"combin":12,"come":5,"command":10,"common":[4,17,22],"compat":[3,6],"compos":22,"concept":[9,20],"config":[7,10],"configur":[1,2,5,6,7,11,12,14,17,19,21,25],"connect":19,"context":17,"convers":[3,10,19],"core":4,"creat":10,"cross":3,"curl":[3,10],"currenc":10,"dashboard":14,"datadog":16,"deepseek":6,"default":6,"defin":[7,11],"demo":7,"deploy":22,"determinist":10,"develop":19,"docker":22,"egress":[2,25],"endpoint":[3,6],"error":3,"exampl":[1,3,7,11,12,16,17],"extern":17,"extract":17,"featur":[4,5,7,11,12,13],"file":10,"filter":1,"first":6,"flask":7,"flow":25,"format":[13,16],"function":11,"gatewai":[3,10,22],"gemini":6,"gener":17,"get":[4,9,20],"googl":6,"grafana":14,"groq":6,"guardrail":18,"guid":[9,20],"guidelin":6,"handl":3,"header":16,"high":25,"how":[13,16,17,18,19],"http":1,"i":11,"implement":[11,17],"inbound":2,"inform":17,"ingress":25,"initi":16,"inner":0,"instanc":6,"instrument":16,"integr":16,"intent":7,"interact":10,"intro":8,"issu":22,"kei":[0,4,7,11,13],"langtrac":16,"let":10,"level":25,"librari":3,"lifecycl":25,"listen":2,"llm":[4,10,12,17,23],"log":13,"logic":0,"loop":0,"mcp":1,"memori":19,"method":12,"metric":14,"minim":10,"mistral":6,"model":[1,2,4,5,6,10,12,22,27],"monitor":14,"moonshot":6,"multi":7,"multipl":6,"name":5,"network":[2,25],"next":[10,17,19],"observ":15,"ollama":6,"openai":[3,6,10],"opentelemetri":16,"orchestr":[0,10,17],"outbound":2,"outer":0,"over":11,"overview":[9,16,19,26],"paramet":7,"plano":[7,8,9,10,11,20],"post":25,"postgresql":19,"practic":[3,5,11,12,16,17,19],"prefer":[6,12],"prepar":17,"prerequisit":[10,19],"process":[7,25],"product":19,"program":1,"prompt":[2,7,10,11],"propag":16,"provid":[4,6],"proxi":[2,10],"python":[3,10,16],"quickstart":10,"qwen":6,"rag":7,"rai":16,"refer":21,"request":[7,25],"requir":6,"resourc":[16,20],"respons":[3,17],"rout":[6,10,12,22],"router":12,"rule":5,"runtim":22,"sdk":3,"see":[3,5,6],"select":6,"send":10,"setup":22,"smoke":22,"solut":22,"soon":5,"stack":22,"staff":24,"start":[4,9,10,20,22],"state":[3,19],"step":[7,10,11,17,19],"storag":19,"string":19,"structur":[6,17],"summari":[7,16],"supabas":19,"support":[3,6],"switch":7,"take":11,"target":[2,7,10,11],"tech":26,"test":[18,22],"thread":27,"tip":11,"togeth":6,"tool":[7,16],"topologi":[2,25],"trace":16,"tracepar":16,"travel":17,"troubleshoot":[19,22],"turn":7,"txt":23,"typic":1,"unsupport":12,"url":6,"us":[1,4,5,10,11,12,16,17],"usag":12,"v":0,"valid":5,"via":14,"welcom":20,"what":11,"why":18,"work":[13,17,18,19],"workflow":11,"x":16,"xai":6,"your":10,"zhipu":6}})
\ No newline at end of file
diff --git a/sitemap.xml b/sitemap.xml
index 0cd358f5..83fe56b4 100755
--- a/sitemap.xml
+++ b/sitemap.xml
@@ -1,2 +1,2 @@
 <?xml version='1.0' encoding='utf-8'?>
-<urlset xmlns="http://www.sitemaps.org/schemas/sitemap/0.9"><url><loc>./docsbuild_with_arch/agent.html</loc></url><url><loc>./docsbuild_with_arch/multi_turn.html</loc></url><url><loc>./docsbuild_with_arch/rag.html</loc></url><url><loc>./docsconcepts/llm_providers/client_libraries.html</loc></url><url><loc>./docsconcepts/llm_providers/llm_providers.html</loc></url><url><loc>./docsconcepts/llm_providers/model_aliases.html</loc></url><url><loc>./docsconcepts/llm_providers/supported_providers.html</loc></url><url><loc>./docsconcepts/prompt_target.html</loc></url><url><loc>./docsconcepts/tech_overview/error_target.html</loc></url><url><loc>./docsconcepts/tech_overview/listener.html</loc></url><url><loc>./docsconcepts/tech_overview/model_serving.html</loc></url><url><loc>./docsconcepts/tech_overview/prompt.html</loc></url><url><loc>./docsconcepts/tech_overview/request_lifecycle.html</loc></url><url><loc>./docsconcepts/tech_overview/tech_overview.html</loc></url><url><loc>./docsconcepts/tech_overview/terminology.html</loc></url><url><loc>./docsconcepts/tech_overview/threading_model.html</loc></url><url><loc>./docsget_started/intro_to_arch.html</loc></url><url><loc>./docsget_started/overview.html</loc></url><url><loc>./docsget_started/quickstart.html</loc></url><url><loc>./docsguides/agent_routing.html</loc></url><url><loc>./docsguides/function_calling.html</loc></url><url><loc>./docsguides/llm_router.html</loc></url><url><loc>./docsguides/observability/access_logging.html</loc></url><url><loc>./docsguides/observability/monitoring.html</loc></url><url><loc>./docsguides/observability/observability.html</loc></url><url><loc>./docsguides/observability/tracing.html</loc></url><url><loc>./docsguides/prompt_guard.html</loc></url><url><loc>./docsindex.html</loc></url><url><loc>./docsresources/configuration_reference.html</loc></url><url><loc>./docsresources/deployment.html</loc></url><url><loc>./docssearch.html</loc></url></urlset>
\ No newline at end of file
+<urlset xmlns="http://www.sitemaps.org/schemas/sitemap/0.9"><url><loc>./docsconcepts/agents.html</loc></url><url><loc>./docsconcepts/filter_chain.html</loc></url><url><loc>./docsconcepts/listeners.html</loc></url><url><loc>./docsconcepts/llm_providers/client_libraries.html</loc></url><url><loc>./docsconcepts/llm_providers/llm_providers.html</loc></url><url><loc>./docsconcepts/llm_providers/model_aliases.html</loc></url><url><loc>./docsconcepts/llm_providers/supported_providers.html</loc></url><url><loc>./docsconcepts/prompt_target.html</loc></url><url><loc>./docsget_started/intro_to_plano.html</loc></url><url><loc>./docsget_started/overview.html</loc></url><url><loc>./docsget_started/quickstart.html</loc></url><url><loc>./docsguides/function_calling.html</loc></url><url><loc>./docsguides/llm_router.html</loc></url><url><loc>./docsguides/observability/access_logging.html</loc></url><url><loc>./docsguides/observability/monitoring.html</loc></url><url><loc>./docsguides/observability/observability.html</loc></url><url><loc>./docsguides/observability/tracing.html</loc></url><url><loc>./docsguides/orchestration.html</loc></url><url><loc>./docsguides/prompt_guard.html</loc></url><url><loc>./docsguides/state.html</loc></url><url><loc>./docsindex.html</loc></url><url><loc>./docsresources/configuration_reference.html</loc></url><url><loc>./docsresources/deployment.html</loc></url><url><loc>./docsresources/llms_txt.html</loc></url><url><loc>./docsresources/tech_overview/model_serving.html</loc></url><url><loc>./docsresources/tech_overview/request_lifecycle.html</loc></url><url><loc>./docsresources/tech_overview/tech_overview.html</loc></url><url><loc>./docsresources/tech_overview/threading_model.html</loc></url><url><loc>./docssearch.html</loc></url></urlset>
\ No newline at end of file