HTTP/1.1 200 OK
Date: Thu, 17 Sep 2026 16:04:29 GMT
Content-Type: text/plain; charset=UTF-8
Transfer-Encoding: chunked
Connection: keep-alive
Access-Control-Allow-Origin: *
X-Kong-Status-Request-ID: swA08FHVHQdk3rSJAqf96TKAWjBpYLCf
X-Kong-Admin-Latency: 2
Server: kong/2.1.0-ai-gateway
# HELP kong_ai_llm_provider_latency_ms LLM response Latency for each AI plugins per ai_provider in Kong
# TYPE kong_ai_llm_provider_latency_ms histogram
kong_ai_llm_provider_latency_ms_bucket{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot",le="1500"} 1
kong_ai_llm_provider_latency_ms_bucket{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot",le="2000"} 1
...
kong_ai_llm_provider_latency_ms_bucket{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot",le="+Inf"} 1
kong_ai_llm_provider_latency_ms_count{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot"} 1
kong_ai_llm_provider_latency_ms_sum{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot"} 1029
# HELP kong_ai_llm_requests_total AI requests total per ai_provider in Kong
# TYPE kong_ai_llm_requests_total counter
kong_ai_llm_requests_total{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot"} 1
# HELP kong_ai_llm_tokens_total AI requests cost per ai_provider/cache in Kong
# TYPE kong_ai_llm_tokens_total counter
kong_ai_llm_tokens_total{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",token_type="completion_tokens",workspace="default",consumer=""} 12
kong_ai_llm_tokens_total{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",token_type="prompt_tokens",workspace="default",consumer=""} 13
kong_ai_llm_tokens_total{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",token_type="total_tokens",workspace="default",consumer=""} 25
# HELP kong_ai_llm_tpot_latency_ms LLM time per token latency for each AI plugins per ai_provider in Kong
# TYPE kong_ai_llm_tpot_latency_ms histogram
kong_ai_llm_tpot_latency_ms_bucket{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot",le="100"} 1
kong_ai_llm_tpot_latency_ms_bucket{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot",le="200"} 1
...
kong_ai_llm_tpot_latency_ms_bucket{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot",le="+Inf"} 1
kong_ai_llm_tpot_latency_ms_count{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot"} 1
kong_ai_llm_tpot_latency_ms_sum{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot"} 85.75
# HELP kong_ai_llm_ttft_latency_ms LLM time to first token latency for each AI plugins per ai_provider in Kong
# TYPE kong_ai_llm_ttft_latency_ms histogram
kong_ai_llm_ttft_latency_ms_bucket{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot",le="1500"} 1
kong_ai_llm_ttft_latency_ms_bucket{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot",le="2000"} 1
...
kong_ai_llm_ttft_latency_ms_bucket{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot",le="+Inf"} 1
kong_ai_llm_ttft_latency_ms_count{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot"} 1
kong_ai_llm_ttft_latency_ms_sum{ai_provider="openai",ai_model="gpt-4o",cache_status="",vector_db="",embeddings_provider="",embeddings_model="",workspace="default",consumer="",request_mode="oneshot"} 1029
# HELP kong_control_plane_connected Kong connected to control plane, 0 is unconnected
# TYPE kong_control_plane_connected gauge
kong_control_plane_connected 1
# HELP kong_data_plane_cluster_cert_expiry_timestamp Unix timestamp of Data Plane's cluster_cert expiry time
# TYPE kong_data_plane_cluster_cert_expiry_timestamp gauge
kong_data_plane_cluster_cert_expiry_timestamp 1792252942
# HELP kong_datastore_reachable Datastore reachable from Kong, 0 is unreachable
# TYPE kong_datastore_reachable gauge
kong_datastore_reachable 1
# HELP kong_http_requests_total HTTP status codes per consumer/service/route in Kong
# TYPE kong_http_requests_total counter
kong_http_requests_total{service="ai-gateway",route="openai-chat",code="200",source="service",type="",workspace="default",consumer=""} 1
# HELP kong_kong_internal_latency_ms Internal latency for each service/route in Kong, excluding the I/O latency
# TYPE kong_kong_internal_latency_ms histogram
kong_kong_internal_latency_ms_bucket{service="ai-gateway",route="openai-chat",workspace="default",le="10"} 1
kong_kong_internal_latency_ms_bucket{service="ai-gateway",route="openai-chat",workspace="default",le="15"} 1
kong_kong_internal_latency_ms_bucket{service="ai-gateway",route="openai-chat",workspace="default",le="20"} 1
...