{"schema_version":"newruntime-agent-readable-v0.2","type":"post","stable_id":"post:parallel-responses-web-research-subagent","slug":"parallel-responses-web-research-subagent","title":"Parallel Packages Web Research As A Responses-Compatible Subagent","description":"Parallel's Responses API offers cited web research behind an OpenAI-compatible endpoint, with bounded effort tiers, streaming, and stateful follow-ups.","retrieval_nugget":"Parallel's Responses API offers cited web research behind an OpenAI-compatible endpoint, with bounded effort tiers, streaming, and stateful follow-ups. Parallel's Responses API turns web research into a specialized agent boundary rather than another search tool that dumps pages into the parent context. The endpoint accepts the OpenAI Responses wire format and returns synthesized answers with citations.","status":"published","published_at":"2026-08-01","updated_at":"2026-08-01","record_date":"2026-08-01","date_kind":"published_at","topics":["research-agents","responses-api","context-engineering","web-search"],"source_urls":["https://parallel.ai/blog/responses-api","https://docs.parallel.ai/responses-api/responses-quickstart"],"visuals":[{"id":"parallel-responses-web-research-subagent","kind":"editorial-diagram","role":"hero","src":"https://newruntime.com/images/posts/parallel-responses-web-research-subagent.webp","alt":"Hand-drawn orchestrator delegating parallel web questions to isolated research workers that return compact cited answers through one compatible response interface.","caption":"Parallel keeps raw web pages inside specialized research workers and returns cited answers to the main agent.","credit":"New Runtime synthesis from Parallel","source_url":"https://parallel.ai/blog/responses-api","generated_with":"gemini-3.1-flash-image","width":1600,"height":900,"legend":[{"label":"Delegate","description":"The orchestrator sends bounded research questions instead of ingesting raw search results."},{"label":"Research","description":"Independent workers search, cross-check, and synthesize in parallel."},{"label":"Return","description":"The parent agent receives a compact answer, citations, and optional structured output."}]}],"telegram_message_id":2898,"telegram_url":"https://t.me/qwgai/2898","telegram_message_ids":[2898,2899],"telegram_delivery_mode":"text_then_media","telegram_media_url":"https://t.me/qwgai/2899","routes":{"html":"https://newruntime.com/posts/parallel-responses-web-research-subagent/","markdown":"https://newruntime.com/posts/parallel-responses-web-research-subagent.md","json":"https://newruntime.com/posts/parallel-responses-web-research-subagent.json"},"source_format":"markdown","next_reads":[{"type":"topic","path":"/topics/context-engineering/","reason":"Explore the context engineering topic hub.","url":"https://newruntime.com/topics/context-engineering/","title":"Context engineering - New Runtime","media_type":"text/html"},{"type":"related_material","path":"/posts/chatgpt-agent-loop-efficiency-stack/","reason":"Shares context engineering.","url":"https://newruntime.com/posts/chatgpt-agent-loop-efficiency-stack/","title":"ChatGPT Cuts Repeated Work Across The Agent Stack","media_type":"text/html"},{"type":"related_material","path":"/posts/contextual-agent-memory-four-layer-system/","reason":"Shares context engineering.","url":"https://newruntime.com/posts/contextual-agent-memory-four-layer-system/","title":"A Vector Store Is Not An Agent Memory System","media_type":"text/html"},{"type":"related_material","path":"/posts/anthropic-tool-search-programmatic-calls/","reason":"Shares context engineering.","url":"https://newruntime.com/posts/anthropic-tool-search-programmatic-calls/","title":"Anthropic Moves Large Tool Libraries Out Of Context","media_type":"text/html"},{"type":"related_material","path":"/posts/drskill-agent-loadout-audit/","reason":"Shares context engineering.","url":"https://newruntime.com/posts/drskill-agent-loadout-audit/","title":"Dr. Skill Audits What An Agent Loads Before It Works","media_type":"text/html"}]}
