{
 "license": "CC-BY-4.0",
 "attribution": "Fide AI, Agentic Cyber Explorer",
 "url": "https://agentic-cyber-explorer.pages.dev/events/openai-pacing-development-cyber-critical-2026/",
 "asOf": "2026-09-26",
 "id": "openai-pacing-development-cyber-critical-2026",
 "date": "2026-08-18",
 "datePrecision": "day",
 "title": "OpenAI pauses RL training and hardens research environments as Astra nears Critical cyber threshold",
 "lane": "policy",
 "kind": "framework",
 "summary": "OpenAI said that the OpenAI-Hugging Face evaluation incident and preliminary evidence that its then-unreleased Astra model may meet the Critical cybersecurity threshold led it to slow scaling, including a two-week pause in reinforcement learning training on deployment models. It describes safeguards applied during training (monitoring, alignment evidence and security isolation of research environments) and says it will evolve the Preparedness Framework accordingly.",
 "whyItMatters": "It is a public case of a lab applying its Critical cyber threshold to development itself, including isolating its own training environments.",
 "actors": [
  "openai"
 ],
 "topics": [
  "capability-thresholds",
  "sandbox-containment",
  "monitoring-and-control"
 ],
 "atlas": [
  "sandbox",
  "monitor",
  "eval-environment"
 ],
 "artifacts": [
  "gpt-6-astra"
 ],
 "sources": [
  {
   "url": "https://openai.com/index/pacing-model-development-cyber-capabilities/",
   "publisher": "OpenAI",
   "title": "Pacing model development in an era of cyber-critical capabilities",
   "date": "2026-08-18",
   "type": "primary",
   "accessed": "2026-09-25"
  }
 ],
 "keyFacts": [
  {
   "fact": "Included a two-week pause in RL training on models intended for deployment; the largest planned frontier RL run remains on hold.",
   "locator": "Opening section"
  },
  {
   "fact": "Security measures include stronger workload sandboxes and network isolation so a single compromise cannot enable unauthorized internet access.",
   "locator": "Security measures"
  },
  {
   "fact": "Monitoring overhead is estimated at roughly 20% of the inference compute being monitored; OpenAI aims to issue an alert within 30 minutes after concerning activity is surfaced by its monitoring system.",
   "locator": "Monitoring section"
  }
 ],
 "significance": 4,
 "fideQuestions": [],
 "methods": [
  "ai-monitoring",
  "sandboxing-egress"
 ],
 "review": "assistant-drafted",
 "addedOn": "2026-09-25"
}