{
  "title": "Agentic RL method and benchmark comparisons",
  "scope": "Qualitative comparisons of methods, information requirements and task design in the cited paper versions.",
  "paper_count": 16,
  "html_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/index.html",
  "json_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/index.json",
  "topics": [
    {
      "title": "Credit assignment methods for multi-turn agents",
      "question": "Which feedback and cross-action comparisons produce the policy's credit signal?",
      "scope": "Qualitative comparisons of methods, information requirements and task design in the cited paper versions.",
      "keywords": [
        "credit assignment",
        "advantage estimation",
        "multi-turn reinforcement learning",
        "sparse rewards"
      ],
      "dimensions": [
        "Credit unit",
        "Feedback",
        "Comparison mechanism",
        "Additional requirements"
      ],
      "html_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/credit-assignment-methods.html",
      "json_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/credit-assignment-methods.json",
      "csv_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/credit-assignment-methods.csv",
      "bibtex_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/credit-assignment-methods.bib",
      "paper_count": 8
    },
    {
      "title": "Memory, context management and selective observation in agents",
      "question": "What information is retained, when is it selected, and how is selection trained?",
      "scope": "Qualitative comparisons of methods, information requirements and task design in the cited paper versions.",
      "keywords": [
        "agent memory",
        "context management",
        "selective observation",
        "partial observability",
        "context budget"
      ],
      "dimensions": [
        "Information retained",
        "Operation and timing",
        "Training or selection rule",
        "Task setting"
      ],
      "html_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/agent-memory-observation.html",
      "json_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/agent-memory-observation.json",
      "csv_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/agent-memory-observation.csv",
      "bibtex_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/agent-memory-observation.bib",
      "paper_count": 4
    },
    {
      "title": "Terminal and software agent benchmark comparison",
      "question": "How do task origin, interaction interface and success criteria differ?",
      "scope": "Qualitative comparisons of methods, information requirements and task design in the cited paper versions.",
      "keywords": [
        "CLI agent benchmark",
        "terminal benchmark",
        "execution feedback",
        "software agent evaluation",
        "ShellOps"
      ],
      "dimensions": [
        "Task source",
        "Interaction interface",
        "Success criteria",
        "Evaluation scope"
      ],
      "html_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/terminal-benchmark-design.html",
      "json_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/terminal-benchmark-design.json",
      "csv_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/terminal-benchmark-design.csv",
      "bibtex_url": "https://hoyant-su-agentic-rl.hf.space/topics/comparisons/terminal-benchmark-design.bib",
      "paper_count": 6
    }
  ]
}
