{
  "schema": "project-experience.writing-task-input.v1",
  "status": "constructed_example_not_evaluated",
  "context_rule": "Only supplied excerpts with Overleaf revision <= cutoff are included. Prompts and selected cutoffs were designed retrospectively; this is not original-context replay or a sealed evaluation.",
  "task": {
    "id": "WA3",
    "title": "Consolidate a mixed draft and preserve the current scope",
    "action": "Revise without regression",
    "cutoff_revision": 2895,
    "prompt": "Edit the current appendix excerpt into a coherent explanation using only the supplied material. Retain the definition of query budget, distinguish saved solver work from the computational cost of model calls, and remove repeated active prose. Keep the scope of the current draft's runtime assertion; do not restore an older broader claim from historical context. Do not infer an optimal budget or create results from a figure that is not supplied. Return a proposed revision, then a short list of claims still needing experimental support.",
    "output_request": "A revised section body and a brief evidence-check note. Preserve useful qualifications; no fixed historical wording is required.",
    "skill_ids": [
      "WS04",
      "WS05",
      "WS06",
      "WS07"
    ],
    "input_note": "Comment-only lines are left out of this task input. The complete manuscript excerpts, including comments, remain available in the source library. The current appendix is at revision 2895; optional earlier context is at revision 2614.",
    "current_source_ids": [
      "C08"
    ],
    "earlier_context_ids": [
      "C06"
    ],
    "sources": [
      {
        "evidence_id": "C08",
        "source_file": "appendix.tex",
        "content_kind": "mixed_prose_and_comments",
        "text": "\\section{Query Budget Analysis}\n\\label{sec:query-budget}\nIn this section, we study how the query budget, the number of times the solver consults the model during search, affects both effectiveness and computational cost. We vary the budget from 1 to 10 calls, and also include an all-calls configuration, to understand how much benefit each additional model query provides.\n\nAs shown in Figure~\\ref{fig:query}, \nthe performance gain exhibits clear diminishing returns: the majority of improvement is obtained from the first three queries, while additional calls beyond six provide only marginal benefits. Because each query introduces non-trivial latency, we interpret the reduction in propagations achieved by the model as a benefit that must be weighed against the computational overhead of issuing more queries.\n\nperformance exhibits diminishing returns as the query budget increases. Most improvement comes from the first three queries. Additional queries beyond six calls offer only marginal gains.\nWe view the reduction in propagations achieved by the model as a benefit that must be traded off against the computational cost of additional queries.\nTo balance this trade-off, we set a default query budget of 3.\nMoreover, as shown in Figure~\\ref{fig:wallclock-methods}, GQSAT also exhibit worse wall-clock performance as the query budget increases from 3 to 5.",
        "text_view": "comment_lines_omitted",
        "presented_text_sha256": "179a678e40edba25a6097dd06ccb965761655e20dc72776775052dacd59f6984",
        "context_status": "A selected manuscript passage, not a complete record of the context available when it was written. The records do not establish who wrote it; scientific claims are unverified.",
        "excerpt_sha256": "2977a8090fad4c8f08af643bce95278a35db9390add20a7a02cef4799343d91c",
        "manuscript_revision": 2895
      },
      {
        "evidence_id": "C06",
        "source_file": "appendix.tex",
        "content_kind": "mixed_prose_and_comments",
        "text": "\\section{Query Budget Analysis}\nWe evaluate our method across different query budgets from 1 to 10 model calls and an all-calls setting.\nAs shown in Figure~\\ref{fig:query}, performance exhibits diminishing returns as the query budget increases. Most improvement comes from the first three queries. Additional queries beyond six calls offer only marginal gains.\nWe view the reduction in propagations achieved by the model as a benefit that must be traded off against the computational cost of additional queries.\nTo balance this trade-off, we set a default query budget of 3.\nMoreover, as shown in Figure~\\ref{fig:wallclock-methods}, both our method and GQSAT exhibit worse wall-clock performance as the query budget increases from 3 to 5.",
        "text_view": "comment_lines_omitted",
        "presented_text_sha256": "0ccf4605824e9f2139a736dbea8c9f47159cc9998b0a9b293bec83f5f1e7d95e",
        "context_status": "A selected manuscript passage, not a complete record of the context available when it was written. The records do not establish who wrote it; scientific claims are unverified.",
        "excerpt_sha256": "798b4f47187458fa64e6c2036d5f4d4f124499b2738336538895cca64bff12fa",
        "manuscript_revision": 2614
      }
    ],
    "prompt_origin": "newly_authored_derived_instruction_not_historical_prompt",
    "historical_note_available": false,
    "reference_completion_included": false,
    "independent_project_group": "ImitSAT-single-project"
  },
  "version": "1.0",
  "use_notice": "For discussing a possible data partnership. Training, redistribution and other uses require a separate agreement."
}
