{"as_of":"2026-08-09T21:35:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:58ce310d09015ad3602d9b9b3f2f3c72898c867a2f31768b1aa24e93251f9153","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":14,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":14,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":14,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":14,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T23:42:26.007060Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T09:39:46.934646Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":"2403.08604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-04T09:39:46.934646Z","title":"DevBench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3","venue":null,"work_id":"a4eb17af-0573-470e-b2b6-600c3e211106","year":2024},"citing_paper":{"arxiv_id":"2404.07972","last_updated":"2024-05-30T08:55:12Z","snapshot_observed_at":"2026-08-09T01:39:11.956106Z","submitted_at":"2024-04-11T17:56:05Z","title":"OSWorld: Benchmarking Multimodal Agents for Open-Ended Tasks in Real Computer Environments","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-13T01:19:32.406859Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2404.07972"},"observation_digest":"sha256:b9839d1214524e9ea4b40dd4b81676d904607dc3219d4267f252f08eaf68e42b","observation_id":"12a5dd74-81c7-41d0-b5c7-26816fe11cc2","resolution":{"observed_at":"2026-05-13T01:19:32.554838Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":"2403.08604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-04T09:39:46.934646Z","title":"DevBench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3","venue":null,"work_id":"a4eb17af-0573-470e-b2b6-600c3e211106","year":2024},"citing_paper":{"arxiv_id":"2502.06556","last_updated":"2026-04-07T02:47:48Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-10T15:24:30Z","title":"MultiFileTest: A Multi-File-Level LLM Unit Test Generation Benchmark and Impact of Error Fixing Mechanisms","version":5},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-23T03:42:36.946141Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2502.06556"},"observation_digest":"sha256:450d20dc2bea923eae7ce3adcfad690a4ff215c36850c8c75a2fcecc77cc89d5","observation_id":"deddc2de-ffbf-4929-8030-a6510c9e708f","resolution":{"observed_at":"2026-05-23T03:45:21.852334Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-08-07T23:42:26.007060Z","title":"Devbench: A comprehensive benchmark for software development, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.08806","last_updated":"2025-02-12T21:42:56Z","snapshot_observed_at":"2026-08-09T13:06:33.955003Z","submitted_at":"2025-02-12T21:42:56Z","title":"CLOVER: A Test Case Generation Benchmark with Coverage, Long-Context, and Verification","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T23:42:26.007060Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2502.08806"},"observation_digest":"sha256:ae472e26f9a123907a811aac32d78c682be26087f4c3eb8848eacfa1b9c7b3c3","observation_id":"afbbbc1c-2504-4bd4-a9d0-115285ec925d","resolution":{"observed_at":"2026-08-07T23:42:26.007060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-08-07T06:00:18.380320Z","title":"Prompting large language models to tackle the full software development lifecycle: A case study,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.11102","last_updated":"2025-06-06T17:52:18Z","snapshot_observed_at":"2026-08-07T05:54:56.593167Z","submitted_at":"2025-06-06T17:52:18Z","title":"Evolutionary Perspectives on the Evaluation of LLM-Based AI Agents: A Comprehensive Survey","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T06:00:18.380320Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2506.11102"},"observation_digest":"sha256:d37c8f599f83bf499698b3f91c214b096b05949e2ed603803440a640b282d061","observation_id":"8794d234-6f6e-4577-af53-7fe360f293b7","resolution":{"observed_at":"2026-08-07T06:00:18.380320Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-08-06T18:08:47.747783Z","title":"Devbench: A comprehensive benchmark for software development","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.09063","last_updated":"2025-07-11T22:45:07Z","snapshot_observed_at":"2026-08-08T09:04:57.537928Z","submitted_at":"2025-07-11T22:45:07Z","title":"SetupBench: Assessing Software Engineering Agents' Ability to Bootstrap Development Environments","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-06T18:08:47.747783Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2507.09063"},"observation_digest":"sha256:0a821f2c00c20e698cee6247c8f803ef1bd274294756be27858fc3ccb22ca28e","observation_id":"2e45cd78-68a0-416c-b7a8-3838f9ccc14d","resolution":{"observed_at":"2026-08-06T18:08:47.747783Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-08-05T18:22:38.950938Z","title":"Devbench: A comprehensive benchmark for software development,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.14703","last_updated":"2025-08-20T13:28:39Z","snapshot_observed_at":"2026-08-09T07:32:17.093278Z","submitted_at":"2025-08-20T13:28:39Z","title":"A Lightweight Incentive-Based Privacy-Preserving Smart Metering Protocol for Value-Added Services","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-05T18:22:38.950938Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2508.14703"},"observation_digest":"sha256:127e48014ca250c34003b3fb5d55f2733c60875526b8a156ce81f3834ae54026","observation_id":"013047fe-7fc1-4db3-8752-041ff5bce123","resolution":{"observed_at":"2026-08-05T18:22:38.950938Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-13T20:49:09.477849Z","title":"Devbench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3, 2024c","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.21489","last_updated":"2026-07-08T16:30:55Z","snapshot_observed_at":"2026-08-03T00:37:21.202769Z","submitted_at":"2026-03-23T02:26:35Z","title":"Effective Strategies for Asynchronous Software Engineering Agents","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-13T20:49:09.477849Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2603.21489"},"observation_digest":"sha256:14da1d5ed68083b990e6bf56a9241c27a05ef1a9b931894b00c84f63c6ede99a","observation_id":"7c8162e6-6e03-4e21-8d57-305d932c3fff","resolution":{"observed_at":"2026-07-13T20:49:09.477849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":"2403.08604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-04T09:39:46.934646Z","title":"DevBench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3","venue":null,"work_id":"a4eb17af-0573-470e-b2b6-600c3e211106","year":2024},"citing_paper":{"arxiv_id":"2604.04009","last_updated":"2026-04-05T07:54:18Z","snapshot_observed_at":"2026-08-05T13:49:50.008471Z","submitted_at":"2026-04-05T07:54:18Z","title":"Benchmarking and Evaluating VLMs for Software Architecture Diagram Understanding","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-13T17:15:31.368264Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2604.04009"},"observation_digest":"sha256:76b3dc0b346b5376f1362e4c38b7c502a658c9d168deb7525064fce27c9001aa","observation_id":"0735e3d5-3728-4c6e-a7af-5ad71e8abe2a","resolution":{"observed_at":"2026-05-13T17:17:57.712918Z","resolver_source":"orphan_title_repair","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":"2403.08604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-04T09:39:46.934646Z","title":"DevBench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3","venue":null,"work_id":"a4eb17af-0573-470e-b2b6-600c3e211106","year":2024},"citing_paper":{"arxiv_id":"2605.05700","last_updated":"2026-05-07T05:44:52Z","snapshot_observed_at":"2026-08-05T12:33:17.159859Z","submitted_at":"2026-05-07T05:44:52Z","title":"An Empirical Study of Proactive Coding Assistants in Real-World Software Development","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-08T09:22:38.776358Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2605.05700"},"observation_digest":"sha256:2a77d2f1154999d5496419a9e7a1757cccebd4d4f095a33f057ccd1278c40533","observation_id":"d37b59c1-2381-4af9-ab0e-7c4745b7fa6c","resolution":{"observed_at":"2026-05-11T20:21:11.811162Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":"2403.08604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-04T09:39:46.934646Z","title":"DevBench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3","venue":null,"work_id":"a4eb17af-0573-470e-b2b6-600c3e211106","year":2024},"citing_paper":{"arxiv_id":"2605.07073","last_updated":"2026-05-08T00:48:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-08T00:48:45Z","title":"TeamBench: Evaluating Agent Coordination under Enforced Role Separation","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-11T00:55:51.358828Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2605.07073"},"observation_digest":"sha256:10e1ef780a838e49686b136add509adafcc6f5116d739a3fa35b2eb8449171d9","observation_id":"09460c20-8bbe-4e87-bd13-2f2b14ed3f72","resolution":{"observed_at":"2026-05-11T05:00:55.450967Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":"2403.08604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-04T09:39:46.934646Z","title":"DevBench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3","venue":null,"work_id":"a4eb17af-0573-470e-b2b6-600c3e211106","year":2024},"citing_paper":{"arxiv_id":"2605.19901","last_updated":"2026-05-19T14:32:50Z","snapshot_observed_at":"2026-07-06T23:30:35.042915Z","submitted_at":"2026-05-19T14:32:50Z","title":"Can LLMs Produce Better Object-Oriented Designs than Human-Involved Development?","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-20T04:10:32.475989Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2605.19901"},"observation_digest":"sha256:cdc6d4b4c7dee9c1a1ba3e909c3db57b353e7033de58dcb8d3a56a2c188b8a98","observation_id":"3b56e7bd-5616-479d-80ad-e0ce8cc6a6f6","resolution":{"observed_at":"2026-05-20T04:13:01.280082Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":"2403.08604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-04T09:39:46.934646Z","title":"DevBench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3","venue":null,"work_id":"a4eb17af-0573-470e-b2b6-600c3e211106","year":2024},"citing_paper":{"arxiv_id":"2605.29277","last_updated":"2026-05-28T02:52:58Z","snapshot_observed_at":"2026-08-02T07:14:00.912490Z","submitted_at":"2026-05-28T02:52:58Z","title":"Code-QA-Bench: Separating Code Reasoning from Documentation Memorization in Repository-Level QA","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-29T07:00:07.102425Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2605.29277"},"observation_digest":"sha256:be5626e6ea1245c4b7ab27090ee1ab27a276a22cf8050438d228adebedebb9cf","observation_id":"3e1880d3-702d-4a89-9fb5-8e012008d4f8","resolution":{"observed_at":"2026-06-29T07:03:12.903482Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":"2403.08604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-04T09:39:46.934646Z","title":"DevBench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3","venue":null,"work_id":"a4eb17af-0573-470e-b2b6-600c3e211106","year":2024},"citing_paper":{"arxiv_id":"2606.22678","last_updated":"2026-06-29T19:02:41Z","snapshot_observed_at":"2026-07-31T20:17:16.405908Z","submitted_at":"2026-06-21T21:41:34Z","title":"RigorBench: Benchmarking Engineering Process Discipline in Autonomous AI Coding Agents","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-26T09:40:00.507836Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2606.22678"},"observation_digest":"sha256:90e45a479a26aee4dcd33b70eeb8acbb44d9292f9b0421746feb9bfce4742a65","observation_id":"48ee2c7e-c4ab-4e01-bc0a-bfb222538f12","resolution":{"observed_at":"2026-07-04T09:39:46.936309Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study","version":3},"cited_work":{"arxiv_id":"2403.08604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.08604","snapshot_observed_at":"2026-07-04T09:39:46.934646Z","title":"DevBench: A comprehensive benchmark for software development.arXiv preprint arXiv:2403.08604, 3","venue":null,"work_id":"a4eb17af-0573-470e-b2b6-600c3e211106","year":2024},"citing_paper":{"arxiv_id":"2606.22678","last_updated":"2026-06-29T19:02:41Z","snapshot_observed_at":"2026-07-31T20:17:16.405908Z","submitted_at":"2026-06-21T21:41:34Z","title":"RigorBench: Benchmarking Engineering Process Discipline in Autonomous AI Coding Agents","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-01T07:12:07.099036Z"},"links":{"cited_paper":"/paper/2403.08604","citing_paper":"/paper/2606.22678"},"observation_digest":"sha256:0d269733fd7c20bb8c5fef696d5485212044de3ab5284bbdbe272fe3fbcc2b33","observation_id":"39a521c5-b9fd-4d12-abcb-c896b04549bc","resolution":{"observed_at":"2026-07-01T08:45:35.420148Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2403.08604/citation-record","integrity":"/paper/2403.08604/integrity","json":"/paper/2403.08604/citation-record.json","paper":"/paper/2403.08604"},"outbound":[],"paper":{"arxiv_id":"2403.08604","last_updated":"2024-12-14T09:45:51Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-06T13:32:36.381010Z","submitted_at":"2024-03-13T15:13:44Z","title":"Prompting Large Language Models to Tackle the Full Software Development Lifecycle: A Case Study"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 14 inbound Pith citation observations for arXiv:2403.08604."}