{"as_of":"2026-08-09T06:59:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:835c6fba3ac1eae04f380e3c9049a314bebae6b1e515305a4cf8bb1d53286715","coverage":[{"denominator":66,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":66,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T19:02:40.696853Z","state":"measured"},{"denominator":66,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":66,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.04682/citation-record","integrity":"/paper/2608.04682/integrity","json":"/paper/2608.04682/citation-record.json","paper":"/paper/2608.04682"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:42.028193Z","title":"ICLR , year=","venue":null,"work_id":"95bc2946-81a7-45db-932e-9eb98778d7cc","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.290931Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:e6411c950c23e1b4d06fc9e91fe6ed1f4a6685cc0d5c6cf35eee0e6efe5e8afa","observation_id":"0565f63a-f9be-4791-8c58-d2f804216905","resolution":{"observed_at":"2026-08-06T19:02:42.033982Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-06T19:02:40.308496Z","title":"arXiv preprint arXiv:2302.13971 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.308496Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:c34cfb2326d5ef9433573c3310f8a56e8cab995a8e67557547212471ab21498a","observation_id":"49516cd8-6ba3-43f7-b0fa-143c03450ce4","resolution":{"observed_at":"2026-08-06T19:02:40.308496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:42.010349Z","title":"ACM transactions on intelligent systems and technology , year=","venue":null,"work_id":"ed67f5a8-4f7f-41cb-a669-0bd940034f47","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.313909Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:325c0d60c3d6d9b613f9929805051532a8aa73e38c9a1ff02c3eaa61e37941c5","observation_id":"b9a61cc0-4534-45ca-80c1-abe1104f57b9","resolution":{"observed_at":"2026-08-06T19:02:42.015769Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.989366Z","title":"Frontiers of Computer Science , year=","venue":null,"work_id":"f3328172-a2c1-406f-8a25-ff859d48a194","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.319687Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:8f0d7a4fef8164590357ed7a31fa0688a6e6632cf6fdfc9986ce69ef61aee1f5","observation_id":"dadae06c-eb74-4022-a9c4-f036a2014744","resolution":{"observed_at":"2026-08-06T19:02:41.994927Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.325309Z","title":"ICSE-FoSE , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.325309Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:4131f1d15917ebf738dfcbde21f5a3056f367138b490ea76b07c039089b5593f","observation_id":"e81df24c-c8ba-4769-a303-c616ca9b4ce5","resolution":{"observed_at":"2026-08-06T19:02:40.325309Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.956015Z","title":"2024 , howpublished =","venue":null,"work_id":"14776416-a9fa-46d5-ae15-1f7e39242e07","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.336945Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:0c8de212b245127ccda861a53fadd393fe3a5e5c8a63ff8f588263782d7595ca","observation_id":"8d662c10-40b9-4050-970d-2e35e528b05c","resolution":{"observed_at":"2026-08-06T19:02:41.964031Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.938374Z","title":"2024 , howpublished =","venue":null,"work_id":"0ee82197-7525-4e6d-86d8-1e8e74f96509","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.342006Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:8fd8e4c393c377c9005a04009a18826905e3878c0b40cfd7feee4c4173e42a11","observation_id":"9115f5ca-b6f8-4f03-b6f3-a2b45d3ebe7d","resolution":{"observed_at":"2026-08-06T19:02:41.943464Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.921314Z","title":"2024 , howpublished =","venue":null,"work_id":"39e6b0b1-876c-456f-958b-86fafdefefe4","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.346986Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:564510e1326e387e456c6879549461eafc5cc5dfdf500292f05bb34f8bdd3634","observation_id":"891edd44-be57-4a68-8a60-b294434e3997","resolution":{"observed_at":"2026-08-06T19:02:41.927060Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.903476Z","title":"2026 , howpublished =","venue":null,"work_id":"ca2daf4f-12f4-4e7c-a6ea-90cfbe418ce1","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.351934Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:a4ebc3c00abda87f629544a25f25be6ff6c266a25d752124c8a48460d703c957","observation_id":"c9d8cead-7ff6-4820-b947-b63811c2c720","resolution":{"observed_at":"2026-08-06T19:02:41.908802Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.884439Z","title":"2026 , howpublished =","venue":null,"work_id":"34e93196-43d0-4ec1-b6c3-f566fe067950","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.356669Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:e555c8f30da06d850d63efc8fc8effa3831475ed28f3cbf47bf00c4990b35983","observation_id":"b30d7ecf-2130-4d80-9f68-d40884c106d2","resolution":{"observed_at":"2026-08-06T19:02:41.890389Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.864849Z","title":"NeurIPS , year=","venue":null,"work_id":"9b8b39d5-ba3b-45d2-b7db-827aeb39a94a","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.361764Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:f7423a2400820d63e2bbae04ce5e7584efec72dcfae04f6ff5d0d4f80cd2d591","observation_id":"85e92ddc-fa13-435f-915e-714e05b68d1c","resolution":{"observed_at":"2026-08-06T19:02:41.870937Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.846962Z","title":"NeurIPS , year=","venue":null,"work_id":"a188d8bf-f235-4ce3-b94a-e10707a713c5","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.372026Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:be338dd64f217a6c158fdf84767a133960f3c473f312227f75063f974629cc4b","observation_id":"0a0e2002-0825-4760-9fee-7dac66b10620","resolution":{"observed_at":"2026-08-06T19:02:41.852951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.06992","last_updated":"2024-10-10T13:13:09Z","snapshot_observed_at":"2026-08-08T05:32:33.029750Z","submitted_at":"2024-10-09T15:38:53Z","title":"SWE-Bench+: Enhanced Coding Benchmark for LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.06992","snapshot_observed_at":"2026-08-06T19:02:40.376652Z","title":"arXiv preprint arXiv:2410.06992 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.376652Z"},"links":{"cited_paper":"/paper/2410.06992","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:64d6dfef79a422a3898663a2f846a647b390041941ec7bc387c043ec40f8b50c","observation_id":"64ba4e39-b21a-4b0a-9c41-5cddd61a9a18","resolution":{"observed_at":"2026-08-06T19:02:40.376652Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.830062Z","title":"ACL , year=","venue":null,"work_id":"1a61c50c-a94c-465b-a57b-120d7809ac04","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.382378Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:0d6de1cde182e56b76a71e78ff836436ac3c2ea0c896db9ef6ce5ce8807b2ffd","observation_id":"6a627b34-1f43-409d-81a5-4827a2338a3c","resolution":{"observed_at":"2026-08-06T19:02:41.835097Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.813383Z","title":"NeurIPS , year=","venue":null,"work_id":"f5b61f48-9a66-4f59-be1c-fd1159cfd6d5","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.403082Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:8fe90d2014e76691c38bb126987763b6c09f385b83851edebecbcf97b5ee8c08","observation_id":"9a0930f9-0958-4d81-a76d-bd09a0406e57","resolution":{"observed_at":"2026-08-06T19:02:41.818320Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.413636Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.413636Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:2165b0ac3670b0e1b119782271dd06d669f77a43f08889a3bc5fc06ed0d8b851","observation_id":"83bd7b1c-0ffa-42f8-85af-8929782a01e3","resolution":{"observed_at":"2026-08-06T19:02:40.413636Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.784335Z","title":"ACL , year=","venue":null,"work_id":"1acb400d-2b03-4766-868a-b0b7bdf05a2c","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.418361Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:801d2fc38ffdcdb4b7f8ba4a8e18b66a50ac27451de8498df4e6aabec7dc09c5","observation_id":"1cfb059a-d1d4-4f63-bb5a-695d361d32a7","resolution":{"observed_at":"2026-08-06T19:02:41.790730Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.766474Z","title":"ICLR , year=","venue":null,"work_id":"eda0e4b2-253d-4e54-9e2a-06bc48726214","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.423733Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:eb36ea6edba5ccddec40d955be09fff1b65d9115dd1b0cb2c53939ede83eb83c","observation_id":"0619187b-1846-4afe-9b3b-eaa0d2de5f5a","resolution":{"observed_at":"2026-08-06T19:02:41.771010Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.749701Z","title":"ACL , year=","venue":null,"work_id":"db474a50-73ac-41a9-8777-85b1f10655fc","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.428427Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:1c27d0cc6cde01081f38353bf1b1c2905ada6ead9dc3db8a0facb61dde801a82","observation_id":"9a479890-d1b4-42f0-a840-d4b7fd6359c7","resolution":{"observed_at":"2026-08-06T19:02:41.754414Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.731244Z","title":"NeurIPS , year=","venue":null,"work_id":"b0a1b423-5145-4b6e-973b-c8c038a168bd","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.433649Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:17b2d57483de384c36ed039b2ef99e785b29b85e08ab562931af9b0f613811a5","observation_id":"d098cacf-1a18-4d77-99eb-6662078e2e95","resolution":{"observed_at":"2026-08-06T19:02:41.737619Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.709353Z","title":"2026 , howpublished =","venue":null,"work_id":"f1e7aa1c-dfa1-4ec8-a99f-0d1fcf700fb9","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.448812Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:00378e3bdba880334950f2076d9bd439f8895261e6acddfb3c38abc30a555a0b","observation_id":"6742c5c2-6174-4d6b-9714-f98aa2d79b7b","resolution":{"observed_at":"2026-08-06T19:02:41.714824Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.690635Z","title":"2026 , howpublished =","venue":null,"work_id":"bafc57fc-44c5-4d85-a3ac-6324ad51e34c","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.459486Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:6c8644ddee43c9e26fce877f358a284a709fcfc86c15b4ca2f0f45f41bcbe2c5","observation_id":"517c0c83-b80f-41fe-aaa4-0e8cc3177cfe","resolution":{"observed_at":"2026-08-06T19:02:41.696949Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.668298Z","title":"Qwen3.5: Accelerating Productivity with Native Multimodal Agents , howpublished =","venue":null,"work_id":"cf732404-8ad0-477e-851d-41bc56191711","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.474687Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:81b05f81bb68610b9e9b207df9564cfcc9b6aabeaec525631a35f66cc0a0e7b9","observation_id":"7fda5a46-3f01-4293-bf9e-341e083e61f5","resolution":{"observed_at":"2026-08-06T19:02:41.678185Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.479235Z","title":"2026 , howpublished =","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.479235Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:5ecb9e57f938a7ab84fd2c3afd72a4e51027fa60e2c102c62871635809dacf04","observation_id":"d437dc11-78b9-4574-80e2-b06b942efee2","resolution":{"observed_at":"2026-08-06T19:02:40.479235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.640573Z","title":"2026 , howpublished =","venue":null,"work_id":"ddf35601-d618-4fc7-841c-666443317641","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.484085Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:dad3b2d3d85f88fdeb24d5551917f32621407e53e7cd930da84ec87e248528a6","observation_id":"ba7c3479-1fa3-49a1-a12a-6b79a7f87839","resolution":{"observed_at":"2026-08-06T19:02:41.645910Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.624139Z","title":"2026 , howpublished =","venue":null,"work_id":"55dee3de-6ff6-4d09-9994-1eaadbe4e419","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.489325Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:0a4f606362721f998ce815279e0267b30da9d44ac7c470bf5543b07c8f71ffc7","observation_id":"dd58faa9-7394-48e2-81ce-f65cf8b8bdf0","resolution":{"observed_at":"2026-08-06T19:02:41.629172Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.509167Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.509167Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:811aaa7e7e556259594f51612f6d9a0479fa452b93d241b5ff653e07af294495","observation_id":"b0c83089-c242-4d7b-bfbd-0785c326eb3b","resolution":{"observed_at":"2026-08-06T19:02:40.509167Z","resolver_source":null,"status":"parse_uncertain"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.596249Z","title":"System card: Claude sonnet 4.6","venue":null,"work_id":"665e9a3e-b727-4567-99ce-265d44c4fbd2","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.513795Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:934cb9b7918c4f8e356b0558bd6921496158c4f4d6e776a605a9a259a1fe6a77","observation_id":"89c6e5dd-035d-4ee2-a1ce-842f3486327d","resolution":{"observed_at":"2026-08-06T19:02:41.601259Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.579587Z","title":"System card: Claude opus 4.8","venue":null,"work_id":"f29257e8-7b29-4df5-9477-a8cbc240a6e0","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.518702Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:2dbd580b6563200a63c19293f124984c39445266ba883aab9cb9e8de2be6cc5c","observation_id":"06eb38fb-7222-444f-b0d7-7abd67c27903","resolution":{"observed_at":"2026-08-06T19:02:41.585132Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.563749Z","title":"A survey on evaluation of large language models","venue":null,"work_id":"eb9c5f75-b36b-4c11-91ae-3b02d80917e8","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.523406Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:c648ce038bd33296b1eeed9a768fbe3d8295d69f372af78cf389e18b10951c85","observation_id":"d0187244-64c2-4c7c-9476-308866a0ec00","resolution":{"observed_at":"2026-08-06T19:02:41.568599Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.26494","last_updated":"2026-07-30T05:04:36Z","snapshot_observed_at":"2026-08-06T03:12:19.331370Z","submitted_at":"2026-05-26T03:16:11Z","title":"The MiniMax-M2 Series: Mini Activations Unleashing Max Real-World Intelligence","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.26494","snapshot_observed_at":"2026-08-06T19:02:40.528271Z","title":"The minimax-m2 series: Mini activations unleashing max real-world intelligence","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.528271Z"},"links":{"cited_paper":"/paper/2605.26494","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:74e2178aa7cc8ff0de9df62192f130da9ee4c9232f78899b080d0e70df8998a3","observation_id":"72be4632-bdf0-49e8-82ed-74b805e6b46b","resolution":{"observed_at":"2026-08-06T19:02:40.528271Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.546902Z","title":"External technical root cause analysis — channel file 291","venue":null,"work_id":"8905016b-c670-4154-877b-1d86541c1a42","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.532761Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:6ed52e0925b349788626b7bf9ac6b24a5a8b0357fa4b4b3504fbe74098f727be","observation_id":"6e8b870a-04a4-444f-b81c-ebaca12c6134","resolution":{"observed_at":"2026-08-06T19:02:41.552095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.529610Z","title":"Deepseek-v4: Towards highly efficient million-token context intelligence","venue":null,"work_id":"0c2a254f-6c21-4d16-b098-633525427a40","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.537453Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:66b3fbb017c2e53e7b4b1ee5da67be69b11987e2872daee27d06d324f15ef736","observation_id":"beb6dfa7-3e20-4817-85df-dfa4e8d756d7","resolution":{"observed_at":"2026-08-06T19:02:41.535235Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.16941","last_updated":"2025-11-14T22:00:03Z","snapshot_observed_at":"2026-08-06T21:34:46.041153Z","submitted_at":"2025-09-21T06:28:17Z","title":"SWE-Bench Pro: Can AI Agents Solve Long-Horizon Software Engineering Tasks?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.16941","snapshot_observed_at":"2026-08-06T19:02:40.542646Z","title":"Swe-bench pro: Can ai agents solve long-horizon software engineering tasks? arXiv preprint arXiv:2509.16941, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.542646Z"},"links":{"cited_paper":"/paper/2509.16941","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:cfc28a9a93a46595c76b1306bab6fa99ca1f6283aa3d375e8be6f3ef1818fd76","observation_id":"49bd99bc-1278-4d1d-ae47-14d10ac73729","resolution":{"observed_at":"2026-08-06T19:02:40.542646Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.512817Z","title":"Large language models for software engineering: Survey and open problems","venue":null,"work_id":"f40d98be-f376-4eb1-a32a-8436b9f21dbc","year":2023},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.547534Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:1e9d287415275110ff69ef15882fbc2fb7502f1836165fa36e7d88ba24293e38","observation_id":"128b5bd0-da9f-491c-bb5a-2f80a28b5e2e","resolution":{"observed_at":"2026-08-06T19:02:41.518388Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.497271Z","title":"Gemini 3.1 pro model card","venue":null,"work_id":"61a2ee32-58e1-4e64-a9b9-2f70f652d209","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.552364Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:9ff3bce71845ac7b26e60527abd388d2fcfd9575bbaa78f4af7a5044adffb231","observation_id":"d4703c88-aaa4-4104-9a54-5d37d6225d5c","resolution":{"observed_at":"2026-08-06T19:02:41.502015Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.13010","last_updated":"2024-05-24T11:47:24Z","snapshot_observed_at":"2026-07-06T17:05:56.282078Z","submitted_at":"2023-12-20T13:22:41Z","title":"AgentCoder: Multi-Agent-based Code Generation with Iterative Testing and Optimisation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.13010","snapshot_observed_at":"2026-08-06T19:02:40.557097Z","title":"Agentcoder: Multi-agent-based code generation with iterative testing and optimisation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.557097Z"},"links":{"cited_paper":"/paper/2312.13010","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:174c6e26d83f4137f185bc1ebac1a326bc3595f9b10f106411be9f6e58752bfc","observation_id":"904c9bb4-aaa4-46e1-9782-cd4833834b9d","resolution":{"observed_at":"2026-08-06T19:02:40.557097Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.481043Z","title":"Mapcoder: Multi-agent code generation for competitive problem solving","venue":null,"work_id":"ea0ec46d-8871-41a2-b1ee-117c9b04ee29","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.561483Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:11703d69e0dde67e94561930675164052fbabedca26c730d5433b6524bce200a","observation_id":"bfed2013-810d-438f-979a-4fccabf85356","resolution":{"observed_at":"2026-08-06T19:02:41.486321Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.462801Z","title":"Swe-bench: Can language models resolve real-world github issues? In ICLR, 2024","venue":null,"work_id":"101e9a78-45b5-4a15-aca7-65515e89fe8c","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.566348Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:60bf682691f8266bb326f6c87c8dfc970a8cd19fbf6e4eab7376f455f83c6575","observation_id":"753f1b08-6ee4-468e-8039-99aa741f11c0","resolution":{"observed_at":"2026-08-06T19:02:41.468934Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02479","last_updated":"2025-04-13T09:42:30Z","snapshot_observed_at":"2026-07-06T18:56:57.369673Z","submitted_at":"2024-08-05T14:01:15Z","title":"From LLMs to LLM-based Agents for Software Engineering: A Survey of Current, Challenges and Future","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02479","snapshot_observed_at":"2026-08-06T19:02:40.571286Z","title":"From llms to llm-based agents for software engineering: A survey of current, challenges and future","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.571286Z"},"links":{"cited_paper":"/paper/2408.02479","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:ae4dbc3b0a92fd80255ce873adc00295f52556485db674905b0eeeb34b887267","observation_id":"1b1f2c14-9151-43ab-bbf6-bc47a76ad8bf","resolution":{"observed_at":"2026-08-06T19:02:40.571286Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2606.15079","last_updated":"2026-06-13T03:21:49Z","snapshot_observed_at":"2026-07-06T23:52:28.557859Z","submitted_at":"2026-06-13T03:21:49Z","title":"Ling and Ring 2.6 Technical Report: Efficient and Instant Agentic Intelligence at Trillion-Parameter Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.15079","snapshot_observed_at":"2026-08-06T19:02:40.575847Z","title":"Ling and ring 2.6 technical report: Efficient and instant agentic intelligence at trillion-parameter scale","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.575847Z"},"links":{"cited_paper":"/paper/2606.15079","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:a06328056285b19f02fd25d00e741c131df0fe4eab38ced7c17e735347baec12","observation_id":"764ae076-5e9f-4c84-a758-364f9bfbb011","resolution":{"observed_at":"2026-08-06T19:02:40.575847Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2602.10471","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.998397Z","title":"Testexplora: Benchmarking llms for proactive bug discovery via repository-level test generation","venue":null,"work_id":"92b91b58-d47f-45bd-b17f-1d4912341b58","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.580855Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:4080267299fd09a45586a4e8837620cb7e47a633fc96c809afbf8032cad11f18","observation_id":"f5745e75-a471-43d7-ba00-7f27b1a049dc","resolution":{"observed_at":"2026-08-06T19:02:41.009629Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.446354Z","title":"Helping our customers through the crowdstrike outage","venue":null,"work_id":"44382069-3647-47db-8052-fb318fac8c5c","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.585756Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:8d5c79f12d3513b4ce4a17791d4c6f24f2258e95e4c4712950ceab36858bca90","observation_id":"60060c8c-6308-4b6d-abe2-43168cc9d958","resolution":{"observed_at":"2026-08-06T19:02:41.451448Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.429637Z","title":"u ndler, Mark N M \\","venue":null,"work_id":"1a022e9b-2db3-456d-b5c6-1370b594a626","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.591162Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:77d052b88599a584cfdd5ceab5a4bc53e531ec4ba6ed93d781c5890913aa084b","observation_id":"24b7c483-1b3f-4c62-a17b-e3ea15790a6d","resolution":{"observed_at":"2026-08-06T19:02:41.435745Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.413281Z","title":"Why swe-bench verified no longer measures frontier coding capabilities","venue":null,"work_id":"61a376c1-015b-437f-8797-304a6160b1a7","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.595627Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:a3ab5c03cc863710f4a9fbbb5146a4e6de8802b9c3be7017c3f6d8b0a04f2541","observation_id":"fac8908f-b69b-4567-8be1-3f35844fff57","resolution":{"observed_at":"2026-08-06T19:02:41.418435Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.396668Z","title":"Separating signal from noise in coding evaluations","venue":null,"work_id":"ba9ab390-b7d1-4b02-b4e7-1e4904eb0560","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.600347Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:2f01a6fb5dfb70cff5f23d1a1a0e3e357d595017058441fa0a6e48314af96b7f","observation_id":"06da3624-803e-471c-8eb2-b598b707635a","resolution":{"observed_at":"2026-08-06T19:02:41.402190Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.377276Z","title":"Crowdstrike to cost fortune 500 \\ 5.4b; insured loss range of \\ 0.54b - \\ 1.08b","venue":null,"work_id":"111ae049-b901-4558-90a9-af626ba8eb7a","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.605278Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:164d00b4f9bdac084c0046d0a6330bcd9a3c33004b6d1f6f6168ad766189db8d","observation_id":"d813e5dd-8a68-49dd-abdd-1a9943392915","resolution":{"observed_at":"2026-08-06T19:02:41.383196Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.360783Z","title":"Qwen3.7 : The agent frontier, May 2026","venue":null,"work_id":"a25b81af-0b60-4c68-b803-3efd631ad40c","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.609842Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:8be757b1b261109e9586f2ead0ab64358bc194e232605c4c03f4918015e75acc","observation_id":"25ba78ca-7e44-4a4c-a537-292814a882fc","resolution":{"observed_at":"2026-08-06T19:02:41.365835Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.08366","last_updated":"2026-05-08T18:21:44Z","snapshot_observed_at":"2026-07-06T23:20:38.385189Z","submitted_at":"2026-05-08T18:21:44Z","title":"SWE Atlas: Benchmarking Coding Agents Beyond Issue Resolution","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.08366","snapshot_observed_at":"2026-08-06T19:02:40.614800Z","title":"Swe atlas: Benchmarking coding agents beyond issue resolution","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.614800Z"},"links":{"cited_paper":"/paper/2605.08366","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:2ddef9a9fa37c549c174d9b52e2ae65777d619633bf65ca1af411851b0790dd1","observation_id":"36983989-f6c6-44e7-9295-216d3effa092","resolution":{"observed_at":"2026-08-06T19:02:40.614800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2607.00248","last_updated":"2026-06-30T22:57:43Z","snapshot_observed_at":"2026-08-03T15:00:12.134373Z","submitted_at":"2026-06-30T22:57:43Z","title":"Seed2.0 Model Card: Towards Intelligence Frontier for Real-World Complexity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2607.00248","snapshot_observed_at":"2026-08-06T19:02:40.619737Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.619737Z"},"links":{"cited_paper":"/paper/2607.00248","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:5aa4da1858255c3aa700dcb3c79d02958768cd8070f22b2bf0f4f1335f4f1073","observation_id":"7e520ca1-e326-40d5-9a5d-33f4815a7fba","resolution":{"observed_at":"2026-08-06T19:02:40.619737Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2601.03267","last_updated":"2026-05-01T23:55:43Z","snapshot_observed_at":"2026-08-02T10:52:10.211700Z","submitted_at":"2025-12-19T07:05:38Z","title":"OpenAI GPT-5 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.03267","snapshot_observed_at":"2026-08-06T19:02:40.624433Z","title":"Openai gpt-5 system card","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.624433Z"},"links":{"cited_paper":"/paper/2601.03267","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:91face91d107bfa7a4af8f8c8f947d240c2dd5cdfc81eb694b442d27c4bbd97a","observation_id":"6bff20fc-1e23-499d-aef2-306784211d9f","resolution":{"observed_at":"2026-08-06T19:02:40.624433Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2602.02276","last_updated":"2026-02-02T16:17:38Z","snapshot_observed_at":"2026-07-06T22:44:09.804048Z","submitted_at":"2026-02-02T16:17:38Z","title":"Kimi K2.5: Visual Agentic Intelligence","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2602.02276","snapshot_observed_at":"2026-08-06T19:02:40.629242Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.629242Z"},"links":{"cited_paper":"/paper/2602.02276","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:737cb01975ac44bcf0cd27c5ae1167d5c2713b87e9681aea2060edb3f8e016c3","observation_id":"4166d0b0-61c3-460c-9d67-305dae3fa442","resolution":{"observed_at":"2026-08-06T19:02:40.629242Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.344226Z","title":"Qwen3.5: Accelerating productivity with native multimodal agents","venue":null,"work_id":"3aa5193a-d358-4e78-bc85-ea2bac52587d","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.633864Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:d824dd529b42d359757d4d746d6bfe632e0de57d4e725fcf149fdb85d587a973","observation_id":"888e5023-ee2f-4712-9181-fc809203c592","resolution":{"observed_at":"2026-08-06T19:02:41.349448Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.326988Z","title":null,"venue":null,"work_id":"537e14a4-282f-41c1-9b47-0d68a63b0819","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.638972Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:7172075b4a03261bdc3df5ef5ca8c9afe97051b92b00d0727bdd69fec036326a","observation_id":"8d751679-314d-461e-9cb2-18c912352ff8","resolution":{"observed_at":"2026-08-06T19:02:41.332028Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.310448Z","title":"Openhands: An open platform for ai software developers as generalist agents","venue":null,"work_id":"e82b0409-6ffa-4a44-972a-2d4db4c08fe9","year":2025},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.643906Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:f01d7a8daf8ce184f1050906af26f0e84037b014304d400bf4cca9d3eaabf19e","observation_id":"4913b48a-561f-4138-9269-3c513df3004f","resolution":{"observed_at":"2026-08-06T19:02:41.315666Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.648923Z","title":"Live-swe-agent: Can software engineering agents self-evolve on the fly? arXiv preprint arXiv:2511.13646, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.648923Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:adf7a30e65388ba748316ed9e76c05c50e203bfdce26d7238e947963b0ef8446","observation_id":"c95108be-bc6a-4c3f-9943-0e1ad7feaf8f","resolution":{"observed_at":"2026-08-06T19:02:40.648923Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.293219Z","title":"Swe-agent: Agent-computer interfaces enable automated software engineering","venue":null,"work_id":"0d6af1be-81d9-4e1b-9817-a87ec64abf3e","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.653538Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:1f231e71e655aadcdf835d4a8932de65c1afa14e2ae8cf5a0714399b4a528be5","observation_id":"70896e99-4ea8-4a36-aaed-baf635d8ea05","resolution":{"observed_at":"2026-08-06T19:02:41.298970Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03859","last_updated":"2024-10-04T18:48:58Z","snapshot_observed_at":"2026-07-06T19:28:08.374354Z","submitted_at":"2024-10-04T18:48:58Z","title":"SWE-bench Multimodal: Do AI Systems Generalize to Visual Software Domains?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.03859","snapshot_observed_at":"2026-08-06T19:02:40.658366Z","title":"Swe-bench multimodal: Do ai systems generalize to visual software domains? arXiv preprint arXiv:2410.03859, 2024 b","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.658366Z"},"links":{"cited_paper":"/paper/2410.03859","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:7c97f4f3424819f2e25253187f7e2101a05b33f98b95ce2cd3b272c87d254d1d","observation_id":"474521f1-6871-4316-9674-9fed3c956337","resolution":{"observed_at":"2026-08-06T19:02:40.658366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.275494Z","title":"Swe-smith: Scaling data for software engineering agents","venue":null,"work_id":"a2bbbae5-0de7-44c2-8cf5-003ff227f24e","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.662708Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:5893a80ccc3e6084f633cffd9213c7d2cc746e3e6e4962e172c22a13acc28385","observation_id":"b3543968-4e74-4956-8813-d4946b3df476","resolution":{"observed_at":"2026-08-06T19:02:41.280645Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.03629","last_updated":"2023-03-10T01:00:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-10-06T01:00:32Z","title":"ReAct: Synergizing Reasoning and Acting in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.03629","snapshot_observed_at":"2026-08-06T19:02:40.667485Z","title":"React: Synergizing reasoning and acting in language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.667485Z"},"links":{"cited_paper":"/paper/2210.03629","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:2e6bf513832c7a33f03e451f6ff6e32660742d2f0b6fa85fdfe17cf865dd3da0","observation_id":"f28180df-8257-44fc-aa16-5f74e3e2d7fe","resolution":{"observed_at":"2026-08-06T19:02:40.667485Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.258952Z","title":"Multi-swe-bench: A multilingual benchmark for issue resolving","venue":null,"work_id":"66172575-dd52-4e0d-aa5f-8af50d5c1163","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.672186Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:dbe90e388d923eee5a5d72793b37663bcd4be397aa83b3dd1d25e9b11aaeadf3","observation_id":"1d73afb6-5baf-41b4-8e2e-f8662083c1a6","resolution":{"observed_at":"2026-08-06T19:02:41.264160Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2602.15763","last_updated":"2026-02-24T10:44:44Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-17T17:50:56Z","title":"GLM-5: from Vibe Coding to Agentic Engineering","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2602.15763","snapshot_observed_at":"2026-08-06T19:02:40.677154Z","title":"Glm-5: from vibe coding to agentic engineering","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.677154Z"},"links":{"cited_paper":"/paper/2602.15763","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:2e4eda11504435c99d2878ff40b157665722d8c2c2f2837533553db06df2c694","observation_id":"9b55618c-b92d-4c42-a61b-1751e94b601e","resolution":{"observed_at":"2026-08-06T19:02:40.677154Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.241313Z","title":"Codeagent: Enhancing code generation with tool-integrated agent systems for real-world repo-level coding challenges","venue":null,"work_id":"169da89d-a4af-459d-a902-e2dd8dd405c3","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.681643Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:b53b776c728c6eb0fdd2e35fb483308f851247566990782385a9b6d4b4c418f1","observation_id":"37d59c2f-702c-447e-a703-8d4c3d9a7959","resolution":{"observed_at":"2026-08-06T19:02:41.247502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.224458Z","title":"Swe-bench goes live! In NeurIPS, 2026","venue":null,"work_id":"830b05dd-1e22-4148-8d5f-eb93ac5163ce","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.686763Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:6477a72bf4eed9c5f39c127b70ba468119e6277db4a9a4971a2ff02581844b3a","observation_id":"772dcec3-c004-4870-91db-1d8927a27f78","resolution":{"observed_at":"2026-08-06T19:02:41.229614Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.207255Z","title":"A survey of large language models","venue":null,"work_id":"8b5fb974-b079-4cc0-a371-57e0f08c9c77","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.691634Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:7b07b5c432998e422acb0b7faf7df1ee86e6f19fe65045079ba0793fa6afcdc1","observation_id":"8eebe9aa-1454-4e4b-b045-82b5beb7fa9f","resolution":{"observed_at":"2026-08-06T19:02:41.212617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.696853Z","title":"Featurebench: Benchmarking agentic coding for complex feature development","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.696853Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:a35a6d2dce2b07b0e1b53c6b9ad10c08db77fecdd47193b60ddddfbf0b67b533","observation_id":"8e9f7764-9d6f-4452-9c46-28fe032d255a","resolution":{"observed_at":"2026-08-06T19:02:40.696853Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","latest_version":1,"primary_category":"cs.SE","snapshot_observed_at":"2026-08-08T23:12:16.855136Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports"},"reference_resolution":{"displayed":66,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":1,"unresolved":20,"verified_exact":1,"verified_fuzzy":44},"total_outbound_references":66},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 66 of 66 outbound references and 0 inbound Pith citation observations for arXiv:2608.04682."}