{"as_of":"2026-08-10T12:46:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:cb0147ed6339b4c60eb4db25760125d8da7641c1cecbf20b06c32e534f6ed879","coverage":[{"denominator":98,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":98,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T23:35:06.211001Z","state":"measured"},{"denominator":98,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":98,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.17644/citation-record","integrity":"/paper/2506.17644/integrity","json":"/paper/2506.17644/citation-record.json","paper":"/paper/2506.17644"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.251099Z","title":"Cybercrime To Cost The World $10.5 Trillion Annually By","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.251099Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:06fc9e313870a69033d27f590fa030c8a995656e5066b64c9f3e83b9d29a2177","observation_id":"9cb6e339-fc42-43e8-9cf1-a0c944bf33a1","resolution":{"observed_at":"2026-08-06T23:34:57.251099Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.504753Z","title":"AI Cyber Challenge Opens Registration, Adds $4 Million in Prizes, Shows Scoring Algorithm and Challenge Exemplar","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.504753Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:5c8c22e48c8bfa9e731e3047e1ea5996058573bc9acacd2df964140b9d318756","observation_id":"ba61912d-7a22-4695-9516-9d12112eb63e","resolution":{"observed_at":"2026-08-06T23:34:57.504753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.662253Z","title":"DEF CON®27 Hacking Conference Contests & Events","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.662253Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:abb12761602173ef57c2b3bb7f4ef0d09d7b8f7730ece7d69bfaf20346d54188","observation_id":"f12cf11d-1a26-4238-a8e1-d2de158ea976","resolution":{"observed_at":"2026-08-06T23:34:57.662253Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.854753Z","title":"0CTF 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.854753Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:55d6fdd6cf5077c89c35c83f14bf97d9bb902a949177da91e1b9e67147cefc26","observation_id":"260f2c39-988b-4727-85b5-939ad76ef832","resolution":{"observed_at":"2026-08-06T23:34:57.854753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.014833Z","title":"All about CTF","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.014833Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:5a4b5365c54ef897ba863e462539fb6b1048a67579cae35d7f34ad8344e48751","observation_id":"19d0120f-3d35-42f7-baf5-04eb405c69bc","resolution":{"observed_at":"2026-08-06T23:34:58.014833Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.064927Z","title":"Assistants API Overview","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.064927Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:631c2e5e61fc838bc8cabb62bd277cd04dec56f36674e063fa0b9bd3d928ff9f","observation_id":"eb70e278-d662-4f74-882f-2477daf33024","resolution":{"observed_at":"2026-08-06T23:34:58.064927Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.184386Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.184386Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:8976cf259d3ca503cc7f48e148b62d24654c4bcaaa1ed43151b7105654d8ec50","observation_id":"3b083d52-2508-49a3-b0de-b6a6a6f3e9d3","resolution":{"observed_at":"2026-08-06T23:34:58.184386Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.301890Z","title":"Capture the Flag","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.301890Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:224c87dcce62799094c6bb07aa8f827a89dbf8d46dae1dc3b120d508270363aa","observation_id":"650dc15d-d2a8-41e6-af98-bbe8750a80c2","resolution":{"observed_at":"2026-08-06T23:34:58.301890Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.476834Z","title":"Capture the Flag for Empowered Cybersecurity Training","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.476834Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:e46d0f6b494a254e3c12b53476e03203c790c0f75a7d2e6ae9fb6c51d56b983b","observation_id":"6a65862f-cbc7-458c-9a5d-dbcf4066deb7","resolution":{"observed_at":"2026-08-06T23:34:58.476834Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.591413Z","title":"CGC: Cyber Grand Challenge","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.591413Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:5ff3f2b9950d44e54157f978203fc348abd8d128bcf0204f37940c126ce8c8dd","observation_id":"e77e7ea6-cd15-434c-ab16-51c38dfd96f3","resolution":{"observed_at":"2026-08-06T23:34:58.591413Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.724750Z","title":"Claude 3.5 Sonnet","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.724750Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d25caf6c39b444b1d012f989758e56dbac4cc243c0544d3abf06390a8ab6a2d7","observation_id":"a79cd1ce-01cf-4ae7-bfe8-bdecf9432220","resolution":{"observed_at":"2026-08-06T23:34:58.724750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:58.866687Z","title":"DeepSeek","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:58.866687Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:166a13e1be35754a62a68ef8c00cd0c25702b6535d9298df510d6cd52275bd84","observation_id":"5273ab2a-785d-4aa5-b7fd-91bd215ac2af","resolution":{"observed_at":"2026-08-06T23:34:58.866687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.000038Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.000038Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:1823b76a86141738fbc873acf96e89b81ae1c3d812c53a48293bb8b501d2615a","observation_id":"8805a6da-5bd5-40b8-86fa-27a117f7be0f","resolution":{"observed_at":"2026-08-06T23:34:59.000038Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.114155Z","title":"Function calling","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.114155Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c8e9690cb172c53be23316dee32f89ffa697c59af386605661d023d0a0b11274","observation_id":"5c07c62d-178a-4c67-ae97-44f8938b5f72","resolution":{"observed_at":"2026-08-06T23:34:59.114155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.248142Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.248142Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d6dafabad3792ee94ff4fbc99e65f3c9bac54809c7b6d004cd387fd959c7f4d0","observation_id":"2e37072a-33d4-4543-92f1-777a2f2ff43e","resolution":{"observed_at":"2026-08-06T23:34:59.248142Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.347005Z","title":"Google CTF","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.347005Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:2a7beacebd961137331e65958f59efb0f307d87b5d193f2d147477e49b7a3fb3","observation_id":"1c1b7b75-ee7f-4ec1-9a59-35bbc5a2d63b","resolution":{"observed_at":"2026-08-06T23:34:59.347005Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:59.451941Z","title":"gpt-3-5-turbo","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.451941Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:fa884e0eb7fcff361952b3d78eb963b056159739aa3cbf1791927b980e060305","observation_id":"79cc5202-a3f4-4b7b-be62-a79c829e95da","resolution":{"observed_at":"2026-08-06T23:34:59.451941Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.676492Z","title":null,"venue":null,"work_id":"da1f0611-f508-47fb-97a0-14353b598398","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.593056Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:47fd80db34290ddf626907a0afe49fe04aac7571a9ed4b43471d3d8a1490d304","observation_id":"5bc6a0a4-3c28-4eed-9afd-a378298a3947","resolution":{"observed_at":"2026-08-06T23:35:09.697153Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.655777Z","title":null,"venue":null,"work_id":"d4eb6bb9-0dfb-46bd-8499-a78d9196b8b2","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.729431Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:a58bcb96a3707d90565cf11f8ba00447eecb3f0fe276381c4fd72def25665b7c","observation_id":"d5d2d54e-f4bf-4255-a601-a8b5a6798d2a","resolution":{"observed_at":"2026-08-06T23:35:09.659284Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.612871Z","title":null,"venue":null,"work_id":"44e0ca21-805e-4132-a426-e8041a2fbda1","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:59.875473Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:ec3dc41a80b521e3c8541707079d8efa93e37a96bf08bd5a423b892f9b0343f8","observation_id":"0a1f6a03-2465-4a06-82ba-b7270204573a","resolution":{"observed_at":"2026-08-06T23:35:09.645896Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.559733Z","title":null,"venue":null,"work_id":"5c45fd35-bb68-4a64-bacf-003dbea1dcb1","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.014973Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:af16c48b3c4d6d95ff7f71e3540b93ee3a8bb13045e581c87ec9ee1d4febe8c9","observation_id":"b68f27cb-b5d1-42a4-a925-ec0c07d4b4de","resolution":{"observed_at":"2026-08-06T23:35:09.588364Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.542639Z","title":"Learning to Reason with LLMs | OpenAI","venue":null,"work_id":"8cf0950c-197b-4e0e-8b50-1704cc92ae2f","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.112630Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:9dc4268ecc34f52461c85861c5d72caedb7ab98461430c229deaba9df75cc21a","observation_id":"1a3b9d3c-cb7b-48b1-a577-74f6f5aaa6c9","resolution":{"observed_at":"2026-08-06T23:35:09.546629Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.493438Z","title":"Meet Llama 3.1","venue":null,"work_id":"d9b38911-715b-404d-8655-8717d3c87640","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.226552Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d9537acd831d67c9a8c6786a5acc119b43dc5ab708dc99bcff614fbf18fa8fef","observation_id":"2b86e511-6854-4305-bf91-9d42170b1dc6","resolution":{"observed_at":"2026-08-06T23:35:09.514750Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.430441Z","title":"Mixtral of experts | Mistral AI | Frontier AI in your hands","venue":null,"work_id":"fac336c1-648b-496f-9265-5c9313ed993c","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.273456Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c11aa009e39324dd58dd768ce5144d68b8545ff758067ffa76847f256c574612","observation_id":"464592a9-cb53-4b38-8bf2-39527696ec49","resolution":{"observed_at":"2026-08-06T23:35:09.446817Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.346944Z","title":"picoCTF - CMU Cybersecurity Competition","venue":null,"work_id":"727d7fa2-5a56-4325-ace9-0004a2d626eb","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.491515Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:910de88c6d349f189276f03ab99db2bc94b87a03f05ebb221a7061b31e205916","observation_id":"19c2bea4-f1c4-43f8-8886-11d25172a832","resolution":{"observed_at":"2026-08-06T23:35:09.374754Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.328547Z","title":"picoCTF2024","venue":null,"work_id":"cf390a1f-3559-4abe-8a75-c61a2ba081ae","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.653650Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:8b56da94866cc1baa0edd6b1c6a0b27a3caecba5152c787b297399696959e987","observation_id":"8896b226-bfba-4220-b710-40b80568e4fa","resolution":{"observed_at":"2026-08-06T23:35:09.333649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.288503Z","title":"Top 10 Cyber Hacking Competitions - Capture the Flag (CTF)","venue":null,"work_id":"63533bac-dd21-4c93-8029-122a8b2c914b","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.761007Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:1473decd4b064dbd558db8cd2ebbc25ecbdbf53b942021d0c107ba9829391878","observation_id":"37147e23-699d-4064-acb6-337f0f9dfada","resolution":{"observed_at":"2026-08-06T23:35:09.306715Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.260365Z","title":"UIUCTF 2024","venue":null,"work_id":"66b2028a-a6c8-497f-bc4f-ebc15553f390","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:00.886621Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d01634812bed55ac9b7e963753b6104530fe4497259ac58cf71a434f81d1e1b2","observation_id":"7cf06032-3ce7-46fa-bb71-bc4b413962b1","resolution":{"observed_at":"2026-08-06T23:35:09.266291Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.242739Z","title":"VicOne & Block Harbor Spearhead Biggest Automotive Cap- ture the Flag Competition for Cybersecurity Enthusiasts World- wide","venue":null,"work_id":"c65f6413-c878-4ebe-a720-f7a932e8b344","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.033315Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:5f460f470acf95565faf2460f2db3c5a68291bd96aa8747fcb7d600b4a444ec4","observation_id":"5bff1ced-39b7-49e3-ac64-2a1d862a56e1","resolution":{"observed_at":"2026-08-06T23:35:09.246250Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.217536Z","title":"Burp Suite - Application Security Testing Software","venue":null,"work_id":"ce584f39-f6db-4c0b-8f1f-73719f03ac89","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.118662Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d6e4851214afc668be6c770d6df42831b916a0e9999a25b496b656e4a2a349f8","observation_id":"31de45ff-2047-459e-bb53-ff29266a8b9a","resolution":{"observed_at":"2026-08-06T23:35:09.222551Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"kql-and-b/4390932","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:07.938631Z","title":"Microsoft Security Copilot Blog","venue":null,"work_id":"e035d54f-1729-40d2-8f2c-12d33508ad12","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.231335Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:930535947c594cbb4eb516a020f9b677d86c168bad3611618775d234b8110793","observation_id":"062a5e70-6c47-46c7-9bff-31b258fbe3c3","resolution":{"observed_at":"2026-08-06T23:35:07.980525Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.184445Z","title":"Proactive Defense: The Role of Offensive Security in Cybersecurity","venue":null,"work_id":"5a584eff-0baf-473e-b9c3-50911b3637bb","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.325726Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:369408ac549f72b01cf9edf1d7dae24047b244fbc522cf39d9354b185a449dbb","observation_id":"a8e0dc8b-8b0b-45b5-bf40-a2c63f1d6465","resolution":{"observed_at":"2026-08-06T23:35:09.197784Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.149674Z","title":"Using AI for Offensive Security","venue":null,"work_id":"2df61746-3392-4ff8-ad98-f18964352f65","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.437033Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:cc2dc065eb571eedaf94baddc380bd09815f3896ed3b35bdce90338a76a25d28","observation_id":"6f9759e6-fe31-4284-95ce-e4e9de9dcdbc","resolution":{"observed_at":"2026-08-06T23:35:09.167578Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.139841Z","title":"What is Automated Vulnerability Remediation? https: //www.sentinelone.com/cybersecurity-101/cybersecurity/what-is-automated- vulnerability-remediation/","venue":null,"work_id":"1c666f32-c9ea-41bd-9160-f6bbf203d1f1","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.552239Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:5536ff7d84a15a77e43d6b48fbab79889e663bbf91ecadeb56bf1dd402d39215","observation_id":"b82a640e-60d4-45f1-abc8-35f5617842ca","resolution":{"observed_at":"2026-08-06T23:35:09.143171Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-06T23:35:01.653994Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.653994Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f4dc0c720bf697f968630fd979ab13a00004294634acb042241891972c660ae5","observation_id":"58b32c76-db2a-4c7c-bf97-38495c49041b","resolution":{"observed_at":"2026-08-06T23:35:01.653994Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.123478Z","title":null,"venue":null,"work_id":"a0a0fc87-6bf1-45f9-9b65-1390bc727cec","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.764451Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:a82cfd752c3ec4a9670e1238b8e8a15f14cf5d7335c296ba754dafc16fa797a0","observation_id":"f8604f7e-7da3-4be3-8670-960caaf3f1b2","resolution":{"observed_at":"2026-08-06T23:35:09.127689Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.13161","last_updated":"2024-04-19T20:11:12Z","snapshot_observed_at":"2026-08-07T13:01:36.699787Z","submitted_at":"2024-04-19T20:11:12Z","title":"CyberSecEval 2: A Wide-Ranging Cybersecurity Evaluation Suite for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.13161","snapshot_observed_at":"2026-08-06T23:35:01.878295Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:01.878295Z"},"links":{"cited_paper":"/paper/2404.13161","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:931df1418b5e418e933b547c1dfe1118fc2236a182b0b854a2df7f0d8aa200b1","observation_id":"0cbcfb12-1d95-4c42-a6f6-0988bfc48c0e","resolution":{"observed_at":"2026-08-06T23:35:01.878295Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.103479Z","title":null,"venue":null,"work_id":"87bb854f-1186-4341-9e83-c60c3028e5e5","year":2018},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.008375Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:e6c13a0f5310d0d902dd76f21111ae00f496c267f204af7571175881207cda9b","observation_id":"4859e62d-cb3b-49a2-91b7-b3416159ae50","resolution":{"observed_at":"2026-08-06T23:35:09.106683Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.092331Z","title":null,"venue":null,"work_id":"4ba4a262-c7d5-4114-84d3-ee8cae9051b4","year":2017},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.138798Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c5c2cc02dbbffa836da55cedfd1a02dce3892d7a1641815c2f1c3be0cde41009","observation_id":"2b9a2b1c-2705-48bf-9dec-bef8021c1d37","resolution":{"observed_at":"2026-08-06T23:35:09.096009Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.06573","last_updated":"2024-07-09T06:07:45Z","snapshot_observed_at":"2026-08-01T19:24:09.434113Z","submitted_at":"2024-07-09T06:07:45Z","title":"LLM for Mobile: An Initial Roadmap","version":1},"cited_work":{"arxiv_id":"2407.06573","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.06573","snapshot_observed_at":"2026-08-06T23:35:07.646513Z","title":"LLM for Mobile: An Initial Roadmap","venue":"cs.SE","work_id":"ba8a34ae-e6a1-4a24-97cf-c6fd93ba7aed","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.325227Z"},"links":{"cited_paper":"/paper/2407.06573","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:6d79234595131f63b804ee41d1cede9adc85e8d58a88bcb6ce14ddebfd77ec13","observation_id":"30f5794f-077c-488f-a7f0-5d7da8de41aa","resolution":{"observed_at":"2026-08-06T23:35:07.675542Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.19470","last_updated":"2025-09-23T03:45:42Z","snapshot_observed_at":"2026-08-10T03:47:52.773511Z","submitted_at":"2025-03-25T09:00:58Z","title":"ReSearch: Learning to Reason with Search for LLMs via Reinforcement Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.19470","snapshot_observed_at":"2026-08-06T23:35:02.407420Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.407420Z"},"links":{"cited_paper":"/paper/2503.19470","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:2ba4df9d6e4c4aba3a157f9e7b3b8f354d95bd4d8d8e60ba39768e37f8cb6885","observation_id":"243b50d8-72a2-4196-b8dc-409c83a4a933","resolution":{"observed_at":"2026-08-06T23:35:02.407420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.072403Z","title":null,"venue":null,"work_id":"5a472660-539a-4b31-9398-02557a3abed7","year":2014},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.488945Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:2205a7b4134d531264de783ce00da40c7d2dc2c76098b41cc8c9e409c9ea6ff0","observation_id":"24fd8075-1986-4e33-a9e0-0f41bff2310d","resolution":{"observed_at":"2026-08-06T23:35:09.084729Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.061883Z","title":null,"venue":null,"work_id":"f435c73f-fdd2-4418-a243-d07b8282fee6","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.559690Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:dfc6053ba7f3f0c84a6d280d2d39c4de724b885b6525df4a83356f58f3e39bc2","observation_id":"bf25768b-f992-4bdd-bf90-a5a351051373","resolution":{"observed_at":"2026-08-06T23:35:09.066019Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:09.039717Z","title":null,"venue":null,"work_id":"2b9d6a1a-9b1b-4293-bcd0-a7075442872e","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.650297Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:aab0670eb0ca1b259bc403b90138827f91b6b466ce859505eeadd98d3accd59c","observation_id":"fb07652c-ed80-4ed3-962e-7e00ff3ae398","resolution":{"observed_at":"2026-08-06T23:35:09.047651Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:02.788629Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.788629Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:b230a896f83a748ad785f149ba0c6a75ab129632d30427d5b92405ba29e600d3","observation_id":"19ccf574-abf8-469d-9c9b-810527772f15","resolution":{"observed_at":"2026-08-06T23:35:02.788629Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.984755Z","title":null,"venue":null,"work_id":"6baaebdb-5f85-4a20-ad3f-760fb78cc723","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.909045Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:de77cdd961ede1bed1cbbd7f3ab3ea494b193879abda517b939b4f46ee2e0fb4","observation_id":"34fb4ab3-b49e-404c-a45a-cff2061fd3cb","resolution":{"observed_at":"2026-08-06T23:35:08.989382Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01637","last_updated":"2025-03-30T00:26:48Z","snapshot_observed_at":"2026-08-08T04:41:33.967516Z","submitted_at":"2024-06-02T16:25:26Z","title":"Teams of LLM Agents can Exploit Zero-Day Vulnerabilities","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01637","snapshot_observed_at":"2026-08-06T23:35:02.972447Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:02.972447Z"},"links":{"cited_paper":"/paper/2406.01637","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:a9183337974835b72d8833e068b7893f31e54775b5044a69d4c7521cfe23f964","observation_id":"52263a2f-b4c7-4518-acc0-075488ea41d0","resolution":{"observed_at":"2026-08-06T23:35:02.972447Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.949958Z","title":null,"venue":null,"work_id":"b869abfd-431b-4826-8c30-d23d09639b01","year":2007},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.105682Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:160e4a18bdfc39e6ccb045a35cef0f5c9c67f79e32c42b83b11d201bc977c60c","observation_id":"c854178b-4f06-428b-90eb-9b4760c2e97c","resolution":{"observed_at":"2026-08-06T23:35:08.953936Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.10997","last_updated":"2024-03-27T09:16:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-18T07:47:33Z","title":"Retrieval-Augmented Generation for Large Language Models: A Survey","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.10997","snapshot_observed_at":"2026-08-06T23:35:03.279545Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.279545Z"},"links":{"cited_paper":"/paper/2312.10997","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:8963461832fd73742ba003c97dfec66d9034723348fa327ab51c48318550d19d","observation_id":"e4b21f42-151b-42ca-8c60-d815a75b1025","resolution":{"observed_at":"2026-08-06T23:35:03.279545Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12420","last_updated":"2023-12-22T14:07:16Z","snapshot_observed_at":"2026-07-06T16:50:26.807430Z","submitted_at":"2023-11-21T08:20:39Z","title":"How Far Have We Gone in Vulnerability Detection Using Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12420","snapshot_observed_at":"2026-08-06T23:35:03.441508Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.441508Z"},"links":{"cited_paper":"/paper/2311.12420","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:35bb4ada4366723c3ac0d54b3eb07842ac98db95f6ee494662cd6bc1a768aea6","observation_id":"8e153b3b-cf27-45dd-9f3f-4cf8fdfec779","resolution":{"observed_at":"2026-08-06T23:35:03.441508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.910692Z","title":null,"venue":null,"work_id":"00092062-f5b4-4583-9d15-3f640165c96f","year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.634822Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:5f16398b06a79a710ba10c0fc8dc2521159b36cc7f77ceff2f1fe0fabbc4b3d7","observation_id":"0786019f-5f60-4474-8a98-0f6d8a838c17","resolution":{"observed_at":"2026-08-06T23:35:08.934745Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.03300","last_updated":"2021-01-12T18:57:11Z","snapshot_observed_at":"2026-08-10T12:35:09.020030Z","submitted_at":"2020-09-07T17:59:25Z","title":"Measuring Massive Multitask Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.03300","snapshot_observed_at":"2026-08-06T23:35:03.753687Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.753687Z"},"links":{"cited_paper":"/paper/2009.03300","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:98b21c01502441a4f3046a9b4412695746dfdef0f01828c8a30ac6a3e63d694d","observation_id":"b6265da3-c39e-4b0c-97d4-da05c695a7b8","resolution":{"observed_at":"2026-08-06T23:35:03.753687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.17788","last_updated":"2024-07-25T05:42:14Z","snapshot_observed_at":"2026-08-05T00:09:41.708678Z","submitted_at":"2024-07-25T05:42:14Z","title":"PenHeal: A Two-Stage LLM Framework for Automated Pentesting and Optimal Remediation","version":1},"cited_work":{"arxiv_id":"2407.17788","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.17788","snapshot_observed_at":"2026-08-06T23:35:07.025550Z","title":"PenHeal: A Two-Stage LLM Framework for Automated Pentesting and Optimal Remediation","venue":"cs.CR","work_id":"9b52dbe8-e248-4ffc-80de-473b1400c14d","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:03.833812Z"},"links":{"cited_paper":"/paper/2407.17788","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:1dd97ae47a32793805ec3438ed9e289c3a3415785417809055075f5b50274359","observation_id":"d326bf29-32f0-4432-a82b-45699265daed","resolution":{"observed_at":"2026-08-06T23:35:07.054889Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.887049Z","title":null,"venue":null,"work_id":"49a665fd-f105-4e8e-90fd-789f0d4fea84","year":2017},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.017101Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:73b8fb28d1dccbfedbdd9690e1e61e1ea19a65cad800ca390288294c0cc50c0b","observation_id":"ef47d230-57d8-47be-b044-1cb5a37ceb36","resolution":{"observed_at":"2026-08-06T23:35:08.890346Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.20787","last_updated":"2025-01-06T07:22:50Z","snapshot_observed_at":"2026-08-05T03:53:59.997985Z","submitted_at":"2024-12-30T08:11:54Z","title":"SecBench: A Comprehensive Multi-Dimensional Benchmarking Dataset for LLMs in Cybersecurity","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.20787","snapshot_observed_at":"2026-08-06T23:35:04.049872Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.049872Z"},"links":{"cited_paper":"/paper/2412.20787","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:06bba29918dcf67e8815b57f11aad38e0ef87a2d6d08175411d1a20da25cf66d","observation_id":"086b98e7-cc7d-49df-9a5d-1232a30abb70","resolution":{"observed_at":"2026-08-06T23:35:04.049872Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.16169","last_updated":"2024-10-23T07:32:15Z","snapshot_observed_at":"2026-07-06T16:53:22.999597Z","submitted_at":"2023-11-16T13:17:20Z","title":"Understanding the Effectiveness of Large Language Models in Detecting Security Vulnerabilities","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.16169","snapshot_observed_at":"2026-08-06T23:35:04.178937Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.178937Z"},"links":{"cited_paper":"/paper/2311.16169","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:14ef98ca0ab9999fc039d7560e911389f109bd616e5ac72f41285093064fecfd","observation_id":"89402b97-58a7-49fc-a238-510ff92fd30d","resolution":{"observed_at":"2026-08-06T23:35:04.178937Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.856949Z","title":null,"venue":null,"work_id":"1596f5bc-e1b4-45f4-83fc-53434346bdef","year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.282967Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:43fa565c2b94c0f99240df0a23418f9b37b87e274bfc0b83834cd7fe084ec542","observation_id":"d131b1f6-b22d-423e-a734-d38acffb53b0","resolution":{"observed_at":"2026-08-06T23:35:08.865416Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.830065Z","title":null,"venue":null,"work_id":"5fd4a2d0-6520-4bc5-ac10-c69a5a722c1d","year":2020},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.420784Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:37486b63d075150c4e3030b544919d09cd794d261bbaceadf7a4c785b9e97c16","observation_id":"1dc2312a-cb0a-4bf8-85a4-7f859317c0c1","resolution":{"observed_at":"2026-08-06T23:35:08.840467Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23383","last_updated":"2025-03-30T10:16:25Z","snapshot_observed_at":"2026-08-07T20:40:27.882593Z","submitted_at":"2025-03-30T10:16:25Z","title":"ToRL: Scaling Tool-Integrated RL","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.23383","snapshot_observed_at":"2026-08-06T23:35:04.518537Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.518537Z"},"links":{"cited_paper":"/paper/2503.23383","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:571e6e347929a1e3d94af2a3d40d9d56cbce3691aa75c90bb0cf9d22c6b301cc","observation_id":"7c9e36bc-fb98-4b7f-90c5-0b5834acc2fa","resolution":{"observed_at":"2026-08-06T23:35:04.518537Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.05459","last_updated":"2024-05-08T06:16:23Z","snapshot_observed_at":"2026-08-02T13:57:57.119489Z","submitted_at":"2024-01-10T09:25:45Z","title":"Personal LLM Agents: Insights and Survey about the Capability, Efficiency and Security","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.05459","snapshot_observed_at":"2026-08-06T23:35:04.607502Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.607502Z"},"links":{"cited_paper":"/paper/2401.05459","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:6155d62f21deb812c45cebe1d3850cf4abcc74d618c733e1f3472bf07e0f94ea","observation_id":"099806aa-a002-40b8-92bd-97eaf28c1dd4","resolution":{"observed_at":"2026-08-06T23:35:04.607502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.774746Z","title":null,"venue":null,"work_id":"aff8d476-c663-40fd-946b-6ac981806faa","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.734768Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:57db184e36d06a3ea711c0fa2af161154bfdadb4aa5634fd5b2c3dff7870c177","observation_id":"8e124330-059f-4511-b28e-b21ab49eebd4","resolution":{"observed_at":"2026-08-06T23:35:08.794746Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.17419","last_updated":"2025-06-25T02:24:46Z","snapshot_observed_at":"2026-08-10T09:42:17.185681Z","submitted_at":"2025-02-24T18:50:52Z","title":"From System 1 to System 2: A Survey of Reasoning Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.17419","snapshot_observed_at":"2026-08-06T23:35:04.775152Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.775152Z"},"links":{"cited_paper":"/paper/2502.17419","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:8642a66962f9d27a84bbdbda207d7cf55d625cc1663cd0047c577c5b91829533","observation_id":"18e460b7-36cd-41a8-9b2d-765927d27ab7","resolution":{"observed_at":"2026-08-06T23:35:04.775152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.01990","last_updated":"2025-08-02T12:44:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-31T18:00:29Z","title":"Advances and Challenges in Foundation Agents: From Brain-Inspired Intelligence to Evolutionary, Collaborative, and Safe Systems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.01990","snapshot_observed_at":"2026-08-06T23:35:04.862557Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.862557Z"},"links":{"cited_paper":"/paper/2504.01990","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:7f9939b26504086a38c3edc97bec9bea8eab7632acc1106879468d2e96aa6c32","observation_id":"26d4119a-f7a0-4190-98f9-31a588c7a61f","resolution":{"observed_at":"2026-08-06T23:35:04.862557Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:04.956694Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:04.956694Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:dba3036566454af68b4d5aed7f1be6061cf4d4f04d859f7940eee8aaef821158","observation_id":"c57679e1-33b8-465b-97d4-1f70a032c624","resolution":{"observed_at":"2026-08-06T23:35:04.956694Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:05.076984Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.076984Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:ef916d04fa5cff39d4f763c075a6f9812dc662c58fa5c290dbe61b5ac4a3b4f6","observation_id":"134025a9-b61d-4fdf-9a4f-bbfc9c00a3b7","resolution":{"observed_at":"2026-08-06T23:35:05.076984Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.596990Z","title":null,"venue":null,"work_id":"168e7a47-2a29-414f-90a7-3d9bc187f071","year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.236867Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:2c2fb50414583312f6a23c21c0dc268e395f601c5eb503412772d3cb6466ae2d","observation_id":"5832686b-7ff9-4d8e-8815-ad7d9eba0abe","resolution":{"observed_at":"2026-08-06T23:35:08.607318Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.627892Z","title":null,"venue":null,"work_id":"41572718-e355-4dc3-8747-833a420baa14","year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.104820Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:cd740e756462a723898272b7dba5469e6aafc0eb2c6500eaf65e95bca561e78b","observation_id":"061e2c6b-20ed-4e55-8616-a217004c1bc6","resolution":{"observed_at":"2026-08-06T23:35:08.634278Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.557612Z","title":null,"venue":null,"work_id":"0bbc35e4-d496-4278-8ab2-e487d2ace1fb","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.435569Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:35be4787ed1eada59873d326d40d8435ba45d70e8445560c195e66bc0e914ebb","observation_id":"28b389a7-918c-48a5-b7d5-85a811bdf43e","resolution":{"observed_at":"2026-08-06T23:35:08.564431Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.583120Z","title":null,"venue":null,"work_id":"d08a32d8-843c-4273-85c7-a3a6069872b5","year":2025},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.320121Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:505edf04ad07b7194860329b32a8a45622d71ad7b1c49330eb34f10d4925e612","observation_id":"595f317a-4f9d-4b7b-bf9d-6156dc4092d4","resolution":{"observed_at":"2026-08-06T23:35:08.586523Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11814","last_updated":"2024-02-19T04:08:44Z","snapshot_observed_at":"2026-08-09T21:05:54.403616Z","submitted_at":"2024-02-19T04:08:44Z","title":"An Empirical Evaluation of LLMs for Solving Offensive Security Challenges","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.11814","snapshot_observed_at":"2026-08-06T23:35:05.704955Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.704955Z"},"links":{"cited_paper":"/paper/2402.11814","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:b59effd73551f8b186c09e8c2b0cd9784309405e0b86a10a1dec60fe4e147887","observation_id":"531bb9b8-a3ff-4145-b1e8-b789c3bd6c80","resolution":{"observed_at":"2026-08-06T23:35:05.704955Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:05.584511Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.584511Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f7321d161e8eeb1172ce517b2683014f9dca4c535f4fc3b842114a7c2bfc23a9","observation_id":"c9f8a3be-43ea-4c3d-811f-12bf00b57524","resolution":{"observed_at":"2026-08-06T23:35:05.584511Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.516212Z","title":null,"venue":null,"work_id":"cdb64cab-0f6b-415d-bfd2-a7891101e89a","year":2021},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.896211Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:165f9d98faf52fff7f8dc97875d61677bf723b4f2358fb8d17f7edb07f8fab06","observation_id":"d30d15d8-7bc0-4777-9a8f-bd27cbc5d928","resolution":{"observed_at":"2026-08-06T23:35:08.524858Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05590","last_updated":"2025-02-18T12:26:33Z","snapshot_observed_at":"2026-07-06T18:27:36.903877Z","submitted_at":"2024-06-08T22:21:42Z","title":"NYU CTF Bench: A Scalable Open-Source Benchmark Dataset for Evaluating LLMs in Offensive Security","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.05590","snapshot_observed_at":"2026-08-06T23:35:05.819135Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.819135Z"},"links":{"cited_paper":"/paper/2406.05590","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:38ce29b0fb45c3efda75630c48e19375668ad0e117da44d48858b82fbc43ccb2","observation_id":"425b3577-9f55-4a6f-9635-ccaf8025aa32","resolution":{"observed_at":"2026-08-06T23:35:05.819135Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.16185","last_updated":"2025-06-07T17:03:11Z","snapshot_observed_at":"2026-08-04T22:35:18.426928Z","submitted_at":"2024-01-29T14:32:27Z","title":"LLM4Vuln: A Unified Evaluation Framework for Decoupling and Enhancing LLMs' Vulnerability Reasoning","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.16185","snapshot_observed_at":"2026-08-06T23:35:05.906614Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.906614Z"},"links":{"cited_paper":"/paper/2401.16185","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c5a5a202083ec50506e63cb707cbdaef7d99c0b7ca1240f95ec6841a84e48356","observation_id":"f0d47db9-8e72-46b0-b5ac-3c7addd97718","resolution":{"observed_at":"2026-08-06T23:35:05.906614Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12566","last_updated":"2021-07-27T03:01:47Z","snapshot_observed_at":"2026-07-06T11:32:47.487435Z","submitted_at":"2021-07-27T03:01:47Z","title":"Thunder CTF: Learning Cloud Security on a Dime","version":1},"cited_work":{"arxiv_id":"2107.12566","doi":null,"metadata_source":"pith","pith_arxiv_id":"2107.12566","snapshot_observed_at":"2026-08-06T23:35:06.652953Z","title":"Thunder CTF: Learning Cloud Security on a Dime","venue":"cs.CR","work_id":"04142027-bf19-44d4-a589-c1a5e5ed429b","year":2021},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.900130Z"},"links":{"cited_paper":"/paper/2107.12566","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:0d5bbadde2920ceaa6d75ba8edf8405e605bf7f4cd1537e5e815b57a7dea0d75","observation_id":"3ed92649-b2e0-40ac-b220-6790d434d41c","resolution":{"observed_at":"2026-08-06T23:35:06.658565Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10443","last_updated":"2023-08-21T03:30:21Z","snapshot_observed_at":"2026-08-05T16:13:40.630996Z","submitted_at":"2023-08-21T03:30:21Z","title":"Using Large Language Models for Cybersecurity Capture-The-Flag Challenges and Certification Questions","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.10443","snapshot_observed_at":"2026-08-06T23:35:05.945362Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.945362Z"},"links":{"cited_paper":"/paper/2308.10443","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:4d4f2112f9d14e24f66bed79637e2a97fdd6153d64235ee3ed5431a87c274233","observation_id":"53081299-40d5-450f-ba3a-f98341845241","resolution":{"observed_at":"2026-08-06T23:35:05.945362Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.447158Z","title":null,"venue":null,"work_id":"da95c98e-ad08-42cd-849f-ea6e5c289fe7","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.926998Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d20005e66b197cc135cd1cec02fab0780b39593fb888b541a54749ed87ba7e36","observation_id":"061eca14-bf72-4530-9a4a-3b9e8c672612","resolution":{"observed_at":"2026-08-06T23:35:08.485010Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.376344Z","title":null,"venue":null,"work_id":"ddc553d6-4733-44c6-b039-1b953ad304fb","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.955326Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:8e3fc673a0e714739d1283262e60a2ef3c9da32c549af31efe658bc72014b17b","observation_id":"e7e2d59b-4718-4bb0-b950-824fad8fc02c","resolution":{"observed_at":"2026-08-06T23:35:08.386518Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.409853Z","title":null,"venue":null,"work_id":"8bdb3ecc-5966-4364-b46c-fac2da5eea6c","year":2022},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.949603Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:0c84b51cd3d5f73f8c46aa43a031f878b2e22cc376e846671bdd138940324644","observation_id":"ecb831ee-ce95-4cf9-8b6a-aea9b41a7e58","resolution":{"observed_at":"2026-08-06T23:35:08.415128Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.12575","last_updated":"2024-07-24T07:49:14Z","snapshot_observed_at":"2026-08-09T16:47:18.281496Z","submitted_at":"2023-12-19T20:19:43Z","title":"LLMs Cannot Reliably Identify and Reason About Security Vulnerabilities (Yet?): A Comprehensive Evaluation, Framework, and Benchmarks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.12575","snapshot_observed_at":"2026-08-06T23:35:05.983497Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.983497Z"},"links":{"cited_paper":"/paper/2312.12575","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:93914230fd791fb1a43b1d19b50dbcf6eb5d1142f34001f53321115e1a4d766c","observation_id":"f450ee49-39b8-4ee6-82a8-6824a1947f21","resolution":{"observed_at":"2026-08-06T23:35:05.983497Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.353257Z","title":null,"venue":null,"work_id":"31bbef40-0815-4926-9a68-96000e41fe40","year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.971759Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:d7a50f9933297e0f7070a297913d5cce51d86f58de5433d5e93488f54dec4dcd","observation_id":"03dbb136-9d8e-43a4-8088-bfcf897e4419","resolution":{"observed_at":"2026-08-06T23:35:08.361894Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:05.991880Z","title":null,"venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.991880Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:eadc760f80805b8685c1dd7835f75a9b6385bab7f3c06359dba28c7cf7de7790","observation_id":"5a9b9dbd-8469-481d-b93b-c746889611af","resolution":{"observed_at":"2026-08-06T23:35:05.991880Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.239120Z","title":null,"venue":null,"work_id":"1248cb14-c093-480b-99c0-455480968d02","year":2020},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.005780Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:5a79ba17920dfbaba40617d16b850a798a045d618dc5255cad461eb377eb17c4","observation_id":"9bae381d-cd8a-47af-887b-cefb8338e859","resolution":{"observed_at":"2026-08-06T23:35:08.245162Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.278904Z","title":"Department of Health and Human Services","venue":null,"work_id":"5b230e85-cc4e-40ec-b0ec-3608946895fe","year":2018},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.988781Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:cbc8b46c664af9d450053138723845f5898047f054b894358ae8dccce65f06af","observation_id":"0c0021e4-5a44-40ac-91e8-a2a920c61026","resolution":{"observed_at":"2026-08-06T23:35:08.288849Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.224753Z","title":null,"venue":null,"work_id":"61b0bbc4-ecd0-4701-b15c-af0f9336cbd8","year":2018},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.057837Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:4563f613fae7e061ca8c75f86dbbd92fa800374af8ebee7a6dd965ce10aa5d5c","observation_id":"600f77bb-75a4-4f1d-ac20-77f11af97fdd","resolution":{"observed_at":"2026-08-06T23:35:08.230058Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.174516Z","title":null,"venue":null,"work_id":"02fda5b3-d423-4d1e-a4d6-b32f683d484f","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.105715Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:fa43a55259947fbb90179ea1dc1f306c5ee7f2b634881e8a32be252a5bbc9f83","observation_id":"ad7539cb-ae77-48be-ae1d-009ccee3cd18","resolution":{"observed_at":"2026-08-06T23:35:08.191001Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.13945","last_updated":"2024-12-21T15:21:27Z","snapshot_observed_at":"2026-07-06T18:03:33.406287Z","submitted_at":"2024-04-22T07:41:41Z","title":"How Multi-Modal LLMs Reshape Visual Deep Learning Testing? A Comprehensive Study Through the Lens of Image Mutation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.13945","snapshot_observed_at":"2026-08-06T23:35:06.031857Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.031857Z"},"links":{"cited_paper":"/paper/2404.13945","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c3823c0237934ba9cae966d224812aa80fe44dc90da144ccce73507cec34cee8","observation_id":"882b0845-7082-451a-a65d-acbe4c3413df","resolution":{"observed_at":"2026-08-06T23:35:06.031857Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.052045Z","title":null,"venue":null,"work_id":"403c3c34-d3fd-44f1-a673-78bed36ae2d9","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.174475Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:192f5de59c5dfcacb9daf8e2b13529925ab2d35b0bea8681dbe23d8e7410f1e9","observation_id":"454d7eff-83d5-4f20-83fb-b4e15faf8d39","resolution":{"observed_at":"2026-08-06T23:35:08.064873Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:06.178897Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.178897Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c42405c68468a41753f95e4c8a47ad8e78fc5d941d17dd11024609fa683cb81e","observation_id":"038fb3f8-62e9-42ac-9fe8-9d1b5928b1bf","resolution":{"observed_at":"2026-08-06T23:35:06.178897Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.109218Z","title":null,"venue":null,"work_id":"84a5ba28-24f3-4aca-937d-36d7b4f9a616","year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.132701Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:8ea35157e0ff4978485a90f7a4134cfb3d79cd1d813cc001fe1f62356074e541","observation_id":"bb0a00d6-6ee1-47c9-a1ff-e900aef2e2d9","resolution":{"observed_at":"2026-08-06T23:35:08.114044Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.03629","last_updated":"2023-03-10T01:00:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-10-06T01:00:32Z","title":"ReAct: Synergizing Reasoning and Acting in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.03629","snapshot_observed_at":"2026-08-06T23:35:06.199616Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.199616Z"},"links":{"cited_paper":"/paper/2210.03629","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:beb46f5dce269d12e444b0ba9d1c4bd27f27a21d41143c60b160251dc1b6ed4b","observation_id":"1d8fa8c4-2b69-43ef-84e7-60eb88c035c4","resolution":{"observed_at":"2026-08-06T23:35:06.199616Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.09210","last_updated":"2024-10-03T04:35:39Z","snapshot_observed_at":"2026-08-10T06:02:30.480313Z","submitted_at":"2023-11-15T18:54:53Z","title":"Chain-of-Note: Enhancing Robustness in Retrieval-Augmented Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.09210","snapshot_observed_at":"2026-08-06T23:35:06.202672Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.202672Z"},"links":{"cited_paper":"/paper/2311.09210","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:cda3ec2b0a47bac9ef41edab0a78e68af99ac83d7335cf4f943652719a481311","observation_id":"fe207e1e-b66f-417f-9ada-cbb87495f96f","resolution":{"observed_at":"2026-08-06T23:35:06.202672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.06838","last_updated":"2025-07-21T05:24:59Z","snapshot_observed_at":"2026-08-04T23:33:44.230238Z","submitted_at":"2024-03-11T15:59:59Z","title":"ACFIX: Guiding LLMs with Mined Common RBAC Practices for Context-Aware Repair of Access Control Vulnerabilities in Smart Contracts","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.06838","snapshot_observed_at":"2026-08-06T23:35:06.207079Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.207079Z"},"links":{"cited_paper":"/paper/2403.06838","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:52871a7a9f4aabbfbdaa91a119b4a1e3e259d4b0465fe4eb8f082159052d8a28","observation_id":"6927f7f0-3050-432d-aeb3-f45efa2e8c02","resolution":{"observed_at":"2026-08-06T23:35:06.207079Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:07.991614Z","title":null,"venue":null,"work_id":"3a00464f-b58a-4237-af7e-77b47055141f","year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.196569Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:b256935779d1a3fe1c9c89f088123eaee379be19882d3f78e40b3c237a2e0a99","observation_id":"056695c2-c176-4ce6-9295-9da74768aaf7","resolution":{"observed_at":"2026-08-06T23:35:07.996021Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.16906","last_updated":"2024-06-06T06:01:41Z","snapshot_observed_at":"2026-08-09T12:26:08.682226Z","submitted_at":"2024-02-25T00:56:27Z","title":"Debug like a Human: A Large Language Model Debugger via Verifying Runtime Execution Step-by-step","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.16906","snapshot_observed_at":"2026-08-06T23:35:06.211001Z","title":"The knowledge fully matches the write- up and accurately reflects its content","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.211001Z"},"links":{"cited_paper":"/paper/2402.16906","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:9c25498ae86bbb1c40502df0d23e54023fc00afc418a9321cd977e394dd662ad","observation_id":"234bcb10-d3a6-47f2-bc4d-fe64c78c9cb2","resolution":{"observed_at":"2026-08-06T23:35:06.211001Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:35:08.014852Z","title":"InMulti-Agent Security Workshop@ NeurIPS’23","venue":null,"work_id":"eda81d94-b4fc-40ed-9193-b07069752afa","year":null},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:06.182493Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:9b90b9978ef7fecf390d1fce0b521a6139d059547360eb5cbb2f9731155cc0ed","observation_id":"89e2fd64-527f-4049-a1b2-eaf4aa002179","resolution":{"observed_at":"2026-08-06T23:35:08.019250Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.07688","last_updated":"2024-06-03T08:14:45Z","snapshot_observed_at":"2026-07-06T17:28:56.092370Z","submitted_at":"2024-02-12T14:53:28Z","title":"CyberMetric: A Benchmark Dataset based on Retrieval-Augmented Generation for Evaluating LLMs in Cybersecurity Knowledge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.07688","snapshot_observed_at":"2026-08-06T23:35:05.975106Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-06T23:35:05.975106Z"},"links":{"cited_paper":"/paper/2402.07688","citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:f8c5bcf3a08c48be6f0a65d8aa4e4ebc449a35b45fe3ebd3926149daa531e1ce","observation_id":"7a79103d-b8de-4a37-8e5e-96ceec3b28e6","resolution":{"observed_at":"2026-08-06T23:35:05.975106Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T23:34:57.394753Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-06T23:34:57.394753Z"},"links":{"citing_paper":"/paper/2506.17644"},"observation_digest":"sha256:c7e2aaec24300c2d488f714c64cf5032aca8f329e4dca2b1c5849d9a8054ab28","observation_id":"1b5fe0dc-2320-4683-80a4-40beeb135602","resolution":{"observed_at":"2026-08-06T23:34:57.394753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.17644","last_updated":"2025-06-21T08:56:20Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-07T23:27:30.726322Z","submitted_at":"2025-06-21T08:56:20Z","title":"Measuring and Augmenting Large Language Models for Solving Capture-the-Flag Challenges"},"reference_resolution":{"displayed":98,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":79,"verified_exact":4,"verified_fuzzy":14},"total_outbound_references":98},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 98 of 98 outbound references and 0 inbound Pith citation observations for arXiv:2506.17644."}