{"as_of":"2026-08-11T01:36:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:16081fa14b244755d0be7f413f357ff4df5e602138ef03ecaf9b6c5138d1d3b7","coverage":[{"denominator":22,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":22,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-26T17:26:22.687010Z","state":"measured"},{"denominator":23,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":23,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-14T07:36:47.336670Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.20093","snapshot_observed_at":"2026-07-14T07:36:47.336670Z","title":"Self-preference is weak or absent in verifiable instruction-following revision: A four-model test under genuine authorship.arXiv preprint arXiv:2606.20093,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11022","last_updated":"2026-07-13T02:41:16Z","snapshot_observed_at":"2026-08-10T16:35:21.863640Z","submitted_at":"2026-07-13T02:41:16Z","title":"When the Reward Suite Is Leaky: A Preregistered Causal Contrast of Natural Verifier False Positives in RLVR","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-14T07:36:47.336670Z"},"links":{"cited_paper":"/paper/2606.20093","citing_paper":"/paper/2607.11022"},"observation_digest":"sha256:60992a629b36ab0977e9bd42bf410b1eca76901fc212851362d9cd170f3618ba","observation_id":"fcdae31f-299a-4fc4-a0fb-31606e3bba6a","resolution":{"observed_at":"2026-07-14T07:36:47.336670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2606.20093/citation-record","integrity":"/paper/2606.20093/integrity","json":"/paper/2606.20093/citation-record.json","paper":"/paper/2606.20093"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"and Feng, Shi , booktitle =","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:7a783fa0b4b6b72cf0bd017b4f7b75783e5d4b175d78d02d787ee09a197292bd","observation_id":"224d18ef-f7c2-4320-91a0-024058ea0f1b","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"Pride and Prejudice:","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:45f935bdd4a01a07a64c8c38e83542b751de30be262dbe74f783e1ba8bd78e5e","observation_id":"2812174a-af35-42ab-a361-bdb0ae7bb4a0","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"Proceedings of the Annual Meeting of the Association for Computational Linguistics (ACL) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:bb7266a15351a290f770cd3dd262989af55ae6ef39ea41726bc9935e92a4f034","observation_id":"a71bbd3a-0199-4994-9dad-ead21b5797b3","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"Challenging the Evaluator:","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:1cd5d733ce46ff01e8af4b031caa7059008c3c522d38db6fa6a87cf3b1f37473","observation_id":"d11f226c-72e3-4099-9375-40a09f647c78","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"Feedback Friction:","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:e5ca4ae2fe98374d8a3169477b8dbfff7b3dffd5535aba4d9da6b8c1a066d9e3","observation_id":"232daf41-669b-4973-847f-7f03d0faa691","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"Cross-Context Review: Improving","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:06812ab59461c5a4b7d474bd93f60b4607047e3fed710ea072b859466de57ea6","observation_id":"d4078201-6eaf-43d0-a696-0059bbbdc73c","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"and Zhang, Hao and Gonzalez, Joseph E","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:acb7859237e60beae6f27ed81ed3ae0086bd3f4e855d6f4eaa2d62381082a9da","observation_id":"1ed36c10-a73c-4f19-b369-ebcf8a2f077a","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:3de09cada151948ea64b1e1927dbbc9545ddce300a1f03b36cbb652f0cfda520","observation_id":"9a36dd8c-09f0-47e4-bbe0-ee0d3efa3884","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"International Conference on Learning Representations (ICLR) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:821e91ee9b8ce631223a20973429b3d4a2a487383e5ec5a06ba0e089119c7847","observation_id":"34ae23b0-d9b7-4884-bd87-b177b198d4cb","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"When Can","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:3827de5c5ec1db12ccb6a67eaa1deee2a10194427353f20cf8c29e83df3c72ad","observation_id":"4caa0cab-36ee-4202-8083-a02db646c764","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-26T17:26:22.687010Z","title":"2023 , note =","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:d67f0d678bef63885269061ca6c23db6c0e6c74cfc121d1b3881167af0f0d9a2","observation_id":"80d19198-67f6-4757-a5d1-debcad56ceee","resolution":{"observed_at":"2026-06-26T17:26:22.687010Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2510.07517","last_updated":"2026-04-09T17:03:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-10-08T20:29:46Z","title":"When Identity Skews Debate: Anonymization for Bias-Reduced Multi-Agent Reasoning","version":5},"cited_work":{"arxiv_id":"2510.07517","doi":null,"metadata_source":"pith","pith_arxiv_id":"2510.07517","snapshot_observed_at":"2026-07-05T08:40:50.485054Z","title":"When Identity Skews Debate: Anonymization for Bias-Reduced Multi-Agent Reasoning","venue":"cs.AI","work_id":"029e5df9-42e4-4358-90e9-83e1f1263aee","year":2025},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"cited_paper":"/paper/2510.07517","citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:b781fe9ca1cc04f3e37995c8279b230a0907f9a59b35a0116420cb4eef4996c5","observation_id":"81dc08ba-b0f6-4ebb-b21a-3c34bebf2c41","resolution":{"observed_at":"2026-07-04T03:59:33.343796Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01798","last_updated":"2024-03-14T04:27:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-03T04:56:12Z","title":"Large Language Models Cannot Self-Correct Reasoning Yet","version":2},"cited_work":{"arxiv_id":"2310.01798","doi":"10.48550/arxiv.2310.01798","metadata_source":"pith","pith_arxiv_id":"2310.01798","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Large Language Models Cannot Self-Correct Reasoning Yet","venue":"cs.CL","work_id":"f63b261b-ef16-40f5-993b-9d37b1a51b92","year":2023},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"cited_paper":"/paper/2310.01798","citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:1e60dfbcec4cc7fb01c5d64aed69e49a62af001ece2b4d932bdafc7d024e1b17","observation_id":"5202651b-0e95-46fd-87df-8a5eedb1a515","resolution":{"observed_at":"2026-07-04T03:59:33.346990Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-07-09T10:48:38.048248+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-09T10:48:38.048248+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2506.11930","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T03:59:33.338691Z","title":"Feedback friction: Llms struggle to fully incorporate external feedback","venue":null,"work_id":"8cc49423-3d1c-469e-8c1c-242ed2e5becd","year":2025},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:23cee7003f5ff20d02f807df6c6eb4a585bc47ccd9f5682218b1103fd180481c","observation_id":"53a1a99b-0185-467d-939a-dad7a98b8ca7","resolution":{"observed_at":"2026-07-04T03:59:33.340449Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01297","last_updated":"2024-12-03T19:14:06Z","snapshot_observed_at":"2026-08-02T06:30:19.601023Z","submitted_at":"2024-06-03T13:05:46Z","title":"When Can LLMs Actually Correct Their Own Mistakes? A Critical Survey of Self-Correction of LLMs","version":3},"cited_work":{"arxiv_id":"2406.01297","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.01297","snapshot_observed_at":"2026-07-04T03:59:33.348485Z","title":"arXiv preprint arXiv:2406.01297 , year=","venue":null,"work_id":"839f05ba-e30d-4789-bf50-b961a76acfc5","year":2024},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"cited_paper":"/paper/2406.01297","citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:811a86b3837d633f778b7d4c395e49e0e2d1cad9ce220032362b42fecf0db4f2","observation_id":"2b771e12-cf91-4116-9dd0-8ca7af3cd7b6","resolution":{"observed_at":"2026-07-04T03:59:33.350376Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.16533","doi":"10.48550/arxiv.2509.16533","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Kim, S.; and Khashabi, D","venue":"ArXiv.org","work_id":"d62bf280-2471-4a9b-b65c-d9a62f1b79dd","year":2025},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:73036394609dfe0234c1830e2e506d1a49f79d3d101c890e0a91cf201c3dbd65","observation_id":"d69827bf-15e8-43c2-8fa9-bc0c07cb0bb1","resolution":{"observed_at":"2026-07-04T03:59:33.343588Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.17651","last_updated":"2023-05-25T19:13:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-30T18:30:01Z","title":"Self-Refine: Iterative Refinement with Self-Feedback","version":2},"cited_work":{"arxiv_id":"2303.17651","doi":"10.1007/s10664-008-","metadata_source":"pith","pith_arxiv_id":"2303.17651","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Self-Refine: Iterative Refinement with Self-Feedback","venue":"cs.CL","work_id":"59181e7f-e58e-45d3-8146-4477a9f53d5a","year":2023},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"cited_paper":"/paper/2303.17651","citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:d79cf1a7529d7afb7e6576598e1e0da0f737507c80882baf994cf51d6156c6a9","observation_id":"fdd6ca02-2b4f-48fe-934e-c2a834417177","resolution":{"observed_at":"2026-07-04T03:59:33.350325Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.13076","last_updated":"2024-04-15T16:49:59Z","snapshot_observed_at":"2026-08-11T01:25:41.522896Z","submitted_at":"2024-04-15T16:49:59Z","title":"LLM Evaluators Recognize and Favor Their Own Generations","version":1},"cited_work":{"arxiv_id":"2404.13076","doi":null,"metadata_source":"pith","pith_arxiv_id":"2404.13076","snapshot_observed_at":"2026-07-10T10:27:02.394079Z","title":"LLM Evaluators Recognize and Favor Their Own Generations","venue":"cs.CL","work_id":"ec9ad2bb-47a2-46c7-b8e9-05f69142decd","year":2024},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"cited_paper":"/paper/2404.13076","citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:42949cad935d604f6bca0c3dc142197c06f4ec0f497496bcd8d16ec538b7b2f9","observation_id":"f1b2ea92-b418-4767-9c43-48bde695430c","resolution":{"observed_at":"2026-07-04T03:59:33.325115Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.12123","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T03:59:33.338389Z","title":"Cross-context review: Improving LLM output quality by separating production and review sessions, 2026","venue":null,"work_id":"5fc6f04c-5877-4db9-8294-aa4f3aad15dc","year":2026},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:d59ec7af182c63323aed6af668f85b01e2e49bc823ebbacf2297ef4f864b44c6","observation_id":"ba22dfe1-cd1f-4b7d-9f97-17ef79ff43f0","resolution":{"observed_at":"2026-07-04T03:59:33.340738Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11436","last_updated":"2024-06-18T04:41:07Z","snapshot_observed_at":"2026-08-10T20:47:31.032186Z","submitted_at":"2024-02-18T03:10:39Z","title":"Pride and Prejudice: LLM Amplifies Self-Bias in Self-Refinement","version":2},"cited_work":{"arxiv_id":"2402.11436","doi":"10.48550/arxiv.2402.11436","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.11436","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Rickard Stureborg, Dimitris Alikaniotis, and Yoshi Suhara","venue":"arXiv (Cornell University)","work_id":"ae09097d-58a5-40bc-8f92-14d333fc4a2d","year":2024},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"cited_paper":"/paper/2402.11436","citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:0d2b13b5f7bd393a275c0b72a51c19caa9add959a6bfb805e3e37afddb310d1d","observation_id":"f80e77ac-d5c0-4868-8fca-9979aaf89230","resolution":{"observed_at":"2026-07-04T03:59:33.333267Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.05685","last_updated":"2023-12-24T02:01:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-09T05:55:52Z","title":"Judging LLM-as-a-Judge with MT-Bench and Chatbot Arena","version":4},"cited_work":{"arxiv_id":"2306.05685","doi":"10.1109/4235.797969","metadata_source":"pith","pith_arxiv_id":"2306.05685","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Judging LLM-as-a-Judge with MT-Bench and Chatbot Arena","venue":"cs.CL","work_id":"d0c30cd7-81e1-4159-a87f-f6adca77ff08","year":2023},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"cited_paper":"/paper/2306.05685","citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:b48bccc57989cf0bc849682c8bfb200a07946a3230166153d55992e8a282164c","observation_id":"7877a5d3-8789-4430-b62b-cbd4d2f16daf","resolution":{"observed_at":"2026-07-04T03:59:33.347259Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.07911","last_updated":"2023-11-14T05:13:55Z","snapshot_observed_at":"2026-07-06T16:47:08.877195Z","submitted_at":"2023-11-14T05:13:55Z","title":"Instruction-Following Evaluation for Large Language Models","version":1},"cited_work":{"arxiv_id":"2311.07911","doi":"10.48550/arxiv.2311.07911","metadata_source":"pith","pith_arxiv_id":"2311.07911","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Instruction-Following Evaluation for Large Language Models","venue":"cs.CL","work_id":"3aa06177-125a-4f5a-8f4a-8070c5986c26","year":2023},"citing_paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-06-26T17:26:22.687010Z"},"links":{"cited_paper":"/paper/2311.07911","citing_paper":"/paper/2606.20093"},"observation_digest":"sha256:5cbbc4031a8119ba3518e06c219fa8fa613c7c860eba53718790a511c5ab27d9","observation_id":"4a3dc0ca-7c40-45cb-a050-a2e6ca9312c3","resolution":{"observed_at":"2026-07-04T03:59:33.335818Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2606.20093","last_updated":"2026-06-18T11:12:46Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T22:10:27.064801Z","submitted_at":"2026-06-18T11:12:46Z","title":"Self-Preference Is Weak or Absent in Verifiable Instruction-Following Revision: A Four-Model Test Under Genuine Authorship"},"reference_resolution":{"displayed":22,"state_counts":{"malformed_identifier":0,"metadata_mismatch":5,"parse_uncertain":0,"unresolved":11,"verified_exact":6,"verified_fuzzy":0},"total_outbound_references":22},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 22 of 22 outbound references and 1 inbound Pith citation observation for arXiv:2606.20093."}