{"as_of":"2026-08-15T14:26:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:85b8ee14054d43e29d3e8d3f3a4d601fd0c0751b4921699615e453b544615a84","coverage":[{"denominator":57,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":57,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T15:02:39.725628Z","state":"measured"},{"denominator":59,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":59,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-15T06:32:42.880941+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T05:23:13.271917Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-10T05:30:23.456663Z","state":"measured"}],"external_citation_measurements":[{"count":3,"observed_at":"2026-08-10T05:30:23.456663Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"cited_work":{"arxiv_id":"2412.11385","doi":"10.48550/arxiv.2412.11385","metadata_source":"pith","pith_arxiv_id":"2412.11385","snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","venue":"cs.CL","work_id":"27a80379-907b-4086-89ef-4e23344e9b2f","year":2024},"citing_paper":{"arxiv_id":"2508.01930","last_updated":"2025-08-03T21:45:37Z","snapshot_observed_at":"2026-08-14T03:42:39.605301Z","submitted_at":"2025-08-03T21:45:37Z","title":"Word Overuse and Alignment in Large Language Models: The Influence of Learning from Human Feedback","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T05:23:13.271917Z"},"links":{"cited_paper":"/paper/2412.11385","citing_paper":"/paper/2508.01930"},"observation_digest":"sha256:8017c542763b1d2048a4759b5cb84a84180850748f6e88797d299e61c4376aa2","observation_id":"c2c65234-96f7-46f1-b363-99954cfc9434","resolution":{"observed_at":"2026-08-06T05:23:13.710345Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.11385","snapshot_observed_at":"2026-08-01T07:19:19.125548Z","title":"S., & Ward, Z","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.21498","last_updated":"2026-07-28T14:30:09Z","snapshot_observed_at":"2026-08-15T08:54:16.185157Z","submitted_at":"2026-07-23T16:47:39Z","title":"Artificial Epanorthosis: Why large language models overuse a classical rhetorical figure, and how to mitigate it","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-01T07:19:19.125548Z"},"links":{"cited_paper":"/paper/2412.11385","citing_paper":"/paper/2607.21498"},"observation_digest":"sha256:22742fd4907985c8a8a06b502c1a5ad410e7be3f4b43bdddedde3a4ef9303992","observation_id":"19fd96cc-7a26-4dfb-92b5-49e260a5df90","resolution":{"observed_at":"2026-08-01T07:19:19.125548Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2412.11385/citation-record","integrity":"/paper/2412.11385/integrity","json":"/paper/2412.11385/citation-record.json","paper":"/paper/2412.11385"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2307.01850","last_updated":"2023-07-04T17:59:31Z","snapshot_observed_at":"2026-08-14T17:55:34.461016Z","submitted_at":"2023-07-04T17:59:31Z","title":"Self-Consuming Generative Models Go MAD","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.01850","snapshot_observed_at":"2026-08-11T15:02:39.464596Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.464596Z"},"links":{"cited_paper":"/paper/2307.01850","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:e01df00170f36e0ba57ec50b6db7c8a29318a5584d68cf95dbc7754ed536bce8","observation_id":"098505f0-9da2-48ad-a95a-d3f91c452ee8","resolution":{"observed_at":"2026-08-11T15:02:39.464596Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.638103Z","title":null,"venue":null,"work_id":"2a9e2141-ef2a-47cb-af18-e53558f88715","year":2017},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.469441Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:6babb2a273bee52f83dcbfb229bd33493dd9e1776b108277c16df4b12159d90d","observation_id":"bbed51b7-49c8-41a7-b3c8-edb19911ef9f","resolution":{"observed_at":"2026-08-11T15:02:40.642934Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.625121Z","title":null,"venue":null,"work_id":"7d8bf37c-6a71-44ab-8f52-436befae139f","year":2020},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.473619Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:d36a8048d2079d11dd17845eec7368668fb2d3721ecb14f0410f09096dc59ff4","observation_id":"699f9400-7932-4715-81ea-dc853c6d7a66","resolution":{"observed_at":"2026-08-11T15:02:40.629102Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.16822","last_updated":"2024-06-17T07:07:30Z","snapshot_observed_at":"2026-08-14T10:26:37.134726Z","submitted_at":"2023-11-28T14:36:43Z","title":"Large Language Models Suffer From Their Own Output: An Analysis of the Self-Consuming Training Loop","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.16822","snapshot_observed_at":"2026-08-11T15:02:39.477913Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.477913Z"},"links":{"cited_paper":"/paper/2311.16822","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:2a8ad05493f6f52a558ebb1a239bd9a861c36feb81aa10d5ffed4a3f87b1dea0","observation_id":"37597351-7ede-412e-bcf7-547cb9633a1c","resolution":{"observed_at":"2026-08-11T15:02:39.477913Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.611422Z","title":null,"venue":null,"work_id":"ae15ccbb-9384-496a-ba79-b9ea31fed19a","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.481510Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:8e6fa6bdc801d8e2d1c55524a6b1dfd24faff963053a794a45269bd27874cad2","observation_id":"4c2a4e93-b684-4864-a43c-a5b44b5a76f4","resolution":{"observed_at":"2026-08-11T15:02:40.616267Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.04132","last_updated":"2024-03-07T01:22:38Z","snapshot_observed_at":"2026-08-02T17:55:33.750637Z","submitted_at":"2024-03-07T01:22:38Z","title":"Chatbot Arena: An Open Platform for Evaluating LLMs by Human Preference","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.04132","snapshot_observed_at":"2026-08-11T15:02:39.485232Z","title":"Gonzalez, and Ion Stoica","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.485232Z"},"links":{"cited_paper":"/paper/2403.04132","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:019334feba2aae4e8d29c41a9a0893f63082c3cb13b66aa9979d41676d063b45","observation_id":"eb7c16a7-f561-49fb-8b37-e4bb51413b9e","resolution":{"observed_at":"2026-08-11T15:02:39.485232Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.492120Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.492120Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:5257add41a805f09b6e126e0bec104f5c782ac470f36686fb128504072349692","observation_id":"718a6a70-bbeb-44ad-a0d9-864506ee3933","resolution":{"observed_at":"2026-08-11T15:02:39.492120Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.584155Z","title":null,"venue":null,"work_id":"6e5afd5d-d982-44c9-8abe-14a1ad08a8c5","year":2018},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.496407Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:b9d48eee630aa61e4a01dc84911055d66fecdad540e115d88e20532008f2d5fc","observation_id":"e86f24b1-d0c6-4710-8b9b-5fe0f120c243","resolution":{"observed_at":"2026-08-11T15:02:40.588675Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.569784Z","title":null,"venue":null,"work_id":"ce6efcab-7154-4a74-aa06-a242c081eecc","year":2018},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.501077Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:a5c2538d98952ba8103e98c4bde6fbed4a89c6b29eb8111562680b741c004f4c","observation_id":"19d4b425-cda5-4544-8c31-77d3b360663d","resolution":{"observed_at":"2026-08-11T15:02:40.574311Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.557495Z","title":null,"venue":null,"work_id":"2f2b33bc-5bb2-4b1a-a7a8-9bc15c5a7fcc","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.505015Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:a01deb8a35a4486186df858490f074e9bb81dc194d0051c96e2994782b490bfb","observation_id":"819bd7ac-16f8-4675-a46f-c4fef14e79fb","resolution":{"observed_at":"2026-08-11T15:02:40.561266Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.13686","last_updated":"2024-10-22T17:06:17Z","snapshot_observed_at":"2026-08-14T22:23:59.128488Z","submitted_at":"2024-09-20T17:54:16Z","title":"The Impact of Large Language Models in Academia: from Writing to Speaking","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.13686","snapshot_observed_at":"2026-08-11T15:02:39.509806Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.509806Z"},"links":{"cited_paper":"/paper/2409.13686","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:b5c4f1d73528ca03a45773a7bb8e478e1249b4e660318c14d903656904648e4c","observation_id":"cce9f4b1-d42b-44c2-802b-16ae67df9e5d","resolution":{"observed_at":"2026-08-11T15:02:39.509806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.514227Z","title":null,"venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.514227Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:a2600ec8f83decbc235c420e714eac446bcb51b1134a9179e56eed9ec42a9545","observation_id":"e26725d2-ef44-4008-ba2f-2932c8ddf88b","resolution":{"observed_at":"2026-08-11T15:02:39.514227Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.16887","last_updated":"2024-03-25T15:56:37Z","snapshot_observed_at":"2026-08-13T00:45:33.293733Z","submitted_at":"2024-03-25T15:56:37Z","title":"ChatGPT \"contamination\": estimating the prevalence of LLMs in the scholarly literature","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.16887","snapshot_observed_at":"2026-08-11T15:02:39.519108Z","title":"contamination","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.519108Z"},"links":{"cited_paper":"/paper/2403.16887","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:6a99d02d9471a7520702cf5a35fc7e9eea0d7ae44aedf9727a84b1427da9fef4","observation_id":"41dce68e-add3-4b04-9581-2ef7c7ba5725","resolution":{"observed_at":"2026-08-11T15:02:39.519108Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.536463Z","title":null,"venue":null,"work_id":"a60ad5da-9079-45f7-b077-f94001aca21e","year":2023},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.524454Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:00715b5a523d32885ace45aad10f6ac0f5b48498da3efa4bbfaf20bc881862cf","observation_id":"407b8f85-323c-47db-8170-d31a7a0c6823","resolution":{"observed_at":"2026-08-11T15:02:40.540219Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.522214Z","title":"a ussler and Tom Juzek. 2017. https://publikationen.uni-tuebingen.de/xmlui/handle/10900/77066 Hot topics surrounding acceptability judgement tasks . In S. Featherston, R. H \\","venue":null,"work_id":"2d08fde7-e289-4821-adb0-630861d5aaee","year":2017},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.530493Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:ab5fcde6f23dbcefc12a335e0c7792857d961f5427ff8bf1b55d46a859ffdeb9","observation_id":"7b5958dd-118f-4318-936a-bea31c1540dd","resolution":{"observed_at":"2026-08-11T15:02:40.527759Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.509412Z","title":null,"venue":null,"work_id":"62224d18-85bf-443d-8dcc-3ace06901666","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.534596Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:fc6d00b27f55ebece1379b133be63deda84085c0bfcd063e8a137914beda5889","observation_id":"140932fe-be73-42a6-a981-16493ebebfff","resolution":{"observed_at":"2026-08-11T15:02:40.513510Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.496375Z","title":null,"venue":null,"work_id":"01b0aa86-bd9f-4aaf-809e-8fecba2f1e9c","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.539176Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:47352b5f6dd96d21d265f0932d972006327d073440be3f7dc03352cc6c8da4aa","observation_id":"c09e8fc1-c15d-425e-99f5-5d5b85fe5339","resolution":{"observed_at":"2026-08-11T15:02:40.500608Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.484464Z","title":null,"venue":null,"work_id":"f143d5fc-919c-4741-991a-fbf9d92cc63a","year":2018},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.543444Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:322fc60c7cd2312f3fe6c58ae38c2a173c9809d10cd812f9b4027941cef4ffff","observation_id":"51c3c825-eb20-44ea-8302-34cbd812e104","resolution":{"observed_at":"2026-08-11T15:02:40.488320Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.471247Z","title":null,"venue":null,"work_id":"d3ac340f-c4d7-496f-ad46-ed847e3dd6b5","year":2018},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.547691Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:edc4619a3802963692dbc8d4ff49f03fb4e052074f59e8a4dd978ba5956eea59","observation_id":"328bb491-0f87-4185-8010-8de6f65f11a5","resolution":{"observed_at":"2026-08-11T15:02:40.475316Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.459562Z","title":null,"venue":null,"work_id":"97a5cdcc-3018-4e91-bc85-8765189882af","year":2017},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.553886Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:63ba93f30cc38cab88dabee3aa015d5314ef929ca59b858199557ac61b912666","observation_id":"6ae2ce1a-2120-4d09-92e9-5752afa75ef9","resolution":{"observed_at":"2026-08-11T15:02:40.463680Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07016","last_updated":"2025-07-03T08:26:13Z","snapshot_observed_at":"2026-08-12T23:45:49.565823Z","submitted_at":"2024-06-11T07:16:34Z","title":"Delving into LLM-assisted writing in biomedical publications through excess vocabulary","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07016","snapshot_observed_at":"2026-08-11T15:02:39.558999Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.558999Z"},"links":{"cited_paper":"/paper/2406.07016","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:9c304963c976cc6f42ad2b5e0531e4ac7d4ed03c033039d5c34a5d38aa12c046","observation_id":"d1f0a5ec-7847-4ee0-8c7c-f3bb297adf53","resolution":{"observed_at":"2026-08-11T15:02:39.558999Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.564857Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.564857Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:318703b1abbabf45037387e964a2b4d996a5ae980529e03ce7f2088ed649e736","observation_id":"90be9cd0-68a7-47a3-8255-77b5fd60f01b","resolution":{"observed_at":"2026-08-11T15:02:39.564857Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.569635Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.569635Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:2dc2e88a67c45819088a936b668207e7c5c4174aea80554618a5b1f95207c615","observation_id":"7801e88f-2f16-42ea-aec3-f4ed3021c75e","resolution":{"observed_at":"2026-08-11T15:02:39.569635Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.441048Z","title":null,"venue":null,"work_id":"6b28b2bd-0d84-4e87-8e3e-a117c1d0291b","year":2019},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.573699Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:bae197d8605d1bc83858186685cd9c415cb0338efe109c588dd658e03433d2a8","observation_id":"cf120133-c47d-4d72-85d7-bd34f9d08b01","resolution":{"observed_at":"2026-08-11T15:02:40.445051Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.429951Z","title":null,"venue":null,"work_id":"5a49581c-e22f-4770-8b42-44f0ba6538df","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.577584Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:95edea89e3da7705a6c14d5bc0cc0c9f3f36e1980a8029d09357095892599e81","observation_id":"bab2c979-b0f1-4205-83b2-7f9f8921201b","resolution":{"observed_at":"2026-08-11T15:02:40.434022Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.417738Z","title":null,"venue":null,"work_id":"c5a39113-06ec-418d-8ee5-443afa32e2b6","year":2020},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.583166Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:e7abbd585c2c8c559238d2e53a59813262b366b67adcfd2728d3fa752c63ef5e","observation_id":"9df96415-65b6-43cb-a3c2-799ddc3279a0","resolution":{"observed_at":"2026-08-11T15:02:40.421590Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.07183","last_updated":"2026-05-19T03:49:32Z","snapshot_observed_at":"2026-08-01T20:53:34.992573Z","submitted_at":"2024-03-11T21:51:39Z","title":"Monitoring AI-Modified Content at Scale: A Case Study on the Impact of ChatGPT on AI Conference Peer Reviews","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.07183","snapshot_observed_at":"2026-08-11T15:02:39.587144Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.587144Z"},"links":{"cited_paper":"/paper/2403.07183","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:e6fa75bbe4f31525de68fb0faffed615f608858ae6f0c89384bece9d5dfc0a40","observation_id":"4da2632e-b2a7-42f7-bdb0-97644cc1ae9d","resolution":{"observed_at":"2026-08-11T15:02:39.587144Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.01268","last_updated":"2024-04-01T17:45:15Z","snapshot_observed_at":"2026-08-14T14:52:07.103780Z","submitted_at":"2024-04-01T17:45:15Z","title":"Mapping the Increasing Use of LLMs in Scientific Papers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.01268","snapshot_observed_at":"2026-08-11T15:02:39.592222Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.592222Z"},"links":{"cited_paper":"/paper/2404.01268","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:4d1c335108e64bd5687d7b0b36e95479f86800daa58b6e2614568adbf55654b3","observation_id":"5e7691b5-e2d1-446a-aa66-776edef4220d","resolution":{"observed_at":"2026-08-11T15:02:39.592222Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.15799","last_updated":"2024-04-24T10:53:39Z","snapshot_observed_at":"2026-08-13T00:23:08.034633Z","submitted_at":"2024-04-24T10:53:39Z","title":"Towards the relationship between AIGC in manuscript writing and author profiles: evidence from preprints in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.15799","snapshot_observed_at":"2026-08-11T15:02:39.597020Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.597020Z"},"links":{"cited_paper":"/paper/2404.15799","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:c6042ad16f7b7c70bf1e8e8869eb1cc29edca21e1015294c626a965a08d95e20","observation_id":"6ebb589c-9c82-4868-b764-4d579785989b","resolution":{"observed_at":"2026-08-11T15:02:39.597020Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.405611Z","title":null,"venue":null,"work_id":"164f851f-94b0-4b12-b0b6-8b8f841e501d","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.604950Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:b42059780cfc03af980952358364f65362df43f7f583dcc67ee61315097669a1","observation_id":"f8b5835f-e729-45c6-a945-b4aa02052cba","resolution":{"observed_at":"2026-08-11T15:02:40.409844Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.392882Z","title":null,"venue":null,"work_id":"bf8d2196-e9eb-4bda-8f29-944b7b2fc60c","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.613373Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:43af64f09c3fa1ba45cfdfe58cae93125dc51a7bbea2dc734f9fe79646c908c4","observation_id":"9dfcd2c4-ee91-4bd0-a72b-ac3ec5bc4c95","resolution":{"observed_at":"2026-08-11T15:02:40.397128Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.379744Z","title":null,"venue":null,"work_id":"57b3a606-6c89-430a-bd3b-a0d84c2103ad","year":2010},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.617280Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:cda76f527505615eba0a8f2511aae81070f1400967a985b8687f6dfcb75c061e","observation_id":"b8325231-224b-4ef2-99ab-7c50909a9d1c","resolution":{"observed_at":"2026-08-11T15:02:40.383954Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.368525Z","title":null,"venue":null,"work_id":"5f4cebe9-9104-4b05-8254-21e379ed79b7","year":2022},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.620846Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:d067bc711424391984328ee38a0fa2136184e89ff2c441fc1298cac9a4c33842","observation_id":"89ddd9d1-7a6d-4fc4-9a2d-c618012471e3","resolution":{"observed_at":"2026-08-11T15:02:40.371845Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.356714Z","title":null,"venue":null,"work_id":"b074a026-52aa-459b-97e4-60e6a680c2ba","year":2023},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.624629Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:8317f2f69a2385dbae352496f25bfccd1dea679c3489a5c05f0b70e6a2bce10b","observation_id":"21f92e3f-f670-44a0-ab04-8d2242cc253d","resolution":{"observed_at":"2026-08-11T15:02:40.360716Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"status/1774021","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.977982Z","title":null,"venue":null,"work_id":"49d0720f-ae19-4186-86dc-0c1e3604aa34","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.628054Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:f2111441702b82fa072d5063e9f1384a38c43042629a1573248929122a7177b1","observation_id":"48035704-52e0-4e1d-804f-bb7488ad8761","resolution":{"observed_at":"2026-08-11T15:02:39.988572Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.344319Z","title":null,"venue":null,"work_id":"4728fc82-4151-42aa-816c-0130a9adb01c","year":2023},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.632075Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:ccd69e4b2ffcb363c2f1856fb5ac1f75139e2cd339dbfcb3edd3ba6fa1e86eae","observation_id":"dfcbd27a-c3e9-4e7f-9a88-ee26bbc58ccb","resolution":{"observed_at":"2026-08-11T15:02:40.348536Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.635166Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.635166Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:f0ffd7bb688b971ff08beb07d05e96720d742f93288331055d12415bea8a3bd7","observation_id":"97841b47-d2a8-4b78-9a80-61ec6b8b4247","resolution":{"observed_at":"2026-08-11T15:02:39.635166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.332456Z","title":null,"venue":null,"work_id":"a7321035-977b-42f3-a547-39bdf244ed50","year":2021},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.639723Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:b41d57abf01dea6d6be9410aa9f985e17354d13bbec9ea2eced7759de1ea697b","observation_id":"3a88b4be-0a91-45e4-8b53-560f9dc9ec96","resolution":{"observed_at":"2026-08-11T15:02:40.336312Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.643998Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.643998Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:70fbc2b7dace4d0117ac67f072ab73ecf7ab9158f05f3a9d25293eff5eb0cc51","observation_id":"f2a85eaf-7c8e-47c4-8fdd-7db308f93dff","resolution":{"observed_at":"2026-08-11T15:02:39.643998Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.312881Z","title":null,"venue":null,"work_id":"d58e0b51-b4a6-4223-8072-3c42397202d1","year":2023},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.648252Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:bebe255893d1520cda885b79955a523e0fc50835783870c8dfa43976b7f2b080","observation_id":"b8977726-f7c7-4434-993f-9c9eb30031c4","resolution":{"observed_at":"2026-08-11T15:02:40.316869Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.300060Z","title":null,"venue":null,"work_id":"e35a90f5-ce1a-4610-a225-070ed0d1376c","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.651932Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:6c3786db4c71b2e9a049054eb7aaea72e20a04f842b72b2fd5820e7d926281c1","observation_id":"2c36d081-b44c-41bd-bf12-ea4597ab7ccb","resolution":{"observed_at":"2026-08-11T15:02:40.304082Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.288756Z","title":null,"venue":null,"work_id":"a28f946a-002b-40cc-9062-5d98adf0f28b","year":2022},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.655897Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:b8110287fe07e86843743dd7740ae2fc8f3111f5f2b33abfca80a11772c28ae8","observation_id":"aa054e70-4ed7-4dba-8bfa-484f2c34ef80","resolution":{"observed_at":"2026-08-11T15:02:40.292416Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.275205Z","title":null,"venue":null,"work_id":"2003d0e7-17a7-428d-9861-fd2df5f92036","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.659794Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:ecd3b9c0ab2e30b450356908363544fb90a2db2c8ef5de322a50f485186e186d","observation_id":"64748305-ae4c-4c01-9f23-fe5fff6febdb","resolution":{"observed_at":"2026-08-11T15:02:40.280163Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.263760Z","title":null,"venue":null,"work_id":"e14e695a-cc0e-4ecd-a22e-ff9396400512","year":2023},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.665217Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:c7f89948d93d2acaeff53c8543a1478ff976c45f097acf17cc1368b807f61a9a","observation_id":"33ca5707-aeb0-4576-af7c-2f1f19b4a31a","resolution":{"observed_at":"2026-08-11T15:02:40.267569Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.251737Z","title":null,"venue":null,"work_id":"7ae8ee48-4d03-4c1a-9d4d-e3a56337c333","year":2015},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.669138Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:aeb8a4b19031e56a0434a82cf69a7384db249ba4911135a4a1804a6a58feb212","observation_id":"d74afef0-34e8-456d-891a-93e827ff4e61","resolution":{"observed_at":"2026-08-11T15:02:40.255835Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.673335Z","title":null,"venue":null,"work_id":null,"year":1948},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.673335Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:863015aadf78001bd66b4f65f68ea75cdae59431a048f9141a995fca3cf228d1","observation_id":"40ee536c-30d1-4597-9821-804a0306f2fd","resolution":{"observed_at":"2026-08-11T15:02:39.673335Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.234937Z","title":null,"venue":null,"work_id":"af9e3810-2ffa-4709-a64d-b5da38d4c0a1","year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.677415Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:5ca51bfc1e7c77cb6002caf8aefcee8c260e61dc4548188f96dc4a935d61149c","observation_id":"8e420eba-c6f8-404a-836e-7f2472383adb","resolution":{"observed_at":"2026-08-11T15:02:40.238179Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17493","last_updated":"2024-04-14T05:20:10Z","snapshot_observed_at":"2026-08-05T16:03:40.517679Z","submitted_at":"2023-05-27T15:10:41Z","title":"The Curse of Recursion: Training on Generated Data Makes Models Forget","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.17493","snapshot_observed_at":"2026-08-11T15:02:39.681519Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.681519Z"},"links":{"cited_paper":"/paper/2305.17493","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:df0711d9bf5b37e6885fe7059cb413462752a75b5d06f11711f32724c0d81603","observation_id":"1dd8f647-4ae9-4f70-a210-cfa6278356a3","resolution":{"observed_at":"2026-08-11T15:02:39.681519Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.685717Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.685717Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:a63f3c7ea5169b59b131307ed5f433e3caa532e557f6fbf0fed65954a8f8658d","observation_id":"b7823f30-20c7-4932-a6fc-9acfbdbf0b5f","resolution":{"observed_at":"2026-08-11T15:02:39.685717Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-07T12:56:43.323460Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-11T15:02:39.690656Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.690656Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:c2f2ae045d98e2951c25838120d40785fc6af034b4124f061fd48dca10a5672b","observation_id":"ea9884f2-6958-4bdb-937e-8cbcfc536b8d","resolution":{"observed_at":"2026-08-11T15:02:39.690656Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:40.215876Z","title":null,"venue":null,"work_id":"fb07dadd-7d62-43dc-9148-3740b1072c33","year":2021},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.695571Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:c85e406796465325be23a810e5d4aaae7b81d16b748f0eaf80570af2263b495b","observation_id":"d1f084ce-fe8a-44c5-b087-01e14354c6a1","resolution":{"observed_at":"2026-08-11T15:02:40.220594Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-15T06:32:42.880941+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.700944Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.700944Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:2e5f136c6135ca535e1311c20cf92855eb5a97261a31e5b6bf83141dc8698515","observation_id":"479e8048-d0d5-43c3-9a0d-ca05c93d2d0a","resolution":{"observed_at":"2026-08-11T15:02:39.700944Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.01754","last_updated":"2026-07-16T15:15:44Z","snapshot_observed_at":"2026-08-12T22:52:29.838259Z","submitted_at":"2024-09-03T10:01:51Z","title":"Empirical evidence of Large Language Model's influence on human spoken communication","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.01754","snapshot_observed_at":"2026-08-11T15:02:39.705996Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.705996Z"},"links":{"cited_paper":"/paper/2409.01754","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:aa3ad242babe58a0c3fa503b1d1c824e11884ffd5241d040528bb97f540f23f6","observation_id":"24228682-6ed1-4983-b481-4a139ac5ef26","resolution":{"observed_at":"2026-08-11T15:02:39.705996Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.10625","last_updated":"2023-04-16T22:08:08Z","snapshot_observed_at":"2026-08-06T09:00:42.886249Z","submitted_at":"2022-05-21T15:34:53Z","title":"Least-to-Most Prompting Enables Complex Reasoning in Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.10625","snapshot_observed_at":"2026-08-11T15:02:39.712292Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.712292Z"},"links":{"cited_paper":"/paper/2205.10625","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:4ec42d93143916a5cce36e6520002e1aafd8659b1de5086d935ef8f2818ae522","observation_id":"60bb11a3-0ad5-42a5-bdf5-a758294ed8aa","resolution":{"observed_at":"2026-08-11T15:02:39.712292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1909.08593","last_updated":"2020-01-08T23:02:36Z","snapshot_observed_at":"2026-08-15T10:46:42.452851Z","submitted_at":"2019-09-18T17:33:39Z","title":"Fine-Tuning Language Models from Human Preferences","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1909.08593","snapshot_observed_at":"2026-08-11T15:02:39.716831Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.716831Z"},"links":{"cited_paper":"/paper/1909.08593","citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:984f48cb8800f1d39937a2de145b23600e01ad8f60a9aee3dca0c20f2c3d4e01","observation_id":"56daf52d-f173-454a-96d1-74895e2c27ac","resolution":{"observed_at":"2026-08-11T15:02:39.716831Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.721028Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.721028Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:fad4ea3512bc68994935bb65021b399b94e1d3c5316a030939bf3a6c30668efa","observation_id":"c4e2b19b-b8b4-4b08-b1ba-06acb199cf84","resolution":{"observed_at":"2026-08-11T15:02:39.721028Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T15:02:39.725628Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-11T15:02:39.725628Z"},"links":{"citing_paper":"/paper/2412.11385"},"observation_digest":"sha256:4b0ffc4c22636251de0a7546f36a4c485c310cbc6de15d7e966ec56aadde3a23","observation_id":"02bd1ccc-3fce-46c6-9c71-ecde52744c36","resolution":{"observed_at":"2026-08-11T15:02:39.725628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2412.11385","last_updated":"2024-12-16T02:27:59Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-13T18:53:13.535644Z","submitted_at":"2024-12-16T02:27:59Z","title":"Why Does ChatGPT \"Delve\" So Much? Exploring the Sources of Lexical Overrepresentation in Large Language Models"},"reference_resolution":{"displayed":57,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":55,"verified_exact":1,"verified_fuzzy":1},"total_outbound_references":57},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-15T06:32:42.880941+00:00","source":"crossref"},{"observed_at":"2026-08-15T06:32:39.529945+00:00","source":"retraction_watch"}],"thesis":"As of 15 August 2026, this Paper Citation Record lists 57 of 57 outbound references and 2 inbound Pith citation observations for arXiv:2412.11385."}