{"as_of":"2026-08-17T14:06:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:63f6968986408c4b9abe36ef4c253e0ec1eb59de1940689de81a8b91cea97501","coverage":[{"denominator":60,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":60,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T01:03:43.999176Z","state":"measured"},{"denominator":61,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":61,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T00:55:32.686738Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.12217","snapshot_observed_at":"2026-08-03T00:55:32.686738Z","title":"arXiv preprint arXiv:2506.12217 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28636","last_updated":"2026-05-19T13:56:13Z","snapshot_observed_at":"2026-08-17T01:28:35.393539Z","submitted_at":"2026-05-19T13:56:13Z","title":"Chain-of-Models: Cross-Model Auditing for Bias-Robust LLM Judges","version":1},"reference_index":201,"source":"arxiv_source","source_observed_at":"2026-08-03T00:55:32.686738Z"},"links":{"cited_paper":"/paper/2506.12217","citing_paper":"/paper/2607.28636"},"observation_digest":"sha256:24032d8adaf30d38fd5deea50b820d4a2724dd55382eac0520a32db18505c365","observation_id":"09e6d2ee-3503-4e81-8bf6-3281c83c64a6","resolution":{"observed_at":"2026-08-03T00:55:32.686738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.12217/citation-record","integrity":"/paper/2506.12217/integrity","json":"/paper/2506.12217/citation-record.json","paper":"/paper/2506.12217"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.09686","last_updated":"2025-01-23T08:44:44Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-16T17:37:58Z","title":"Towards Large Reasoning Models: A Survey of Reinforced Reasoning with Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09686","snapshot_observed_at":"2026-08-07T01:03:38.561365Z","title":"Towards large reasoning models: A survey of reinforced reasoning with large language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.561365Z"},"links":{"cited_paper":"/paper/2501.09686","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:b7ef7c3e3f6b85991d593ed96d4490f67f1dc358412a4806033c933522dd8227","observation_id":"1629fedd-c51a-443c-8c86-a9e44716013d","resolution":{"observed_at":"2026-08-07T01:03:38.561365Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:50.006721Z","title":"Self-reasoning language models: Unfold hidden reasoning chains with few reasoning catalyst,","venue":null,"work_id":"27856238-d4e3-4a6d-87a4-556895271826","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.644078Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:ac74e418d3ad033bb951569ffd563653a6e193d3611fce2223a04d956c7fd6f3","observation_id":"c3f47b45-1cf7-4040-b36c-88f44fb55353","resolution":{"observed_at":"2026-08-07T01:03:50.091457Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:38.746258Z","title":"Reinforcement learning with verifiable rewards: Grpo’s effective loss, dynamics, and success amplification,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.746258Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:55ea27bfcf40c3d2a052f403788476a5a444af0d5ec941d428ddb8bc20a4b451","observation_id":"d39c0232-4edf-4f1f-93f0-77c3ec0e014d","resolution":{"observed_at":"2026-08-07T01:03:38.746258Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05379","last_updated":"2025-03-10T07:11:14Z","snapshot_observed_at":"2026-08-16T12:52:11.523420Z","submitted_at":"2025-03-07T12:46:42Z","title":"R1-Omni: Explainable Omni-Multimodal Emotion Recognition with Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.05379","snapshot_observed_at":"2026-08-07T01:03:38.824120Z","title":"R1-omni: Explainable omni-multimodal emotion recognition with rein- forcement learning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.824120Z"},"links":{"cited_paper":"/paper/2503.05379","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:111aa4a469202c081d83aaa29b2241b5c8b2fc55532343e318d5b2e59646c0b2","observation_id":"72ccaf44-fc65-488a-a1cd-10b9b7ca1794","resolution":{"observed_at":"2026-08-07T01:03:38.824120Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:38.946446Z","title":"Reasoning beyond limits: Advances and open problems for llms,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:38.946446Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:097ec3d2291c883ea81bc3c5fc220b87ad7d62ae6417c6952793070a8c77c7b3","observation_id":"95f02528-e180-46f1-b36e-a332307da8d4","resolution":{"observed_at":"2026-08-07T01:03:38.946446Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23829","last_updated":"2025-04-01T14:48:02Z","snapshot_observed_at":"2026-08-16T12:45:26.696220Z","submitted_at":"2025-03-31T08:22:49Z","title":"Crossing the Reward Bridge: Expanding RL with Verifiable Rewards Across Diverse Domains","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.23829","snapshot_observed_at":"2026-08-07T01:03:39.063087Z","title":"Expanding rl with verifiable rewards across diverse domains,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.063087Z"},"links":{"cited_paper":"/paper/2503.23829","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:aebe44de7977246ce6c4d751886ad12322101d44779c120bd05e21cb49a954cb","observation_id":"45d137f3-ca87-4e2f-b01a-e084297c6e34","resolution":{"observed_at":"2026-08-07T01:03:39.063087Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-08-15T12:33:55.451951Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-07T01:03:39.150691Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.150691Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:064362cf64ed6456fd9793173a83c43b04f38fee56f5faefe74298983b79516e","observation_id":"6ce3f4ec-e201-4c21-8cfa-de0702166cd9","resolution":{"observed_at":"2026-08-07T01:03:39.150691Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.20783","last_updated":"2025-10-06T09:30:03Z","snapshot_observed_at":"2026-08-13T12:34:54.476684Z","submitted_at":"2025-03-26T17:59:14Z","title":"Understanding R1-Zero-Like Training: A Critical Perspective","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.20783","snapshot_observed_at":"2026-08-07T01:03:39.248657Z","title":"Understanding r1-zero-like training: A critical perspective,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.248657Z"},"links":{"cited_paper":"/paper/2503.20783","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:a5b406832bde6ab230f0fcc36721cb370e591e1da8f9d862d6a56f6ce3611b92","observation_id":"00266b23-2c47-4af4-a0f7-ac7c0164fc40","resolution":{"observed_at":"2026-08-07T01:03:39.248657Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.18892","last_updated":"2025-08-06T08:42:32Z","snapshot_observed_at":"2026-07-06T20:57:57.039376Z","submitted_at":"2025-03-24T17:06:10Z","title":"SimpleRL-Zoo: Investigating and Taming Zero Reinforcement Learning for Open Base Models in the Wild","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.18892","snapshot_observed_at":"2026-08-07T01:03:39.348782Z","title":"Simplerl-zoo: Investigating and taming zero reinforcement learning for open base models in the wild,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.348782Z"},"links":{"cited_paper":"/paper/2503.18892","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:9283b36f27c66debe4cd71330bb00857f9e1c01696da47eb88fbf6e183106c90","observation_id":"fec69ab1-664a-4a23-b22f-5e9923cbb2d1","resolution":{"observed_at":"2026-08-07T01:03:39.348782Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16084","last_updated":"2025-06-30T15:59:26Z","snapshot_observed_at":"2026-08-15T13:55:17.887932Z","submitted_at":"2025-04-22T17:59:56Z","title":"TTRL: Test-Time Reinforcement Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.16084","snapshot_observed_at":"2026-08-07T01:03:39.439450Z","title":"Ttrl: Test-time reinforcement learning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.439450Z"},"links":{"cited_paper":"/paper/2504.16084","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:4eff79975b36f8488d6e35658408dc73921d59933504d753835b501c1af80c28","observation_id":"39da37bb-30e3-44ac-a74b-d2899675fa2e","resolution":{"observed_at":"2026-08-07T01:03:39.439450Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13837","last_updated":"2025-11-24T06:11:04Z","snapshot_observed_at":"2026-08-14T14:03:15.178702Z","submitted_at":"2025-04-18T17:59:56Z","title":"Does Reinforcement Learning Really Incentivize Reasoning Capacity in LLMs Beyond the Base Model?","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.13837","snapshot_observed_at":"2026-08-07T01:03:39.555965Z","title":"Does reinforcement learning really incentivize reasoning capacity in llms beyond the base model?,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.555965Z"},"links":{"cited_paper":"/paper/2504.13837","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:9839c17697a29c8c9657055cb1b8b5536b37af547315a518059f58c65a86f203","observation_id":"6e48c6ec-e596-4e7a-a061-ee7584b866ec","resolution":{"observed_at":"2026-08-07T01:03:39.555965Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10400","last_updated":"2025-02-16T09:33:15Z","snapshot_observed_at":"2026-08-16T13:42:49.208149Z","submitted_at":"2024-06-14T20:07:11Z","title":"Self-Reflection Makes Large Language Models Safer, Less Biased, and Ideologically Neutral","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10400","snapshot_observed_at":"2026-08-07T01:03:39.646395Z","title":"Self-reflection outcome is sensitive to prompt construction,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.646395Z"},"links":{"cited_paper":"/paper/2406.10400","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:1c7574fe782d318bf43f79c1f12ad5ee807944be7e33d8c824ec37d507797724","observation_id":"5012415e-37d7-479a-b840-2af01d93001e","resolution":{"observed_at":"2026-08-07T01:03:39.646395Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:39.763786Z","title":"Dynamic early exit in reasoning models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.763786Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:b654476b7e4c290e398f04301c2d72b34ead2b1e712e9e108223cdc6dad57b9f","observation_id":"3a583017-4b0b-4ce9-a51f-739836db59d9","resolution":{"observed_at":"2026-08-07T01:03:39.763786Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.06682","last_updated":"2024-10-16T23:19:46Z","snapshot_observed_at":"2026-08-16T13:55:15.115349Z","submitted_at":"2024-05-05T18:56:46Z","title":"Self-Reflection in LLM Agents: Effects on Problem-Solving Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.06682","snapshot_observed_at":"2026-08-07T01:03:39.875381Z","title":"Self-reflection in llm agents: Effects on problem-solving performance,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.875381Z"},"links":{"cited_paper":"/paper/2405.06682","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:b6f1b5028c99ea2ee75872390ca078de3fe750d027d4b3849e945c94c98f8d49","observation_id":"0ff7c1eb-0746-4af3-a970-991237b9ad2b","resolution":{"observed_at":"2026-08-07T01:03:39.875381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.16419","last_updated":"2025-08-21T19:14:40Z","snapshot_observed_at":"2026-08-11T13:10:23.709172Z","submitted_at":"2025-03-20T17:59:38Z","title":"Stop Overthinking: A Survey on Efficient Reasoning for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.16419","snapshot_observed_at":"2026-08-07T01:03:39.949009Z","title":"Stop overthinking: Asurveyonefficientreasoningforlargelanguagemodels,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:39.949009Z"},"links":{"cited_paper":"/paper/2503.16419","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:3bc77d3ab9dc4131f8a50c2d7bafaa281abd4b77d05774d19dd157160d94f1d4","observation_id":"e75301d6-8fc0-4c2b-8fd8-258fd41755a5","resolution":{"observed_at":"2026-08-07T01:03:39.949009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.808216Z","title":"Steering llama 2 via contrastive activation addition,","venue":null,"work_id":"0745a100-d217-4ccd-87e1-d5d74358f3b1","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.075252Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:867d7f1befcb25acfe8897bb439433fd608d2ed794f188c8d96f27605d7624a4","observation_id":"1c3c4d59-7537-49e7-87a4-48675c718d1f","resolution":{"observed_at":"2026-08-07T01:03:49.915796Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1801.10198","last_updated":"2018-01-30T20:07:01Z","snapshot_observed_at":"2026-08-14T19:50:32.053437Z","submitted_at":"2018-01-30T20:07:01Z","title":"Generating Wikipedia by Summarizing Long Sequences","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1801.10198","snapshot_observed_at":"2026-08-07T01:03:40.150110Z","title":"Generating wikipedia by summarizing long sequences,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.150110Z"},"links":{"cited_paper":"/paper/1801.10198","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:9c1e5fdace0ad4c453ed93e57b9cd0558ff5f233d8f50e09ab55253e7aa7816c","observation_id":"dfcd3ad3-4410-464d-bd53-7cafe43242bd","resolution":{"observed_at":"2026-08-07T01:03:40.150110Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.607196Z","title":"Beyond accuracy: Evaluating the reasoning behavior of large language models - a survey,","venue":null,"work_id":"cbbc5d79-a6f5-4984-90a3-4b1982b41863","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.258094Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:86ff3f01f499b53a5d864ec56c3ed6e49ebf92e894dd10fe899b13d10b387c39","observation_id":"d7c5a580-fc9f-4a8e-afd2-9006c9ee36f6","resolution":{"observed_at":"2026-08-07T01:03:49.713626Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.409413Z","title":"Oat: A research-friendly framework for llm online alignment,","venue":null,"work_id":"aed4a729-51c1-4061-977d-474fb660b103","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.356585Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:7b23a4c0762497cb8b9797e283c670012e4abd0c98966cc3db294a7cd4864435","observation_id":"6edd9d22-89a5-4dc0-9c6f-af10c5a203a7","resolution":{"observed_at":"2026-08-07T01:03:49.490397Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.221893Z","title":"Self- consistency improves chain of thought reasoning in language models,","venue":null,"work_id":"54b239e4-7c40-4e94-9625-effc4d5b436e","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.427956Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d3fd1fd582702c7115641fb081c7b5664d8161bd7d42c3ed07f9c441fdb02eae","observation_id":"aa276d97-fc20-4b0b-9055-911fb8bdbae9","resolution":{"observed_at":"2026-08-07T01:03:49.324127Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:49.065199Z","title":"X-reasoner: Towards generalizable reasoning across modalities and domains,","venue":null,"work_id":"d7fd41cd-9359-4c55-a606-0653fcea7162","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.536878Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:94eb796508c39035d2cccf5a1cffd0685aa251403ee33ef2450089d7d29a71c7","observation_id":"f3ca3234-ffce-4943-a683-0ebd222aaa98","resolution":{"observed_at":"2026-08-07T01:03:49.145880Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.10400","last_updated":"2025-02-24T08:57:10Z","snapshot_observed_at":"2026-08-15T21:48:04.229100Z","submitted_at":"2024-12-05T16:10:42Z","title":"Reinforcement Learning Enhanced LLMs: A Survey","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.10400","snapshot_observed_at":"2026-08-07T01:03:40.631750Z","title":"Reinforcement learning enhanced llms: A survey,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.631750Z"},"links":{"cited_paper":"/paper/2412.10400","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:74dbbb91377926b4b3a1e1c8ad8919a83f2561889a65f76e1b762f404b50d35d","observation_id":"271c3ef1-143a-4feb-a781-27cf896f307c","resolution":{"observed_at":"2026-08-07T01:03:40.631750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.03335","last_updated":"2025-10-16T08:23:36Z","snapshot_observed_at":"2026-08-11T20:12:00.210052Z","submitted_at":"2025-05-06T09:08:00Z","title":"Absolute Zero: Reinforced Self-play Reasoning with Zero Data","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.03335","snapshot_observed_at":"2026-08-07T01:03:40.693792Z","title":"Absolute zero: Reinforced self-play reasoning with zero data,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.693792Z"},"links":{"cited_paper":"/paper/2505.03335","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:242e647a2a8d4e7528075603cf9899a2fbd6a708d15a99ebb3339136512e9b55","observation_id":"a66ec951-743b-4545-99c4-8e4e8a841cbf","resolution":{"observed_at":"2026-08-07T01:03:40.693792Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.09129","last_updated":"2024-04-14T02:47:32Z","snapshot_observed_at":"2026-08-16T14:01:09.543696Z","submitted_at":"2024-04-14T02:47:32Z","title":"When Hindsight is Not 20/20: Testing Limits on Reflective Thinking in Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.09129","snapshot_observed_at":"2026-08-07T01:03:40.835799Z","title":"When hindsight is not 20/20: Testing limits on reflective thinking in large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.835799Z"},"links":{"cited_paper":"/paper/2404.09129","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:089f8d59b64ec3281d6cb5f749e86063bcf45837faef9a29db3e734bcbd37fd9","observation_id":"f9f122bc-638e-4812-94d2-a741f75f391c","resolution":{"observed_at":"2026-08-07T01:03:40.835799Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.892164Z","title":"Demystifyinglongchain-of-thoughtreasoninginLLMs,","venue":null,"work_id":"593e40ec-1d05-43bb-b7cf-55f9fd549fd6","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:40.918114Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d9033bda1f6eb1cbe178fc91c326d8298e1a436e7819511e2720e3638221e347","observation_id":"d7ed41c2-594d-4bc7-8560-4063ddfa933d","resolution":{"observed_at":"2026-08-07T01:03:48.981885Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16720","last_updated":"2026-04-30T02:46:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-21T18:04:31Z","title":"OpenAI o1 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16720","snapshot_observed_at":"2026-08-07T01:03:41.000635Z","title":"Openai o1 system card,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.000635Z"},"links":{"cited_paper":"/paper/2412.16720","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:a9bbf0983080d109bb079235f09fb03451b5bb5efa8988e5b487e0614f1a006f","observation_id":"9646293b-ee21-4f86-9fc0-d89a56d2de0f","resolution":{"observed_at":"2026-08-07T01:03:41.000635Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.00656","last_updated":"2025-10-08T07:50:45Z","snapshot_observed_at":"2026-08-17T13:58:40.683829Z","submitted_at":"2024-12-31T21:55:10Z","title":"2 OLMo 2 Furious","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.00656","snapshot_observed_at":"2026-08-07T01:03:41.101622Z","title":"2 olmo 2 furious,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.101622Z"},"links":{"cited_paper":"/paper/2501.00656","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:e9935a2b734f0d9ced00bddca7ef132d6e8501bb81c4576c6e908a622c9dd7e1","observation_id":"469303ad-d1e6-4663-a969-09b9a3aad043","resolution":{"observed_at":"2026-08-07T01:03:41.101622Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.708810Z","title":"Qwen3technical report,","venue":null,"work_id":"3e1375b7-3354-43fb-ab92-0eaab029689c","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.187541Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:605461fee412140d22eb53fbfbb77090e13ae335408e1fdf1c59fbfb8a0fdc0d","observation_id":"a6bebd1d-2878-441b-b47a-317a95e26e40","resolution":{"observed_at":"2026-08-07T01:03:48.803809Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.524846Z","title":"Measuring mathematical problem solving with the math dataset,","venue":null,"work_id":"c0c06683-3b2b-4300-ad76-be6b82fc90fc","year":2021},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.281529Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:01fa41f65f720eaa7b26eb462f0b2d5c02e3aad14a6b0e60ef44a7d44f66492e","observation_id":"7c38a833-c42f-4bdf-b9e3-6d709d13d656","resolution":{"observed_at":"2026-08-07T01:03:48.604954Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-07T01:03:41.367515Z","title":"Qwen2.5 technical report,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.367515Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:ab0a7edbb07d7eb8e8fd4e89649238538152afb9e1c250ac705a725126aa14c0","observation_id":"c3b0ce6b-1dcd-4094-a847-9ac37c243daf","resolution":{"observed_at":"2026-08-07T01:03:41.367515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.305684Z","title":"Umap: Uniform manifold approximation and projection,","venue":null,"work_id":"4c53cb44-2ed2-4de6-bb60-4224652c4f14","year":2018},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.455519Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:9dcf10bae198029c9413cc6a74a569791e0e7a66a1835156194b5a454a71cffc","observation_id":"5f2a934e-8188-4b61-b082-cfc0a740d4f2","resolution":{"observed_at":"2026-08-07T01:03:48.402203Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:48.105369Z","title":"Refusal in language models is mediated by a single direction,","venue":null,"work_id":"22c5c963-05d1-4048-be22-af5458561e6b","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.581041Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d433bcda851fdeb00a48ad2e77a81b0d3c02d60707e98002ae6052ca9ece0e05","observation_id":"c48b7e33-6880-4a41-9bcc-4e2415ae0261","resolution":{"observed_at":"2026-08-07T01:03:48.223985Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.17148","last_updated":"2025-03-03T21:15:30Z","snapshot_observed_at":"2026-08-16T12:58:32.947185Z","submitted_at":"2025-01-28T18:51:24Z","title":"AxBench: Steering LLMs? Even Simple Baselines Outperform Sparse Autoencoders","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.17148","snapshot_observed_at":"2026-08-07T01:03:41.633267Z","title":"Axbench: Steering llms? even simple baselines outperform sparse autoencoders,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.633267Z"},"links":{"cited_paper":"/paper/2501.17148","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:48b21e5c202172bc0f5a4d6d8a787e0d04717fc2df80c3adb421790f4ccecd73","observation_id":"aa3cfd68-4461-42d7-9c7b-067819b1df5e","resolution":{"observed_at":"2026-08-07T01:03:41.633267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.927657Z","title":"GPQA: A graduate-level google-proof q&a benchmark,","venue":null,"work_id":"b4a52e1e-fe79-414e-b1fd-2d6c6bcb1404","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.688003Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:9f38f52197f8276a146939187fafc105d6db32b4d5925fbd16def0e4fec1590b","observation_id":"574728a9-8514-40e5-808d-9ebd8bffb441","resolution":{"observed_at":"2026-08-07T01:03:48.007208Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T01:03:41.772289Z","title":"The llama 3 herd of models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.772289Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:4025093aa08e128ecddda8ded5f0a12d524c029c47cc09c423da011d1f6641f1","observation_id":"663e3187-2685-436c-8ec6-a625da5ab9ca","resolution":{"observed_at":"2026-08-07T01:03:41.772289Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.775398Z","title":"s1: Simple test-time scaling,","venue":null,"work_id":"4c7242c8-6b95-4250-a836-309efa758e88","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.883084Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:0f183805996771633d47afe4145f0d552716e120cb4245d792c37eb7aa22c73a","observation_id":"176a923f-3bc6-4556-bb92-d9d3e3b7a73e","resolution":{"observed_at":"2026-08-07T01:03:47.843540Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.609327Z","title":"Discovering latent knowledge in language models without supervision,","venue":null,"work_id":"d9ecbae2-a49a-4b2c-92b8-cab0d4ac86d1","year":2022},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:41.999426Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:aa8f8e5bfdf43d0631df61fb3552ee1c326a8fb11e9b06fc347706dcef21716a","observation_id":"455bfbd3-f1f3-4569-9fb0-13256cdc8c65","resolution":{"observed_at":"2026-08-07T01:03:47.687530Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-08-12T09:06:50.363435Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15043","snapshot_observed_at":"2026-08-07T01:03:42.097166Z","title":"Universal and transferable adversarial attacks on aligned language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.097166Z"},"links":{"cited_paper":"/paper/2307.15043","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:b6526d270c054462f0a8eda26f7acc473ea6eaf8bd4a9f536c426e1db09c7b59","observation_id":"e080e39d-6fdc-46bd-b468-be482af58d7d","resolution":{"observed_at":"2026-08-07T01:03:42.097166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14433","last_updated":"2024-02-22T10:25:14Z","snapshot_observed_at":"2026-08-16T14:16:18.185492Z","submitted_at":"2024-02-22T10:25:14Z","title":"A Language Model's Guide Through Latent Space","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14433","snapshot_observed_at":"2026-08-07T01:03:42.165517Z","title":"A language model’s guide through latent space,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.165517Z"},"links":{"cited_paper":"/paper/2402.14433","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:ce337ac20efd4b60e7d027bf4cc501d3378108d512ac2daee82a81f316e0d263","observation_id":"066fae06-1e46-4c3e-943f-2ef5c7231f6c","resolution":{"observed_at":"2026-08-07T01:03:42.165517Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.03813","last_updated":"2023-12-06T18:27:07Z","snapshot_observed_at":"2026-08-16T14:37:04.314654Z","submitted_at":"2023-12-06T18:27:07Z","title":"Improving Activation Steering in Language Models with Mean-Centring","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.03813","snapshot_observed_at":"2026-08-07T01:03:42.231617Z","title":"Improvingactivationsteeringinlanguagemodels with mean-centring,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.231617Z"},"links":{"cited_paper":"/paper/2312.03813","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:a98f72805f9462ef7da54a68d115b2a1f278fae28fa9a0d444b483584d2197c3","observation_id":"ccda8ba5-f1c8-4351-b338-06c32a8b5839","resolution":{"observed_at":"2026-08-07T01:03:42.231617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.443594Z","title":"Finding alignments between interpretable causal variables and distributed neural representations,","venue":null,"work_id":"704a07e8-8912-4fd7-a283-7781f17fdfa2","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.293333Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:5429360af3b39f81f1a44f959c74b5cef878d195d9007d23f7c1585acdbb8591","observation_id":"4b8fa7cc-b175-46cf-a4bc-2e07327b259b","resolution":{"observed_at":"2026-08-07T01:03:47.540585Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.231503Z","title":"Generative agents: Interactivesimulacraofhumanbehavior,","venue":null,"work_id":"58604775-c758-4c0d-b3a3-8065e31b59a3","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.349684Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:80ffb4d15f043cf34da8fdb5c51ba571161c619b05a8cb1e6b06c75739f808c4","observation_id":"75b62a26-f9ce-45a9-bf6d-25a79750aa6c","resolution":{"observed_at":"2026-08-07T01:03:47.341340Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:47.070162Z","title":"Man is to computer programmer as woman is to homemaker? debiasing word embeddings,","venue":null,"work_id":"951afe25-0bda-4db5-98a1-fa91255a26b8","year":2016},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.442869Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:36ecee178cb9ba523f1253663b21108b0552135d5f681173b6d3233dc7d35778","observation_id":"2e2ee699-326c-43f4-9ce6-e929d40d1b2a","resolution":{"observed_at":"2026-08-07T01:03:47.140600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.909600Z","title":"Sparseautoencodersfindhighlyinter- pretable features in language models,","venue":null,"work_id":"0cf0e31a-b14b-4305-b6eb-9b42fe8f49e7","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.512310Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:904f6a16f1cf12b99f7925013852b33b569c3948ddbfeeda9bb107e20b0dd223","observation_id":"5646bbed-0a47-48ee-92c0-a8de0ca8edb5","resolution":{"observed_at":"2026-08-07T01:03:46.980720Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.760950Z","title":"Gold doesn‘t always glitter: Spectral removal of linear and nonlinear guardedattributeinformation,","venue":null,"work_id":"2e587e80-176b-42ad-9a49-df3b7049bd34","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.632456Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:28600e8f4df309bfee910c583cf5f9cad0ecdbd822c8d49cbb3f84654cfef2a0","observation_id":"f7461889-ce63-41fe-89a9-fbd38f2db7d1","resolution":{"observed_at":"2026-08-07T01:03:46.832967Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.595599Z","title":"LEACE: Perfect linear concept erasure in closed form,","venue":null,"work_id":"4df176e9-c6c9-4cca-a5dd-d001179ab7c2","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.748817Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:751da90a843ce6ff23e596854d344d0c3d6d9b87fd2968db1168c58d5b0c2806","observation_id":"ee86d914-0601-459b-b18b-37f5071d3aff","resolution":{"observed_at":"2026-08-07T01:03:46.659998Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.416338Z","title":"Monitoring latent world states in language models with proposi- tional probes,","venue":null,"work_id":"f193c926-0c04-4ff0-95f2-031e38902e3a","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.820828Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:cfef62eb1d0ec1a3b30ec30ab7359d75d64533aab919e17faaed5c09b05eb999","observation_id":"4876e0ea-6e93-4259-9ab9-3a14a1faeb66","resolution":{"observed_at":"2026-08-07T01:03:46.506000Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.276366Z","title":"Let’s verify step by step,","venue":null,"work_id":"6a2867a7-344a-4b77-aaf8-667407e0867b","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.869100Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:6850614830459264375987940c339ae3f7cdaea278d607c6be5691a5faa8b0fe","observation_id":"14c8a776-3a4d-498f-8938-059650e3aa86","resolution":{"observed_at":"2026-08-07T01:03:46.336643Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:46.141934Z","title":"Self-refine: Iterative refinement with self-feedback,","venue":null,"work_id":"ff0f3286-ad0f-4c29-bc52-e2d444f3d0e8","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:42.995770Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:dce546156336caf78a467eccdd242a5f566a19d8ebc8bc91ed6fb4ea41571ca0","observation_id":"45964bdf-c1b5-4127-be86-22247b70986f","resolution":{"observed_at":"2026-08-07T01:03:46.194754Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.954249Z","title":"Fine-tuning with divergent chains of thought boosts reasoning through self-correction in language models,","venue":null,"work_id":"a6694914-fc18-421e-9cdc-66de285902fc","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.095901Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:f4ba9df5ca071c502614285f98a9f7581d256837cdc82ab7f13f39f3e6fae545","observation_id":"4f77234f-d40c-4b67-ac9a-698fee7368e6","resolution":{"observed_at":"2026-08-07T01:03:46.041393Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.807342Z","title":"STar: Bootstrapping reasoning with reasoning,","venue":null,"work_id":"f2da2629-3a02-45ed-b049-a19ebdf94b3f","year":2022},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.195392Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:9d81a2f60418de5bcfb5340099387a60689f2d25a311e1cd09109072da54c1d1","observation_id":"b0fac060-c3ff-45bf-be9d-73d9cfe73d47","resolution":{"observed_at":"2026-08-07T01:03:45.858720Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.594857Z","title":"Let’s verify step by step,","venue":null,"work_id":"6e0ab76d-2a2b-465f-ba08-d41449af2328","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.254472Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:9d911b4c392109619c2d65d3b35e7433958853d0890f515d34407c8193f2c56a","observation_id":"2912de11-1a0d-4b74-84d1-31f3dffcfaae","resolution":{"observed_at":"2026-08-07T01:03:45.699379Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.07374","last_updated":"2025-02-18T05:20:33Z","snapshot_observed_at":"2026-08-15T09:05:14.806013Z","submitted_at":"2025-02-11T08:48:48Z","title":"LLMs Can Easily Learn to Reason from Demonstrations Structure, not content, is what matters!","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.07374","snapshot_observed_at":"2026-08-07T01:03:43.321333Z","title":"Llms can easily learn to reason from demonstrations structure, not content, is what matters!,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.321333Z"},"links":{"cited_paper":"/paper/2502.07374","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:4b8e03a065cd5f7a494f5ece4bf75a65d142ddfbfe17bd1d390d34eca4ba7673","observation_id":"b5856e66-e3bf-4112-b2c0-a838138b03f0","resolution":{"observed_at":"2026-08-07T01:03:43.321333Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:43.439679Z","title":"Shorterbetter: Guidingreasoningmodelstofindoptimalinferencelengthforefficient reasoning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.439679Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:67fe015966e9f8b5d710c6e64eade95aeb03d8be503443b3a02a70e9678ce499","observation_id":"1af91da7-bec5-4818-8733-323b5f4b5934","resolution":{"observed_at":"2026-08-07T01:03:43.439679Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.332299Z","title":"Unlocking the capabilities of thought: A reasoning boundary framework to quantify and optimize chain-of-thought,","venue":null,"work_id":"08dbf383-56b3-4cdb-bcf3-380ae1905c70","year":2024},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.510697Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:1403a881eb76a9d68d96db2273f1553822aee513fae62a30edfe5cc599963a4f","observation_id":"ee2f8254-4724-4d15-b0f4-e5a60cf4666e","resolution":{"observed_at":"2026-08-07T01:03:45.472186Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12599","last_updated":"2025-06-03T02:14:54Z","snapshot_observed_at":"2026-08-15T22:38:53.825110Z","submitted_at":"2025-01-22T02:48:14Z","title":"Kimi k1.5: Scaling Reinforcement Learning with LLMs","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12599","snapshot_observed_at":"2026-08-07T01:03:43.598058Z","title":"Kimi k1. 5: Scaling reinforcement learning with llms,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.598058Z"},"links":{"cited_paper":"/paper/2501.12599","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:ea0a5e4e7151f72945fe88361807ca6f90feca21ff8e19f31652dcc1e7ce6940","observation_id":"1e93143a-2574-4054-be9a-79baf7e88ac1","resolution":{"observed_at":"2026-08-07T01:03:43.598058Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.04022","last_updated":"2025-04-05T02:24:07Z","snapshot_observed_at":"2026-08-16T12:43:48.435094Z","submitted_at":"2025-04-05T02:24:07Z","title":"Rethinking Reflection in Pre-Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.04022","snapshot_observed_at":"2026-08-07T01:03:43.678277Z","title":"Rethinking reflection in pre-training,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.678277Z"},"links":{"cited_paper":"/paper/2504.04022","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:22842b4092cf00fe26c27b732a65ba2d7300a849a17731a8e7887a7e68410aba","observation_id":"c86d77b5-78b6-4952-baaf-a77e3547715d","resolution":{"observed_at":"2026-08-07T01:03:43.678277Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:45.103569Z","title":"Reflexion: Languageagentswithverbal reinforcement learning,","venue":null,"work_id":"c73c2aaf-8ec5-465e-8aaf-39a6e521d6fe","year":2023},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.774554Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:a7fb2baa8ee57a7832297abf9afc403b45af0dda8e40d21db2ead60dbb286b51","observation_id":"ab5573a1-46ee-49f2-b945-5d27a0a59f08","resolution":{"observed_at":"2026-08-07T01:03:45.213263Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.20571","last_updated":"2025-10-24T10:02:36Z","snapshot_observed_at":"2026-08-15T17:20:54.840134Z","submitted_at":"2025-04-29T09:24:30Z","title":"Reinforcement Learning for Reasoning in Large Language Models with One Training Example","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.20571","snapshot_observed_at":"2026-08-07T01:03:43.838361Z","title":"Rein- forcement learning for reasoning in large language models with one training example,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.838361Z"},"links":{"cited_paper":"/paper/2504.20571","citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:d8091922865c2142df504b0cdf706f3ace6dd52c19832fb8c94df75941f09bab","observation_id":"3fa9b162-1ea8-45d9-b1a8-faa954e327d7","resolution":{"observed_at":"2026-08-07T01:03:43.838361Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:03:44.974472Z","title":"There may not be aha moment in r1-zero-like training — a pilot study","venue":null,"work_id":"e5444a1f-56e7-4686-bfe8-778a818f14fd","year":2025},"citing_paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-07T01:03:43.999176Z"},"links":{"citing_paper":"/paper/2506.12217"},"observation_digest":"sha256:39db1b5aee319a8cddf9b4d6c3f71ef1eaf52ad0a2991a6c37fd46a85fbdcc54","observation_id":"33193c25-3851-45af-84ae-f7a83be9ac32","resolution":{"observed_at":"2026-08-07T01:03:45.036245Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.12217","last_updated":"2025-06-13T20:40:13Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-13T15:09:57.842743Z","submitted_at":"2025-06-13T20:40:13Z","title":"From Emergence to Control: Probing and Modulating Self-Reflection in Language Models"},"reference_resolution":{"displayed":60,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":31,"verified_exact":0,"verified_fuzzy":29},"total_outbound_references":60},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 17 August 2026, this Paper Citation Record lists 60 of 60 outbound references and 1 inbound Pith citation observation for arXiv:2506.12217."}