{"as_of":"2026-08-18T00:44:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:fd0d3c8521e8a477f8d0e45b63266654daa529f757169db372cf8bb9ea4ff0e6","coverage":[{"denominator":19,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":19,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:28:34.870429Z","state":"measured"},{"denominator":21,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":21,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:27:50.270831Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-07T05:27:50.650205Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"cited_work":{"arxiv_id":"2506.02522","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.02522","snapshot_observed_at":"2026-08-07T05:27:50.650205Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","venue":"cs.AI","work_id":"1b4cd731-8e74-4650-aaa2-43f65c0e858c","year":2025},"citing_paper":{"arxiv_id":"2506.07976","last_updated":"2025-06-10T12:50:18Z","snapshot_observed_at":"2026-08-16T07:40:03.364615Z","submitted_at":"2025-06-09T17:50:02Z","title":"Thinking vs. Doing: Agents that Reason by Scaling Test-Time Interaction","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-07T05:27:50.270831Z"},"links":{"cited_paper":"/paper/2506.02522","citing_paper":"/paper/2506.07976"},"observation_digest":"sha256:1a337cccac98cb87ed5dc7e317f5ea8731b534b2f0984ccc4f7e95673f02b224","observation_id":"220a18c6-fc7e-4132-af06-b6ec616c7f4b","resolution":{"observed_at":"2026-08-07T05:27:50.656331Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.02522","snapshot_observed_at":"2026-08-01T18:32:47.516639Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.17281","last_updated":"2026-07-19T14:59:12Z","snapshot_observed_at":"2026-08-08T22:36:45.673041Z","submitted_at":"2026-07-19T14:59:12Z","title":"AIGB-R1: Self-Evolving Generative Auto-Bidding via Hierarchical Planner-Executor Optimization","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-01T18:32:47.516639Z"},"links":{"cited_paper":"/paper/2506.02522","citing_paper":"/paper/2607.17281"},"observation_digest":"sha256:64575b0a5dc8becac1a569101093d05fa0fc669e3f1e4283d03b7e3f2d23341d","observation_id":"875246e5-b29c-49ab-8784-3ef4efaaa397","resolution":{"observed_at":"2026-08-01T18:32:47.516639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.02522/citation-record","integrity":"/paper/2506.02522/integrity","json":"/paper/2506.02522/citation-record.json","paper":"/paper/2506.02522"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:38.238554Z","title":null,"venue":null,"work_id":"a238f5c8-5524-4ce2-a2f8-9e68ad948d8a","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:33.415893Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:03953bc4e0a224399b25f02b79b5847b24d169370c3bc23c067469f387fe1b1e","observation_id":"1f335743-a4d8-44b0-bd21-9d7f708683aa","resolution":{"observed_at":"2026-08-07T11:28:38.357713Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:38.034545Z","title":"This indirect approach may help redistribute the load and reduce stress on overloaded lines","venue":null,"work_id":"aaffffba-025d-4e1d-9270-3b08b18e0994","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:33.483874Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:ab7cdaa03a4a90793c88f7b5cf1b5a16363a10b7f180e774c27d81771c42c0f7","observation_id":"8d2ed40a-3e31-477d-9b0f-8469af8c790d","resolution":{"observed_at":"2026-08-07T11:28:38.124839Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:37.781480Z","title":"Response Format: Please analyze the situation and provide your response in the following format:","venue":null,"work_id":"ae49f599-a98b-4a7f-a692-b0885b14ca63","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:33.591029Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:b0b426bfe0b6b86b56ce28bba9399c2768a0f33f050fab23e6d79072c3293e00","observation_id":"83ebd474-b496-4467-96b7-74e72f135892","resolution":{"observed_at":"2026-08-07T11:28:37.908860Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:36.452344Z","title":null,"venue":null,"work_id":"c12e8cac-fdd7-427d-8335-507e29ad44ce","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.112315Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:2f1df0c2b1b8c834fa8e91a5ab9877dd032a793ffc32a7165334e671fc3f5e87","observation_id":"ba3ff72b-f7c0-4db4-b668-238ed4a2d7df","resolution":{"observed_at":"2026-08-07T11:28:36.569688Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:36.270687Z","title":null,"venue":null,"work_id":"2bfde8d7-129f-463b-af9e-d2058cc64ab4","year":2012},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.215049Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:69749a39f6747523b99737e21c4e32cd422d889cd66a15619f60bb846ab43185","observation_id":"b3353fcf-a473-4010-bfb1-7ee956f448f0","resolution":{"observed_at":"2026-08-07T11:28:36.344039Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:37.544938Z","title":null,"venue":null,"work_id":"1d2268cf-fdee-4936-bb87-7dd81d7bf75e","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:33.648087Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:ce03a56ae7c5dbd84b0b9894065e3a836f11650952397b72301765db9c48936e","observation_id":"b3c5939d-5bca-4081-b108-f355c06890b0","resolution":{"observed_at":"2026-08-07T11:28:37.661405Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:37.311132Z","title":null,"venue":null,"work_id":"14d5ddcb-34f2-4684-8bd6-72dad633a660","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:33.736506Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:26f4b57ec72fabfd483c3102f80424a2499cff27019d4d1aec5a026147bc5145","observation_id":"ef894c80-0f56-4784-94a3-e87a99e0d046","resolution":{"observed_at":"2026-08-07T11:28:37.449814Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:37.089742Z","title":null,"venue":null,"work_id":"cd5f2b15-e86c-48c3-b6ce-0249347dd13f","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:33.832598Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:151ecc56c34a420f961c9f2e6559018e817b25553fc8784c317043ef71860bf3","observation_id":"9e545372-8807-420d-a531-45a550df92b2","resolution":{"observed_at":"2026-08-07T11:28:37.200160Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:36.868608Z","title":"Do not add any bold formatting, asterisks, or other special characters","venue":null,"work_id":"54044217-d3f0-406a-850b-95415e8e8ba0","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:33.937756Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:67c9ee77956f6bc8f73021b4afebe5aab68904ad9be46d3e1d72e33ca313df76","observation_id":"004c9d1f-9415-4a8f-88c5-34f5e51c1b34","resolution":{"observed_at":"2026-08-07T11:28:36.957770Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:36.669914Z","title":null,"venue":null,"work_id":"56fc2935-232a-4638-a1db-2d452f8248e4","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.038427Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:18b1402f26aa60baf1bb1a8ceeac7104ccf5d33d2d5c1b84d43fe9f7dbeaefb7","observation_id":"9170f947-dd51-4270-8980-6c61086a9334","resolution":{"observed_at":"2026-08-07T11:28:36.783099Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:36.062980Z","title":null,"venue":null,"work_id":"be92c9b2-2e8a-4fed-b5d0-4a02c40455f3","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.332683Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:d783475f3ea9ef5a31ee073eecd88a3aa61847c64a47540776c0f47e592e87cb","observation_id":"b5d03cd0-9ae5-4cec-9791-f1a290bd870b","resolution":{"observed_at":"2026-08-07T11:28:36.175249Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:35.832763Z","title":null,"venue":null,"work_id":"e8982c04-0d8a-45bd-9c46-57e437c6a3c7","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.444238Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:43911fbc63658210757128d354979a4007ef76a4b20b5ad575d5bbf31bf7ba96","observation_id":"ab48466a-aee0-44ca-91b3-6d6e086c48ba","resolution":{"observed_at":"2026-08-07T11:28:35.940234Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:35.692350Z","title":"Provide your analysis results in the following format: Key Decision Point Indices: [X, Y , A, B] Reward Adjustments: [W, V , T, S]","venue":null,"work_id":"c7e84bba-2660-4e78-b8f8-16dae651b9aa","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.537903Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:cab6437c686f7907f1dae4e51faa6bd8da77087143537c59a76ad7a2cea0a0f9","observation_id":"e2cce5ea-204f-4798-9eec-640e0f3fdb2a","resolution":{"observed_at":"2026-08-07T11:28:35.774192Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:35.472513Z","title":null,"venue":null,"work_id":"a746ae47-b578-422f-bff4-a9ecc1b12fd3","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.579666Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:907f41d7d2892dab5e4371f45affec103e78ae1d67e6c391f31c9c206bc86b75","observation_id":"81282aff-970e-4ba9-a56b-736856669363","resolution":{"observed_at":"2026-08-07T11:28:35.587029Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:35.368266Z","title":null,"venue":null,"work_id":"7dba4927-c229-49f2-8d85-e2d3014c8208","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.676695Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:e70f83f652fb71f4264b6c48374f9dcf177825ea8e3e4df576be7ed1bfcbfaa4","observation_id":"964ea50d-072e-42e5-a21c-90ac4fca0db0","resolution":{"observed_at":"2026-08-07T11:28:35.395167Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:35.212557Z","title":null,"venue":null,"work_id":"a6363f6c-7aff-4288-a304-7b69bda856a1","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.780169Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:28a42d9c2f3520b6daed603d825dacd626daeffdcec08e32b16318af43b7dc0e","observation_id":"0bbff5db-78fb-4fed-890d-3682a4a99a4a","resolution":{"observed_at":"2026-08-07T11:28:35.284666Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:28:35.032745Z","title":"For fLLM, this includes power grid topology rules and operational constraints, while for gLLM, it focuses on reward assessment criteria and safety standards","venue":null,"work_id":"74ebb029-5432-4930-aa3e-2415e77b10eb","year":null},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:34.870429Z"},"links":{"citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:48468a38c1cea0ffa5e438b5004369a170b5634fe8b5c0a38ee99c8c30f99340","observation_id":"fb7859fd-c488-4285-858d-19cfd71402f7","resolution":{"observed_at":"2026-08-07T11:28:35.111568Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.08958","last_updated":"2023-12-14T14:07:41Z","snapshot_observed_at":"2026-08-16T14:34:46.573369Z","submitted_at":"2023-12-14T14:07:41Z","title":"LiFT: Unsupervised Reinforcement Learning with Foundation Models as Teachers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.08958","snapshot_observed_at":"2026-08-07T11:28:33.304802Z","title":"Nam, T., Lee, J., Zhang, J., Hwang, S","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":6299,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:33.304802Z"},"links":{"cited_paper":"/paper/2312.08958","citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:8a2e8a595478fbb1e2048a51bec383d6823a7fd95649b771630422bc223303f2","observation_id":"9536d212-42d4-4a51-ad15-448fd7227b50","resolution":{"observed_at":"2026-08-07T11:28:33.304802Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1710.11248","last_updated":"2018-08-13T18:33:24Z","snapshot_observed_at":"2026-08-14T20:18:41.596978Z","submitted_at":"2017-10-30T21:22:28Z","title":"Learning Robust Rewards with Adversarial Inverse Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1710.11248","snapshot_observed_at":"2026-08-07T11:28:33.254559Z","title":"Fu, J., Luo, K., and Levine, S","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making","version":1},"reference_index":8677,"source":"pdf_text","source_observed_at":"2026-08-07T11:28:33.254559Z"},"links":{"cited_paper":"/paper/1710.11248","citing_paper":"/paper/2506.02522"},"observation_digest":"sha256:cd3a4c604b17df210bbd4795e928ede1175256c01e4108e372a97d95a12de70a","observation_id":"a812d602-5a9f-462b-bee5-07be222a85d0","resolution":{"observed_at":"2026-08-07T11:28:33.254559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.02522","last_updated":"2025-06-03T06:52:37Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-17T18:34:56.722469Z","submitted_at":"2025-06-03T06:52:37Z","title":"Think Twice, Act Once: A Co-Evolution Framework of LLM and RL for Large-Scale Decision Making"},"reference_resolution":{"displayed":19,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":14,"verified_exact":0,"verified_fuzzy":5},"total_outbound_references":19},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 19 of 19 outbound references and 2 inbound Pith citation observations for arXiv:2506.02522."}