{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:KP6EGVPA6F5UBH6473RRXPREIP","short_pith_number":"pith:KP6EGVPA","schema_version":"1.0","canonical_sha256":"53fc4355e0f17b409fdcfee31bbe2443dc3dbf71c320961d26de8aac02a91402","source":{"kind":"arxiv","id":"2309.01458","version":1},"attestation_state":"computed","paper":{"title":"Leveraging Reward Consistency for Interpretable Feature Discovery in Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.LG","authors_text":"Gao Huang, Huanqian Wang, Mukun Tong, Qisen Yang, Shiji Song, Wenjie Shi","submitted_at":"2023-09-04T09:09:54Z","abstract_excerpt":"The black-box nature of deep reinforcement learning (RL) hinders them from real-world applications. Therefore, interpreting and explaining RL agents have been active research topics in recent years. Existing methods for post-hoc explanations usually adopt the action matching principle to enable an easy understanding of vision-based RL agents. In this paper, it is argued that the commonly used action matching principle is more like an explanation of deep neural networks (DNNs) than the interpretation of RL agents. It may lead to irrelevant or misplaced feature attribution when different DNNs' o"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2309.01458","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2023-09-04T09:09:54Z","cross_cats_sorted":[],"title_canon_sha256":"2e0e4f512aa8cbad6d0fbe19a55be92423afd81d4bc6b465bed9b705fbb4cf65","abstract_canon_sha256":"b257ffe8b701c13201a01f099816cf8aa4e26b6f58a2b24baf250fc8fcd6b5ca"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T06:47:38.215354Z","signature_b64":"v/bxhHjNdJQY5tm5viHhQZNwE8XnHgbA7LqQxnYxtshuNgyccvCrEJB+vNyWPIJ85pCrEJsb91WEILIeBhY/Cg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"53fc4355e0f17b409fdcfee31bbe2443dc3dbf71c320961d26de8aac02a91402","last_reissued_at":"2026-07-05T06:47:38.214914Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T06:47:38.214914Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Leveraging Reward Consistency for Interpretable Feature Discovery in Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.LG","authors_text":"Gao Huang, Huanqian Wang, Mukun Tong, Qisen Yang, Shiji Song, Wenjie Shi","submitted_at":"2023-09-04T09:09:54Z","abstract_excerpt":"The black-box nature of deep reinforcement learning (RL) hinders them from real-world applications. Therefore, interpreting and explaining RL agents have been active research topics in recent years. Existing methods for post-hoc explanations usually adopt the action matching principle to enable an easy understanding of vision-based RL agents. In this paper, it is argued that the commonly used action matching principle is more like an explanation of deep neural networks (DNNs) than the interpretation of RL agents. It may lead to irrelevant or misplaced feature attribution when different DNNs' o"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2309.01458","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2309.01458/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2309.01458","created_at":"2026-07-05T06:47:38.214962+00:00"},{"alias_kind":"arxiv_version","alias_value":"2309.01458v1","created_at":"2026-07-05T06:47:38.214962+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2309.01458","created_at":"2026-07-05T06:47:38.214962+00:00"},{"alias_kind":"pith_short_12","alias_value":"KP6EGVPA6F5U","created_at":"2026-07-05T06:47:38.214962+00:00"},{"alias_kind":"pith_short_16","alias_value":"KP6EGVPA6F5UBH64","created_at":"2026-07-05T06:47:38.214962+00:00"},{"alias_kind":"pith_short_8","alias_value":"KP6EGVPA","created_at":"2026-07-05T06:47:38.214962+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/KP6EGVPA6F5UBH6473RRXPREIP","json":"https://pith.science/pith/KP6EGVPA6F5UBH6473RRXPREIP.json","graph_json":"https://pith.science/api/pith-number/KP6EGVPA6F5UBH6473RRXPREIP/graph.json","events_json":"https://pith.science/api/pith-number/KP6EGVPA6F5UBH6473RRXPREIP/events.json","paper":"https://pith.science/paper/KP6EGVPA"},"agent_actions":{"view_html":"https://pith.science/pith/KP6EGVPA6F5UBH6473RRXPREIP","download_json":"https://pith.science/pith/KP6EGVPA6F5UBH6473RRXPREIP.json","view_paper":"https://pith.science/paper/KP6EGVPA","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2309.01458&json=true","fetch_graph":"https://pith.science/api/pith-number/KP6EGVPA6F5UBH6473RRXPREIP/graph.json","fetch_events":"https://pith.science/api/pith-number/KP6EGVPA6F5UBH6473RRXPREIP/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/KP6EGVPA6F5UBH6473RRXPREIP/action/timestamp_anchor","attest_storage":"https://pith.science/pith/KP6EGVPA6F5UBH6473RRXPREIP/action/storage_attestation","attest_author":"https://pith.science/pith/KP6EGVPA6F5UBH6473RRXPREIP/action/author_attestation","sign_citation":"https://pith.science/pith/KP6EGVPA6F5UBH6473RRXPREIP/action/citation_signature","submit_replication":"https://pith.science/pith/KP6EGVPA6F5UBH6473RRXPREIP/action/replication_record"}},"created_at":"2026-07-05T06:47:38.214962+00:00","updated_at":"2026-07-05T06:47:38.214962+00:00"}