Files
hermes-agent/tests/agent/test_ghost_skill_pruning.py
ethernet 890bbbda1f Merge remote-tracking branch 'origin/main' into ethie/pm-clean
# Conflicts:
#	apps/desktop/e2e/archived-hidden-session-recoverable.spec.ts
#	apps/desktop/e2e/bot-chat-message-agent-friendly-name.spec.ts
#	apps/desktop/e2e/bot-mailbox-unreadable-ticket.spec.ts
#	apps/desktop/e2e/bot-mode-roster-localized.spec.ts
#	apps/desktop/e2e/bot-mode-row-click-mirrors-registry.spec.ts
#	apps/desktop/e2e/bot-mode-tab-shows-bot-name.spec.ts
#	apps/desktop/e2e/bot-roster-group-row-organisation.spec.ts
#	apps/desktop/e2e/bot-roster-ignores-infra-dirs.spec.ts
#	apps/desktop/e2e/bot-roster-timestamp-meta.spec.ts
#	apps/desktop/e2e/bot-roster-user-sections.spec.ts
#	apps/desktop/e2e/bot-routines-pane-narrow.spec.ts
#	apps/desktop/e2e/bot-row-open-recent-session.spec.ts
#	apps/desktop/e2e/bot-tile-ignores-ambient-composer-model.spec.ts
#	apps/desktop/e2e/group-composer-auto-grow.spec.ts
#	apps/desktop/e2e/group-create-gate-remote-roster.spec.ts
#	apps/desktop/e2e/group-prompt-renamed-primary-handle.spec.ts
#	apps/desktop/e2e/hosted-room-backend-continuity.spec.ts
#	apps/desktop/e2e/hosted-room-legacy-store-migration.spec.ts
#	apps/desktop/e2e/settings-scope-chips-bot-title.spec.ts
#	apps/desktop/e2e/worktree-branch-status.spec.ts
#	apps/desktop/electron/backend-probes.test.ts
#	apps/desktop/electron/connection-apply.test.ts
#	apps/desktop/electron/desktop-electron-pin.test.ts
#	apps/desktop/electron/desktop-uninstall.test.ts
#	apps/desktop/electron/gateway-file-download-transport.test.ts
#	apps/desktop/electron/gateway-stop-before-update.test.ts
#	apps/desktop/electron/github-api-auth.test.ts
#	apps/desktop/electron/registry-primary-profile-scope.test.ts
#	apps/desktop/electron/update-api-check.test.ts
#	apps/desktop/electron/update-handoff-marker.test.ts
#	apps/desktop/electron/venv-blocker-scan.test.ts
#	apps/desktop/scripts/after-extract.test.mjs
#	apps/desktop/scripts/local-pack-publish.test.mjs
#	apps/desktop/scripts/tasks-scroll.test.mjs
#	apps/desktop/src/app/settings/model-settings.test.tsx
#	apps/desktop/src/app/updates-overlay.blockers.test.tsx
#	apps/desktop/src/components/desktop-install-overlay.test.tsx
#	apps/desktop/src/lib/update-copy.test.ts
#	scripts/ci/check_os_marker_fakes.py
#	tests-js/desktop-mac-usage-descriptions.test.ts
#	tests-js/node-engine-alignment.test.ts
#	tests/agent/lsp/test_install_and_lint_fixes.py
#	tests/agent/test_command_token_source.py
#	tests/agent/test_compression_boundary_hook.py
#	tests/agent/test_create_openai_client_ssl_verify.py
#	tests/agent/test_custom_provider_ca_probes.py
#	tests/agent/test_endpoint_blackhole.py
#	tests/agent/test_estimator_parity.py
#	tests/agent/test_in_place_compaction.py
#	tests/agent/test_moa_loop_mode.py
#	tests/agent/test_model_metadata.py
#	tests/agent/test_skill_session_platform_gate.py
#	tests/agent/test_skill_utils.py
#	tests/agent/test_ssl_ca_guard.py
#	tests/computer_use/test_doctor.py
#	tests/cron/test_codex_execution_paths.py
#	tests/cron/test_cron_bot_chat_delivery.py
#	tests/cron/test_cron_script.py
#	tests/cron/test_media_delivery_parity.py
#	tests/cron/test_misfire_catchup.py
#	tests/cron/test_parallel_pool.py
#	tests/cron/test_recurring_eagain_redispatch.py
#	tests/gateway/test_choice_picker.py
#	tests/gateway/test_control_socket_windows_live.py
#	tests/gateway/test_dingtalk.py
#	tests/gateway/test_feishu.py
#	tests/gateway/test_feishu_onboard.py
#	tests/gateway/test_gateway_shutdown.py
#	tests/gateway/test_matrix.py
#	tests/gateway/test_model_command_custom_providers.py
#	tests/gateway/test_reasoning_command.py
#	tests/gateway/test_runtime_footer.py
#	tests/gateway/test_session.py
#	tests/gateway/test_session_hygiene.py
#	tests/gateway/test_status.py
#	tests/gateway/test_teams.py
#	tests/gateway/test_turn_lease.py
#	tests/gateway/test_whatsapp_connect.py
#	tests/hermes_cli/test_approvals_command.py
#	tests/hermes_cli/test_auth_store_lock_concurrent.py
#	tests/hermes_cli/test_backup.py
#	tests/hermes_cli/test_banner_git_state.py
#	tests/hermes_cli/test_certifi_repair.py
#	tests/hermes_cli/test_cmd_update.py
#	tests/hermes_cli/test_compat_manifest_targets.py
#	tests/hermes_cli/test_computer_use_cli.py
#	tests/hermes_cli/test_cpr_local_leak.py
#	tests/hermes_cli/test_dashboard_auth_gate.py
#	tests/hermes_cli/test_dashboard_procs_kill_grace.py
#	tests/hermes_cli/test_desktop_lifecycle_windows_live.py
#	tests/hermes_cli/test_doctor.py
#	tests/hermes_cli/test_doctor_command_install.py
#	tests/hermes_cli/test_fleet_config_migration_windows_live.py
#	tests/hermes_cli/test_gateway.py
#	tests/hermes_cli/test_gateway_platform_gating.py
#	tests/hermes_cli/test_gateway_restart_loop.py
#	tests/hermes_cli/test_gateway_task_probe.py
#	tests/hermes_cli/test_gateway_wsl.py
#	tests/hermes_cli/test_gui_command.py
#	tests/hermes_cli/test_install_cua_driver.py
#	tests/hermes_cli/test_kanban_db.py
#	tests/hermes_cli/test_lazy_command_exports.py
#	tests/hermes_cli/test_lazy_refresh_venv_repair.py
#	tests/hermes_cli/test_linux_desktop_entry.py
#	tests/hermes_cli/test_local_runtime.py
#	tests/hermes_cli/test_local_runtime_updates.py
#	tests/hermes_cli/test_managed_uv.py
#	tests/hermes_cli/test_mcp_reload_confirm_gate.py
#	tests/hermes_cli/test_nous_subscription.py
#	tests/hermes_cli/test_npm_engine.py
#	tests/hermes_cli/test_personality_none.py
#	tests/hermes_cli/test_pet_toggle.py
#	tests/hermes_cli/test_plan_reconciliation_windows_live.py
#	tests/hermes_cli/test_plugin_event_bus.py
#	tests/hermes_cli/test_plugin_manifest_v2.py
#	tests/hermes_cli/test_plugin_packs.py
#	tests/hermes_cli/test_plugins_cmd.py
#	tests/hermes_cli/test_plugins_cmd_enable_disable_nested.py
#	tests/hermes_cli/test_process_identity.py
#	tests/hermes_cli/test_profiles.py
#	tests/hermes_cli/test_profiles_sidebar_cache.py
#	tests/hermes_cli/test_pty_bridge.py
#	tests/hermes_cli/test_resolve_turn_limit.py
#	tests/hermes_cli/test_serve_runtime_inventory.py
#	tests/hermes_cli/test_session_vacuum_config.py
#	tests/hermes_cli/test_set_config_value.py
#	tests/hermes_cli/test_signal_handler_kanban_worker.py
#	tests/hermes_cli/test_slash_confirm_windows.py
#	tests/hermes_cli/test_stale_pid_guard.py
#	tests/hermes_cli/test_startup_fast_guards.py
#	tests/hermes_cli/test_status.py
#	tests/hermes_cli/test_telegram_managed_bot.py
#	tests/hermes_cli/test_tools_config.py
#	tests/hermes_cli/test_update_apply_shallow_count.py
#	tests/hermes_cli/test_update_autostash.py
#	tests/hermes_cli/test_update_concurrent_quarantine.py
#	tests/hermes_cli/test_update_fetch_failure_classifier.py
#	tests/hermes_cli/test_update_fleet_probe_resume_token.py
#	tests/hermes_cli/test_update_handoff_backend_reap.py
#	tests/hermes_cli/test_update_handoff_desktop_rebuild.py
#	tests/hermes_cli/test_update_head_moved_gate.py
#	tests/hermes_cli/test_update_host_obligation.py
#	tests/hermes_cli/test_update_import_guard.py
#	tests/hermes_cli/test_update_interrupted_recovery.py
#	tests/hermes_cli/test_update_inventory.py
#	tests/hermes_cli/test_update_launchd_unloaded_gateway.py
#	tests/hermes_cli/test_update_missing_configured_deps.py
#	tests/hermes_cli/test_update_modified_notice.py
#	tests/hermes_cli/test_update_multiplex_migration_hook.py
#	tests/hermes_cli/test_update_no_gateway_restart.py
#	tests/hermes_cli/test_update_orphan_backend_reap.py
#	tests/hermes_cli/test_update_parked_branch_guard.py
#	tests/hermes_cli/test_update_post_pull_syntax_guard.py
#	tests/hermes_cli/test_update_receipt.py
#	tests/hermes_cli/test_update_self_lock.py
#	tests/hermes_cli/test_update_shim_fail_closed.py
#	tests/hermes_cli/test_update_shim_self_lock.py
#	tests/hermes_cli/test_update_sqlite_remediation.py
#	tests/hermes_cli/test_update_stale_dashboard.py
#	tests/hermes_cli/test_update_stale_virtualenv.py
#	tests/hermes_cli/test_update_venv_health.py
#	tests/hermes_cli/test_update_venv_ownership_preflight.py
#	tests/hermes_cli/test_update_wedged_gateway.py
#	tests/hermes_cli/test_update_yes_flag.py
#	tests/hermes_cli/test_update_zip_two_phase.py
#	tests/hermes_cli/test_urllib_security.py
#	tests/hermes_cli/test_ux_messages_auth_config.py
#	tests/hermes_cli/test_ux_messages_startup.py
#	tests/hermes_cli/test_venv_holder_classifier.py
#	tests/hermes_cli/test_verify_console_scripts.py
#	tests/hermes_cli/test_verify_core_dependencies.py
#	tests/hermes_cli/test_web_server.py
#	tests/hermes_cli/test_web_server_console_ws.py
#	tests/hermes_cli/test_web_server_ws_ping.py
#	tests/hermes_cli/test_web_ui_build.py
#	tests/hermes_state/test_fts_rebuild_admission.py
#	tests/hermes_state/test_hermes_state.py
#	tests/plugins/memory/test_memory_lazy_install.py
#	tests/plugins/test_google_meet_plugin.py
#	tests/plugins/test_langfuse_plugin.py
#	tests/plugins/test_security_guidance_plugin.py
#	tests/plugins/test_transform_llm_output_hook.py
#	tests/scripts/desktop_update/test_desktop_update_windows_gateway_flag.py
#	tests/scripts/desktop_update/test_desktop_update_windows_python_handoff.py
#	tests/scripts/desktop_update/test_desktop_update_windows_timestamp.py
#	tests/scripts/install/test_install_clone_throttle_fallback.py
#	tests/scripts/install/test_install_lockfile_churn.py
#	tests/scripts/install/test_install_no_initial_commit.py
#	tests/scripts/install/test_install_sh_browser_install.py
#	tests/scripts/install/test_install_sh_node_prerelease.py
#	tests/scripts/install/test_install_sh_symlink_stomp.py
#	tests/scripts/install/test_install_sh_uv_lock_config.py
#	tests/scripts/install/test_install_unmerged_index.py
#	tests/scripts/test_contributor_map.py
#	tests/scripts/test_run_tests_parallel.py
#	tests/skills/test_competitor_news_monitor_skill.py
#	tests/skills/test_document_to_action_items_skill.py
#	tests/skills/test_google_workspace_setup.py
#	tests/skills/test_google_workspace_setup_deps.py
#	tests/skills/test_grounded_citations_skill.py
#	tests/skills/test_ip_as_logo_skill.py
#	tests/skills/test_live_dashboard_skill.py
#	tests/skills/test_mcp_oauth_remote_gateway_skill.py
#	tests/skills/test_office_document_skills.py
#	tests/skills/test_openclaw_migration.py
#	tests/skills/test_product_price_monitor_skill.py
#	tests/skills/test_scrollcraft_skill.py
#	tests/skills/test_setup_wizard_generator_skill.py
#	tests/skills/test_weekly_review_planning_skill.py
#	tests/test_engines_satisfiable.py
#	tests/test_fast_safe_load.py
#	tests/test_hermes_bootstrap.py
#	tests/test_hermes_constants.py
#	tests/test_hermes_logging.py
#	tests/test_managed_runtime_resolution.py
#	tests/test_model_tools_async_bridge.py
#	tests/test_packaging_build_guard.py
#	tests/test_packaging_metadata.py
#	tests/test_yaml_indent_consistency.py
#	tests/tools/test_approval_timeout_overflow.py
#	tests/tools/test_base_environment.py
#	tests/tools/test_bot_mode_dm.py
#	tests/tools/test_browser_chromium_check.py
#	tests/tools/test_browser_hardening.py
#	tests/tools/test_browser_homebrew_paths.py
#	tests/tools/test_browser_npx_warmup.py
#	tests/tools/test_browser_orphan_reaper.py
#	tests/tools/test_browser_real_profile.py
#	tests/tools/test_browser_use_cli.py
#	tests/tools/test_clipboard.py
#	tests/tools/test_code_execution.py
#	tests/tools/test_code_execution_modes.py
#	tests/tools/test_code_execution_windows_env.py
#	tests/tools/test_computer_use.py
#	tests/tools/test_delegate_liveness_timeout.py
#	tests/tools/test_execute_code_approval_cluster.py
#	tests/tools/test_execution_flag_detection.py
#	tests/tools/test_fal_common.py
#	tests/tools/test_file_operations.py
#	tests/tools/test_file_tools.py
#	tests/tools/test_file_tools_cwd_resolution.py
#	tests/tools/test_file_tools_live.py
#	tests/tools/test_lazy_deps.py
#	tests/tools/test_lazy_deps_durable_target.py
#	tests/tools/test_lazy_deps_managed.py
#	tests/tools/test_local_env_blocklist.py
#	tests/tools/test_local_tempdir.py
#	tests/tools/test_macos_protected_search.py
#	tests/tools/test_mcp_npx_cached_bin.py
#	tests/tools/test_oneshot_completion_linger.py
#	tests/tools/test_process_registry.py
#	tests/tools/test_read_file_schema_gating.py
#	tests/tools/test_skill_improvements.py
#	tests/tools/test_skills_sync.py
#	tests/tools/test_termux_api_detection.py
#	tests/tools/test_tirith_security.py
#	tests/tools/test_transcription_tools.py
#	tests/tools/test_tts_streaming.py
#	tests/tools/test_wake_word.py
#	tests/tui_gateway/test_compute_host_borrowed_lease.py
#	tests/tui_gateway/test_compute_host_turn_protocol.py
#	tests/tui_gateway/test_isolated_orphan_activity.py
#	tests/tui_gateway/test_protocol.py
#	tests/tui_gateway/test_slash_worker_profile_home.py
#	tests/tui_gateway/test_subprocess_encoding.py
#	tests/tui_gateway/test_tui_gateway_server.py
#	ui-tui/src/__tests__/terminalParity.test.ts
#	ui-tui/src/__tests__/termuxComposerLayout.test.ts
#	ui-tui/src/__tests__/textInputFastEcho.test.ts
2026-09-23 07:02:44 -04:00

263 lines
10 KiB
Python

"""Ghost-skill defense tests (#32106, salvage of PR #44166).
When compaction reduces an old ``skill_view`` result to a metadata-only
summary, the model still believes the skill is loaded even though its
instructions are gone. The defense has three layers:
- P0/P1: the pruned tool-result summary carries a canonical
``[SKILL_PRUNED: ...]`` marker with the exact reload call, and the
system prompt (SKILLS_GUIDANCE) tells the model how to react to it.
- Phase-1 protection: a skill loaded just before compaction (or actively
referenced in the protected tail) keeps its full body through the
ordinary prune passes.
- P2: markers entering the summarizer are extracted BEFORE the aux LLM
call and deterministically re-injected if the model paraphrased them
away — including on the static fallback path.
Test patterns for the marker emit checks adapted from PR #32375
(@LeonSGP43) with credit.
"""
from unittest.mock import MagicMock, patch
from agent.context_compressor import (
SKILL_PRUNED_MARKER_PREFIX,
SUMMARY_PREFIX,
ContextCompressor,
_extract_pruned_skill_names,
_reinject_pruned_skill_markers,
_skill_pruned_marker,
_summarize_tool_result,
)
def _make_compressor(**overrides):
kwargs = dict(
model="test/model",
quiet_mode=True,
protect_first_n=1,
protect_last_n=2,
)
kwargs.update(overrides)
with patch(
"agent.context_compressor.get_model_context_length", return_value=100000
):
return ContextCompressor(**kwargs)
def _skill_view_pair(call_id, skill_name, size=6000):
return [
{
"role": "assistant",
"content": "",
"tool_calls": [{
"id": call_id,
"type": "function",
"function": {
"name": "skill_view",
"arguments": f'{{"name":"{skill_name}"}}',
},
}],
},
{
"role": "tool",
"tool_call_id": call_id,
"content": f"# {skill_name} instructions\n" + "x" * size,
},
]
class TestSkillPrunedMarkerEmit:
"""Marker emit — patterns adapted from PR #32375 (@LeonSGP43)."""
def test_skill_view_summary_marks_pruned_content(self):
summary = _summarize_tool_result(
"skill_view", '{"name":"docker-management"}', "x" * 6000
)
assert summary.startswith("[skill_view] name=docker-management (6,000 chars)")
assert _skill_pruned_marker("docker-management") in summary
assert "reload with skill_view(name='docker-management')" in summary
def test_small_skill_view_summary_not_marked(self):
summary = _summarize_tool_result(
"skill_view", '{"name":"docker-management"}', "x" * 1234
)
assert summary == "[skill_view] name=docker-management (1,234 chars)"
assert SKILL_PRUNED_MARKER_PREFIX not in summary
def test_marker_extractor_round_trips_the_emitted_marker(self):
"""Emit and check sides share one canonical string.
The original PR #44166 emitted ``[SKILL_PRUNED:`` but presence-
checked ``[SKILL_PRUNED]`` — re-injection fired even when the
marker had survived. Pin the round trip.
"""
summary = _summarize_tool_result("skill_view", '{"name":"pdf"}', "x" * 6000)
assert _extract_pruned_skill_names(summary) == ["pdf"]
assert _skill_pruned_marker("pdf") in summary
class TestReinjectPrunedSkillMarkers:
def test_reinjects_when_marker_paraphrased_away(self):
out = _reinject_pruned_skill_markers(
"The pdf skill was loaded earlier but its content was summarized.",
["pdf"],
)
assert _skill_pruned_marker("pdf") in out
assert "## Pruned Skills" in out
def test_partial_survival_reinjects_only_missing(self):
marker_a = _skill_pruned_marker("alpha")
out = _reinject_pruned_skill_markers("body\n" + marker_a, ["alpha", "beta"])
assert out.count(_skill_pruned_marker("alpha")) == 1
assert out.count(_skill_pruned_marker("beta")) == 1
def test_reinjected_summary_still_classifies_standalone(self):
body = _reinject_pruned_skill_markers("## Goal\nwork\n", ["pdf"])
full = SUMMARY_PREFIX + "\n\n" + body
assert ContextCompressor.classify_summary_content(full) == "standalone"
class TestProtectedSkillPrune:
"""Phase-1 prune must not demote a just-loaded skill (#32106)."""
def _filler(self, n, start=0):
out = []
for i in range(n):
role = "user" if (start + i) % 2 == 0 else "assistant"
out.append({"role": role, "content": f"filler {start + i} " + "y" * 400})
return out
def test_recently_loaded_skill_survives_prune(self):
c = _make_compressor()
# skill loaded within the last 10 messages, but OUTSIDE the
# protected tail count — without the guard it would be demoted.
msgs = (
self._filler(10)
+ _skill_view_pair("call_s", "fresh-skill")
+ self._filler(6, start=10)
)
result, _ = c._prune_old_tool_results(msgs, protect_tail_count=4)
skill_row = result[11]
assert skill_row["content"].startswith("# fresh-skill instructions")
assert SKILL_PRUNED_MARKER_PREFIX not in skill_row["content"]
def test_pressure_demotion_overrides_skill_protection(self):
"""Pass-4 must still demote protected skill bodies (#61932 guard)."""
c = _make_compressor()
msgs = (
self._filler(2)
+ _skill_view_pair("call_s", "fresh-skill", size=60000)
+ [{"role": "user", "content": "active ask"}]
)
# Tiny token budget → protected region exceeds the soft ceiling and
# the pressure pass must reclaim the skill body despite protection.
result, pruned = c._prune_old_tool_results(
msgs, protect_tail_count=4, protect_tail_tokens=100
)
skill_row = result[3]
assert pruned >= 1
assert _skill_pruned_marker("fresh-skill") in skill_row["content"]
class TestMarkerSurvivesRealCompress:
"""P2 layer: markers survive a real compress() with a mocked aux LLM."""
def _mock_response(self, text):
response = MagicMock()
response.choices = [MagicMock()]
response.choices[0].message.content = text
return response
def _messages_with_pruned_skill_in_middle(self):
"""Transcript whose compressed middle carries a prune marker row."""
pruned_row_content = (
"[skill_view] name=pdf (48,201 chars) " + _skill_pruned_marker("pdf")
)
msgs = [
{"role": "system", "content": "System prompt"},
{"role": "user", "content": "Build the PDF report"},
{
"role": "assistant",
"content": "",
"tool_calls": [{
"id": "call_pdf",
"type": "function",
"function": {
"name": "skill_view",
"arguments": '{"name":"pdf"}',
},
}],
},
{"role": "tool", "tool_call_id": "call_pdf", "content": pruned_row_content},
{"role": "assistant", "content": "Loaded the skill, working."},
{"role": "user", "content": "continue"},
{"role": "assistant", "content": "more work " + "z" * 500},
{"role": "user", "content": "latest ask"},
{"role": "assistant", "content": "ack"},
]
return msgs
def _summary_text_of(self, result):
for msg in result:
if ContextCompressor.classify_summary_content(msg.get("content")):
return msg["content"]
raise AssertionError(f"no summary message found in {result!r}")
def test_marker_reinjected_when_summarizer_drops_it(self):
c = _make_compressor(protect_first_n=1, protect_last_n=2)
msgs = self._messages_with_pruned_skill_in_middle()
drop_response = self._mock_response(
"## Goal\nBuild the PDF report.\n\n## Completed Actions\n"
"1. Loaded some skills and worked on the report."
)
with (
patch.object(c, "_find_tail_cut_by_tokens", return_value=7),
patch(
"agent.context_compressor.call_llm", return_value=drop_response
) as mock_call,
):
result = c.compress(msgs, force=True)
assert mock_call.called
summary_text = self._summary_text_of(result)
assert _skill_pruned_marker("pdf") in summary_text
# Stored iterative-update state carries the marker too.
assert _skill_pruned_marker("pdf") in c._previous_summary
def test_marker_survives_iterative_recompression(self):
"""Markers in a rehydrated handoff summary survive iterative rewrites.
On re-compression the previous handoff (carrying the marker) is
rehydrated into ``_previous_summary``; even when the summarizer's
iterative update drops the marker, re-injection restores it.
"""
c = _make_compressor(protect_first_n=1, protect_last_n=2)
prior_handoff = (
SUMMARY_PREFIX
+ "\n\n## Goal\nOld work.\n\n## Pruned Skills\n"
+ _skill_pruned_marker("pdf")
)
msgs = [
{"role": "system", "content": "System prompt"},
{"role": "user", "content": prior_handoff},
] + [
{"role": "assistant" if i % 2 == 0 else "user", "content": f"turn {i} " + "q" * 300}
for i in range(8)
]
drop_response = self._mock_response("## Goal\nNext task in flight.")
with (
patch.object(c, "_find_tail_cut_by_tokens", return_value=8),
patch("agent.context_compressor.call_llm", return_value=drop_response),
):
result = c.compress(msgs, force=True)
summary_text = self._summary_text_of(result)
assert _skill_pruned_marker("pdf") in summary_text
class TestReinjectionBoundsAndRedaction:
# The cap is applied at the collection sites in _generate_summary /
# _build_static_fallback_summary; the helper itself is mechanical.
def test_reinjection_block_is_redacted(self, monkeypatch):
import agent.redact as redact_mod
# force=True redaction must win even when redaction is disabled.
monkeypatch.setattr(redact_mod, "_REDACT_ENABLED", False, raising=False)
secret = "ghp_" + "a1B2" * 6
out = _reinject_pruned_skill_markers("body", [f"x {secret}"])
assert secret not in out