# Conflicts: # apps/desktop/e2e/archived-hidden-session-recoverable.spec.ts # apps/desktop/e2e/bot-chat-message-agent-friendly-name.spec.ts # apps/desktop/e2e/bot-mailbox-unreadable-ticket.spec.ts # apps/desktop/e2e/bot-mode-roster-localized.spec.ts # apps/desktop/e2e/bot-mode-row-click-mirrors-registry.spec.ts # apps/desktop/e2e/bot-mode-tab-shows-bot-name.spec.ts # apps/desktop/e2e/bot-roster-group-row-organisation.spec.ts # apps/desktop/e2e/bot-roster-ignores-infra-dirs.spec.ts # apps/desktop/e2e/bot-roster-timestamp-meta.spec.ts # apps/desktop/e2e/bot-roster-user-sections.spec.ts # apps/desktop/e2e/bot-routines-pane-narrow.spec.ts # apps/desktop/e2e/bot-row-open-recent-session.spec.ts # apps/desktop/e2e/bot-tile-ignores-ambient-composer-model.spec.ts # apps/desktop/e2e/group-composer-auto-grow.spec.ts # apps/desktop/e2e/group-create-gate-remote-roster.spec.ts # apps/desktop/e2e/group-prompt-renamed-primary-handle.spec.ts # apps/desktop/e2e/hosted-room-backend-continuity.spec.ts # apps/desktop/e2e/hosted-room-legacy-store-migration.spec.ts # apps/desktop/e2e/settings-scope-chips-bot-title.spec.ts # apps/desktop/e2e/worktree-branch-status.spec.ts # apps/desktop/electron/backend-probes.test.ts # apps/desktop/electron/connection-apply.test.ts # apps/desktop/electron/desktop-electron-pin.test.ts # apps/desktop/electron/desktop-uninstall.test.ts # apps/desktop/electron/gateway-file-download-transport.test.ts # apps/desktop/electron/gateway-stop-before-update.test.ts # apps/desktop/electron/github-api-auth.test.ts # apps/desktop/electron/registry-primary-profile-scope.test.ts # apps/desktop/electron/update-api-check.test.ts # apps/desktop/electron/update-handoff-marker.test.ts # apps/desktop/electron/venv-blocker-scan.test.ts # apps/desktop/scripts/after-extract.test.mjs # apps/desktop/scripts/local-pack-publish.test.mjs # apps/desktop/scripts/tasks-scroll.test.mjs # apps/desktop/src/app/settings/model-settings.test.tsx # apps/desktop/src/app/updates-overlay.blockers.test.tsx # apps/desktop/src/components/desktop-install-overlay.test.tsx # apps/desktop/src/lib/update-copy.test.ts # scripts/ci/check_os_marker_fakes.py # tests-js/desktop-mac-usage-descriptions.test.ts # tests-js/node-engine-alignment.test.ts # tests/agent/lsp/test_install_and_lint_fixes.py # tests/agent/test_command_token_source.py # tests/agent/test_compression_boundary_hook.py # tests/agent/test_create_openai_client_ssl_verify.py # tests/agent/test_custom_provider_ca_probes.py # tests/agent/test_endpoint_blackhole.py # tests/agent/test_estimator_parity.py # tests/agent/test_in_place_compaction.py # tests/agent/test_moa_loop_mode.py # tests/agent/test_model_metadata.py # tests/agent/test_skill_session_platform_gate.py # tests/agent/test_skill_utils.py # tests/agent/test_ssl_ca_guard.py # tests/computer_use/test_doctor.py # tests/cron/test_codex_execution_paths.py # tests/cron/test_cron_bot_chat_delivery.py # tests/cron/test_cron_script.py # tests/cron/test_media_delivery_parity.py # tests/cron/test_misfire_catchup.py # tests/cron/test_parallel_pool.py # tests/cron/test_recurring_eagain_redispatch.py # tests/gateway/test_choice_picker.py # tests/gateway/test_control_socket_windows_live.py # tests/gateway/test_dingtalk.py # tests/gateway/test_feishu.py # tests/gateway/test_feishu_onboard.py # tests/gateway/test_gateway_shutdown.py # tests/gateway/test_matrix.py # tests/gateway/test_model_command_custom_providers.py # tests/gateway/test_reasoning_command.py # tests/gateway/test_runtime_footer.py # tests/gateway/test_session.py # tests/gateway/test_session_hygiene.py # tests/gateway/test_status.py # tests/gateway/test_teams.py # tests/gateway/test_turn_lease.py # tests/gateway/test_whatsapp_connect.py # tests/hermes_cli/test_approvals_command.py # tests/hermes_cli/test_auth_store_lock_concurrent.py # tests/hermes_cli/test_backup.py # tests/hermes_cli/test_banner_git_state.py # tests/hermes_cli/test_certifi_repair.py # tests/hermes_cli/test_cmd_update.py # tests/hermes_cli/test_compat_manifest_targets.py # tests/hermes_cli/test_computer_use_cli.py # tests/hermes_cli/test_cpr_local_leak.py # tests/hermes_cli/test_dashboard_auth_gate.py # tests/hermes_cli/test_dashboard_procs_kill_grace.py # tests/hermes_cli/test_desktop_lifecycle_windows_live.py # tests/hermes_cli/test_doctor.py # tests/hermes_cli/test_doctor_command_install.py # tests/hermes_cli/test_fleet_config_migration_windows_live.py # tests/hermes_cli/test_gateway.py # tests/hermes_cli/test_gateway_platform_gating.py # tests/hermes_cli/test_gateway_restart_loop.py # tests/hermes_cli/test_gateway_task_probe.py # tests/hermes_cli/test_gateway_wsl.py # tests/hermes_cli/test_gui_command.py # tests/hermes_cli/test_install_cua_driver.py # tests/hermes_cli/test_kanban_db.py # tests/hermes_cli/test_lazy_command_exports.py # tests/hermes_cli/test_lazy_refresh_venv_repair.py # tests/hermes_cli/test_linux_desktop_entry.py # tests/hermes_cli/test_local_runtime.py # tests/hermes_cli/test_local_runtime_updates.py # tests/hermes_cli/test_managed_uv.py # tests/hermes_cli/test_mcp_reload_confirm_gate.py # tests/hermes_cli/test_nous_subscription.py # tests/hermes_cli/test_npm_engine.py # tests/hermes_cli/test_personality_none.py # tests/hermes_cli/test_pet_toggle.py # tests/hermes_cli/test_plan_reconciliation_windows_live.py # tests/hermes_cli/test_plugin_event_bus.py # tests/hermes_cli/test_plugin_manifest_v2.py # tests/hermes_cli/test_plugin_packs.py # tests/hermes_cli/test_plugins_cmd.py # tests/hermes_cli/test_plugins_cmd_enable_disable_nested.py # tests/hermes_cli/test_process_identity.py # tests/hermes_cli/test_profiles.py # tests/hermes_cli/test_profiles_sidebar_cache.py # tests/hermes_cli/test_pty_bridge.py # tests/hermes_cli/test_resolve_turn_limit.py # tests/hermes_cli/test_serve_runtime_inventory.py # tests/hermes_cli/test_session_vacuum_config.py # tests/hermes_cli/test_set_config_value.py # tests/hermes_cli/test_signal_handler_kanban_worker.py # tests/hermes_cli/test_slash_confirm_windows.py # tests/hermes_cli/test_stale_pid_guard.py # tests/hermes_cli/test_startup_fast_guards.py # tests/hermes_cli/test_status.py # tests/hermes_cli/test_telegram_managed_bot.py # tests/hermes_cli/test_tools_config.py # tests/hermes_cli/test_update_apply_shallow_count.py # tests/hermes_cli/test_update_autostash.py # tests/hermes_cli/test_update_concurrent_quarantine.py # tests/hermes_cli/test_update_fetch_failure_classifier.py # tests/hermes_cli/test_update_fleet_probe_resume_token.py # tests/hermes_cli/test_update_handoff_backend_reap.py # tests/hermes_cli/test_update_handoff_desktop_rebuild.py # tests/hermes_cli/test_update_head_moved_gate.py # tests/hermes_cli/test_update_host_obligation.py # tests/hermes_cli/test_update_import_guard.py # tests/hermes_cli/test_update_interrupted_recovery.py # tests/hermes_cli/test_update_inventory.py # tests/hermes_cli/test_update_launchd_unloaded_gateway.py # tests/hermes_cli/test_update_missing_configured_deps.py # tests/hermes_cli/test_update_modified_notice.py # tests/hermes_cli/test_update_multiplex_migration_hook.py # tests/hermes_cli/test_update_no_gateway_restart.py # tests/hermes_cli/test_update_orphan_backend_reap.py # tests/hermes_cli/test_update_parked_branch_guard.py # tests/hermes_cli/test_update_post_pull_syntax_guard.py # tests/hermes_cli/test_update_receipt.py # tests/hermes_cli/test_update_self_lock.py # tests/hermes_cli/test_update_shim_fail_closed.py # tests/hermes_cli/test_update_shim_self_lock.py # tests/hermes_cli/test_update_sqlite_remediation.py # tests/hermes_cli/test_update_stale_dashboard.py # tests/hermes_cli/test_update_stale_virtualenv.py # tests/hermes_cli/test_update_venv_health.py # tests/hermes_cli/test_update_venv_ownership_preflight.py # tests/hermes_cli/test_update_wedged_gateway.py # tests/hermes_cli/test_update_yes_flag.py # tests/hermes_cli/test_update_zip_two_phase.py # tests/hermes_cli/test_urllib_security.py # tests/hermes_cli/test_ux_messages_auth_config.py # tests/hermes_cli/test_ux_messages_startup.py # tests/hermes_cli/test_venv_holder_classifier.py # tests/hermes_cli/test_verify_console_scripts.py # tests/hermes_cli/test_verify_core_dependencies.py # tests/hermes_cli/test_web_server.py # tests/hermes_cli/test_web_server_console_ws.py # tests/hermes_cli/test_web_server_ws_ping.py # tests/hermes_cli/test_web_ui_build.py # tests/hermes_state/test_fts_rebuild_admission.py # tests/hermes_state/test_hermes_state.py # tests/plugins/memory/test_memory_lazy_install.py # tests/plugins/test_google_meet_plugin.py # tests/plugins/test_langfuse_plugin.py # tests/plugins/test_security_guidance_plugin.py # tests/plugins/test_transform_llm_output_hook.py # tests/scripts/desktop_update/test_desktop_update_windows_gateway_flag.py # tests/scripts/desktop_update/test_desktop_update_windows_python_handoff.py # tests/scripts/desktop_update/test_desktop_update_windows_timestamp.py # tests/scripts/install/test_install_clone_throttle_fallback.py # tests/scripts/install/test_install_lockfile_churn.py # tests/scripts/install/test_install_no_initial_commit.py # tests/scripts/install/test_install_sh_browser_install.py # tests/scripts/install/test_install_sh_node_prerelease.py # tests/scripts/install/test_install_sh_symlink_stomp.py # tests/scripts/install/test_install_sh_uv_lock_config.py # tests/scripts/install/test_install_unmerged_index.py # tests/scripts/test_contributor_map.py # tests/scripts/test_run_tests_parallel.py # tests/skills/test_competitor_news_monitor_skill.py # tests/skills/test_document_to_action_items_skill.py # tests/skills/test_google_workspace_setup.py # tests/skills/test_google_workspace_setup_deps.py # tests/skills/test_grounded_citations_skill.py # tests/skills/test_ip_as_logo_skill.py # tests/skills/test_live_dashboard_skill.py # tests/skills/test_mcp_oauth_remote_gateway_skill.py # tests/skills/test_office_document_skills.py # tests/skills/test_openclaw_migration.py # tests/skills/test_product_price_monitor_skill.py # tests/skills/test_scrollcraft_skill.py # tests/skills/test_setup_wizard_generator_skill.py # tests/skills/test_weekly_review_planning_skill.py # tests/test_engines_satisfiable.py # tests/test_fast_safe_load.py # tests/test_hermes_bootstrap.py # tests/test_hermes_constants.py # tests/test_hermes_logging.py # tests/test_managed_runtime_resolution.py # tests/test_model_tools_async_bridge.py # tests/test_packaging_build_guard.py # tests/test_packaging_metadata.py # tests/test_yaml_indent_consistency.py # tests/tools/test_approval_timeout_overflow.py # tests/tools/test_base_environment.py # tests/tools/test_bot_mode_dm.py # tests/tools/test_browser_chromium_check.py # tests/tools/test_browser_hardening.py # tests/tools/test_browser_homebrew_paths.py # tests/tools/test_browser_npx_warmup.py # tests/tools/test_browser_orphan_reaper.py # tests/tools/test_browser_real_profile.py # tests/tools/test_browser_use_cli.py # tests/tools/test_clipboard.py # tests/tools/test_code_execution.py # tests/tools/test_code_execution_modes.py # tests/tools/test_code_execution_windows_env.py # tests/tools/test_computer_use.py # tests/tools/test_delegate_liveness_timeout.py # tests/tools/test_execute_code_approval_cluster.py # tests/tools/test_execution_flag_detection.py # tests/tools/test_fal_common.py # tests/tools/test_file_operations.py # tests/tools/test_file_tools.py # tests/tools/test_file_tools_cwd_resolution.py # tests/tools/test_file_tools_live.py # tests/tools/test_lazy_deps.py # tests/tools/test_lazy_deps_durable_target.py # tests/tools/test_lazy_deps_managed.py # tests/tools/test_local_env_blocklist.py # tests/tools/test_local_tempdir.py # tests/tools/test_macos_protected_search.py # tests/tools/test_mcp_npx_cached_bin.py # tests/tools/test_oneshot_completion_linger.py # tests/tools/test_process_registry.py # tests/tools/test_read_file_schema_gating.py # tests/tools/test_skill_improvements.py # tests/tools/test_skills_sync.py # tests/tools/test_termux_api_detection.py # tests/tools/test_tirith_security.py # tests/tools/test_transcription_tools.py # tests/tools/test_tts_streaming.py # tests/tools/test_wake_word.py # tests/tui_gateway/test_compute_host_borrowed_lease.py # tests/tui_gateway/test_compute_host_turn_protocol.py # tests/tui_gateway/test_isolated_orphan_activity.py # tests/tui_gateway/test_protocol.py # tests/tui_gateway/test_slash_worker_profile_home.py # tests/tui_gateway/test_subprocess_encoding.py # tests/tui_gateway/test_tui_gateway_server.py # ui-tui/src/__tests__/terminalParity.test.ts # ui-tui/src/__tests__/termuxComposerLayout.test.ts # ui-tui/src/__tests__/textInputFastEcho.test.ts
258 lines
12 KiB
Python
258 lines
12 KiB
Python
"""Tests for the infinite compaction loop fix (issue #40803).
|
|
|
|
When summary_target_ratio is large enough that the entire transcript fits
|
|
within soft_ceiling, the backward walk in _find_tail_cut_by_tokens never
|
|
breaks early. Without the fix this produces either a no-op compression
|
|
(compress_start >= compress_end) or a single-message compression whose
|
|
summary-of-one overhead saves 0 tokens — both of which cause the
|
|
compressor to fire on every subsequent turn with no progress.
|
|
|
|
The fix adds two safeguards:
|
|
1. _find_tail_cut_by_tokens: when the whole transcript fits in soft_ceiling,
|
|
re-walk with the raw (non-inflated) budget to find a meaningful cut.
|
|
2. compress(): when compress_start >= compress_end, defer retries via the
|
|
structural no-op backoff (#93022) so should_compress() anti-thrashing
|
|
fires without burning anti-thrash strikes on transcript-shape facts.
|
|
"""
|
|
|
|
from unittest.mock import patch
|
|
|
|
import time
|
|
|
|
from agent.context_compressor import ContextCompressor
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Helpers
|
|
# ---------------------------------------------------------------------------
|
|
|
|
def _make_compressor(**kwargs) -> ContextCompressor:
|
|
defaults = dict(
|
|
model="test-model",
|
|
threshold_percent=0.65,
|
|
protect_first_n=2,
|
|
protect_last_n=3,
|
|
quiet_mode=True,
|
|
)
|
|
defaults.update(kwargs)
|
|
# NOTE: 96K < 512K, so the small-context floor raises the effective
|
|
# threshold_percent to 0.75 → threshold_tokens = 72_000. Tests use
|
|
# 73_000 as the "over threshold" probe value.
|
|
with patch("agent.context_compressor.get_model_context_length", return_value=96000):
|
|
return ContextCompressor(**defaults)
|
|
|
|
def _build_session(n_turns: int, words_per_turn: int = 20) -> list:
|
|
"""Build a multi-turn conversation with a system prompt."""
|
|
base_text = " ".join(["a"] * words_per_turn)
|
|
messages = [{"role": "system", "content": "You are a helpful agent."}]
|
|
for i in range(n_turns):
|
|
messages.append({"role": "user", "content": f"{base_text} (user turn {i})"})
|
|
messages.append({"role": "assistant", "content": f"{base_text} (assistant turn {i})"})
|
|
return messages
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Test: compress_start >= compress_end registers as ineffective
|
|
# ---------------------------------------------------------------------------
|
|
|
|
class TestCompressNoOpRegistersIneffective:
|
|
"""When compress_start >= compress_end, the fix defers further attempts
|
|
via the structural no-op backoff (#93022) so the anti-thrashing guard
|
|
fires.
|
|
|
|
We trigger this path by having _find_tail_cut_by_tokens return
|
|
head_end (which makes compress_end = head_end + 1, same as
|
|
compress_start after alignment)."""
|
|
|
|
def test_no_op_arms_structural_backoff(self):
|
|
"""compress_start >= compress_end -> backoff armed, strikes untouched."""
|
|
comp = _make_compressor(
|
|
summary_target_ratio=0.45,
|
|
config_context_length=96000,
|
|
)
|
|
# A large session that passes the min_for_compress check
|
|
messages = _build_session(10, words_per_turn=10)
|
|
comp.last_prompt_tokens = 73_000
|
|
|
|
# Mock _find_tail_cut_by_tokens to return head_end,
|
|
# causing compress_start >= compress_end
|
|
original = comp._find_tail_cut_by_tokens
|
|
comp._find_tail_cut_by_tokens = lambda msgs, he: he # force no-op
|
|
|
|
result = comp.compress(messages, current_tokens=73_000)
|
|
|
|
assert len(result) == len(messages), (
|
|
"no-op compression must return the transcript unchanged"
|
|
)
|
|
assert comp._ineffective_compression_count == 0, (
|
|
"a structural impossibility is not an ineffective strike (#93022)"
|
|
)
|
|
assert comp._structural_no_op_backoff_until > time.monotonic(), (
|
|
"structural no-op must arm the retry backoff"
|
|
)
|
|
|
|
def test_two_no_ops_block_should_compress(self):
|
|
"""After 2 no-op compressions, should_compress returns False."""
|
|
comp = _make_compressor(
|
|
summary_target_ratio=0.45,
|
|
config_context_length=96000,
|
|
)
|
|
messages = _build_session(10, words_per_turn=10)
|
|
comp.last_prompt_tokens = 73_000
|
|
comp._find_tail_cut_by_tokens = lambda msgs, he: he # force no-op
|
|
|
|
comp.compress(messages, current_tokens=73_000)
|
|
comp.compress(messages, current_tokens=73_000)
|
|
|
|
assert comp._ineffective_compression_count == 0, (
|
|
"structural no-ops defer via backoff instead of striking (#93022)"
|
|
)
|
|
assert not comp.should_compress(73_000), (
|
|
"should_compress should return False while the structural backoff holds"
|
|
)
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Test: _find_tail_cut_by_tokens raw-budget fallback
|
|
# ---------------------------------------------------------------------------
|
|
|
|
class TestTailCutRawBudgetFallback:
|
|
"""When the entire transcript fits within soft_ceiling, the fix
|
|
re-walks with the raw budget to find a meaningful cut point."""
|
|
|
|
def test_meaningful_cut_with_large_ratio(self):
|
|
"""With summary_target_ratio=0.45, _find_tail_cut_by_tokens still
|
|
leaves a meaningful compressable region."""
|
|
comp = _make_compressor(
|
|
summary_target_ratio=0.45,
|
|
config_context_length=96000,
|
|
)
|
|
messages = _build_session(20, words_per_turn=20)
|
|
head_end = comp._protect_head_size(messages)
|
|
head_end = comp._align_boundary_forward(messages, head_end)
|
|
|
|
cut = comp._find_tail_cut_by_tokens(messages, head_end)
|
|
|
|
n = len(messages)
|
|
middle_size = cut - head_end
|
|
assert middle_size >= 3, (
|
|
f"Expected at least 3 messages in compressable region, got {middle_size} "
|
|
f"(cut={cut}, head_end={head_end}, n={n})"
|
|
)
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Test: Effective compression resets counter
|
|
# ---------------------------------------------------------------------------
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Test: anti-thrashing in should_compress
|
|
# ---------------------------------------------------------------------------
|
|
|
|
class TestAntiThrashing:
|
|
"""Directly test the should_compress anti-thrashing guard."""
|
|
|
|
def test_ineffective_count_2_blocks(self):
|
|
"""_ineffective_compression_count >= 2 -> should_compress returns False."""
|
|
comp = _make_compressor(config_context_length=96000)
|
|
comp.last_prompt_tokens = 73_000
|
|
comp._ineffective_compression_count = 2
|
|
assert not comp.should_compress(73_000)
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Test: summary-LLM cooldown guard in should_compress (#11529)
|
|
# ---------------------------------------------------------------------------
|
|
|
|
class TestCooldownGuard:
|
|
"""should_compress() must skip compression while the summary LLM is in
|
|
cooldown, otherwise a 429/transient failure re-fires _compress_context()
|
|
every turn (inserting a fallback marker repeatedly) and freezes the CLI.
|
|
"""
|
|
|
|
def test_active_cooldown_blocks(self):
|
|
"""A future cooldown deadline -> should_compress returns False even
|
|
when tokens are over threshold."""
|
|
comp = _make_compressor(config_context_length=96000)
|
|
comp.last_prompt_tokens = 73_000
|
|
comp._summary_failure_cooldown_until = time.monotonic() + 60
|
|
assert not comp.should_compress(73_000)
|
|
|
|
# ---------------------------------------------------------------------------
|
|
# Test: #48621 — gpt-5.3-codex-spark short-session boundary
|
|
#
|
|
# Issue #48621 Bug 2 claims that a short high-token session (15-20 messages,
|
|
# ~90k tokens on a 128k model with protect_last_n=20) hits
|
|
# compress_start >= compress_end, causing a silent context wipe. The
|
|
# raw-budget fallback added in the #40803 fix already mitigates this: the
|
|
# boundary logic always exposes a minimal compressible window. This test
|
|
# locks that behavior in for the exact #48621 parameters.
|
|
# ---------------------------------------------------------------------------
|
|
|
|
class TestCodexSparkShortSessionBoundary:
|
|
"""Verify that gpt-5.3-codex-spark's short-session scenario always yields
|
|
a non-empty compressible window (no silent wipe)."""
|
|
|
|
def test_short_high_token_session_has_compressible_window(self):
|
|
"""16 messages with large tool outputs on a 128k model must leave
|
|
a compressible middle (compress_start < compress_end)."""
|
|
comp = _make_compressor(
|
|
model="gpt-5.3-codex-spark",
|
|
threshold_percent=0.70,
|
|
protect_first_n=3,
|
|
protect_last_n=20,
|
|
config_context_length=128000,
|
|
)
|
|
# Build system + 3 head pairs + 3 tool groups (large outputs) + tail pair
|
|
big_tool = "x" * 20000 # ~5k tokens each
|
|
messages = [{"role": "system", "content": "You are a helpful agent."}]
|
|
for i in range(3):
|
|
messages.append({"role": "user", "content": f"Question {i}"})
|
|
messages.append({"role": "assistant", "content": f"Answer {i}"})
|
|
for i in range(3):
|
|
messages.append({"role": "user", "content": f"Run command {i}"})
|
|
messages.append({
|
|
"role": "assistant", "content": "",
|
|
"tool_calls": [{
|
|
"id": f"tc{i}", "type": "function",
|
|
"function": {"name": "terminal", "arguments": "{}"},
|
|
}],
|
|
})
|
|
messages.append({"role": "tool", "tool_call_id": f"tc{i}", "content": big_tool})
|
|
messages.append({"role": "user", "content": "Final question"})
|
|
messages.append({"role": "assistant", "content": "Final answer"})
|
|
|
|
head = comp._protect_head_size(messages)
|
|
compress_start = comp._align_boundary_forward(messages, head)
|
|
compress_end = comp._find_tail_cut_by_tokens(messages, compress_start)
|
|
|
|
assert compress_start < compress_end, (
|
|
f"No compressible window: start={compress_start}, end={compress_end}. "
|
|
f"This would cause the silent context wipe described in #48621."
|
|
)
|
|
assert comp.has_content_to_compress(messages) is True
|
|
|
|
class TestPressureRealFloor:
|
|
"""Regression: Cyrillic-heavy sessions under-count in the rough estimate,
|
|
letting real prompts ride the provider window (64,842→64,995 observed)
|
|
while the pre-API gate saw sub-threshold pressure."""
|
|
|
|
def _compressor(self, last_real, awaiting=False):
|
|
class _C:
|
|
last_real_prompt_tokens = last_real
|
|
awaiting_real_usage_after_compression = awaiting
|
|
return _C()
|
|
|
|
def test_real_floor_lifts_undercounted_rough(self):
|
|
from agent.conversation_loop import _pressure_with_real_floor
|
|
assert _pressure_with_real_floor(self._compressor(64_842), 45_000) == 64_842
|
|
|
|
def test_rough_wins_when_larger(self):
|
|
from agent.conversation_loop import _pressure_with_real_floor
|
|
assert _pressure_with_real_floor(self._compressor(30_000), 45_000) == 45_000
|
|
|
|
def test_stale_real_ignored_right_after_compaction(self):
|
|
from agent.conversation_loop import _pressure_with_real_floor
|
|
compressor = self._compressor(64_842, awaiting=True)
|
|
assert _pressure_with_real_floor(compressor, 20_000) == 20_000
|
|
|
|
def test_zero_and_missing_real_are_safe(self):
|
|
from agent.conversation_loop import _pressure_with_real_floor
|
|
assert _pressure_with_real_floor(self._compressor(0), 10_000) == 10_000
|
|
assert _pressure_with_real_floor(object(), 10_000) == 10_000
|