mirror of
https://ghfast.top/https://github.com/aeroxw/tick-stock-panel.git
synced 2026-09-12 20:14:16 +08:00
* fix(concurrency): 共享缓存/任务表加锁, 全局限速, depth 原子写, 认证热路径缓存 修复多线程下的竞态与阻塞: - overview/strategy_cache/PanelCache/StrategyMonitor._watching 四处共享状态加锁, 消除 "dict/OrderedDict mutated" 与丢更新/半写读取 - strategy_cache/depth parquet 改临时文件 + os.replace 原子写 - rate_limits 改进程级共享时间轴限速, 并发同步不再聚合超过单能力 rpm; scheduler 令牌账目与 sleep 分离, sleep 不再独占锁串行化其他请求 - auth.is_configured() 内存缓存, 认证中间件不再每请求读盘阻塞事件循环 - api/backtest 任务清理/取消全程持 _jobs_lock, 并用 Semaphore(2) 限并发重回测 Co-Authored-By: Claude Opus 4.8 (1M context) <noreply@anthropic.com> * perf(data): limit_ladder 去 N+1 全市场重算, 指标裁剪, factor 向量化 - limit_ladder 前一日 consecutive 改窄读单日 parquet 存储列 (谓词/投影下推), 替代 range(1,10) 逐日 _load_enriched_for_date 全市场指标重算 (最坏 9x) - compute_indicators 新增可选 needed 裁剪 (默认 None 行为逐位不变, 已对照验证), factor 只算所需因子列 - factor._calc_period_return 用 Polars join 替代 Python 逐行 price_map 循环, _add_groups 去 map_elements 改纯表达式 (输出逐位一致) - screener ext value_map 按 parquet mtime 记忆化, 免每请求磁盘重读 (DuckDB 过滤仍用隔离 :memory: 连接, 不扩大注入面) Co-Authored-By: Claude Opus 4.8 (1M context) <noreply@anthropic.com> * refactor(backend): 报表存储去重, 删死代码, DuckDB 视图重建收敛, 管道失败如实标记 - 三份近乎逐字复制的 *_reports.py 收敛到共享 JsonReportStore (原子写 + 锁), 各模块公有 API/id 格式/上限/落盘 schema 完全保持不变 - 删除 ext_pull.py 中字节相同的死 _run_loop (Python 只绑第二个) 及无用 import - 13 张 DuckDB 视图重建收敛为唯一权威 repository.rebuild_views(), daily_pipeline 与 /api/data/clear 改为调用 (修好 clear 路径漏挂视图的漂移) - daily_pipeline 累积 stage_errors 并在末尾抛出, 部分失败不再误报成功; free/None 模式的能力门控跳过不计入失败 Co-Authored-By: Claude Opus 4.8 (1M context) <noreply@anthropic.com> * feat(frontend): SSE 连接态, 路由代码分割, 查询失效修复, 三态与无障碍 - 实时行情 SSE: 连接态 store + 指数退避 + 断线徽标/toast (避免静默丢告警); 回测 SSE 断线有界重连 + 可重试, 不再永久卡住进度条 - router 全部 React.lazy + Suspense, vite manualChunks 拆图表库 (echarts 变独立 1MB 懒加载 chunk, 首屏包显著减小) - 修 Data 清库后其它页显示旧数据 (改回广域失效); 修 Watchlist kline 失效键 永不匹配; query key 收敛到 QK 工厂 (新增 strategyDetail) - Monitor/Analysis/StockAnalysis/ExtPages/CustomSignals 补 loading 门控与 error/empty 三态区分 - 新增共享 Modal 原语 (焦点陷阱/ESC/焦点还原/aria), 改造 3 个高频弹窗; Toast/AlertToast 加 aria-live 与键盘可达; Watchlist/LimitUpLadder 卡片 memo 修复本轮 review 发现的缺陷: - Modal 焦点 effect 依赖 onClose 致每次输入抢焦点 → 改 ref 只装一次 - StrategySettingsDialog 删除确认框被 Modal 面板裁剪 → 移出作兄弟节点 Co-Authored-By: Claude Opus 4.8 (1M context) <noreply@anthropic.com> * fix(quant): 修正 ST 板块限价套错 与 因子 Sharpe 年化频率 两个不报错但会算错数的领域 bug: 1. ST 5% 涨跌停限幅被无条件套到创业板/科创板 ST 股: 注册制改革后 创业板(300/301)、科创板(688/689) 的风险警示股仍执行 20%, 北交所 30%, 只有主板 ST 才是 5%。原代码 _is_st 先判且覆盖板块限幅, 导致 创业板/科创板 ST 的涨停价按 5% 计算 → +5% 被误报涨停、真 +20% 涨停被漏报, 污染 signal_limit_up / consecutive_limit_ups / 连板梯队 / near_limit_up。 修正: ST 5% 仅在 ~(创业板|科创板|北交所) 时生效 (EOD + 盘中两条路径 + near_limit_up)。 2. 因子回测 Sharpe 一律乘 √252, 但 group_nav 每点是一个调仓周期收益: 月频调仓下是月收益, 乘 √252 会把 Sharpe 高估 √(252/12) ≈ 4.6x (周频 ≈2.2x), 使无效因子显示成明星因子, 废掉"先筛无效指标"的用途。 修正: 年化系数按 config.rebalance 取 √252/√52/√12。 新增 tests/test_st_limit_and_sharpe.py (5 例) 覆盖两处修正。 Co-Authored-By: Claude Opus 4.8 (1M context) <noreply@anthropic.com> --------- Co-authored-by: Claude Opus 4.8 (1M context) <noreply@anthropic.com>
178 lines
6.6 KiB
Python
178 lines
6.6 KiB
Python
"""策略结果缓存 — 写入本地文件,供策略页面秒加载。
|
|
|
|
缓存结构:
|
|
{
|
|
"as_of": "2024-01-15",
|
|
"results": { strategy_id: { total, as_of, rows } },
|
|
"today_ever_matched": { strategy_id: [symbol, ...] }, // 今日曾命中 symbol 并集
|
|
"today_ever_rows": { strategy_id: { symbol: row_data } },// 今日曾命中的完整行数据
|
|
"updated_at": 1705324800000 # Unix ms
|
|
}
|
|
|
|
文件路径: data/user_data/strategy_cache.json
|
|
"""
|
|
from __future__ import annotations
|
|
|
|
import json
|
|
import logging
|
|
import os
|
|
import threading
|
|
import time
|
|
from datetime import date, datetime
|
|
from pathlib import Path
|
|
from typing import Any
|
|
|
|
|
|
def _json_default(obj: Any) -> Any:
|
|
"""处理 date/datetime 等 JSON 不认识的类型。"""
|
|
if isinstance(obj, date):
|
|
return obj.isoformat()
|
|
if isinstance(obj, datetime):
|
|
return obj.isoformat()
|
|
raise TypeError(f"Object of type {type(obj).__name__} is not JSON serializable")
|
|
|
|
|
|
logger = logging.getLogger(__name__)
|
|
|
|
_CACHE_FILENAME = "strategy_cache.json"
|
|
|
|
# 读写同一 JSON 文件的进程内锁: write_cache 的 read-modify-write 与并发 read_cache
|
|
# 无锁会丢更新/读到半写文件。read_cache 与 write_cache 共用此锁; write 内部复用
|
|
# _read_cache_unlocked 避免自死锁。写入用临时文件 + os.replace 做到原子替换。
|
|
_file_lock = threading.Lock()
|
|
|
|
|
|
def _cache_path(data_dir: Path) -> Path:
|
|
return data_dir / "user_data" / _CACHE_FILENAME
|
|
|
|
|
|
def _enriched_parquet_path(data_dir: Path, as_of: str) -> Path:
|
|
"""返回 enriched parquet 文件路径。"""
|
|
return data_dir / "kline_daily_enriched" / f"date={as_of}" / "part.parquet"
|
|
|
|
|
|
def _get_enriched_mtime(data_dir: Path, as_of: str) -> float | None:
|
|
"""返回 enriched parquet 文件的 mtime (秒)。文件不存在返回 None。"""
|
|
p = _enriched_parquet_path(data_dir, as_of)
|
|
try:
|
|
return p.stat().st_mtime
|
|
except FileNotFoundError:
|
|
return None
|
|
|
|
|
|
def read_cache(data_dir: Path) -> dict | None:
|
|
"""读取策略缓存文件。返回 None 表示无缓存或读取失败。
|
|
|
|
说明: 原先有 enriched mtime 过期校验 (数据文件变化 → 判过期返回 None),
|
|
但在有实时行情的系统里, enriched parquet 每轮被刷新 → mtime 必然变化 →
|
|
缓存被永久判死, 策略页读不到数据。且判过期后不触发重算, 只能让用户手动重跑,
|
|
保护价值有限。故移除: 盘后缓存总能读出, 实时新鲜度由 /api/screener/cached
|
|
端点叠加监控引擎的内存实时结果 (latest_strategy_results) 来保证。
|
|
"""
|
|
with _file_lock:
|
|
return _read_cache_unlocked(data_dir)
|
|
|
|
|
|
def _read_cache_unlocked(data_dir: Path) -> dict | None:
|
|
"""实际读取逻辑 (不持锁)。供 read_cache 与 write_cache 复用, 避免重入死锁。"""
|
|
path = _cache_path(data_dir)
|
|
if not path.exists():
|
|
return None
|
|
try:
|
|
text = path.read_text(encoding="utf-8")
|
|
if not text.strip():
|
|
return None
|
|
cached = json.loads(text)
|
|
except Exception as e: # noqa: BLE001
|
|
logger.warning("读取策略缓存失败: %s", e)
|
|
return None
|
|
|
|
return cached
|
|
|
|
|
|
def _rows_to_symbol_map(rows: list[dict]) -> dict[str, dict]:
|
|
"""将 rows 列表转为 {symbol: row_data} 映射。"""
|
|
result: dict[str, dict] = {}
|
|
for row in rows:
|
|
sym = row.get("symbol")
|
|
if sym:
|
|
result[sym] = row
|
|
return result
|
|
|
|
|
|
def write_cache(
|
|
data_dir: Path,
|
|
as_of: str,
|
|
results: dict[str, Any],
|
|
) -> None:
|
|
"""将策略结果写入缓存文件,同时更新今日曾命中集合。
|
|
|
|
- 日期变更时重置 today_ever_matched 和 today_ever_rows
|
|
- 同一天内合并 (并集) 之前曾命中的 symbol,并用最新行数据更新
|
|
"""
|
|
path = _cache_path(data_dir)
|
|
path.parent.mkdir(parents=True, exist_ok=True)
|
|
|
|
# 整个 read-modify-write 持锁: 避免并发 write 丢更新, 也避免与 read_cache 撕裂
|
|
with _file_lock:
|
|
_write_cache_locked(path, data_dir, as_of, results)
|
|
|
|
|
|
def _write_cache_locked(
|
|
path: Path,
|
|
data_dir: Path,
|
|
as_of: str,
|
|
results: dict[str, Any],
|
|
) -> None:
|
|
"""持 _file_lock 后的实际写入逻辑 (read-merge-write + 原子替换)。"""
|
|
# 读取旧缓存 (已持锁, 走不重入的 _read_cache_unlocked)
|
|
old = _read_cache_unlocked(data_dir)
|
|
old_as_of = old.get("as_of") if old else None
|
|
old_ever_rows: dict[str, dict[str, dict]] = old.get("today_ever_rows", {}) if old else {}
|
|
|
|
# 当前命中的行数据 → symbol 映射
|
|
current_row_maps: dict[str, dict[str, dict]] = {}
|
|
for sid, r in results.items():
|
|
current_row_maps[sid] = _rows_to_symbol_map(r.get("rows", []))
|
|
|
|
if old_as_of and old_as_of == as_of and old_ever_rows:
|
|
# 同一天: 合并 — 用当前行数据更新旧数据 (保持最新价格等)
|
|
merged_rows: dict[str, dict[str, dict]] = {}
|
|
all_keys = set(old_ever_rows.keys()) | set(current_row_maps.keys())
|
|
for sid in all_keys:
|
|
old_map = old_ever_rows.get(sid, {})
|
|
cur_map = current_row_maps.get(sid, {})
|
|
# 以旧数据为基础,用当前数据覆盖 (当前数据更新鲜)
|
|
combined = {**old_map, **cur_map}
|
|
merged_rows[sid] = combined
|
|
today_ever_rows = merged_rows
|
|
else:
|
|
# 新的一天或首次写入
|
|
today_ever_rows = current_row_maps
|
|
|
|
# 从 ever_rows 提取 symbol 列表 (用于快速计数)
|
|
today_ever_matched = {sid: sorted(maps.keys()) for sid, maps in today_ever_rows.items()}
|
|
|
|
# enriched_mtime: 盘后缓存写入时记录 (向后兼容旧字段)。read_cache 已不再用它
|
|
# 做过期校验, 实时新鲜度改由 /cached 端点叠加监控引擎内存结果保证。
|
|
enriched_mtime = _get_enriched_mtime(data_dir, as_of)
|
|
|
|
payload = {
|
|
"as_of": as_of,
|
|
"results": results,
|
|
"today_ever_matched": today_ever_matched,
|
|
"today_ever_rows": today_ever_rows,
|
|
"enriched_mtime": enriched_mtime,
|
|
"updated_at": int(time.time() * 1000),
|
|
}
|
|
try:
|
|
# 原子写: 先写临时文件再 os.replace, 避免读侧读到半写的 JSON
|
|
tmp = path.with_name(path.name + ".tmp")
|
|
tmp.write_text(json.dumps(payload, ensure_ascii=False, default=_json_default), encoding="utf-8")
|
|
os.replace(tmp, path)
|
|
total_rows = sum(len(r.get("rows", [])) for r in results.values())
|
|
total_ever = sum(len(v) for v in today_ever_matched.values())
|
|
logger.info("策略缓存已写入: %s, %d 策略, %d 命中, %d 曾命中", as_of, len(results), total_rows, total_ever)
|
|
except Exception as e: # noqa: BLE001
|
|
logger.warning("写入策略缓存失败: %s", e)
|