리뷰 두 건을 반영했다. 계약 - tech-log-tree.json 하나가 분해 계약이자 색인이다. 사람이 읽는 트리·Node Specification· 후보 대장은 없어졌고, 문서에 남아 있던 그 개념을 걷어냈다 - candidateScope — 후보를 찾는 SSOT 범위. 접어 넣은 제2부·제3부는 근거이지 후보가 아니다 - sourceRepository — 분석한 저장소의 경로·리비전·판단 근거. 리비전을 모르면 null 로 두고 지어내지 않는다. 갈래가 여럿이면 revisions - 검사기: 계약 미채택·PENDING·PROMOTE↔글감 양방향·candidateScope·sourceRepository 를 error/warn 으로 센다. 옛 스키마도 검사를 피하지 못한다. 테스트 22 → 31 기록 쓰기 - 템플릿 5종에 source·sourceRevision·topicName, Question 에 닫는 조건, 본문 없는 종류에서 assets 제거. 고정 절 개수 삭제 - check_evidence.mjs — 인용한 코드가 SSOT 에 있는지, 앵커가 SSOT 를 가리키는지, 제목이 계약과 같은지, 리비전이 저장소에 있는지. 게시된 기록에서 SSOT 와 다른 URL 을 잡았다 문체 - 문체 규칙의 정본을 ai-tells.md 로. explaining.md 의 질문체 제목·절 끝 대조 반복·그림 예고 규칙을 삭제해 충돌을 없앴다. 첫 절 「설명 뒤에 평가를 붙이지 않는다」에 지우는 사례 네 유형 - voice 스킬의 「독자 쪽을 본다」를 자료에 오독 기록이 있을 때로 좁히고, 평가만 더한 예시를 교체 - check_prose: 안내 문장을 요구하던 경고 제거, 문장이 끝나지 않은 채 문단이 끝나는 조각 검사 추가 Co-Authored-By: Claude Fable 5.1 <noreply@anthropic.com>
293 lines
12 KiB
Python
293 lines
12 KiB
Python
#!/usr/bin/env python3
|
|
"""분해 작업 재료를 `tech-log-tree.json` 으로 옮기고 폴더에서 지운다.
|
|
|
|
`analysis/` 를 `final/document.md` 로 옮긴 것과 같은 일을 `tech-log-studio/` 에서 한다.
|
|
끝난 프로젝트의 `tech-log-studio/` 에는 `tech-log-tree.json` 과 기록 폴더만 있다.
|
|
|
|
옮기는 것
|
|
|
|
root-tree.md 사람이 읽는 트리 + Node Specifications → topics
|
|
candidate-ledger.json 후보와 처분 → candidates
|
|
root-tree-source-manifest.json 원본 해시 → ssotSha256
|
|
_meta/state.json 생성·편집·검증 이력 → history
|
|
_meta/** 편집 과정 기록 → 지운다 (git 에 남는다)
|
|
|
|
python3 scripts/fold-studio-contract-into-index.py <프로젝트> [--dry-run] [--keep]
|
|
|
|
**옮기는 것이지 요약하는 것이 아니다.** 노드의 칸은 하나도 버리지 않는다.
|
|
"""
|
|
from __future__ import annotations
|
|
|
|
import argparse
|
|
import datetime
|
|
import hashlib
|
|
import json
|
|
import os
|
|
import re
|
|
import shutil
|
|
import sys
|
|
|
|
sys.path.insert(0, os.path.dirname(os.path.abspath(__file__)))
|
|
import techlog # noqa: E402
|
|
from techlog import KINDS # noqa: E402
|
|
|
|
ROOT = os.path.dirname(os.path.dirname(os.path.abspath(__file__)))
|
|
|
|
KIND_LABELS = {"CASE": "case", "CONCEPT": "concept", "REFERENCE": "reference",
|
|
"OPEN QUESTION": "question", "QUESTION": "question", "DECISION": "decision"}
|
|
FIELD_NAMES = ("slug", "readiness", "disposition", "source", "code", "evidence",
|
|
"classification", "missing-verification", "relations", "scope", "exceptions",
|
|
"known", "unknown", "next-verification", "decision-criterion",
|
|
"decision-status", "decision-evidence", "grounds", "basis-version")
|
|
MULTI = {"source", "code", "evidence", "relations", "grounds", "decision-evidence",
|
|
"known", "unknown", "scope", "exceptions"}
|
|
|
|
_TOPIC_HEAD = re.compile(r"^##\s+TOPIC(?:\s+\d+)?\s+[—-]\s+(.+?)\s*$")
|
|
_SPEC_HEAD = re.compile(r"^###\s+(CASE|CONCEPT|REFERENCE|OPEN QUESTION|QUESTION|DECISION)\s+[—-]\s+(.+?)\s*$")
|
|
_BRANCH = re.compile(r"^[├└]──\s+(CASE|CONCEPT|REFERENCE|OPEN QUESTION|QUESTION|DECISION)\s*$")
|
|
_ITEM = re.compile(r"^(?:│|\s)\s{2,}[├└]──\s+(.+?)\s*$")
|
|
_FIELD = re.compile(r"^-\s+([a-z][a-z-]*):\s*(.*)$")
|
|
_CONT = re.compile(r"^\s{2,}-\s+(.+?)\s*$")
|
|
_INLINE = re.compile(r"\s+·\s+(" + "|".join(FIELD_NAMES) + r"):\s*")
|
|
_READER_Q = re.compile(r"^독자\s*질문\s*[—:-]\s*(.+?)\s*$")
|
|
EMPTY = ("(추가 없음)", "(없음)", "(none)", "-")
|
|
|
|
|
|
def _set(fields: dict, key: str, value: str) -> None:
|
|
value = value.strip()
|
|
if not value:
|
|
fields.setdefault(key, [])
|
|
return
|
|
items = [v.strip() for v in value.split(" · ")] if key in MULTI else [value]
|
|
fields[key] = [v for v in items if v]
|
|
|
|
|
|
def parse_root_tree(path: str) -> dict:
|
|
"""사람이 읽는 트리와 Node Specifications 를 하나로 읽는다."""
|
|
text = open(path, encoding="utf-8").read()
|
|
header, body = {}, text
|
|
if text.startswith("---"):
|
|
end = text.find("\n---", 3)
|
|
for line in text[3:end].splitlines():
|
|
m = re.match(r"^([A-Za-z][A-Za-z0-9_]*):\s*(.*)$", line)
|
|
if m:
|
|
header[m.group(1)] = m.group(2).strip().strip('"')
|
|
body = text[end + 4:]
|
|
|
|
topics, specs, prose = [], [], []
|
|
in_specs = False
|
|
topic = branch = spec = None
|
|
spec_topic = ""
|
|
pending = None
|
|
expect = 0
|
|
|
|
for raw in body.splitlines():
|
|
line = raw.rstrip()
|
|
stripped = line.strip()
|
|
if line.startswith("# Node Specifications"):
|
|
in_specs = True
|
|
topic = branch = None
|
|
continue
|
|
if not in_specs:
|
|
if stripped.startswith(">"):
|
|
prose.append(stripped.lstrip("> ").rstrip())
|
|
continue
|
|
if stripped == "---":
|
|
continue
|
|
if stripped == "PROJECT":
|
|
expect = -1
|
|
continue
|
|
if expect == -1:
|
|
if stripped:
|
|
expect = 0
|
|
continue
|
|
if stripped == "TOPIC":
|
|
topic = {"topic": "", "title": "", "readerQuestion": "", "nodes": []}
|
|
topics.append(topic)
|
|
branch, expect = None, 1
|
|
continue
|
|
if topic is not None and expect in (1, 2, 3):
|
|
if expect == 1 and stripped:
|
|
topic["title"] = stripped
|
|
expect = 2
|
|
continue
|
|
if expect == 2 and stripped:
|
|
topic["topic"] = stripped
|
|
expect = 3
|
|
continue
|
|
if expect == 3:
|
|
m = _READER_Q.match(stripped)
|
|
if m:
|
|
topic["readerQuestion"] = m.group(1)
|
|
expect = 0
|
|
continue
|
|
if stripped:
|
|
expect = 0
|
|
m = _BRANCH.match(stripped)
|
|
if m:
|
|
branch = KIND_LABELS[m.group(1)]
|
|
continue
|
|
m = _ITEM.match(line)
|
|
if m and topic is not None and branch:
|
|
title = m.group(1).strip()
|
|
if title not in EMPTY:
|
|
topic["nodes"].append({"kind": branch, "title": title})
|
|
continue
|
|
|
|
m = _TOPIC_HEAD.match(line)
|
|
if m:
|
|
spec_topic, spec, pending = m.group(1).strip(), None, None
|
|
continue
|
|
m = _SPEC_HEAD.match(line)
|
|
if m:
|
|
spec = {"topic": spec_topic, "kind": KIND_LABELS[m.group(1)],
|
|
"title": m.group(2).strip(), "fields": {}}
|
|
specs.append(spec)
|
|
pending = None
|
|
continue
|
|
if spec is None:
|
|
continue
|
|
m = _FIELD.match(line)
|
|
if m:
|
|
key, value = m.group(1), m.group(2).strip()
|
|
parts = _INLINE.split(value)
|
|
_set(spec["fields"], key, parts[0])
|
|
rest = parts[1:]
|
|
while rest:
|
|
_set(spec["fields"], rest[0], rest[1])
|
|
rest = rest[2:]
|
|
pending = key if not rest else None
|
|
continue
|
|
m = _CONT.match(line)
|
|
if m and pending:
|
|
spec["fields"].setdefault(pending, []).append(m.group(1).strip())
|
|
continue
|
|
if not stripped:
|
|
pending = None
|
|
|
|
return {"header": header, "topics": topics, "specs": specs,
|
|
"prose": [p for p in prose if p]}
|
|
|
|
|
|
def node_from_spec(spec: dict) -> dict:
|
|
node = {"title": spec["title"], "kind": spec["kind"]}
|
|
for key in ("slug", "readiness"):
|
|
values = spec["fields"].get(key) or []
|
|
node[key] = values[0].strip("`").strip() if values else ""
|
|
node["readiness"] = node["readiness"].upper()
|
|
for key, values in spec["fields"].items():
|
|
if key in ("slug", "readiness") or not values:
|
|
continue
|
|
node[key] = values if key in MULTI or len(values) > 1 else values[0]
|
|
return node
|
|
|
|
|
|
def main() -> int:
|
|
ap = argparse.ArgumentParser(description="분해 작업 재료를 tech-log-tree.json 으로 옮긴다.")
|
|
ap.add_argument("project")
|
|
ap.add_argument("--dry-run", action="store_true")
|
|
ap.add_argument("--keep", action="store_true", help="옮기기만 하고 지우지 않는다")
|
|
args = ap.parse_args()
|
|
|
|
base = os.path.join(ROOT, "docs", args.project)
|
|
studio = os.path.join(base, "tech-log-studio")
|
|
tree_path = os.path.join(studio, "root-tree.md")
|
|
index_path = os.path.join(studio, "tech-log-tree.json")
|
|
if not os.path.exists(tree_path):
|
|
print(f"{args.project}: root-tree.md 가 없다 — 이미 옮겼거나 분해 계약을 쓴 적이 없다")
|
|
return 0
|
|
|
|
parsed = parse_root_tree(tree_path)
|
|
header = parsed["header"]
|
|
by_key = {(s["topic"], s["kind"], s["title"]): s for s in parsed["specs"]}
|
|
|
|
topics = {}
|
|
used = set()
|
|
for t in parsed["topics"]:
|
|
entry = {"topic": t["topic"], "title": t["title"],
|
|
"readerQuestion": t["readerQuestion"], "kinds": {k: [] for k in KINDS}}
|
|
for n in t["nodes"]:
|
|
key = (t["topic"], n["kind"], n["title"])
|
|
spec = by_key.get(key)
|
|
entry["kinds"][n["kind"]].append(
|
|
node_from_spec(spec) if spec else {"title": n["title"], "kind": n["kind"],
|
|
"slug": "", "readiness": ""})
|
|
used.add(key)
|
|
topics[t["topic"]] = entry
|
|
# 사람이 읽는 트리에 줄이 없던 Node Specification 도 잃지 않는다
|
|
orphans = 0
|
|
for key, spec in by_key.items():
|
|
if key in used:
|
|
continue
|
|
orphans += 1
|
|
entry = topics.setdefault(spec["topic"], {
|
|
"topic": spec["topic"], "title": "", "readerQuestion": "",
|
|
"kinds": {k: [] for k in KINDS}})
|
|
node = node_from_spec(spec)
|
|
node["listedInTree"] = False
|
|
entry["kinds"][spec["kind"]].append(node)
|
|
|
|
ledger_path = os.path.join(studio, "candidate-ledger.json")
|
|
ledger = json.load(open(ledger_path, encoding="utf-8")) if os.path.exists(ledger_path) else {}
|
|
meta_state_path = os.path.join(studio, "_meta", "state.json")
|
|
meta_state = json.load(open(meta_state_path, encoding="utf-8")) \
|
|
if os.path.exists(meta_state_path) else {}
|
|
|
|
index = json.load(open(index_path, encoding="utf-8")) if os.path.exists(index_path) else {}
|
|
ssot = header.get("sourceDocument", "final/document.md")
|
|
out = {
|
|
"schemaVersion": 4,
|
|
"project": args.project,
|
|
"ssot": ssot,
|
|
"ssotSha256": techlog.sha256_of(os.path.join(base, ssot)),
|
|
"sourceRevision": header.get("sourceRevision", ""),
|
|
"generatedAt": datetime.date.today().isoformat(),
|
|
"note": ("이 프로젝트의 글감 전부다. 분해 계약이자 색인이고, 이 파일이 정본이다. "
|
|
"노드의 칸(readiness·source·classification·relations…)은 사람이 적고, "
|
|
"file·publication·status 는 기록 파일에서 읽어 채운다 — "
|
|
"python3 scripts/build-tech-log-tree.py <프로젝트>"),
|
|
"contract": {
|
|
"decomposition": parsed["prose"],
|
|
"readinessValues": techlog.READINESS,
|
|
"dispositionValues": ledger.get("dispositionValues", {}),
|
|
},
|
|
"counts": {},
|
|
"topics": topics,
|
|
"candidates": ledger.get("candidates", []),
|
|
"history": {k: v for k, v in meta_state.items()
|
|
if k not in ("schemaVersion", "project", "rootTreePath")},
|
|
}
|
|
for key in ("counts", "unmapped", "explicitAnalysisCandidates", "cycle2", "cycle3",
|
|
"conceptRecall", "migration", "fold"):
|
|
if key in ledger:
|
|
out["history"].setdefault("ledger", {})[key] = ledger[key]
|
|
|
|
total = sum(len(v) for t in topics.values() for v in t["kinds"].values())
|
|
out["counts"] = {"topics": len(topics), "nodes": total,
|
|
"candidates": len(out["candidates"])}
|
|
|
|
print(f"{args.project}: 주제 {len(topics)} · 글감 {total} · 후보 {len(out['candidates'])}")
|
|
if orphans:
|
|
print(f" 사람이 읽는 트리에 줄이 없던 노드 {orphans}건은 listedInTree=false 로 옮겼다")
|
|
if args.dry_run:
|
|
print(" (--dry-run: 쓰지 않았다)")
|
|
return 0
|
|
|
|
with open(index_path, "w", encoding="utf-8") as fh:
|
|
json.dump(out, fh, ensure_ascii=False, indent=2)
|
|
fh.write("\n")
|
|
print(f" tech-log-tree.json {os.path.getsize(index_path):,} bytes")
|
|
|
|
if not args.keep:
|
|
for name in ("root-tree.md", "candidate-ledger.json", "root-tree-source-manifest.json"):
|
|
path = os.path.join(studio, name)
|
|
if os.path.exists(path):
|
|
os.remove(path)
|
|
shutil.rmtree(os.path.join(studio, "_meta"), ignore_errors=True)
|
|
print(" root-tree.md · candidate-ledger.json · manifest · _meta/ 를 지웠다")
|
|
return 0
|
|
|
|
|
|
if __name__ == "__main__":
|
|
raise SystemExit(main())
|