-
Notifications
You must be signed in to change notification settings - Fork 4
Expand file tree
/
Copy pathopenclaw.plugin.json
More file actions
93 lines (93 loc) · 2.57 KB
/
Copy pathopenclaw.plugin.json
File metadata and controls
93 lines (93 loc) · 2.57 KB
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
76
77
78
79
80
81
82
83
84
85
86
87
88
89
90
91
92
93
{
"id": "tokenranger",
"name": "TokenRanger",
"description": "Compresses session context via a local SLM (Ollama) before sending to cloud LLMs, reducing token costs by 50-80%.",
"configSchema": {
"type": "object",
"additionalProperties": false,
"properties": {
"serviceUrl": {
"type": "string",
"default": "http://127.0.0.1:8100"
},
"timeoutMs": {
"type": "number",
"default": 10000
},
"minPromptLength": {
"type": "number",
"default": 500
},
"ollamaUrl": {
"type": "string",
"default": "http://127.0.0.1:11434"
},
"preferredModel": {
"type": "string",
"default": "qwen3:8b"
},
"compressionStrategy": {
"type": "string",
"enum": ["auto", "full", "light", "passthrough"],
"default": "auto"
},
"inferenceMode": {
"type": "string",
"enum": ["auto", "cpu", "gpu", "remote"],
"default": "auto"
},
"metricsEnabled": {
"type": "boolean",
"default": false
},
"metricsUrl": {
"type": "string",
"default": "http://192.168.1.203:8101"
}
}
},
"uiHints": {
"serviceUrl": {
"label": "Compression Service URL",
"placeholder": "http://127.0.0.1:8100",
"help": "URL of the TokenRanger FastAPI service"
},
"timeoutMs": {
"label": "Timeout (ms)",
"help": "Max time to wait for compression before falling through",
"advanced": true
},
"minPromptLength": {
"label": "Min Prompt Length",
"help": "Only compress when session history exceeds this many characters",
"advanced": true
},
"ollamaUrl": {
"label": "Ollama URL",
"placeholder": "http://127.0.0.1:11434",
"advanced": true
},
"preferredModel": {
"label": "Preferred SLM Model",
"help": "Ollama model for context compression (e.g. qwen3:8b, qwen3:1.7b)"
},
"compressionStrategy": {
"label": "Compression Strategy",
"help": "auto: GPU detection selects full/light; passthrough: no compression"
},
"inferenceMode": {
"label": "Inference Mode",
"help": "auto: probe GPU; cpu: force light; gpu: force full; remote: use remote Ollama"
},
"metricsEnabled": {
"label": "Metrics",
"help": "Send compression events to the centralized metrics collector",
"advanced": true
},
"metricsUrl": {
"label": "Metrics Collector URL",
"placeholder": "http://192.168.1.203:8101",
"advanced": true
}
}
}