Public Access
Compare commits
1
Commits
| Author | SHA1 | Date | |
|---|---|---|---|
|
|
20eed38477 |
@@ -18,7 +18,7 @@ jobs:
|
|||||||
uses: actions/checkout@v7
|
uses: actions/checkout@v7
|
||||||
|
|
||||||
- name: Setup .NET
|
- name: Setup .NET
|
||||||
uses: actions/setup-dotnet@v5
|
uses: actions/setup-dotnet@v6
|
||||||
with:
|
with:
|
||||||
dotnet-version: "10.0.x"
|
dotnet-version: "10.0.x"
|
||||||
|
|
||||||
|
|||||||
@@ -13,50 +13,21 @@ namespace MeetingAssistant.Tests;
|
|||||||
|
|
||||||
public sealed class LiteLlmScreenshotOcrClientTests
|
public sealed class LiteLlmScreenshotOcrClientTests
|
||||||
{
|
{
|
||||||
[Fact]
|
|
||||||
public async Task ExtractUsesAgentStreamingTransportByDefault()
|
|
||||||
{
|
|
||||||
var screenshotPath = await CreateScreenshotAsync([1, 2, 3]);
|
|
||||||
var handler = new RecordingHandler(
|
|
||||||
CreateStreamedTextResponse("Streamed OCR text"),
|
|
||||||
"text/event-stream");
|
|
||||||
var client = new LiteLlmScreenshotOcrClient(
|
|
||||||
() => handler,
|
|
||||||
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
|
||||||
var options = new MeetingAssistantOptions
|
|
||||||
{
|
|
||||||
Agent =
|
|
||||||
{
|
|
||||||
Endpoint = "https://summary.local",
|
|
||||||
Model = "vision-model",
|
|
||||||
Key = "agent-key",
|
|
||||||
UseStreaming = true
|
|
||||||
}
|
|
||||||
};
|
|
||||||
|
|
||||||
var result = await client.ExtractAsync(
|
|
||||||
screenshotPath,
|
|
||||||
"Extract screenshot.",
|
|
||||||
options,
|
|
||||||
CancellationToken.None);
|
|
||||||
|
|
||||||
Assert.Equal("Streamed OCR text", result.Text);
|
|
||||||
using var payload = JsonDocument.Parse(handler.RequestBody!);
|
|
||||||
Assert.True(payload.RootElement.GetProperty("stream").GetBoolean());
|
|
||||||
var message = Assert.Single(payload.RootElement.GetProperty("input").EnumerateArray());
|
|
||||||
Assert.Equal("message", message.GetProperty("type").GetString());
|
|
||||||
var content = message.GetProperty("content").EnumerateArray().ToArray();
|
|
||||||
Assert.Equal("input_text", content[0].GetProperty("type").GetString());
|
|
||||||
Assert.Equal("Extract screenshot.", content[0].GetProperty("text").GetString());
|
|
||||||
Assert.Equal("input_image", content[1].GetProperty("type").GetString());
|
|
||||||
Assert.Equal("data:image/png;base64,AQID", content[1].GetProperty("image_url").GetString());
|
|
||||||
}
|
|
||||||
|
|
||||||
[Fact]
|
[Fact]
|
||||||
public async Task ExtractUsesAgentEndpointAndModelWhenOcrEndpointAndModelAreBlank()
|
public async Task ExtractUsesAgentEndpointAndModelWhenOcrEndpointAndModelAreBlank()
|
||||||
{
|
{
|
||||||
var screenshotPath = await CreateScreenshotAsync([1, 2, 3]);
|
var screenshotPath = await CreateScreenshotAsync([1, 2, 3]);
|
||||||
var handler = new RecordingHandler(CreateNonStreamingTextResponse("Visible slide text"));
|
var handler = new RecordingHandler("""
|
||||||
|
{
|
||||||
|
"output": [
|
||||||
|
{
|
||||||
|
"content": [
|
||||||
|
{ "type": "output_text", "text": "Visible slide text" }
|
||||||
|
]
|
||||||
|
}
|
||||||
|
]
|
||||||
|
}
|
||||||
|
""");
|
||||||
var client = new LiteLlmScreenshotOcrClient(
|
var client = new LiteLlmScreenshotOcrClient(
|
||||||
() => handler,
|
() => handler,
|
||||||
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
||||||
@@ -66,8 +37,7 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
{
|
{
|
||||||
Endpoint = "https://summary.local",
|
Endpoint = "https://summary.local",
|
||||||
Model = "summary-model",
|
Model = "summary-model",
|
||||||
Key = "agent-key",
|
Key = "agent-key"
|
||||||
UseStreaming = false
|
|
||||||
},
|
},
|
||||||
Screenshots =
|
Screenshots =
|
||||||
{
|
{
|
||||||
@@ -92,7 +62,6 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
Assert.Equal("Bearer", handler.Authorization?.Scheme);
|
Assert.Equal("Bearer", handler.Authorization?.Scheme);
|
||||||
Assert.Equal("ocr-key", handler.Authorization?.Parameter);
|
Assert.Equal("ocr-key", handler.Authorization?.Parameter);
|
||||||
using var payload = JsonDocument.Parse(handler.RequestBody!);
|
using var payload = JsonDocument.Parse(handler.RequestBody!);
|
||||||
Assert.False(payload.RootElement.GetProperty("stream").GetBoolean());
|
|
||||||
Assert.Equal("summary-model", payload.RootElement.GetProperty("model").GetString());
|
Assert.Equal("summary-model", payload.RootElement.GetProperty("model").GetString());
|
||||||
var content = payload.RootElement
|
var content = payload.RootElement
|
||||||
.GetProperty("input")[0]
|
.GetProperty("input")[0]
|
||||||
@@ -105,7 +74,7 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
public async Task ExtractUsesScreenshotOcrEndpointAndModelWhenConfigured()
|
public async Task ExtractUsesScreenshotOcrEndpointAndModelWhenConfigured()
|
||||||
{
|
{
|
||||||
var screenshotPath = await CreateScreenshotAsync([4, 5, 6]);
|
var screenshotPath = await CreateScreenshotAsync([4, 5, 6]);
|
||||||
var handler = new RecordingHandler(CreateNonStreamingTextResponse("OCR result"));
|
var handler = new RecordingHandler("""{ "output_text": "OCR result" }""");
|
||||||
var client = new LiteLlmScreenshotOcrClient(
|
var client = new LiteLlmScreenshotOcrClient(
|
||||||
() => handler,
|
() => handler,
|
||||||
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
||||||
@@ -115,8 +84,7 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
{
|
{
|
||||||
Endpoint = "https://summary.local",
|
Endpoint = "https://summary.local",
|
||||||
Model = "summary-model",
|
Model = "summary-model",
|
||||||
Key = "agent-key",
|
Key = "agent-key"
|
||||||
UseStreaming = false
|
|
||||||
},
|
},
|
||||||
Screenshots =
|
Screenshots =
|
||||||
{
|
{
|
||||||
@@ -145,14 +113,11 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
public async Task ExtractParsesCropMetadataAndOmitsMetadataFromReturnedText()
|
public async Task ExtractParsesCropMetadataAndOmitsMetadataFromReturnedText()
|
||||||
{
|
{
|
||||||
var screenshotPath = await CreateScreenshotAsync(CreatePngBytes(8, 6));
|
var screenshotPath = await CreateScreenshotAsync(CreatePngBytes(8, 6));
|
||||||
var handler = new RecordingHandler(CreateNonStreamingTextResponse(
|
var handler = new RecordingHandler("""
|
||||||
"""
|
{
|
||||||
Slide text
|
"output_text": "Slide text\n\n```json\n{ \"crop\": { \"x\": 1, \"y\": 2, \"width\": 3, \"height\": 4 } }\n```"
|
||||||
|
}
|
||||||
```json
|
""");
|
||||||
{ "crop": { "x": 1, "y": 2, "width": 3, "height": 4 } }
|
|
||||||
```
|
|
||||||
"""));
|
|
||||||
var client = new LiteLlmScreenshotOcrClient(
|
var client = new LiteLlmScreenshotOcrClient(
|
||||||
() => handler,
|
() => handler,
|
||||||
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
||||||
@@ -160,8 +125,7 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
{
|
{
|
||||||
Agent =
|
Agent =
|
||||||
{
|
{
|
||||||
Key = "agent-key",
|
Key = "agent-key"
|
||||||
UseStreaming = false
|
|
||||||
}
|
}
|
||||||
};
|
};
|
||||||
|
|
||||||
@@ -184,14 +148,11 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
public async Task ExtractParsesAttendeeMetadataAndOmitsMetadataFromReturnedText()
|
public async Task ExtractParsesAttendeeMetadataAndOmitsMetadataFromReturnedText()
|
||||||
{
|
{
|
||||||
var screenshotPath = await CreateScreenshotAsync([1, 2, 3]);
|
var screenshotPath = await CreateScreenshotAsync([1, 2, 3]);
|
||||||
var handler = new RecordingHandler(CreateNonStreamingTextResponse(
|
var handler = new RecordingHandler("""
|
||||||
"""
|
{
|
||||||
Visible participant tiles: Ada and Grace.
|
"output_text": "Visible participant tiles: Ada and Grace.\n\n```json\n{ \"crop\": null, \"attendees\": [\"Ada Lovelace\", \"Grace Hopper\"] }\n```"
|
||||||
|
}
|
||||||
```json
|
""");
|
||||||
{ "crop": null, "attendees": ["Ada Lovelace", "Grace Hopper"] }
|
|
||||||
```
|
|
||||||
"""));
|
|
||||||
var client = new LiteLlmScreenshotOcrClient(
|
var client = new LiteLlmScreenshotOcrClient(
|
||||||
() => handler,
|
() => handler,
|
||||||
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
||||||
@@ -199,8 +160,7 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
{
|
{
|
||||||
Agent =
|
Agent =
|
||||||
{
|
{
|
||||||
Key = "agent-key",
|
Key = "agent-key"
|
||||||
UseStreaming = false
|
|
||||||
}
|
}
|
||||||
};
|
};
|
||||||
|
|
||||||
@@ -218,14 +178,11 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
public async Task ExtractIgnoresMalformedAttendeesMetadataAndStillParsesCrop()
|
public async Task ExtractIgnoresMalformedAttendeesMetadataAndStillParsesCrop()
|
||||||
{
|
{
|
||||||
var screenshotPath = await CreateScreenshotAsync(CreatePngBytes(8, 6));
|
var screenshotPath = await CreateScreenshotAsync(CreatePngBytes(8, 6));
|
||||||
var handler = new RecordingHandler(CreateNonStreamingTextResponse(
|
var handler = new RecordingHandler("""
|
||||||
"""
|
{
|
||||||
Slide text
|
"output_text": "Slide text\n\n```json\n{ \"crop\": { \"x\": 1, \"y\": 2, \"width\": 3, \"height\": 4 }, \"attendees\": \"Ada\" }\n```"
|
||||||
|
}
|
||||||
```json
|
""");
|
||||||
{ "crop": { "x": 1, "y": 2, "width": 3, "height": 4 }, "attendees": "Ada" }
|
|
||||||
```
|
|
||||||
"""));
|
|
||||||
var client = new LiteLlmScreenshotOcrClient(
|
var client = new LiteLlmScreenshotOcrClient(
|
||||||
() => handler,
|
() => handler,
|
||||||
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
NullLogger<LiteLlmScreenshotOcrClient>.Instance);
|
||||||
@@ -233,8 +190,7 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
{
|
{
|
||||||
Agent =
|
Agent =
|
||||||
{
|
{
|
||||||
Key = "agent-key",
|
Key = "agent-key"
|
||||||
UseStreaming = false
|
|
||||||
}
|
}
|
||||||
};
|
};
|
||||||
|
|
||||||
@@ -252,12 +208,10 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
private sealed class RecordingHandler : HttpMessageHandler
|
private sealed class RecordingHandler : HttpMessageHandler
|
||||||
{
|
{
|
||||||
private readonly string responseBody;
|
private readonly string responseBody;
|
||||||
private readonly string mediaType;
|
|
||||||
|
|
||||||
public RecordingHandler(string responseBody, string mediaType = "application/json")
|
public RecordingHandler(string responseBody)
|
||||||
{
|
{
|
||||||
this.responseBody = responseBody;
|
this.responseBody = responseBody;
|
||||||
this.mediaType = mediaType;
|
|
||||||
}
|
}
|
||||||
|
|
||||||
public Uri? RequestUri { get; private set; }
|
public Uri? RequestUri { get; private set; }
|
||||||
@@ -277,7 +231,7 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
: await request.Content.ReadAsStringAsync(cancellationToken);
|
: await request.Content.ReadAsStringAsync(cancellationToken);
|
||||||
return new HttpResponseMessage(HttpStatusCode.OK)
|
return new HttpResponseMessage(HttpStatusCode.OK)
|
||||||
{
|
{
|
||||||
Content = new StringContent(responseBody, Encoding.UTF8, mediaType)
|
Content = new StringContent(responseBody, Encoding.UTF8, "application/json")
|
||||||
};
|
};
|
||||||
}
|
}
|
||||||
}
|
}
|
||||||
@@ -295,55 +249,6 @@ public sealed class LiteLlmScreenshotOcrClientTests
|
|||||||
return stream.ToArray();
|
return stream.ToArray();
|
||||||
}
|
}
|
||||||
|
|
||||||
private static string CreateNonStreamingTextResponse(string text)
|
|
||||||
{
|
|
||||||
return JsonSerializer.Serialize(new
|
|
||||||
{
|
|
||||||
id = "resp_ocr",
|
|
||||||
created_at = 1779147100,
|
|
||||||
model = "vision-model",
|
|
||||||
@object = "response",
|
|
||||||
output = new[]
|
|
||||||
{
|
|
||||||
new
|
|
||||||
{
|
|
||||||
id = "msg_ocr",
|
|
||||||
type = "message",
|
|
||||||
status = "completed",
|
|
||||||
content = new[]
|
|
||||||
{
|
|
||||||
new
|
|
||||||
{
|
|
||||||
type = "output_text",
|
|
||||||
annotations = Array.Empty<object>(),
|
|
||||||
text
|
|
||||||
}
|
|
||||||
},
|
|
||||||
role = "assistant"
|
|
||||||
}
|
|
||||||
},
|
|
||||||
parallel_tool_calls = true,
|
|
||||||
status = "completed",
|
|
||||||
store = false
|
|
||||||
});
|
|
||||||
}
|
|
||||||
|
|
||||||
private static string CreateStreamedTextResponse(string text)
|
|
||||||
{
|
|
||||||
var delta = JsonSerializer.Serialize(new
|
|
||||||
{
|
|
||||||
type = "response.output_text.delta",
|
|
||||||
item_id = "msg_ocr",
|
|
||||||
output_index = 0,
|
|
||||||
content_index = 0,
|
|
||||||
delta = text
|
|
||||||
});
|
|
||||||
return
|
|
||||||
$"data: {delta}{Environment.NewLine}{Environment.NewLine}" +
|
|
||||||
"""data: {"type":"response.completed","response":{"id":"resp_ocr","created_at":1779147100,"model":"vision-model","object":"response","output":[],"parallel_tool_calls":true,"status":"completed","store":false}}""" +
|
|
||||||
$"{Environment.NewLine}{Environment.NewLine}data: [DONE]{Environment.NewLine}{Environment.NewLine}";
|
|
||||||
}
|
|
||||||
|
|
||||||
private static async Task<string> CreateScreenshotAsync(byte[] bytes)
|
private static async Task<string> CreateScreenshotAsync(byte[] bytes)
|
||||||
{
|
{
|
||||||
var screenshotPath = Path.Combine(Path.GetTempPath(), "meeting-assistant-tests", Guid.NewGuid().ToString("N") + ".png");
|
var screenshotPath = Path.Combine(Path.GetTempPath(), "meeting-assistant-tests", Guid.NewGuid().ToString("N") + ".png");
|
||||||
|
|||||||
@@ -13,23 +13,22 @@ public sealed class TaskbarIconTests
|
|||||||
{
|
{
|
||||||
var menu = MeetingTaskbarMenuBuilder.Build(
|
var menu = MeetingTaskbarMenuBuilder.Build(
|
||||||
Status(),
|
Status(),
|
||||||
[Profile("default", "Ctrl+Alt+M"), Profile("english", "Ctrl+Alt+L")],
|
[Profile("default", "Ctrl+Alt+M"), Profile("english", "Ctrl+Alt+L")]);
|
||||||
[new MicrophoneDevice("integrated", "integrated microphone")],
|
|
||||||
"integrated");
|
|
||||||
|
|
||||||
Assert.Equal(RecordingProcessState.Idle, menu.State);
|
Assert.Equal(RecordingProcessState.Idle, menu.State);
|
||||||
AssertMenuLayout(
|
Assert.Contains(menu.Items, item =>
|
||||||
menu,
|
item.Action == MeetingTaskbarAction.EditRules &&
|
||||||
("Open agent", MeetingTaskbarAction.EditRules, false),
|
item.Text == "Open agent");
|
||||||
("Microphone", MeetingTaskbarAction.OpenSubmenu, true),
|
Assert.Contains(menu.Items, item =>
|
||||||
("Start meeting recording (default)\tCtrl+Alt+M", MeetingTaskbarAction.StartRecording, false),
|
item.Action == MeetingTaskbarAction.StartRecording &&
|
||||||
("Start meeting recording (english)\tCtrl+Alt+L", MeetingTaskbarAction.StartRecording, false),
|
item.ProfileName == "default" &&
|
||||||
("Exit", MeetingTaskbarAction.Exit, true));
|
item.Text == "Start meeting recording (default)\tCtrl+Alt+M");
|
||||||
Assert.Equal(
|
Assert.Contains(menu.Items, item =>
|
||||||
["default", "english"],
|
item.Action == MeetingTaskbarAction.StartRecording &&
|
||||||
menu.Items
|
item.ProfileName == "english" &&
|
||||||
.Where(item => item.Action == MeetingTaskbarAction.StartRecording)
|
item.Text == "Start meeting recording (english)\tCtrl+Alt+L");
|
||||||
.Select(item => item.ProfileName));
|
Assert.DoesNotContain(menu.Items, item => item.Action == MeetingTaskbarAction.StopRecording);
|
||||||
|
Assert.DoesNotContain(menu.Items, item => item.Action == MeetingTaskbarAction.AbortRecording);
|
||||||
}
|
}
|
||||||
|
|
||||||
[Fact]
|
[Fact]
|
||||||
@@ -61,41 +60,27 @@ public sealed class TaskbarIconTests
|
|||||||
}
|
}
|
||||||
|
|
||||||
[Fact]
|
[Fact]
|
||||||
public void RecordingMenuPrioritizesFinishMeetingInDedicatedSection()
|
public void RecordingMenuOffersStopAbortAndOtherProfileSwitches()
|
||||||
{
|
{
|
||||||
var menu = MeetingTaskbarMenuBuilder.Build(
|
var menu = MeetingTaskbarMenuBuilder.Build(
|
||||||
Status(isRecording: true, state: RecordingProcessState.Recording, profile: "default"),
|
Status(isRecording: true, state: RecordingProcessState.Recording, profile: "default"),
|
||||||
[Profile("default", "Ctrl+Alt+M"), Profile("english", "Ctrl+Alt+L")],
|
[Profile("default", "Ctrl+Alt+M"), Profile("english", "Ctrl+Alt+L"), Profile("french", "Ctrl+Alt+F")]);
|
||||||
[new MicrophoneDevice("integrated", "integrated microphone")],
|
|
||||||
"integrated");
|
|
||||||
|
|
||||||
Assert.Equal(RecordingProcessState.Recording, menu.State);
|
Assert.Equal(RecordingProcessState.Recording, menu.State);
|
||||||
AssertMenuLayout(
|
Assert.Contains(menu.Items, item => item.Action == MeetingTaskbarAction.StopRecording);
|
||||||
menu,
|
Assert.Contains(menu.Items, item => item.Action == MeetingTaskbarAction.AbortRecording);
|
||||||
("Open agent", MeetingTaskbarAction.EditRules, false),
|
Assert.Contains(menu.Items, item =>
|
||||||
("Finish meeting", MeetingTaskbarAction.StopRecording, true),
|
item.Action == MeetingTaskbarAction.SwitchProfile &&
|
||||||
("Microphone", MeetingTaskbarAction.OpenSubmenu, true),
|
item.ProfileName == "english" &&
|
||||||
("Cancel meeting recording and discard", MeetingTaskbarAction.AbortRecording, false),
|
item.Text == "Switch to english\tCtrl+Alt+L");
|
||||||
("Switch to english\tCtrl+Alt+L", MeetingTaskbarAction.SwitchProfile, false),
|
Assert.Contains(menu.Items, item =>
|
||||||
("Exit", MeetingTaskbarAction.Exit, true));
|
item.Action == MeetingTaskbarAction.SwitchProfile &&
|
||||||
Assert.Equal(
|
item.ProfileName == "french" &&
|
||||||
"english",
|
item.Text == "Switch to french\tCtrl+Alt+F");
|
||||||
Assert.Single(menu.Items, item => item.Action == MeetingTaskbarAction.SwitchProfile).ProfileName);
|
Assert.DoesNotContain(menu.Items, item =>
|
||||||
}
|
item.Action == MeetingTaskbarAction.SwitchProfile &&
|
||||||
|
item.ProfileName == "default");
|
||||||
[Fact]
|
Assert.DoesNotContain(menu.Items, item => item.Action == MeetingTaskbarAction.StartRecording);
|
||||||
public void RecordingMenuKeepsFinishMeetingIsolatedWithoutMicrophones()
|
|
||||||
{
|
|
||||||
var menu = MeetingTaskbarMenuBuilder.Build(
|
|
||||||
Status(isRecording: true, state: RecordingProcessState.Recording, profile: "default"),
|
|
||||||
[Profile("default")]);
|
|
||||||
|
|
||||||
AssertMenuLayout(
|
|
||||||
menu,
|
|
||||||
("Open agent", MeetingTaskbarAction.EditRules, false),
|
|
||||||
("Finish meeting", MeetingTaskbarAction.StopRecording, true),
|
|
||||||
("Cancel meeting recording and discard", MeetingTaskbarAction.AbortRecording, true),
|
|
||||||
("Exit", MeetingTaskbarAction.Exit, true));
|
|
||||||
}
|
}
|
||||||
|
|
||||||
[Fact]
|
[Fact]
|
||||||
@@ -188,15 +173,6 @@ public sealed class TaskbarIconTests
|
|||||||
});
|
});
|
||||||
}
|
}
|
||||||
|
|
||||||
private static void AssertMenuLayout(
|
|
||||||
MeetingTaskbarMenu menu,
|
|
||||||
params (string Text, MeetingTaskbarAction Action, bool StartsSection)[] expected)
|
|
||||||
{
|
|
||||||
Assert.Equal(
|
|
||||||
expected,
|
|
||||||
menu.Items.Select(item => (item.Text, item.Action, item.StartsSection)));
|
|
||||||
}
|
|
||||||
|
|
||||||
private static RecordingStatus Status(
|
private static RecordingStatus Status(
|
||||||
bool isRecording = false,
|
bool isRecording = false,
|
||||||
RecordingProcessState state = RecordingProcessState.Idle,
|
RecordingProcessState state = RecordingProcessState.Idle,
|
||||||
|
|||||||
@@ -26,7 +26,7 @@
|
|||||||
<PackageReference Include="Microsoft.CognitiveServices.Speech.Extension.MAS" Version="$(MicrosoftSpeechVersion)" ExcludeAssets="build" />
|
<PackageReference Include="Microsoft.CognitiveServices.Speech.Extension.MAS" Version="$(MicrosoftSpeechVersion)" ExcludeAssets="build" />
|
||||||
<PackageReference Include="Microsoft.EntityFrameworkCore.Sqlite" Version="10.0.10" />
|
<PackageReference Include="Microsoft.EntityFrameworkCore.Sqlite" Version="10.0.10" />
|
||||||
<PackageReference Include="NAudio" Version="2.3.0" />
|
<PackageReference Include="NAudio" Version="2.3.0" />
|
||||||
<PackageReference Include="NCalcSync" Version="7.0.2" />
|
<PackageReference Include="NCalcSync" Version="6.4.0" />
|
||||||
<PackageReference Include="RazorLight" Version="2.3.1" />
|
<PackageReference Include="RazorLight" Version="2.3.1" />
|
||||||
<PackageReference Include="SQLitePCLRaw.bundle_e_sqlite3" Version="3.0.5" />
|
<PackageReference Include="SQLitePCLRaw.bundle_e_sqlite3" Version="3.0.5" />
|
||||||
<PackageReference Include="System.Drawing.Common" Version="10.0.10" />
|
<PackageReference Include="System.Drawing.Common" Version="10.0.10" />
|
||||||
|
|||||||
@@ -1,13 +1,15 @@
|
|||||||
|
using System.Net.Http.Headers;
|
||||||
|
using System.Text;
|
||||||
using System.Text.Json;
|
using System.Text.Json;
|
||||||
|
using System.Text.Json.Nodes;
|
||||||
using System.Text.RegularExpressions;
|
using System.Text.RegularExpressions;
|
||||||
using MeetingAssistant.MeetingNotes;
|
using MeetingAssistant.MeetingNotes;
|
||||||
using MeetingAssistant.Summary;
|
|
||||||
using Microsoft.Extensions.AI;
|
|
||||||
|
|
||||||
namespace MeetingAssistant.Screenshots;
|
namespace MeetingAssistant.Screenshots;
|
||||||
|
|
||||||
public sealed partial class LiteLlmScreenshotOcrClient : IScreenshotOcrClient
|
public sealed partial class LiteLlmScreenshotOcrClient : IScreenshotOcrClient
|
||||||
{
|
{
|
||||||
|
private static readonly JsonSerializerOptions JsonOptions = new(JsonSerializerDefaults.Web);
|
||||||
private readonly ILogger<LiteLlmScreenshotOcrClient> logger;
|
private readonly ILogger<LiteLlmScreenshotOcrClient> logger;
|
||||||
private readonly Func<HttpMessageHandler>? httpMessageHandlerFactory;
|
private readonly Func<HttpMessageHandler>? httpMessageHandlerFactory;
|
||||||
|
|
||||||
@@ -38,30 +40,20 @@ public sealed partial class LiteLlmScreenshotOcrClient : IScreenshotOcrClient
|
|||||||
: options.Agent.Model;
|
: options.Agent.Model;
|
||||||
var key = ResolveApiKey(options);
|
var key = ResolveApiKey(options);
|
||||||
var imageBytes = await File.ReadAllBytesAsync(screenshotPath, cancellationToken);
|
var imageBytes = await File.ReadAllBytesAsync(screenshotPath, cancellationToken);
|
||||||
var httpClient = CreateHttpClient();
|
using var httpClient = CreateHttpClient();
|
||||||
httpClient.BaseAddress = LiteLlmResponsesChatClient.NormalizeEndpoint(new Uri(endpoint));
|
httpClient.BaseAddress = NormalizeEndpoint(new Uri(endpoint));
|
||||||
using var chatClient = new LiteLlmResponsesChatClient(
|
httpClient.DefaultRequestHeaders.Authorization = new AuthenticationHeaderValue("Bearer", key);
|
||||||
httpClient,
|
var payload = CreatePayload(model, CreatePrompt(prompt, imageBytes), imageBytes);
|
||||||
key,
|
using var content = new StringContent(payload.ToJsonString(JsonOptions), Encoding.UTF8, "application/json");
|
||||||
model,
|
using var response = await httpClient.PostAsync("responses", content, cancellationToken);
|
||||||
enableThinking: false,
|
var responseJson = await response.Content.ReadAsStringAsync(cancellationToken);
|
||||||
reasoningEffort: "none",
|
if (!response.IsSuccessStatusCode)
|
||||||
reconnectionAttempts: options.Agent.ReconnectionAttempts,
|
{
|
||||||
reconnectionDelay: options.Agent.ReconnectionDelay,
|
throw new InvalidOperationException(
|
||||||
logger: logger,
|
$"Screenshot OCR request failed with {(int)response.StatusCode} {response.ReasonPhrase}: {responseJson}");
|
||||||
firstRequestIsUser: false,
|
}
|
||||||
useStreaming: options.Agent.UseStreaming);
|
|
||||||
var response = await chatClient.GetResponseAsync(
|
var text = ParseOutputText(responseJson);
|
||||||
[
|
|
||||||
new ChatMessage(
|
|
||||||
ChatRole.User,
|
|
||||||
[
|
|
||||||
new TextContent(CreatePrompt(prompt, imageBytes)),
|
|
||||||
new DataContent(imageBytes, "image/png")
|
|
||||||
])
|
|
||||||
],
|
|
||||||
cancellationToken: cancellationToken);
|
|
||||||
var text = response.Text ?? string.Empty;
|
|
||||||
logger.LogInformation("Screenshot OCR completed for {ScreenshotPath}", screenshotPath);
|
logger.LogInformation("Screenshot OCR completed for {ScreenshotPath}", screenshotPath);
|
||||||
return ParseOcrResult(text);
|
return ParseOcrResult(text);
|
||||||
}
|
}
|
||||||
@@ -73,6 +65,35 @@ public sealed partial class LiteLlmScreenshotOcrClient : IScreenshotOcrClient
|
|||||||
: new HttpClient(httpMessageHandlerFactory());
|
: new HttpClient(httpMessageHandlerFactory());
|
||||||
}
|
}
|
||||||
|
|
||||||
|
private static JsonObject CreatePayload(string model, string prompt, byte[] imageBytes)
|
||||||
|
{
|
||||||
|
return new JsonObject
|
||||||
|
{
|
||||||
|
["model"] = model,
|
||||||
|
["store"] = false,
|
||||||
|
["input"] = new JsonArray
|
||||||
|
{
|
||||||
|
new JsonObject
|
||||||
|
{
|
||||||
|
["role"] = "user",
|
||||||
|
["content"] = new JsonArray
|
||||||
|
{
|
||||||
|
new JsonObject
|
||||||
|
{
|
||||||
|
["type"] = "input_text",
|
||||||
|
["text"] = prompt
|
||||||
|
},
|
||||||
|
new JsonObject
|
||||||
|
{
|
||||||
|
["type"] = "input_image",
|
||||||
|
["image_url"] = "data:image/png;base64," + Convert.ToBase64String(imageBytes)
|
||||||
|
}
|
||||||
|
}
|
||||||
|
}
|
||||||
|
}
|
||||||
|
};
|
||||||
|
}
|
||||||
|
|
||||||
private static string CreatePrompt(string prompt, byte[] imageBytes)
|
private static string CreatePrompt(string prompt, byte[] imageBytes)
|
||||||
{
|
{
|
||||||
return TryReadPngDimensions(imageBytes, out var width, out var height)
|
return TryReadPngDimensions(imageBytes, out var width, out var height)
|
||||||
@@ -189,6 +210,46 @@ public sealed partial class LiteLlmScreenshotOcrClient : IScreenshotOcrClient
|
|||||||
bytes[offset + 3];
|
bytes[offset + 3];
|
||||||
}
|
}
|
||||||
|
|
||||||
|
private static string ParseOutputText(string responseJson)
|
||||||
|
{
|
||||||
|
using var document = JsonDocument.Parse(responseJson);
|
||||||
|
var root = document.RootElement;
|
||||||
|
var parts = new List<string>();
|
||||||
|
if (root.TryGetProperty("output_text", out var outputText) &&
|
||||||
|
outputText.ValueKind == JsonValueKind.String &&
|
||||||
|
!string.IsNullOrWhiteSpace(outputText.GetString()))
|
||||||
|
{
|
||||||
|
parts.Add(outputText.GetString()!);
|
||||||
|
}
|
||||||
|
|
||||||
|
if (root.TryGetProperty("output", out var output) && output.ValueKind == JsonValueKind.Array)
|
||||||
|
{
|
||||||
|
foreach (var item in output.EnumerateArray())
|
||||||
|
{
|
||||||
|
if (!item.TryGetProperty("content", out var content) || content.ValueKind != JsonValueKind.Array)
|
||||||
|
{
|
||||||
|
continue;
|
||||||
|
}
|
||||||
|
|
||||||
|
foreach (var block in content.EnumerateArray())
|
||||||
|
{
|
||||||
|
if (block.TryGetProperty("type", out var type) &&
|
||||||
|
type.GetString() == "output_text" &&
|
||||||
|
block.TryGetProperty("text", out var text) &&
|
||||||
|
text.ValueKind == JsonValueKind.String &&
|
||||||
|
!string.IsNullOrWhiteSpace(text.GetString()))
|
||||||
|
{
|
||||||
|
parts.Add(text.GetString()!);
|
||||||
|
}
|
||||||
|
}
|
||||||
|
}
|
||||||
|
}
|
||||||
|
|
||||||
|
return parts.Count == 0
|
||||||
|
? ""
|
||||||
|
: string.Join(Environment.NewLine + Environment.NewLine, parts);
|
||||||
|
}
|
||||||
|
|
||||||
private static string ResolveApiKey(MeetingAssistantOptions options)
|
private static string ResolveApiKey(MeetingAssistantOptions options)
|
||||||
{
|
{
|
||||||
if (!string.IsNullOrWhiteSpace(options.Screenshots.Ocr.Key))
|
if (!string.IsNullOrWhiteSpace(options.Screenshots.Ocr.Key))
|
||||||
@@ -217,6 +278,17 @@ public sealed partial class LiteLlmScreenshotOcrClient : IScreenshotOcrClient
|
|||||||
$"No screenshot OCR API key configured. Set MeetingAssistant:Screenshots:Ocr:Key or environment variable '{options.Screenshots.Ocr.KeyEnv}'.");
|
$"No screenshot OCR API key configured. Set MeetingAssistant:Screenshots:Ocr:Key or environment variable '{options.Screenshots.Ocr.KeyEnv}'.");
|
||||||
}
|
}
|
||||||
|
|
||||||
|
private static Uri NormalizeEndpoint(Uri endpoint)
|
||||||
|
{
|
||||||
|
var value = endpoint.ToString().TrimEnd('/');
|
||||||
|
if (!value.EndsWith("/v1", StringComparison.OrdinalIgnoreCase))
|
||||||
|
{
|
||||||
|
value += "/v1";
|
||||||
|
}
|
||||||
|
|
||||||
|
return new Uri(value + "/");
|
||||||
|
}
|
||||||
|
|
||||||
[GeneratedRegex("```json\\s*(?<json>.*?)\\s*```", RegexOptions.Singleline | RegexOptions.IgnoreCase)]
|
[GeneratedRegex("```json\\s*(?<json>.*?)\\s*```", RegexOptions.Singleline | RegexOptions.IgnoreCase)]
|
||||||
private static partial Regex JsonCodeBlockRegex();
|
private static partial Regex JsonCodeBlockRegex();
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -659,9 +659,9 @@ public sealed class LiteLlmResponsesChatClient : IChatClient
|
|||||||
private static void AddInputItem(JsonArray input, StringBuilder instructions, ChatMessage message)
|
private static void AddInputItem(JsonArray input, StringBuilder instructions, ChatMessage message)
|
||||||
{
|
{
|
||||||
var role = message.Role.Value;
|
var role = message.Role.Value;
|
||||||
|
var text = message.Text;
|
||||||
if (role == ChatRole.System.Value)
|
if (role == ChatRole.System.Value)
|
||||||
{
|
{
|
||||||
var text = message.Text;
|
|
||||||
if (!string.IsNullOrWhiteSpace(text))
|
if (!string.IsNullOrWhiteSpace(text))
|
||||||
{
|
{
|
||||||
instructions.AppendLine(text);
|
instructions.AppendLine(text);
|
||||||
@@ -670,37 +670,21 @@ public sealed class LiteLlmResponsesChatClient : IChatClient
|
|||||||
return;
|
return;
|
||||||
}
|
}
|
||||||
|
|
||||||
var isAssistant = role == ChatRole.Assistant.Value;
|
if (!string.IsNullOrWhiteSpace(text))
|
||||||
var messageContent = new JsonArray();
|
|
||||||
foreach (var content in message.Contents)
|
|
||||||
{
|
|
||||||
if (content is TextContent textContent && !string.IsNullOrWhiteSpace(textContent.Text))
|
|
||||||
{
|
|
||||||
messageContent.Add(new JsonObject
|
|
||||||
{
|
|
||||||
["type"] = isAssistant ? "output_text" : "input_text",
|
|
||||||
["text"] = textContent.Text
|
|
||||||
});
|
|
||||||
}
|
|
||||||
else if (!isAssistant &&
|
|
||||||
content is DataContent dataContent &&
|
|
||||||
dataContent.HasTopLevelMediaType("image"))
|
|
||||||
{
|
|
||||||
messageContent.Add(new JsonObject
|
|
||||||
{
|
|
||||||
["type"] = "input_image",
|
|
||||||
["image_url"] = dataContent.Uri.ToString()
|
|
||||||
});
|
|
||||||
}
|
|
||||||
}
|
|
||||||
|
|
||||||
if (messageContent.Count > 0)
|
|
||||||
{
|
{
|
||||||
|
var isAssistant = role == ChatRole.Assistant.Value;
|
||||||
input.Add(new JsonObject
|
input.Add(new JsonObject
|
||||||
{
|
{
|
||||||
["type"] = "message",
|
["type"] = "message",
|
||||||
["role"] = isAssistant ? "assistant" : "user",
|
["role"] = isAssistant ? "assistant" : "user",
|
||||||
["content"] = messageContent
|
["content"] = new JsonArray
|
||||||
|
{
|
||||||
|
new JsonObject
|
||||||
|
{
|
||||||
|
["type"] = isAssistant ? "output_text" : "input_text",
|
||||||
|
["text"] = text
|
||||||
|
}
|
||||||
|
}
|
||||||
});
|
});
|
||||||
}
|
}
|
||||||
|
|
||||||
@@ -777,7 +761,7 @@ public sealed class LiteLlmResponsesChatClient : IChatClient
|
|||||||
return Math.Max(1, (int)Math.Ceiling(json.Length / 4.0));
|
return Math.Max(1, (int)Math.Ceiling(json.Length / 4.0));
|
||||||
}
|
}
|
||||||
|
|
||||||
internal static Uri NormalizeEndpoint(Uri endpoint)
|
private static Uri NormalizeEndpoint(Uri endpoint)
|
||||||
{
|
{
|
||||||
var value = endpoint.ToString().TrimEnd('/');
|
var value = endpoint.ToString().TrimEnd('/');
|
||||||
if (!value.EndsWith("/v1", StringComparison.OrdinalIgnoreCase))
|
if (!value.EndsWith("/v1", StringComparison.OrdinalIgnoreCase))
|
||||||
|
|||||||
@@ -26,8 +26,7 @@ public sealed record MeetingTaskbarMenuItem(
|
|||||||
string? ProfileName = null,
|
string? ProfileName = null,
|
||||||
string? MicrophoneDeviceId = null,
|
string? MicrophoneDeviceId = null,
|
||||||
bool IsChecked = false,
|
bool IsChecked = false,
|
||||||
IReadOnlyList<MeetingTaskbarMenuItem>? Items = null,
|
IReadOnlyList<MeetingTaskbarMenuItem>? Items = null);
|
||||||
bool StartsSection = false);
|
|
||||||
|
|
||||||
public static class MeetingTaskbarMenuBuilder
|
public static class MeetingTaskbarMenuBuilder
|
||||||
{
|
{
|
||||||
@@ -42,29 +41,23 @@ public static class MeetingTaskbarMenuBuilder
|
|||||||
new("Open agent", MeetingTaskbarAction.EditRules)
|
new("Open agent", MeetingTaskbarAction.EditRules)
|
||||||
};
|
};
|
||||||
|
|
||||||
|
if (microphones is { Count: > 0 })
|
||||||
|
{
|
||||||
|
items.Add(BuildMicrophoneMenu(microphones, currentMicrophoneDeviceId));
|
||||||
|
}
|
||||||
|
|
||||||
if (status.IsRecording)
|
if (status.IsRecording)
|
||||||
{
|
{
|
||||||
items.Add(new MeetingTaskbarMenuItem(
|
items.Add(new MeetingTaskbarMenuItem(
|
||||||
"Finish meeting",
|
"Stop meeting recording and transcribe",
|
||||||
MeetingTaskbarAction.StopRecording,
|
MeetingTaskbarAction.StopRecording));
|
||||||
StartsSection: true));
|
items.Add(new MeetingTaskbarMenuItem(
|
||||||
}
|
|
||||||
|
|
||||||
var secondaryControls = new List<MeetingTaskbarMenuItem>();
|
|
||||||
if (microphones is { Count: > 0 })
|
|
||||||
{
|
|
||||||
secondaryControls.Add(BuildMicrophoneMenu(microphones, currentMicrophoneDeviceId));
|
|
||||||
}
|
|
||||||
|
|
||||||
if (status.IsRecording)
|
|
||||||
{
|
|
||||||
secondaryControls.Add(new MeetingTaskbarMenuItem(
|
|
||||||
"Cancel meeting recording and discard",
|
"Cancel meeting recording and discard",
|
||||||
MeetingTaskbarAction.AbortRecording));
|
MeetingTaskbarAction.AbortRecording));
|
||||||
|
|
||||||
foreach (var profile in launchProfiles.Where(profile => !IsActiveProfile(profile, status)))
|
foreach (var profile in launchProfiles.Where(profile => !IsActiveProfile(profile, status)))
|
||||||
{
|
{
|
||||||
secondaryControls.Add(new MeetingTaskbarMenuItem(
|
items.Add(new MeetingTaskbarMenuItem(
|
||||||
AppendHotkey($"Switch to {profile.Name}", profile.Options.Hotkey.Toggle),
|
AppendHotkey($"Switch to {profile.Name}", profile.Options.Hotkey.Toggle),
|
||||||
MeetingTaskbarAction.SwitchProfile,
|
MeetingTaskbarAction.SwitchProfile,
|
||||||
profile.Name));
|
profile.Name));
|
||||||
@@ -74,18 +67,16 @@ public static class MeetingTaskbarMenuBuilder
|
|||||||
{
|
{
|
||||||
foreach (var profile in launchProfiles)
|
foreach (var profile in launchProfiles)
|
||||||
{
|
{
|
||||||
secondaryControls.Add(new MeetingTaskbarMenuItem(
|
items.Add(new MeetingTaskbarMenuItem(
|
||||||
AppendHotkey($"Start meeting recording ({profile.Name})", profile.Options.Hotkey.Toggle),
|
AppendHotkey($"Start meeting recording ({profile.Name})", profile.Options.Hotkey.Toggle),
|
||||||
MeetingTaskbarAction.StartRecording,
|
MeetingTaskbarAction.StartRecording,
|
||||||
profile.Name));
|
profile.Name));
|
||||||
}
|
}
|
||||||
}
|
}
|
||||||
|
|
||||||
AddSection(items, secondaryControls);
|
|
||||||
items.Add(new MeetingTaskbarMenuItem(
|
items.Add(new MeetingTaskbarMenuItem(
|
||||||
"Exit",
|
"Exit",
|
||||||
MeetingTaskbarAction.Exit,
|
MeetingTaskbarAction.Exit));
|
||||||
StartsSection: true));
|
|
||||||
|
|
||||||
return new MeetingTaskbarMenu(
|
return new MeetingTaskbarMenu(
|
||||||
status.State,
|
status.State,
|
||||||
@@ -111,19 +102,6 @@ public static class MeetingTaskbarMenuBuilder
|
|||||||
Items: microphoneItems);
|
Items: microphoneItems);
|
||||||
}
|
}
|
||||||
|
|
||||||
private static void AddSection(
|
|
||||||
List<MeetingTaskbarMenuItem> items,
|
|
||||||
IReadOnlyList<MeetingTaskbarMenuItem> section)
|
|
||||||
{
|
|
||||||
if (section.Count == 0)
|
|
||||||
{
|
|
||||||
return;
|
|
||||||
}
|
|
||||||
|
|
||||||
items.Add(section[0] with { StartsSection = true });
|
|
||||||
items.AddRange(section.Skip(1));
|
|
||||||
}
|
|
||||||
|
|
||||||
private static string BuildTooltip(RecordingStatus status)
|
private static string BuildTooltip(RecordingStatus status)
|
||||||
{
|
{
|
||||||
return status.State switch
|
return status.State switch
|
||||||
|
|||||||
@@ -196,12 +196,14 @@ public sealed class UnoTaskbarIconService : IHostedService, IDisposable
|
|||||||
var popupMenu = new PopupMenu();
|
var popupMenu = new PopupMenu();
|
||||||
for (var index = 0; index < menu.Items.Count; index++)
|
for (var index = 0; index < menu.Items.Count; index++)
|
||||||
{
|
{
|
||||||
var menuItem = menu.Items[index];
|
if (index == 1 ||
|
||||||
if (index > 0 && menuItem.StartsSection)
|
(menu.Items[index].Action == MeetingTaskbarAction.Exit &&
|
||||||
|
menu.Items[index - 1].Action != MeetingTaskbarAction.EditRules))
|
||||||
{
|
{
|
||||||
popupMenu.Items.Add(new PopupMenuSeparator());
|
popupMenu.Items.Add(new PopupMenuSeparator());
|
||||||
}
|
}
|
||||||
|
|
||||||
|
var menuItem = menu.Items[index];
|
||||||
popupMenu.Items.Add(BuildPopupItem(menuItem));
|
popupMenu.Items.Add(BuildPopupItem(menuItem));
|
||||||
}
|
}
|
||||||
|
|
||||||
@@ -288,7 +290,7 @@ public sealed class UnoTaskbarIconService : IHostedService, IDisposable
|
|||||||
return string.Join(
|
return string.Join(
|
||||||
"|",
|
"|",
|
||||||
FlattenMenuItems(menu.Items).Select(item =>
|
FlattenMenuItems(menu.Items).Select(item =>
|
||||||
$"{item.Action}:{item.ProfileName}:{item.MicrophoneDeviceId}:{item.IsChecked}:{item.StartsSection}:{item.Text}"));
|
$"{item.Action}:{item.ProfileName}:{item.MicrophoneDeviceId}:{item.IsChecked}:{item.Text}"));
|
||||||
}
|
}
|
||||||
|
|
||||||
private static IEnumerable<MeetingTaskbarMenuItem> FlattenMenuItems(
|
private static IEnumerable<MeetingTaskbarMenuItem> FlattenMenuItems(
|
||||||
|
|||||||
@@ -317,7 +317,7 @@ When enabled on Windows, Meeting Assistant periodically syncs today's Outlook Cl
|
|||||||
|
|
||||||
`Screenshots:Hotkey` configures a global hotkey that captures the currently active window during an active meeting. Screenshots are written under `Screenshots:AttachmentsFolder`, which defaults to an `Attachments` folder beside the assistant context note, and each capture appends a timestamped markdown image link to the assistant context.
|
`Screenshots:Hotkey` configures a global hotkey that captures the currently active window during an active meeting. Screenshots are written under `Screenshots:AttachmentsFolder`, which defaults to an `Attachments` folder beside the assistant context note, and each capture appends a timestamped markdown image link to the assistant context.
|
||||||
|
|
||||||
`Screenshots:Ocr` optionally enables vision extraction for screenshots. Blank `Endpoint` or `Model` values fall back to the summary `Agent` endpoint and model. `Key` or `KeyEnv` can be set specifically for OCR; otherwise the summary agent key configuration is used. Screenshot OCR also inherits `Agent:UseStreaming`, using the Responses SSE transport when it is `true` and the non-streaming Responses transport when it is `false`. Automatic summarization scans the meeting note for user-added Obsidian image embeds such as `![[whiteboard.png]]` and Markdown image embeds such as ``, adds resolvable local images to the assistant context without copying them or changing the meeting note, runs OCR without crop or attendee updates, and waits for all pending OCR work to complete or hit `Timeout` before the assistant context moves to `summarizing`. Failed or timed-out screenshot OCR writes a retry link that targets `/meetings/screenshot-ocr/retry` with the exact saved screenshot and OCR block id.
|
`Screenshots:Ocr` optionally enables vision extraction for screenshots. Blank `Endpoint` or `Model` values fall back to the summary `Agent` endpoint and model. `Key` or `KeyEnv` can be set specifically for OCR; otherwise the summary agent key configuration is used. Automatic summarization scans the meeting note for user-added Obsidian image embeds such as `![[whiteboard.png]]` and Markdown image embeds such as ``, adds resolvable local images to the assistant context without copying them or changing the meeting note, runs OCR without crop or attendee updates, and waits for all pending OCR work to complete or hit `Timeout` before the assistant context moves to `summarizing`. Failed or timed-out screenshot OCR writes a retry link that targets `/meetings/screenshot-ocr/retry` with the exact saved screenshot and OCR block id.
|
||||||
|
|
||||||
| Setting | Purpose |
|
| Setting | Purpose |
|
||||||
| --- | --- |
|
| --- | --- |
|
||||||
|
|||||||
@@ -1,2 +0,0 @@
|
|||||||
schema: spec-driven
|
|
||||||
created: 2026-07-29
|
|
||||||
@@ -1,34 +0,0 @@
|
|||||||
## Context
|
|
||||||
|
|
||||||
`LiteLlmScreenshotOcrClient` currently builds and posts a raw Responses JSON payload, then parses the successful body as one JSON document. It inherits endpoint, model, and key values from `AgentOptions`, but never reads `AgentOptions.UseStreaming`. The summary and workflow agents already use `LiteLlmResponsesChatClient`, which selects the OpenAI SDK streaming or non-streaming Responses method and maps both through the Microsoft.Extensions.AI adapter.
|
|
||||||
|
|
||||||
## Goals / Non-Goals
|
|
||||||
|
|
||||||
**Goals:**
|
|
||||||
|
|
||||||
- Make screenshot OCR use `Agent:UseStreaming` without adding another setting.
|
|
||||||
- Reuse the supported Responses SDK transport and response adapter.
|
|
||||||
- Preserve screenshot-specific endpoint, model, key, prompt, image, crop, attendee, and timeout behavior.
|
|
||||||
- Keep non-streaming screenshot OCR working when streaming is disabled.
|
|
||||||
|
|
||||||
**Non-Goals:**
|
|
||||||
|
|
||||||
- Expose OCR token deltas to the UI or assistant context.
|
|
||||||
- Change screenshot OCR retry, crop, attendee, or note-block semantics.
|
|
||||||
- Add file upload or remote image URL support.
|
|
||||||
|
|
||||||
## Decisions
|
|
||||||
|
|
||||||
1. Route screenshot OCR through `LiteLlmResponsesChatClient` instead of maintaining a second Responses parser. The screenshot client will construct one user chat message containing prompt text and PNG `DataContent`, then consume the buffered `ChatResponse.Text`. This keeps transport selection, SDK request creation, SSE assembly, and non-streaming mapping in one client.
|
|
||||||
|
|
||||||
2. Extend the shared Responses message translator to map image `DataContent` to an `input_image` content block using its data URI. Text and image blocks remain in one `type: message` input item, matching the existing OCR payload shape.
|
|
||||||
|
|
||||||
3. Use `AgentOptions.UseStreaming` for screenshot OCR even when the screenshot-specific endpoint or model overrides are set. Endpoint/model/key remain independently overrideable; transport is an agent-wide behavior setting.
|
|
||||||
|
|
||||||
4. Disable reasoning and compaction for the one-turn OCR request, preserving the existing screenshot client behavior while still using the agent reconnection settings and selected transport.
|
|
||||||
|
|
||||||
## Risks / Trade-offs
|
|
||||||
|
|
||||||
- **[Shared-client diagnostics mention summary context]** Some low-level logs are named for the summary pipeline. → Avoid passing summary compaction state and keep screenshot-specific completion/failure logs at the screenshot client boundary.
|
|
||||||
- **[Multimodal translation expands shared client scope]** Incorrect content mapping could affect summary requests. → Add request-body behavior coverage proving prompt and PNG data URI are preserved, while existing summary message tests protect text translation.
|
|
||||||
- **[Provider image support varies]** A configured model may reject image input. → Preserve the provider error and existing screenshot OCR failure/retry behavior.
|
|
||||||
@@ -1,24 +0,0 @@
|
|||||||
## Why
|
|
||||||
|
|
||||||
Screenshot OCR inherits its endpoint, model, and key from the summary agent, but it bypasses the shared Responses client and ignores `Agent:UseStreaming`. With streaming enabled, the screenshot client still expects one JSON document and cannot consume the configured LiteLLM Responses event stream.
|
|
||||||
|
|
||||||
## What Changes
|
|
||||||
|
|
||||||
- Make screenshot OCR inherit the summary agent's streaming transport selection.
|
|
||||||
- Send screenshot image input through the same OpenAI Responses SDK and Agent Framework adapter used by the summary client.
|
|
||||||
- Preserve the existing non-streaming screenshot OCR path when streaming is disabled.
|
|
||||||
- Add behavior coverage for streamed and non-streamed screenshot OCR responses.
|
|
||||||
|
|
||||||
## Capabilities
|
|
||||||
|
|
||||||
### New Capabilities
|
|
||||||
|
|
||||||
- None.
|
|
||||||
|
|
||||||
### Modified Capabilities
|
|
||||||
|
|
||||||
- `meeting-summary`: Require screenshot OCR to honor the configured agent streaming transport while preserving image input and OCR metadata parsing.
|
|
||||||
|
|
||||||
## Impact
|
|
||||||
|
|
||||||
The change affects the screenshot OCR client, the shared LiteLLM Responses message translation, focused tests, and agent configuration documentation. It does not change the local HTTP API or screenshot note format.
|
|
||||||
@@ -1,103 +0,0 @@
|
|||||||
## MODIFIED Requirements
|
|
||||||
|
|
||||||
### Requirement: Meeting screenshots are captured into assistant context
|
|
||||||
Meeting Assistant SHALL expose a configurable screenshot hotkey.
|
|
||||||
|
|
||||||
When a meeting is active and the screenshot hotkey is pressed, Meeting Assistant SHALL capture the currently active window.
|
|
||||||
|
|
||||||
The screenshot image SHALL be saved into a configurable attachments folder for the assistant context note. By default, the folder SHALL be `Attachments` beside the assistant context note.
|
|
||||||
|
|
||||||
After the image is saved, Meeting Assistant SHALL append a markdown image link to the assistant context note with a meeting-relative timestamp that correlates to transcript timestamps.
|
|
||||||
|
|
||||||
Meeting Assistant SHALL allow optional screenshot OCR configuration with endpoint URL, API key or key environment variable, model, prompt, and timeout.
|
|
||||||
|
|
||||||
When screenshot OCR is configured, Meeting Assistant SHALL send the screenshot and prompt to the configured OpenAI-compatible Responses endpoint and append the model result after the screenshot link in the assistant context note.
|
|
||||||
|
|
||||||
Screenshot OCR SHALL honor the configured `MeetingAssistant:Agent:UseStreaming` transport selection.
|
|
||||||
|
|
||||||
When streaming is enabled, screenshot OCR SHALL consume the Responses result through the supported OpenAI Responses and Agent Framework Server-Sent Events adapter.
|
|
||||||
|
|
||||||
When streaming is disabled, screenshot OCR SHALL consume the result through the supported non-streaming OpenAI Responses client and adapter.
|
|
||||||
|
|
||||||
The screenshot OCR prompt SHALL ask the model to return pixel crop coordinates when it can confidently isolate only the presentation, shared screen, or similarly relevant meeting content.
|
|
||||||
|
|
||||||
When OCR returns valid crop coordinates within the original image bounds, Meeting Assistant SHALL save a cropped PNG beside the original screenshot and SHALL link the cropped image before the OCR result in the assistant context note.
|
|
||||||
|
|
||||||
When OCR returns no crop coordinates or invalid crop coordinates, Meeting Assistant SHALL keep the original screenshot link and OCR result without writing a cropped image.
|
|
||||||
|
|
||||||
After transcription finishes and before summarization starts, Meeting Assistant SHALL scan the meeting note for user-authored Obsidian image embeds and Markdown image embeds.
|
|
||||||
|
|
||||||
When configured screenshot OCR is enabled and the meeting note contains image embeds, Meeting Assistant SHALL append each resolvable image to the assistant context note, state that the image came from the meeting note, preserve the original embed text for cross-reference, and run OCR for the linked image.
|
|
||||||
|
|
||||||
Meeting-note image OCR SHALL NOT copy the image file, SHALL NOT write crop images, SHALL NOT add attendees from OCR metadata, and SHALL NOT modify the meeting note.
|
|
||||||
|
|
||||||
Meeting Assistant SHALL wait for meeting-note image OCR to finish or time out before transitioning the assistant context to summarizing.
|
|
||||||
|
|
||||||
When screenshot OCR fails or times out, Meeting Assistant SHALL write the failure status into the assistant context note with a retry link for that exact screenshot.
|
|
||||||
|
|
||||||
When the screenshot OCR retry link is activated, Meeting Assistant SHALL rerun OCR for the saved screenshot and replace that screenshot's OCR block in the assistant context note.
|
|
||||||
|
|
||||||
When screenshot OCR is not configured, Meeting Assistant SHALL skip OCR and keep the screenshot link.
|
|
||||||
|
|
||||||
The default OCR prompt SHALL explain that the image is from a meeting and ask the model to identify who is talking, who is presenting, what is presented, capture slide text in markdown, convert diagrams to Mermaid when possible, indicate whether visible people are clearly the exact meeting participants or only a partial result, return crop coordinates only for confidently isolated presentation/shared-screen content, and otherwise describe the scene.
|
|
||||||
|
|
||||||
#### Scenario: Screenshot is linked with meeting timestamp
|
|
||||||
- **GIVEN** a meeting started at `10:00:00`
|
|
||||||
- **WHEN** the user captures a screenshot at `10:03:05`
|
|
||||||
- **THEN** Meeting Assistant saves the screenshot under the configured attachments folder
|
|
||||||
- **AND** appends a markdown image link to assistant context with timestamp `[00:03:05]`
|
|
||||||
|
|
||||||
#### Scenario: OCR result is appended after screenshot
|
|
||||||
- **GIVEN** screenshot OCR is configured
|
|
||||||
- **WHEN** the user captures a screenshot
|
|
||||||
- **THEN** Meeting Assistant appends the screenshot link to assistant context
|
|
||||||
- **AND** appends the OCR result for that screenshot after the link when processing completes
|
|
||||||
|
|
||||||
#### Scenario: Streaming screenshot OCR is assembled
|
|
||||||
- **GIVEN** screenshot OCR is configured
|
|
||||||
- **AND** `MeetingAssistant:Agent:UseStreaming` is `true`
|
|
||||||
- **WHEN** the Responses endpoint returns screenshot OCR output as Server-Sent Events
|
|
||||||
- **THEN** Meeting Assistant sends the prompt and screenshot as one multimodal Responses message
|
|
||||||
- **AND** appends the assembled OCR text without a JSON document parse failure
|
|
||||||
|
|
||||||
#### Scenario: Screenshot OCR streaming can be disabled
|
|
||||||
- **GIVEN** screenshot OCR is configured
|
|
||||||
- **AND** `MeetingAssistant:Agent:UseStreaming` is `false`
|
|
||||||
- **WHEN** Meeting Assistant requests screenshot OCR
|
|
||||||
- **THEN** it uses the supported non-streaming Responses client and adapter
|
|
||||||
- **AND** preserves the prompt and screenshot image input
|
|
||||||
|
|
||||||
#### Scenario: OCR crop is saved and linked before OCR text
|
|
||||||
- **GIVEN** screenshot OCR is configured
|
|
||||||
- **AND** OCR returns valid crop coordinates for a shared screen
|
|
||||||
- **WHEN** OCR processing completes
|
|
||||||
- **THEN** Meeting Assistant saves a cropped screenshot beside the original image
|
|
||||||
- **AND** links the cropped screenshot before the OCR text in assistant context
|
|
||||||
|
|
||||||
#### Scenario: Meeting note image embeds are OCRed before summarization
|
|
||||||
- **GIVEN** screenshot OCR is configured
|
|
||||||
- **AND** the meeting note contains `![[whiteboard.png]]`
|
|
||||||
- **AND** the meeting note contains ``
|
|
||||||
- **WHEN** transcription finishes
|
|
||||||
- **THEN** Meeting Assistant appends both images to the assistant context as images from the meeting note
|
|
||||||
- **AND** includes the original embed text for each image
|
|
||||||
- **AND** runs OCR for each image without copying files, writing crop images, adding attendees, or modifying the meeting note
|
|
||||||
- **AND** waits for this OCR to finish or time out before transitioning to summarizing
|
|
||||||
|
|
||||||
#### Scenario: OCR failure can be retried for the same screenshot
|
|
||||||
- **GIVEN** screenshot OCR is configured
|
|
||||||
- **AND** OCR fails or times out for a captured screenshot
|
|
||||||
- **WHEN** Meeting Assistant writes the OCR failure status
|
|
||||||
- **THEN** the assistant context includes a retry link for that exact screenshot
|
|
||||||
- **WHEN** the retry link is activated
|
|
||||||
- **THEN** Meeting Assistant reruns OCR against the saved screenshot
|
|
||||||
- **AND** replaces that screenshot's OCR block with the retry result
|
|
||||||
|
|
||||||
#### Scenario: OCR is skipped when not configured
|
|
||||||
- **GIVEN** screenshot OCR is not configured
|
|
||||||
- **WHEN** the user captures a screenshot
|
|
||||||
- **THEN** Meeting Assistant saves and links the screenshot without calling a model endpoint
|
|
||||||
|
|
||||||
#### Scenario: OCR reports whether visible people are complete or partial
|
|
||||||
- **WHEN** Meeting Assistant uses the built-in screenshot OCR prompt
|
|
||||||
- **THEN** the prompt asks the model to state whether the screenshot clearly shows exactly who is in the meeting or only a partial participant result
|
|
||||||
@@ -1,15 +0,0 @@
|
|||||||
## 1. Streaming screenshot OCR
|
|
||||||
|
|
||||||
- [x] 1.1 Add a failing screenshot-client behavior test for an SSE response with prompt and image input.
|
|
||||||
- [x] 1.2 Route screenshot OCR through the shared Responses client and add multimodal message translation.
|
|
||||||
|
|
||||||
## 2. Non-streaming compatibility
|
|
||||||
|
|
||||||
- [x] 2.1 Add behavior coverage proving `Agent:UseStreaming=false` preserves non-streaming screenshot OCR and image input.
|
|
||||||
- [x] 2.2 Document that screenshot OCR inherits the agent streaming setting.
|
|
||||||
|
|
||||||
## 3. Verification
|
|
||||||
|
|
||||||
- [x] 3.1 Refactor the touched screenshot and shared client paths for DRYness, SOLID boundaries, and simplicity while preserving behavior.
|
|
||||||
- [x] 3.2 Run focused tests, the full solution tests, and strict OpenSpec validation.
|
|
||||||
- [x] 3.3 Restart Meeting Assistant only while idle and verify screenshot OCR against the deployed LiteLLM endpoint.
|
|
||||||
@@ -1,2 +0,0 @@
|
|||||||
schema: spec-driven
|
|
||||||
created: 2026-08-04
|
|
||||||
@@ -1,45 +0,0 @@
|
|||||||
## Context
|
|
||||||
|
|
||||||
The tray-menu builder currently returns a flat list of semantic actions, while the Windows renderer infers separators from item indexes and the Exit action. During an active recording, the normal stop action is added after the microphone submenu and uses a long implementation-oriented label. This makes the primary meeting-completion action look equivalent to cancel, profile switching, and device selection.
|
|
||||||
|
|
||||||
## Goals / Non-Goals
|
|
||||||
|
|
||||||
**Goals:**
|
|
||||||
|
|
||||||
- Give normal meeting completion the concise label `Finish meeting`.
|
|
||||||
- Make that action the only item in the section immediately below `Open agent` while recording.
|
|
||||||
- Keep fine-grained recording controls in a distinct following section.
|
|
||||||
- Make section boundaries observable in platform-independent menu behavior tests.
|
|
||||||
|
|
||||||
**Non-Goals:**
|
|
||||||
|
|
||||||
- Change what normal stop, abort, profile switching, or microphone selection does.
|
|
||||||
- Change idle-menu actions, hotkeys, endpoints, or recording state transitions.
|
|
||||||
- Add icons, confirmation prompts, or nested submenus.
|
|
||||||
|
|
||||||
## Decisions
|
|
||||||
|
|
||||||
### Represent section starts in the menu model
|
|
||||||
|
|
||||||
Add a section-start flag to `MeetingTaskbarMenuItem`. The Windows renderer will insert a separator before items carrying the flag instead of deriving layout from array indexes and action types.
|
|
||||||
|
|
||||||
This keeps layout intent in the platform-independent builder where behavior tests can observe it. Keeping another renderer-only special case was rejected because it would leave the requested prominence untestable without Windows UI automation.
|
|
||||||
|
|
||||||
### Build prioritized and fine-grained controls as separate groups
|
|
||||||
|
|
||||||
While recording, the builder will add `Open agent`, then `Finish meeting` as a new section, then collect microphone, cancel/discard, and profile-switch actions into a fine-grained group whose first item starts another section. Exit remains the final section.
|
|
||||||
|
|
||||||
The action continues to use the existing normal-stop command so transcription, speaker processing, OCR, and summarization semantics do not change.
|
|
||||||
|
|
||||||
## Risks / Trade-offs
|
|
||||||
|
|
||||||
- **A section flag could produce adjacent separators if assigned carelessly** → The builder marks only the first item of each non-empty group, and the renderer follows those explicit starts.
|
|
||||||
- **Menu ordering changes while recording** → Limit reordering to the active-recording state; idle and processing actions retain their existing relative order.
|
|
||||||
|
|
||||||
## Migration Plan
|
|
||||||
|
|
||||||
No configuration or data migration is required. Deploying the updated executable changes only tray-menu presentation. Rollback restores the previous label and grouping.
|
|
||||||
|
|
||||||
## Open Questions
|
|
||||||
|
|
||||||
None.
|
|
||||||
@@ -1,25 +0,0 @@
|
|||||||
## Why
|
|
||||||
|
|
||||||
The active-recording tray menu labels its most important completion action as the verbose `Stop meeting recording and transcribe` and groups it with rarely used controls. Finishing a meeting should be immediately recognizable and visually prioritized during normal use.
|
|
||||||
|
|
||||||
## What Changes
|
|
||||||
|
|
||||||
- Rename the active-recording stop action to `Finish meeting` without changing its normal stop, transcription, or summary behavior.
|
|
||||||
- Place `Finish meeting` by itself in the section immediately below `Open agent`.
|
|
||||||
- Place microphone selection, cancel/discard, and profile-switch controls in a separate lower-priority section.
|
|
||||||
- Represent tray-menu section boundaries explicitly so ordering and prominence are behavior-testable.
|
|
||||||
|
|
||||||
## Capabilities
|
|
||||||
|
|
||||||
### New Capabilities
|
|
||||||
|
|
||||||
None.
|
|
||||||
|
|
||||||
### Modified Capabilities
|
|
||||||
|
|
||||||
- `meeting-recording`: Prioritize the normal meeting completion action in the Windows tray menu with a concise label and dedicated section.
|
|
||||||
|
|
||||||
## Impact
|
|
||||||
|
|
||||||
- Affects the platform-independent tray-menu model/builder, Windows tray-menu rendering, and taskbar behavior tests.
|
|
||||||
- Does not change recording lifecycle semantics, hotkeys, endpoints, or generated meeting artifacts.
|
|
||||||
-62
@@ -1,62 +0,0 @@
|
|||||||
## MODIFIED Requirements
|
|
||||||
|
|
||||||
### Requirement: Windows taskbar icon controls recording
|
|
||||||
Meeting Assistant SHALL show a Windows taskbar notification icon when running on Windows.
|
|
||||||
|
|
||||||
The taskbar icon SHALL indicate whether the newest meeting process is idle, actively recording, or post-recording processing/summarizing.
|
|
||||||
|
|
||||||
When a new meeting is actively recording while an older stopped meeting is still transcribing, recognizing speakers, or summarizing, the taskbar icon SHALL show the new active recording state.
|
|
||||||
|
|
||||||
The taskbar icon right-click menu SHALL expose recording controls based on the current state and configured launch profiles.
|
|
||||||
|
|
||||||
The taskbar icon right-click menu SHALL expose an Exit action in every recording state.
|
|
||||||
|
|
||||||
When Meeting Assistant is idle or only processing older stopped meetings, the menu SHALL allow starting a meeting recording for each configured launch profile.
|
|
||||||
|
|
||||||
When a meeting is actively recording, the menu SHALL allow stopping the recording and continuing transcription/summary generation.
|
|
||||||
|
|
||||||
During an active recording, the normal stop action SHALL be labeled `Finish meeting` and SHALL be the only action in a dedicated menu section immediately below the `Open agent` section.
|
|
||||||
|
|
||||||
During an active recording, microphone selection, cancel/discard, and profile-switch actions SHALL appear in a separate fine-grained controls section below `Finish meeting`.
|
|
||||||
|
|
||||||
When a meeting is actively recording, the menu SHALL allow canceling the recording and discarding that run's artifacts.
|
|
||||||
|
|
||||||
When a meeting is actively recording, the menu SHALL allow switching to each configured launch profile other than the current active profile.
|
|
||||||
|
|
||||||
Selecting Exit while Meeting Assistant is idle SHALL stop the application without an additional confirmation prompt.
|
|
||||||
|
|
||||||
Selecting Exit while Meeting Assistant is recording, transcribing, recognizing speakers, or summarizing SHALL show a confirmation dialog before stopping the application.
|
|
||||||
|
|
||||||
#### Scenario: Idle tray menu can start configured profiles
|
|
||||||
- **GIVEN** launch profiles `default` and `english` are configured
|
|
||||||
- **AND** no meeting recording is active
|
|
||||||
- **WHEN** the taskbar menu is opened
|
|
||||||
- **THEN** it offers start recording actions for `default` and `english`
|
|
||||||
|
|
||||||
#### Scenario: Recording tray menu prioritizes finishing the meeting
|
|
||||||
- **GIVEN** launch profiles `default` and `english` are configured
|
|
||||||
- **AND** a meeting is actively recording with profile `default`
|
|
||||||
- **WHEN** the taskbar menu is opened
|
|
||||||
- **THEN** `Finish meeting` is the only action in the section immediately below `Open agent`
|
|
||||||
- **AND** microphone selection, cancel/discard, and switching to `english` appear in a separate following section
|
|
||||||
- **AND** the menu does not offer switching to `default`
|
|
||||||
|
|
||||||
#### Scenario: Active recording has priority over older summarizing runs
|
|
||||||
- **GIVEN** an older meeting is still summarizing
|
|
||||||
- **WHEN** a newer meeting is actively recording
|
|
||||||
- **THEN** the taskbar icon indicates recording
|
|
||||||
|
|
||||||
#### Scenario: Tray menu always exposes Exit
|
|
||||||
- **GIVEN** Meeting Assistant is running
|
|
||||||
- **WHEN** the taskbar menu is opened
|
|
||||||
- **THEN** it offers an Exit action
|
|
||||||
|
|
||||||
#### Scenario: Idle Exit stops immediately
|
|
||||||
- **GIVEN** no recording, transcription, speaker recognition, or summary work is running
|
|
||||||
- **WHEN** the user selects Exit from the taskbar menu
|
|
||||||
- **THEN** Meeting Assistant stops the application without an additional confirmation prompt
|
|
||||||
|
|
||||||
#### Scenario: In-progress Exit asks for confirmation
|
|
||||||
- **GIVEN** Meeting Assistant is recording, transcribing, recognizing speakers, or summarizing
|
|
||||||
- **WHEN** the user selects Exit from the taskbar menu
|
|
||||||
- **THEN** Meeting Assistant asks for confirmation before stopping the application
|
|
||||||
@@ -1,15 +0,0 @@
|
|||||||
## 1. Tray Menu Behavior
|
|
||||||
|
|
||||||
- [x] 1.1 Add a failing behavior test proving that an active recording labels the normal stop action `Finish meeting`, places it alone immediately below `Open agent`, and keeps fine-grained controls in the following section.
|
|
||||||
- [x] 1.2 Add explicit section metadata to the tray-menu model, reorder the active-recording actions, and render separators from that metadata.
|
|
||||||
|
|
||||||
## 2. Verification
|
|
||||||
|
|
||||||
- [x] 2.1 Review the touched menu builder and renderer for DRYness, SOLID design, and simplicity while preserving behavior.
|
|
||||||
- [x] 2.2 Run focused taskbar-menu tests, the Windows application build, the full solution tests, and strict OpenSpec validation.
|
|
||||||
|
|
||||||
## 3. Refactor Follow-up
|
|
||||||
|
|
||||||
- [x] 3.1 Lock down idle section boundaries and active-recording layout when no microphone is available.
|
|
||||||
- [x] 3.2 Remove the tray-menu section helper's hidden input mutation without changing rendered behavior.
|
|
||||||
- [x] 3.3 Run focused and full verification, then validate the OpenSpec change strictly.
|
|
||||||
Reference in New Issue
Block a user