2022-12-01 05:33:20 +00:00
|
|
|
# Copyright (c) 2022 Kyle Schouviller (https://github.com/kyle0654)
|
|
|
|
|
2023-11-21 02:57:10 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
from typing import TYPE_CHECKING, Optional
|
|
|
|
|
|
|
|
from invokeai.app.services.events.events_common import (
|
|
|
|
BatchEnqueuedEvent,
|
|
|
|
BulkDownloadCompleteEvent,
|
|
|
|
BulkDownloadErrorEvent,
|
|
|
|
BulkDownloadStartedEvent,
|
|
|
|
DownloadCancelledEvent,
|
|
|
|
DownloadCompleteEvent,
|
|
|
|
DownloadErrorEvent,
|
|
|
|
DownloadProgressEvent,
|
|
|
|
DownloadStartedEvent,
|
2024-03-10 12:23:11 +00:00
|
|
|
EventBase,
|
2024-03-14 08:04:19 +00:00
|
|
|
InvocationCompleteEvent,
|
|
|
|
InvocationDenoiseProgressEvent,
|
|
|
|
InvocationErrorEvent,
|
|
|
|
InvocationStartedEvent,
|
|
|
|
ModelInstallCancelledEvent,
|
|
|
|
ModelInstallCompleteEvent,
|
|
|
|
ModelInstallDownloadProgressEvent,
|
|
|
|
ModelInstallDownloadsCompleteEvent,
|
|
|
|
ModelInstallErrorEvent,
|
|
|
|
ModelInstallStartedEvent,
|
|
|
|
ModelLoadCompleteEvent,
|
|
|
|
ModelLoadStartedEvent,
|
|
|
|
QueueClearedEvent,
|
|
|
|
QueueItemStatusChangedEvent,
|
|
|
|
SessionCanceledEvent,
|
|
|
|
SessionCompleteEvent,
|
|
|
|
SessionStartedEvent,
|
2023-10-09 00:04:03 +00:00
|
|
|
)
|
feat(nodes,ui): fix soft locks on session/invocation retrieval
When a queue item is popped for processing, we need to retrieve its session from the DB. Pydantic serializes the graph at this stage.
It's possible for a graph to have been made invalid during the graph preparation stage (e.g. an ancestor node executes, and its output is not valid for its successor node's input field).
When this occurs, the session in the DB will fail validation, but we don't have a chance to find out until it is retrieved and parsed by pydantic.
This logic was previously not wrapped in any exception handling.
Just after retrieving a session, we retrieve the specific invocation to execute from the session. It's possible that this could also have some sort of error, though it should be impossible for it to be a pydantic validation error (that would have been caught during session validation). There was also no exception handling here.
When either of these processes fail, the processor gets soft-locked because the processor's cleanup logic is never run. (I didn't dig deeper into exactly what cleanup is not happening, because the fix is to just handle the exceptions.)
This PR adds exception handling to both the session retrieval and node retrieval and events for each: `session_retrieval_error` and `invocation_retrieval_error`.
These events are caught and displayed in the UI as toasts, along with the type of the python exception (e.g. `Validation Error`). The events are also logged to the browser console.
2023-07-23 02:27:59 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
if TYPE_CHECKING:
|
|
|
|
from invokeai.app.invocations.baseinvocation import BaseInvocation, BaseInvocationOutput
|
2024-03-10 12:23:11 +00:00
|
|
|
from invokeai.app.services.events.events_common import EventBase
|
2024-03-14 08:04:19 +00:00
|
|
|
from invokeai.app.services.model_install.model_install_common import ModelInstallJob
|
|
|
|
from invokeai.app.services.session_processor.session_processor_common import ProgressImage
|
|
|
|
from invokeai.app.services.session_queue.session_queue_common import (
|
|
|
|
BatchStatus,
|
|
|
|
EnqueueBatchResult,
|
|
|
|
SessionQueueItem,
|
|
|
|
SessionQueueStatus,
|
|
|
|
)
|
|
|
|
from invokeai.backend.model_manager.config import AnyModelConfig, SubModelType
|
2022-12-01 05:33:20 +00:00
|
|
|
|
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
class EventServiceBase:
|
2022-12-01 05:33:20 +00:00
|
|
|
"""Basic event bus, to have an empty stand-in when not needed"""
|
2023-03-03 06:02:00 +00:00
|
|
|
|
2024-03-10 12:23:11 +00:00
|
|
|
def dispatch(self, event: "EventBase") -> None:
|
2022-12-01 05:33:20 +00:00
|
|
|
pass
|
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
# region: Invocation
|
2023-12-22 17:35:57 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_invocation_started(self, queue_item: "SessionQueueItem", invocation: "BaseInvocation") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when an invocation is started"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(InvocationStartedEvent.build(queue_item, invocation))
|
2023-11-26 02:45:59 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_invocation_denoise_progress(
|
2023-03-03 06:02:00 +00:00
|
|
|
self,
|
2024-03-14 08:04:19 +00:00
|
|
|
queue_item: "SessionQueueItem",
|
|
|
|
invocation: "BaseInvocation",
|
2022-12-01 05:33:20 +00:00
|
|
|
step: int,
|
2023-03-15 12:50:26 +00:00
|
|
|
total_steps: int,
|
2024-03-14 08:04:19 +00:00
|
|
|
progress_image: "ProgressImage",
|
2022-12-01 05:33:20 +00:00
|
|
|
) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted at each step during denoising of an invocation."""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(InvocationDenoiseProgressEvent.build(queue_item, invocation, step, total_steps, progress_image))
|
2022-12-01 05:33:20 +00:00
|
|
|
|
2023-03-03 06:02:00 +00:00
|
|
|
def emit_invocation_complete(
|
2024-03-14 08:04:19 +00:00
|
|
|
self, queue_item: "SessionQueueItem", invocation: "BaseInvocation", output: "BaseInvocationOutput"
|
2022-12-01 05:33:20 +00:00
|
|
|
) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when an invocation is complete"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(InvocationCompleteEvent.build(queue_item, invocation, output))
|
2023-03-03 06:02:00 +00:00
|
|
|
|
|
|
|
def emit_invocation_error(
|
2024-03-14 08:04:19 +00:00
|
|
|
self, queue_item: "SessionQueueItem", invocation: "BaseInvocation", error_type: str, error: str
|
2023-02-27 18:01:07 +00:00
|
|
|
) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when an invocation encounters an error"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(InvocationErrorEvent.build(queue_item, invocation, error_type, error))
|
2022-12-01 05:33:20 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
# endregion
|
2022-12-01 05:33:20 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
# region Session
|
2023-05-11 04:09:19 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_session_started(self, queue_item: "SessionQueueItem") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a session has started"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(SessionStartedEvent.build(queue_item))
|
2023-05-11 04:09:19 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_session_complete(self, queue_item: "SessionQueueItem") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a session has completed all invocations"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(SessionCompleteEvent.build(queue_item))
|
feat(nodes,ui): fix soft locks on session/invocation retrieval
When a queue item is popped for processing, we need to retrieve its session from the DB. Pydantic serializes the graph at this stage.
It's possible for a graph to have been made invalid during the graph preparation stage (e.g. an ancestor node executes, and its output is not valid for its successor node's input field).
When this occurs, the session in the DB will fail validation, but we don't have a chance to find out until it is retrieved and parsed by pydantic.
This logic was previously not wrapped in any exception handling.
Just after retrieving a session, we retrieve the specific invocation to execute from the session. It's possible that this could also have some sort of error, though it should be impossible for it to be a pydantic validation error (that would have been caught during session validation). There was also no exception handling here.
When either of these processes fail, the processor gets soft-locked because the processor's cleanup logic is never run. (I didn't dig deeper into exactly what cleanup is not happening, because the fix is to just handle the exceptions.)
This PR adds exception handling to both the session retrieval and node retrieval and events for each: `session_retrieval_error` and `invocation_retrieval_error`.
These events are caught and displayed in the UI as toasts, along with the type of the python exception (e.g. `Validation Error`). The events are also logged to the browser console.
2023-07-23 02:27:59 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_session_canceled(self, queue_item: "SessionQueueItem") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a session is canceled"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(SessionCanceledEvent.build(queue_item))
|
|
|
|
|
|
|
|
# endregion
|
|
|
|
|
|
|
|
# region Queue
|
feat: queued generation (#4502)
* fix(config): fix typing issues in `config/`
`config/invokeai_config.py`:
- use `Optional` for things that are optional
- fix typing of `ram_cache_size()` and `vram_cache_size()`
- remove unused and incorrectly typed method `autoconvert_path`
- fix types and logic for `parse_args()`, in which `InvokeAIAppConfig.initconf` *must* be a `DictConfig`, but function would allow it to be set as a `ListConfig`, which presumably would cause issues elsewhere
`config/base.py`:
- use `cls` for first arg of class methods
- use `Optional` for things that are optional
- fix minor type issue related to setting of `env_prefix`
- remove unused `add_subparser()` method, which calls `add_parser()` on an `ArgumentParser` (method only available on the `_SubParsersAction` object, which is returned from ArgumentParser.add_subparsers()`)
* feat: queued generation and batches
Due to a very messy branch with broad addition of `isort` on `main` alongside it, some git surgery was needed to get an agreeable git history. This commit represents all of the work on queued generation. See PR for notes.
* chore: flake8, isort, black
* fix(nodes): fix incorrect service stop() method
* fix(nodes): improve names of a few variables
* fix(tests): fix up tests after changes to batches/queue
* feat(tests): add unit tests for session queue helper functions
* feat(ui): dynamic prompts is always enabled
* feat(queue): add queue_status_changed event
* feat(ui): wip queue graphs
* feat(nodes): move cleanup til after invoker startup
* feat(nodes): add cancel_by_batch_ids
* feat(ui): wip batch graphs & UI
* fix(nodes): remove `Batch.batch_id` from required
* fix(ui): cleanup and use fixedCacheKey for all mutations
* fix(ui): remove orphaned nodes from canvas graphs
* fix(nodes): fix cancel_by_batch_ids result count
* fix(ui): only show cancel batch tooltip when batches were canceled
* chore: isort
* fix(api): return `[""]` when dynamic prompts generates no prompts
Just a simple fallback so we always have a prompt.
* feat(ui): dynamicPrompts.combinatorial is always on
There seems to be little purpose in using the combinatorial generation for dynamic prompts. I've disabled it by hiding it from the UI and defaulting combinatorial to true. If we want to enable it again in the future it's straightforward to do so.
* feat: add queue_id & support logic
* feat(ui): fix upscale button
It prepends the upscale operation to queue
* feat(nodes): return queue item when enqueuing a single graph
This facilitates one-off graph async workflows in the client.
* feat(ui): move controlnet autoprocess to queue
* fix(ui): fix non-serializable DOMRect in redux state
* feat(ui): QueueTable performance tweaks
* feat(ui): update queue list
Queue items expand to show the full queue item. Just as JSON for now.
* wip threaded session_processor
* feat(nodes,ui): fully migrate queue to session_processor
* feat(nodes,ui): add processor events
* feat(ui): ui tweaks
* feat(nodes,ui): consolidate events, reduce network requests
* feat(ui): cleanup & abstract queue hooks
* feat(nodes): optimize batch permutation
Use a generator to do only as much work as is needed.
Previously, though we only ended up creating exactly as many queue items as was needed, there was still some intermediary work that calculated *all* permutations. When that number was very high, the system had a very hard time and used a lot of memory.
The logic has been refactored to use a generator. Additionally, the batch validators are optimized to return early and use less memory.
* feat(ui): add seed behaviour parameter
This dynamic prompts parameter allows the seed to be randomized per prompt or per iteration:
- Per iteration: Use the same seed for all prompts in a single dynamic prompt expansion
- Per prompt: Use a different seed for every single prompt
"Per iteration" is appropriate for exploring a the latents space with a stable starting noise, while "Per prompt" provides more variation.
* fix(ui): remove extraneous random seed nodes from linear graphs
* fix(ui): fix controlnet autoprocess not working when queue is running
* feat(queue): add timestamps to queue status updates
Also show execution time in queue list
* feat(queue): change all execution-related events to use the `queue_id` as the room, also include `queue_item_id` in InvocationQueueItem
This allows for much simpler handling of queue items.
* feat(api): deprecate sessions router
* chore(backend): tidy logging in `dependencies.py`
* fix(backend): respect `use_memory_db`
* feat(backend): add `config.log_sql` (enables sql trace logging)
* feat: add invocation cache
Supersedes #4574
The invocation cache provides simple node memoization functionality. Nodes that use the cache are memoized and not re-executed if their inputs haven't changed. Instead, the stored output is returned.
## Results
This feature provides anywhere some significant to massive performance improvement.
The improvement is most marked on large batches of generations where you only change a couple things (e.g. different seed or prompt for each iteration) and low-VRAM systems, where skipping an extraneous model load is a big deal.
## Overview
A new `invocation_cache` service is added to handle the caching. There's not much to it.
All nodes now inherit a boolean `use_cache` field from `BaseInvocation`. This is a node field and not a class attribute, because specific instances of nodes may want to opt in or out of caching.
The recently-added `invoke_internal()` method on `BaseInvocation` is used as an entrypoint for the cache logic.
To create a cache key, the invocation is first serialized using pydantic's provided `json()` method, skipping the unique `id` field. Then python's very fast builtin `hash()` is used to create an integer key. All implementations of `InvocationCacheBase` must provide a class method `create_key()` which accepts an invocation and outputs a string or integer key.
## In-Memory Implementation
An in-memory implementation is provided. In this implementation, the node outputs are stored in memory as python classes. The in-memory cache does not persist application restarts.
Max node cache size is added as `node_cache_size` under the `Generation` config category.
It defaults to 512 - this number is up for discussion, but given that these are relatively lightweight pydantic models, I think it's safe to up this even higher.
Note that the cache isn't storing the big stuff - tensors and images are store on disk, and outputs include only references to them.
## Node Definition
The default for all nodes is to use the cache. The `@invocation` decorator now accepts an optional `use_cache: bool` argument to override the default of `True`.
Non-deterministic nodes, however, should set this to `False`. Currently, all random-stuff nodes, including `dynamic_prompt`, are set to `False`.
The field name `use_cache` is now effectively a reserved field name and possibly a breaking change if any community nodes use this as a field name. In hindsight, all our reserved field names should have been prefixed with underscores or something.
## One Gotcha
Leaf nodes probably want to opt out of the cache, because if they are not cached, their outputs are not saved again.
If you run the same graph multiple times, you only end up with a single image output, because the image storage side-effects are in the `invoke()` method, which is bypassed if we have a cache hit.
## Linear UI
The linear graphs _almost_ just work, but due to the gotcha, we need to be careful about the final image-outputting node. To resolve this, a `SaveImageInvocation` node is added and used in the linear graphs.
This node is similar to `ImagePrimitive`, except it saves a copy of its input image, and has `use_cache` set to `False` by default.
This is now the leaf node in all linear graphs, and is the only node in those graphs with `use_cache == False` _and_ the only node with `is_intermedate == False`.
## Workflow Editor
All nodes now have a footer with a new `Use Cache [ ]` checkbox. It defaults to the value set by the invocation in its python definition, but can be changed by the user.
The workflow/node validation logic has been updated to migrate old workflows to use the new default values for `use_cache`. Users may still want to review the settings that have been chosen. In the event of catastrophic failure when running this migration, the default value of `True` is applied, as this is correct for most nodes.
Users should consider saving their workflows after loading them in and having them updated.
## Future Enhancements - Callback
A future enhancement would be to provide a callback to the `use_cache` flag that would be run as the node is executed to determine, based on its own internal state, if the cache should be used or not.
This would be useful for `DynamicPromptInvocation`, where the deterministic behaviour is determined by the `combinatorial: bool` field.
## Future Enhancements - Persisted Cache
Similar to how the latents storage is backed by disk, the invocation cache could be persisted to the database or disk. We'd need to be very careful about deserializing outputs, but it's perhaps worth exploring in the future.
* fix(ui): fix queue list item width
* feat(nodes): do not send the whole node on every generator progress
* feat(ui): strip out old logic related to sessions
Things like `isProcessing` are no longer relevant with queue. Removed them all & updated everything be appropriate for queue. May be a few little quirks I've missed...
* feat(ui): fix up param collapse labels
* feat(ui): click queue count to go to queue tab
* tidy(queue): update comment, query format
* feat(ui): fix progress bar when canceling
* fix(ui): fix circular dependency
* feat(nodes): bail on node caching logic if `node_cache_size == 0`
* feat(nodes): handle KeyError on node cache pop
* feat(nodes): bypass cache codepath if caches is disabled
more better no do thing
* fix(ui): reset api cache on connect/disconnect
* feat(ui): prevent enqueue when no prompts generated
* feat(ui): add queue controls to workflow editor
* feat(ui): update floating buttons & other incidental UI tweaks
* fix(ui): fix missing/incorrect translation keys
* fix(tests): add config service to mock invocation services
invoking needs access to `node_cache_size` to occur
* optionally remove pause/resume buttons from queue UI
* option to disable prepending
* chore(ui): remove unused file
* feat(queue): remove `order_id` entirely, `item_id` is now an autoinc pk
---------
Co-authored-by: Mary Hipp <maryhipp@Marys-MacBook-Air.local>
2023-09-20 05:09:24 +00:00
|
|
|
|
2023-10-09 00:04:03 +00:00
|
|
|
def emit_queue_item_status_changed(
|
2024-03-14 08:04:19 +00:00
|
|
|
self, queue_item: "SessionQueueItem", batch_status: "BatchStatus", queue_status: "SessionQueueStatus"
|
2023-10-09 00:04:03 +00:00
|
|
|
) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a queue item's status changes"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(QueueItemStatusChangedEvent.build(queue_item, batch_status, queue_status))
|
feat: queued generation (#4502)
* fix(config): fix typing issues in `config/`
`config/invokeai_config.py`:
- use `Optional` for things that are optional
- fix typing of `ram_cache_size()` and `vram_cache_size()`
- remove unused and incorrectly typed method `autoconvert_path`
- fix types and logic for `parse_args()`, in which `InvokeAIAppConfig.initconf` *must* be a `DictConfig`, but function would allow it to be set as a `ListConfig`, which presumably would cause issues elsewhere
`config/base.py`:
- use `cls` for first arg of class methods
- use `Optional` for things that are optional
- fix minor type issue related to setting of `env_prefix`
- remove unused `add_subparser()` method, which calls `add_parser()` on an `ArgumentParser` (method only available on the `_SubParsersAction` object, which is returned from ArgumentParser.add_subparsers()`)
* feat: queued generation and batches
Due to a very messy branch with broad addition of `isort` on `main` alongside it, some git surgery was needed to get an agreeable git history. This commit represents all of the work on queued generation. See PR for notes.
* chore: flake8, isort, black
* fix(nodes): fix incorrect service stop() method
* fix(nodes): improve names of a few variables
* fix(tests): fix up tests after changes to batches/queue
* feat(tests): add unit tests for session queue helper functions
* feat(ui): dynamic prompts is always enabled
* feat(queue): add queue_status_changed event
* feat(ui): wip queue graphs
* feat(nodes): move cleanup til after invoker startup
* feat(nodes): add cancel_by_batch_ids
* feat(ui): wip batch graphs & UI
* fix(nodes): remove `Batch.batch_id` from required
* fix(ui): cleanup and use fixedCacheKey for all mutations
* fix(ui): remove orphaned nodes from canvas graphs
* fix(nodes): fix cancel_by_batch_ids result count
* fix(ui): only show cancel batch tooltip when batches were canceled
* chore: isort
* fix(api): return `[""]` when dynamic prompts generates no prompts
Just a simple fallback so we always have a prompt.
* feat(ui): dynamicPrompts.combinatorial is always on
There seems to be little purpose in using the combinatorial generation for dynamic prompts. I've disabled it by hiding it from the UI and defaulting combinatorial to true. If we want to enable it again in the future it's straightforward to do so.
* feat: add queue_id & support logic
* feat(ui): fix upscale button
It prepends the upscale operation to queue
* feat(nodes): return queue item when enqueuing a single graph
This facilitates one-off graph async workflows in the client.
* feat(ui): move controlnet autoprocess to queue
* fix(ui): fix non-serializable DOMRect in redux state
* feat(ui): QueueTable performance tweaks
* feat(ui): update queue list
Queue items expand to show the full queue item. Just as JSON for now.
* wip threaded session_processor
* feat(nodes,ui): fully migrate queue to session_processor
* feat(nodes,ui): add processor events
* feat(ui): ui tweaks
* feat(nodes,ui): consolidate events, reduce network requests
* feat(ui): cleanup & abstract queue hooks
* feat(nodes): optimize batch permutation
Use a generator to do only as much work as is needed.
Previously, though we only ended up creating exactly as many queue items as was needed, there was still some intermediary work that calculated *all* permutations. When that number was very high, the system had a very hard time and used a lot of memory.
The logic has been refactored to use a generator. Additionally, the batch validators are optimized to return early and use less memory.
* feat(ui): add seed behaviour parameter
This dynamic prompts parameter allows the seed to be randomized per prompt or per iteration:
- Per iteration: Use the same seed for all prompts in a single dynamic prompt expansion
- Per prompt: Use a different seed for every single prompt
"Per iteration" is appropriate for exploring a the latents space with a stable starting noise, while "Per prompt" provides more variation.
* fix(ui): remove extraneous random seed nodes from linear graphs
* fix(ui): fix controlnet autoprocess not working when queue is running
* feat(queue): add timestamps to queue status updates
Also show execution time in queue list
* feat(queue): change all execution-related events to use the `queue_id` as the room, also include `queue_item_id` in InvocationQueueItem
This allows for much simpler handling of queue items.
* feat(api): deprecate sessions router
* chore(backend): tidy logging in `dependencies.py`
* fix(backend): respect `use_memory_db`
* feat(backend): add `config.log_sql` (enables sql trace logging)
* feat: add invocation cache
Supersedes #4574
The invocation cache provides simple node memoization functionality. Nodes that use the cache are memoized and not re-executed if their inputs haven't changed. Instead, the stored output is returned.
## Results
This feature provides anywhere some significant to massive performance improvement.
The improvement is most marked on large batches of generations where you only change a couple things (e.g. different seed or prompt for each iteration) and low-VRAM systems, where skipping an extraneous model load is a big deal.
## Overview
A new `invocation_cache` service is added to handle the caching. There's not much to it.
All nodes now inherit a boolean `use_cache` field from `BaseInvocation`. This is a node field and not a class attribute, because specific instances of nodes may want to opt in or out of caching.
The recently-added `invoke_internal()` method on `BaseInvocation` is used as an entrypoint for the cache logic.
To create a cache key, the invocation is first serialized using pydantic's provided `json()` method, skipping the unique `id` field. Then python's very fast builtin `hash()` is used to create an integer key. All implementations of `InvocationCacheBase` must provide a class method `create_key()` which accepts an invocation and outputs a string or integer key.
## In-Memory Implementation
An in-memory implementation is provided. In this implementation, the node outputs are stored in memory as python classes. The in-memory cache does not persist application restarts.
Max node cache size is added as `node_cache_size` under the `Generation` config category.
It defaults to 512 - this number is up for discussion, but given that these are relatively lightweight pydantic models, I think it's safe to up this even higher.
Note that the cache isn't storing the big stuff - tensors and images are store on disk, and outputs include only references to them.
## Node Definition
The default for all nodes is to use the cache. The `@invocation` decorator now accepts an optional `use_cache: bool` argument to override the default of `True`.
Non-deterministic nodes, however, should set this to `False`. Currently, all random-stuff nodes, including `dynamic_prompt`, are set to `False`.
The field name `use_cache` is now effectively a reserved field name and possibly a breaking change if any community nodes use this as a field name. In hindsight, all our reserved field names should have been prefixed with underscores or something.
## One Gotcha
Leaf nodes probably want to opt out of the cache, because if they are not cached, their outputs are not saved again.
If you run the same graph multiple times, you only end up with a single image output, because the image storage side-effects are in the `invoke()` method, which is bypassed if we have a cache hit.
## Linear UI
The linear graphs _almost_ just work, but due to the gotcha, we need to be careful about the final image-outputting node. To resolve this, a `SaveImageInvocation` node is added and used in the linear graphs.
This node is similar to `ImagePrimitive`, except it saves a copy of its input image, and has `use_cache` set to `False` by default.
This is now the leaf node in all linear graphs, and is the only node in those graphs with `use_cache == False` _and_ the only node with `is_intermedate == False`.
## Workflow Editor
All nodes now have a footer with a new `Use Cache [ ]` checkbox. It defaults to the value set by the invocation in its python definition, but can be changed by the user.
The workflow/node validation logic has been updated to migrate old workflows to use the new default values for `use_cache`. Users may still want to review the settings that have been chosen. In the event of catastrophic failure when running this migration, the default value of `True` is applied, as this is correct for most nodes.
Users should consider saving their workflows after loading them in and having them updated.
## Future Enhancements - Callback
A future enhancement would be to provide a callback to the `use_cache` flag that would be run as the node is executed to determine, based on its own internal state, if the cache should be used or not.
This would be useful for `DynamicPromptInvocation`, where the deterministic behaviour is determined by the `combinatorial: bool` field.
## Future Enhancements - Persisted Cache
Similar to how the latents storage is backed by disk, the invocation cache could be persisted to the database or disk. We'd need to be very careful about deserializing outputs, but it's perhaps worth exploring in the future.
* fix(ui): fix queue list item width
* feat(nodes): do not send the whole node on every generator progress
* feat(ui): strip out old logic related to sessions
Things like `isProcessing` are no longer relevant with queue. Removed them all & updated everything be appropriate for queue. May be a few little quirks I've missed...
* feat(ui): fix up param collapse labels
* feat(ui): click queue count to go to queue tab
* tidy(queue): update comment, query format
* feat(ui): fix progress bar when canceling
* fix(ui): fix circular dependency
* feat(nodes): bail on node caching logic if `node_cache_size == 0`
* feat(nodes): handle KeyError on node cache pop
* feat(nodes): bypass cache codepath if caches is disabled
more better no do thing
* fix(ui): reset api cache on connect/disconnect
* feat(ui): prevent enqueue when no prompts generated
* feat(ui): add queue controls to workflow editor
* feat(ui): update floating buttons & other incidental UI tweaks
* fix(ui): fix missing/incorrect translation keys
* fix(tests): add config service to mock invocation services
invoking needs access to `node_cache_size` to occur
* optionally remove pause/resume buttons from queue UI
* option to disable prepending
* chore(ui): remove unused file
* feat(queue): remove `order_id` entirely, `item_id` is now an autoinc pk
---------
Co-authored-by: Mary Hipp <maryhipp@Marys-MacBook-Air.local>
2023-09-20 05:09:24 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_batch_enqueued(self, enqueue_result: "EnqueueBatchResult") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a batch is enqueued"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(BatchEnqueuedEvent.build(enqueue_result))
|
feat: queued generation (#4502)
* fix(config): fix typing issues in `config/`
`config/invokeai_config.py`:
- use `Optional` for things that are optional
- fix typing of `ram_cache_size()` and `vram_cache_size()`
- remove unused and incorrectly typed method `autoconvert_path`
- fix types and logic for `parse_args()`, in which `InvokeAIAppConfig.initconf` *must* be a `DictConfig`, but function would allow it to be set as a `ListConfig`, which presumably would cause issues elsewhere
`config/base.py`:
- use `cls` for first arg of class methods
- use `Optional` for things that are optional
- fix minor type issue related to setting of `env_prefix`
- remove unused `add_subparser()` method, which calls `add_parser()` on an `ArgumentParser` (method only available on the `_SubParsersAction` object, which is returned from ArgumentParser.add_subparsers()`)
* feat: queued generation and batches
Due to a very messy branch with broad addition of `isort` on `main` alongside it, some git surgery was needed to get an agreeable git history. This commit represents all of the work on queued generation. See PR for notes.
* chore: flake8, isort, black
* fix(nodes): fix incorrect service stop() method
* fix(nodes): improve names of a few variables
* fix(tests): fix up tests after changes to batches/queue
* feat(tests): add unit tests for session queue helper functions
* feat(ui): dynamic prompts is always enabled
* feat(queue): add queue_status_changed event
* feat(ui): wip queue graphs
* feat(nodes): move cleanup til after invoker startup
* feat(nodes): add cancel_by_batch_ids
* feat(ui): wip batch graphs & UI
* fix(nodes): remove `Batch.batch_id` from required
* fix(ui): cleanup and use fixedCacheKey for all mutations
* fix(ui): remove orphaned nodes from canvas graphs
* fix(nodes): fix cancel_by_batch_ids result count
* fix(ui): only show cancel batch tooltip when batches were canceled
* chore: isort
* fix(api): return `[""]` when dynamic prompts generates no prompts
Just a simple fallback so we always have a prompt.
* feat(ui): dynamicPrompts.combinatorial is always on
There seems to be little purpose in using the combinatorial generation for dynamic prompts. I've disabled it by hiding it from the UI and defaulting combinatorial to true. If we want to enable it again in the future it's straightforward to do so.
* feat: add queue_id & support logic
* feat(ui): fix upscale button
It prepends the upscale operation to queue
* feat(nodes): return queue item when enqueuing a single graph
This facilitates one-off graph async workflows in the client.
* feat(ui): move controlnet autoprocess to queue
* fix(ui): fix non-serializable DOMRect in redux state
* feat(ui): QueueTable performance tweaks
* feat(ui): update queue list
Queue items expand to show the full queue item. Just as JSON for now.
* wip threaded session_processor
* feat(nodes,ui): fully migrate queue to session_processor
* feat(nodes,ui): add processor events
* feat(ui): ui tweaks
* feat(nodes,ui): consolidate events, reduce network requests
* feat(ui): cleanup & abstract queue hooks
* feat(nodes): optimize batch permutation
Use a generator to do only as much work as is needed.
Previously, though we only ended up creating exactly as many queue items as was needed, there was still some intermediary work that calculated *all* permutations. When that number was very high, the system had a very hard time and used a lot of memory.
The logic has been refactored to use a generator. Additionally, the batch validators are optimized to return early and use less memory.
* feat(ui): add seed behaviour parameter
This dynamic prompts parameter allows the seed to be randomized per prompt or per iteration:
- Per iteration: Use the same seed for all prompts in a single dynamic prompt expansion
- Per prompt: Use a different seed for every single prompt
"Per iteration" is appropriate for exploring a the latents space with a stable starting noise, while "Per prompt" provides more variation.
* fix(ui): remove extraneous random seed nodes from linear graphs
* fix(ui): fix controlnet autoprocess not working when queue is running
* feat(queue): add timestamps to queue status updates
Also show execution time in queue list
* feat(queue): change all execution-related events to use the `queue_id` as the room, also include `queue_item_id` in InvocationQueueItem
This allows for much simpler handling of queue items.
* feat(api): deprecate sessions router
* chore(backend): tidy logging in `dependencies.py`
* fix(backend): respect `use_memory_db`
* feat(backend): add `config.log_sql` (enables sql trace logging)
* feat: add invocation cache
Supersedes #4574
The invocation cache provides simple node memoization functionality. Nodes that use the cache are memoized and not re-executed if their inputs haven't changed. Instead, the stored output is returned.
## Results
This feature provides anywhere some significant to massive performance improvement.
The improvement is most marked on large batches of generations where you only change a couple things (e.g. different seed or prompt for each iteration) and low-VRAM systems, where skipping an extraneous model load is a big deal.
## Overview
A new `invocation_cache` service is added to handle the caching. There's not much to it.
All nodes now inherit a boolean `use_cache` field from `BaseInvocation`. This is a node field and not a class attribute, because specific instances of nodes may want to opt in or out of caching.
The recently-added `invoke_internal()` method on `BaseInvocation` is used as an entrypoint for the cache logic.
To create a cache key, the invocation is first serialized using pydantic's provided `json()` method, skipping the unique `id` field. Then python's very fast builtin `hash()` is used to create an integer key. All implementations of `InvocationCacheBase` must provide a class method `create_key()` which accepts an invocation and outputs a string or integer key.
## In-Memory Implementation
An in-memory implementation is provided. In this implementation, the node outputs are stored in memory as python classes. The in-memory cache does not persist application restarts.
Max node cache size is added as `node_cache_size` under the `Generation` config category.
It defaults to 512 - this number is up for discussion, but given that these are relatively lightweight pydantic models, I think it's safe to up this even higher.
Note that the cache isn't storing the big stuff - tensors and images are store on disk, and outputs include only references to them.
## Node Definition
The default for all nodes is to use the cache. The `@invocation` decorator now accepts an optional `use_cache: bool` argument to override the default of `True`.
Non-deterministic nodes, however, should set this to `False`. Currently, all random-stuff nodes, including `dynamic_prompt`, are set to `False`.
The field name `use_cache` is now effectively a reserved field name and possibly a breaking change if any community nodes use this as a field name. In hindsight, all our reserved field names should have been prefixed with underscores or something.
## One Gotcha
Leaf nodes probably want to opt out of the cache, because if they are not cached, their outputs are not saved again.
If you run the same graph multiple times, you only end up with a single image output, because the image storage side-effects are in the `invoke()` method, which is bypassed if we have a cache hit.
## Linear UI
The linear graphs _almost_ just work, but due to the gotcha, we need to be careful about the final image-outputting node. To resolve this, a `SaveImageInvocation` node is added and used in the linear graphs.
This node is similar to `ImagePrimitive`, except it saves a copy of its input image, and has `use_cache` set to `False` by default.
This is now the leaf node in all linear graphs, and is the only node in those graphs with `use_cache == False` _and_ the only node with `is_intermedate == False`.
## Workflow Editor
All nodes now have a footer with a new `Use Cache [ ]` checkbox. It defaults to the value set by the invocation in its python definition, but can be changed by the user.
The workflow/node validation logic has been updated to migrate old workflows to use the new default values for `use_cache`. Users may still want to review the settings that have been chosen. In the event of catastrophic failure when running this migration, the default value of `True` is applied, as this is correct for most nodes.
Users should consider saving their workflows after loading them in and having them updated.
## Future Enhancements - Callback
A future enhancement would be to provide a callback to the `use_cache` flag that would be run as the node is executed to determine, based on its own internal state, if the cache should be used or not.
This would be useful for `DynamicPromptInvocation`, where the deterministic behaviour is determined by the `combinatorial: bool` field.
## Future Enhancements - Persisted Cache
Similar to how the latents storage is backed by disk, the invocation cache could be persisted to the database or disk. We'd need to be very careful about deserializing outputs, but it's perhaps worth exploring in the future.
* fix(ui): fix queue list item width
* feat(nodes): do not send the whole node on every generator progress
* feat(ui): strip out old logic related to sessions
Things like `isProcessing` are no longer relevant with queue. Removed them all & updated everything be appropriate for queue. May be a few little quirks I've missed...
* feat(ui): fix up param collapse labels
* feat(ui): click queue count to go to queue tab
* tidy(queue): update comment, query format
* feat(ui): fix progress bar when canceling
* fix(ui): fix circular dependency
* feat(nodes): bail on node caching logic if `node_cache_size == 0`
* feat(nodes): handle KeyError on node cache pop
* feat(nodes): bypass cache codepath if caches is disabled
more better no do thing
* fix(ui): reset api cache on connect/disconnect
* feat(ui): prevent enqueue when no prompts generated
* feat(ui): add queue controls to workflow editor
* feat(ui): update floating buttons & other incidental UI tweaks
* fix(ui): fix missing/incorrect translation keys
* fix(tests): add config service to mock invocation services
invoking needs access to `node_cache_size` to occur
* optionally remove pause/resume buttons from queue UI
* option to disable prepending
* chore(ui): remove unused file
* feat(queue): remove `order_id` entirely, `item_id` is now an autoinc pk
---------
Co-authored-by: Mary Hipp <maryhipp@Marys-MacBook-Air.local>
2023-09-20 05:09:24 +00:00
|
|
|
|
|
|
|
def emit_queue_cleared(self, queue_id: str) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a queue is cleared"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(QueueClearedEvent.build(queue_id))
|
|
|
|
|
|
|
|
# endregion
|
|
|
|
|
|
|
|
# region Download
|
2023-11-21 02:57:10 +00:00
|
|
|
|
2023-12-22 17:35:57 +00:00
|
|
|
def emit_download_started(self, source: str, download_path: str) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a download is started"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(DownloadStartedEvent.build(source, download_path))
|
2023-12-22 17:35:57 +00:00
|
|
|
|
|
|
|
def emit_download_progress(self, source: str, download_path: str, current_bytes: int, total_bytes: int) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted at intervals during a download"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(DownloadProgressEvent.build(source, download_path, current_bytes, total_bytes))
|
2023-12-22 17:35:57 +00:00
|
|
|
|
|
|
|
def emit_download_complete(self, source: str, download_path: str, total_bytes: int) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a download is completed"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(DownloadCompleteEvent.build(source, download_path, total_bytes))
|
2023-12-22 17:35:57 +00:00
|
|
|
|
|
|
|
def emit_download_cancelled(self, source: str) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a download is cancelled"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(DownloadCancelledEvent.build(source))
|
2023-12-22 17:35:57 +00:00
|
|
|
|
|
|
|
def emit_download_error(self, source: str, error_type: str, error: str) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a download encounters an error"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(DownloadErrorEvent.build(source, error_type, error))
|
2023-12-22 17:35:57 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
# endregion
|
2024-01-14 19:54:53 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
# region Model loading
|
2024-03-19 20:54:49 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_model_load_started(self, config: "AnyModelConfig", submodel_type: Optional["SubModelType"] = None) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a model load is started."""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(ModelLoadStartedEvent.build(config, submodel_type))
|
2024-03-19 20:54:49 +00:00
|
|
|
|
2024-03-14 08:29:57 +00:00
|
|
|
def emit_model_load_complete(
|
|
|
|
self, config: "AnyModelConfig", submodel_type: Optional["SubModelType"] = None
|
|
|
|
) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a model load is complete."""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(ModelLoadCompleteEvent.build(config, submodel_type))
|
2023-11-21 02:57:10 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
# endregion
|
2023-11-21 02:57:10 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
# region Model install
|
2023-11-21 02:57:10 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_model_install_download_progress(self, job: "ModelInstallJob") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted at intervals while the install job is in progress (remote models only)."""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(ModelInstallDownloadProgressEvent.build(job))
|
2023-11-26 18:18:21 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_model_install_downloads_complete(self, job: "ModelInstallJob") -> None:
|
|
|
|
self.dispatch(ModelInstallDownloadsCompleteEvent.build(job))
|
2023-11-26 18:18:21 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_model_install_started(self, job: "ModelInstallJob") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted once when an install job is started (after any download)."""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(ModelInstallStartedEvent.build(job))
|
|
|
|
|
|
|
|
def emit_model_install_complete(self, job: "ModelInstallJob") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when an install job is completed successfully."""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(ModelInstallCompleteEvent.build(job))
|
|
|
|
|
|
|
|
def emit_model_install_cancelled(self, job: "ModelInstallJob") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when an install job is cancelled."""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(ModelInstallCancelledEvent.build(job))
|
|
|
|
|
|
|
|
def emit_model_install_error(self, job: "ModelInstallJob") -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when an install job encounters an exception."""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(ModelInstallErrorEvent.build(job))
|
|
|
|
|
|
|
|
# endregion
|
|
|
|
|
|
|
|
# region Bulk image download
|
2024-01-08 00:55:59 +00:00
|
|
|
|
2024-02-19 18:54:48 +00:00
|
|
|
def emit_bulk_download_started(
|
|
|
|
self, bulk_download_id: str, bulk_download_item_id: str, bulk_download_item_name: str
|
|
|
|
) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a bulk image download is started"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(BulkDownloadStartedEvent.build(bulk_download_id, bulk_download_item_id, bulk_download_item_name))
|
2024-01-08 03:17:03 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_bulk_download_complete(
|
2024-01-14 04:35:33 +00:00
|
|
|
self, bulk_download_id: str, bulk_download_item_id: str, bulk_download_item_name: str
|
|
|
|
) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a bulk image download is complete"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(BulkDownloadCompleteEvent.build(bulk_download_id, bulk_download_item_id, bulk_download_item_name))
|
2024-01-08 03:17:03 +00:00
|
|
|
|
2024-03-14 08:04:19 +00:00
|
|
|
def emit_bulk_download_error(
|
2024-02-19 18:54:48 +00:00
|
|
|
self, bulk_download_id: str, bulk_download_item_id: str, bulk_download_item_name: str, error: str
|
|
|
|
) -> None:
|
2024-03-14 08:25:55 +00:00
|
|
|
"""Emitted when a bulk image download has an error"""
|
2024-03-14 08:04:19 +00:00
|
|
|
self.dispatch(
|
|
|
|
BulkDownloadErrorEvent.build(bulk_download_id, bulk_download_item_id, bulk_download_item_name, error)
|
2024-01-08 00:55:59 +00:00
|
|
|
)
|
2024-03-14 08:04:19 +00:00
|
|
|
|
|
|
|
# endregion
|