xref: /libCEED/backends/opt/ceed-opt-operator.c (revision 82946b17cb1e5ff4e735292ffaa6377862cf1d29)
189c6efa4Sjeremylt // Copyright (c) 2017-2018, Lawrence Livermore National Security, LLC.
289c6efa4Sjeremylt // Produced at the Lawrence Livermore National Laboratory. LLNL-CODE-734707.
389c6efa4Sjeremylt // All Rights reserved. See files LICENSE and NOTICE for details.
489c6efa4Sjeremylt //
589c6efa4Sjeremylt // This file is part of CEED, a collection of benchmarks, miniapps, software
689c6efa4Sjeremylt // libraries and APIs for efficient high-order finite element and spectral
789c6efa4Sjeremylt // element discretizations for exascale applications. For more information and
889c6efa4Sjeremylt // source code availability see http://github.com/ceed.
989c6efa4Sjeremylt //
1089c6efa4Sjeremylt // The CEED research is supported by the Exascale Computing Project 17-SC-20-SC,
1189c6efa4Sjeremylt // a collaborative effort of two U.S. Department of Energy organizations (Office
1289c6efa4Sjeremylt // of Science and the National Nuclear Security Administration) responsible for
1389c6efa4Sjeremylt // the planning and preparation of a capable exascale ecosystem, including
1489c6efa4Sjeremylt // software, applications, hardware, advanced system engineering and early
1589c6efa4Sjeremylt // testbed platforms, in support of the nation's exascale computing imperative.
1689c6efa4Sjeremylt 
1789c6efa4Sjeremylt #include <string.h>
1889c6efa4Sjeremylt #include "ceed-opt.h"
1989c6efa4Sjeremylt #include "../ref/ceed-ref.h"
2089c6efa4Sjeremylt 
21f10650afSjeremylt //------------------------------------------------------------------------------
22f10650afSjeremylt // Setup Input/Output Fields
23f10650afSjeremylt //------------------------------------------------------------------------------
2489c6efa4Sjeremylt static int CeedOperatorSetupFields_Opt(CeedQFunction qf, CeedOperator op,
2589c6efa4Sjeremylt                                        bool inOrOut, const CeedInt blksize,
2689c6efa4Sjeremylt                                        CeedElemRestriction *blkrestr,
2789c6efa4Sjeremylt                                        CeedVector *fullevecs, CeedVector *evecs,
2889c6efa4Sjeremylt                                        CeedVector *qvecs, CeedInt starte,
2989c6efa4Sjeremylt                                        CeedInt numfields, CeedInt Q) {
304d537eeaSYohann   CeedInt dim, ierr, ncomp, size, P;
3189c6efa4Sjeremylt   Ceed ceed;
3289c6efa4Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
3389c6efa4Sjeremylt   CeedBasis basis;
3489c6efa4Sjeremylt   CeedElemRestriction r;
3589c6efa4Sjeremylt   CeedOperatorField *opfields;
3689c6efa4Sjeremylt   CeedQFunctionField *qffields;
3789c6efa4Sjeremylt   if (inOrOut) {
3889c6efa4Sjeremylt     ierr = CeedOperatorGetFields(op, NULL, &opfields);
3989c6efa4Sjeremylt     CeedChk(ierr);
4089c6efa4Sjeremylt     ierr = CeedQFunctionGetFields(qf, NULL, &qffields);
4189c6efa4Sjeremylt     CeedChk(ierr);
4289c6efa4Sjeremylt   } else {
4389c6efa4Sjeremylt     ierr = CeedOperatorGetFields(op, &opfields, NULL);
4489c6efa4Sjeremylt     CeedChk(ierr);
4589c6efa4Sjeremylt     ierr = CeedQFunctionGetFields(qf, &qffields, NULL);
4689c6efa4Sjeremylt     CeedChk(ierr);
4789c6efa4Sjeremylt   }
4889c6efa4Sjeremylt 
4989c6efa4Sjeremylt   // Loop over fields
5089c6efa4Sjeremylt   for (CeedInt i=0; i<numfields; i++) {
5189c6efa4Sjeremylt     CeedEvalMode emode;
5289c6efa4Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qffields[i], &emode); CeedChk(ierr);
5389c6efa4Sjeremylt 
5489c6efa4Sjeremylt     if (emode != CEED_EVAL_WEIGHT) {
5589c6efa4Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opfields[i], &r);
5689c6efa4Sjeremylt       CeedChk(ierr);
5789c6efa4Sjeremylt       CeedElemRestriction_Ref *data;
5889c6efa4Sjeremylt       ierr = CeedElemRestrictionGetData(r, (void *)&data); CeedChk(ierr);
5989c6efa4Sjeremylt       Ceed ceed;
6089c6efa4Sjeremylt       ierr = CeedElemRestrictionGetCeed(r, &ceed); CeedChk(ierr);
618795c945Sjeremylt       CeedInt nelem, elemsize, nnodes;
6289c6efa4Sjeremylt       ierr = CeedElemRestrictionGetNumElements(r, &nelem); CeedChk(ierr);
6389c6efa4Sjeremylt       ierr = CeedElemRestrictionGetElementSize(r, &elemsize); CeedChk(ierr);
648795c945Sjeremylt       ierr = CeedElemRestrictionGetNumNodes(r, &nnodes); CeedChk(ierr);
6589c6efa4Sjeremylt       ierr = CeedElemRestrictionGetNumComponents(r, &ncomp); CeedChk(ierr);
6689c6efa4Sjeremylt       ierr = CeedElemRestrictionCreateBlocked(ceed, nelem, elemsize,
678795c945Sjeremylt                                               blksize, nnodes, ncomp,
6889c6efa4Sjeremylt                                               CEED_MEM_HOST, CEED_COPY_VALUES,
6989c6efa4Sjeremylt                                               data->indices, &blkrestr[i+starte]);
7089c6efa4Sjeremylt       CeedChk(ierr);
7189c6efa4Sjeremylt       ierr = CeedElemRestrictionCreateVector(blkrestr[i+starte], NULL,
7289c6efa4Sjeremylt                                              &fullevecs[i+starte]);
7389c6efa4Sjeremylt       CeedChk(ierr);
7489c6efa4Sjeremylt     }
7589c6efa4Sjeremylt 
7689c6efa4Sjeremylt     switch(emode) {
7789c6efa4Sjeremylt     case CEED_EVAL_NONE:
784d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChk(ierr);
794d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &evecs[i]); CeedChk(ierr);
804d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChk(ierr);
8189c6efa4Sjeremylt       break;
8289c6efa4Sjeremylt     case CEED_EVAL_INTERP:
834d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChk(ierr);
8489c6efa4Sjeremylt       ierr = CeedElemRestrictionGetElementSize(r, &P);
8589c6efa4Sjeremylt       CeedChk(ierr);
864d537eeaSYohann       ierr = CeedVectorCreate(ceed, P*size*blksize, &evecs[i]); CeedChk(ierr);
874d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChk(ierr);
8889c6efa4Sjeremylt       break;
8989c6efa4Sjeremylt     case CEED_EVAL_GRAD:
9089c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
914d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChk(ierr);
9289c6efa4Sjeremylt       ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
9389c6efa4Sjeremylt       ierr = CeedElemRestrictionGetElementSize(r, &P);
9489c6efa4Sjeremylt       CeedChk(ierr);
954d537eeaSYohann       ierr = CeedVectorCreate(ceed, P*size/dim*blksize, &evecs[i]); CeedChk(ierr);
964d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChk(ierr);
9789c6efa4Sjeremylt       break;
9889c6efa4Sjeremylt     case CEED_EVAL_WEIGHT: // Only on input fields
9989c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
10089c6efa4Sjeremylt       ierr = CeedVectorCreate(ceed, Q*blksize, &qvecs[i]); CeedChk(ierr);
10189c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
102a7b7f929Sjeremylt                             CEED_EVAL_WEIGHT, CEED_VECTOR_NONE, qvecs[i]);
103a7b7f929Sjeremylt       CeedChk(ierr);
10489c6efa4Sjeremylt 
10589c6efa4Sjeremylt       break;
10689c6efa4Sjeremylt     case CEED_EVAL_DIV:
10789c6efa4Sjeremylt       break; // Not implimented
10889c6efa4Sjeremylt     case CEED_EVAL_CURL:
10989c6efa4Sjeremylt       break; // Not implimented
11089c6efa4Sjeremylt     }
11189c6efa4Sjeremylt   }
11289c6efa4Sjeremylt   return 0;
11389c6efa4Sjeremylt }
11489c6efa4Sjeremylt 
115f10650afSjeremylt //------------------------------------------------------------------------------
116f10650afSjeremylt // Setup Operator
117f10650afSjeremylt //------------------------------------------------------------------------------
11889c6efa4Sjeremylt static int CeedOperatorSetup_Opt(CeedOperator op) {
11989c6efa4Sjeremylt   int ierr;
12089c6efa4Sjeremylt   bool setupdone;
12189c6efa4Sjeremylt   ierr = CeedOperatorGetSetupStatus(op, &setupdone); CeedChk(ierr);
12289c6efa4Sjeremylt   if (setupdone) return 0;
12389c6efa4Sjeremylt   Ceed ceed;
12489c6efa4Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
12589c6efa4Sjeremylt   Ceed_Opt *ceedimpl;
12689c6efa4Sjeremylt   ierr = CeedGetData(ceed, (void *)&ceedimpl); CeedChk(ierr);
12789c6efa4Sjeremylt   const CeedInt blksize = ceedimpl->blksize;
12889c6efa4Sjeremylt   CeedOperator_Opt *impl;
12989c6efa4Sjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
13089c6efa4Sjeremylt   CeedQFunction qf;
13189c6efa4Sjeremylt   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
13289c6efa4Sjeremylt   CeedInt Q, numinputfields, numoutputfields;
13389c6efa4Sjeremylt   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
13416911fdaSjeremylt   ierr = CeedQFunctionGetIdentityStatus(qf, &impl->identityqf); CeedChk(ierr);
13589c6efa4Sjeremylt   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
13689c6efa4Sjeremylt   CeedChk(ierr);
13789c6efa4Sjeremylt   CeedOperatorField *opinputfields, *opoutputfields;
13889c6efa4Sjeremylt   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
13989c6efa4Sjeremylt   CeedChk(ierr);
14089c6efa4Sjeremylt   CeedQFunctionField *qfinputfields, *qfoutputfields;
14189c6efa4Sjeremylt   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
14289c6efa4Sjeremylt   CeedChk(ierr);
14389c6efa4Sjeremylt 
14489c6efa4Sjeremylt   // Allocate
14589c6efa4Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->blkrestr);
14689c6efa4Sjeremylt   CeedChk(ierr);
14789c6efa4Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->evecs);
14889c6efa4Sjeremylt   CeedChk(ierr);
14989c6efa4Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->edata);
15089c6efa4Sjeremylt   CeedChk(ierr);
15189c6efa4Sjeremylt 
15289c6efa4Sjeremylt   ierr = CeedCalloc(16, &impl->inputstate); CeedChk(ierr);
15389c6efa4Sjeremylt   ierr = CeedCalloc(16, &impl->evecsin); CeedChk(ierr);
15489c6efa4Sjeremylt   ierr = CeedCalloc(16, &impl->evecsout); CeedChk(ierr);
15589c6efa4Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsin); CeedChk(ierr);
15689c6efa4Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsout); CeedChk(ierr);
15789c6efa4Sjeremylt 
15889c6efa4Sjeremylt   impl->numein = numinputfields; impl->numeout = numoutputfields;
15989c6efa4Sjeremylt 
16089c6efa4Sjeremylt   // Set up infield and outfield pointer arrays
16189c6efa4Sjeremylt   // Infields
16289c6efa4Sjeremylt   ierr = CeedOperatorSetupFields_Opt(qf, op, 0, blksize, impl->blkrestr,
16389c6efa4Sjeremylt                                      impl->evecs, impl->evecsin,
16489c6efa4Sjeremylt                                      impl->qvecsin, 0,
16589c6efa4Sjeremylt                                      numinputfields, Q);
16689c6efa4Sjeremylt   CeedChk(ierr);
16789c6efa4Sjeremylt   // Outfields
16889c6efa4Sjeremylt   ierr = CeedOperatorSetupFields_Opt(qf, op, 1, blksize, impl->blkrestr,
16989c6efa4Sjeremylt                                      impl->evecs, impl->evecsout,
17089c6efa4Sjeremylt                                      impl->qvecsout, numinputfields,
17189c6efa4Sjeremylt                                      numoutputfields, Q);
17289c6efa4Sjeremylt   CeedChk(ierr);
17389c6efa4Sjeremylt 
17416911fdaSjeremylt   // Identity QFunctions
17516911fdaSjeremylt   if (impl->identityqf) {
17616911fdaSjeremylt     CeedEvalMode inmode, outmode;
17716911fdaSjeremylt     CeedQFunctionField *infields, *outfields;
17816911fdaSjeremylt     ierr = CeedQFunctionGetFields(qf, &infields, &outfields); CeedChk(ierr);
17916911fdaSjeremylt 
18016911fdaSjeremylt     for (CeedInt i=0; i<numinputfields; i++) {
18116911fdaSjeremylt       ierr = CeedQFunctionFieldGetEvalMode(infields[i], &inmode);
18216911fdaSjeremylt       CeedChk(ierr);
18316911fdaSjeremylt       ierr = CeedQFunctionFieldGetEvalMode(outfields[i], &outmode);
18416911fdaSjeremylt       CeedChk(ierr);
18516911fdaSjeremylt 
18616911fdaSjeremylt       ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChk(ierr);
18716911fdaSjeremylt       impl->qvecsout[i] = impl->qvecsin[i];
18855e4cc5bSjeremylt       ierr = CeedVectorAddReference(impl->qvecsin[i]); CeedChk(ierr);
18916911fdaSjeremylt     }
19016911fdaSjeremylt   }
19116911fdaSjeremylt 
19289c6efa4Sjeremylt   ierr = CeedOperatorSetSetupDone(op); CeedChk(ierr);
19389c6efa4Sjeremylt 
19489c6efa4Sjeremylt   return 0;
19589c6efa4Sjeremylt }
19689c6efa4Sjeremylt 
197f10650afSjeremylt //------------------------------------------------------------------------------
198f10650afSjeremylt // Setup Input Fields
199f10650afSjeremylt //------------------------------------------------------------------------------
2001d102b48SJeremy L Thompson static inline int CeedOperatorSetupInputs_Opt(CeedInt numinputfields,
2011d102b48SJeremy L Thompson     CeedQFunctionField *qfinputfields, CeedOperatorField *opinputfields,
2021d102b48SJeremy L Thompson     CeedVector invec, CeedOperator_Opt *impl, CeedRequest *request) {
2031d102b48SJeremy L Thompson   CeedInt ierr;
20489c6efa4Sjeremylt   CeedEvalMode emode;
20589c6efa4Sjeremylt   CeedVector vec;
2061d102b48SJeremy L Thompson   CeedTransposeMode lmode;
20789c6efa4Sjeremylt   uint64_t state;
20889c6efa4Sjeremylt 
20989c6efa4Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
21089c6efa4Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
21189c6efa4Sjeremylt     CeedChk(ierr);
21289c6efa4Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
21389c6efa4Sjeremylt     } else {
21489c6efa4Sjeremylt       // Get input vector
21589c6efa4Sjeremylt       ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
21689c6efa4Sjeremylt       if (vec != CEED_VECTOR_ACTIVE) {
21789c6efa4Sjeremylt         // Restrict
21889c6efa4Sjeremylt         ierr = CeedVectorGetState(vec, &state); CeedChk(ierr);
21989c6efa4Sjeremylt         if (state != impl->inputstate[i]) {
22089c6efa4Sjeremylt           ierr = CeedOperatorFieldGetLMode(opinputfields[i], &lmode);
22189c6efa4Sjeremylt           CeedChk(ierr);
22289c6efa4Sjeremylt           ierr = CeedElemRestrictionApply(impl->blkrestr[i], CEED_NOTRANSPOSE,
22389c6efa4Sjeremylt                                           lmode, vec, impl->evecs[i], request);
22489c6efa4Sjeremylt           CeedChk(ierr);
22589c6efa4Sjeremylt           impl->inputstate[i] = state;
22689c6efa4Sjeremylt         }
22789c6efa4Sjeremylt       } else {
22889c6efa4Sjeremylt         // Set Qvec for CEED_EVAL_NONE
22989c6efa4Sjeremylt         if (emode == CEED_EVAL_NONE) {
23089c6efa4Sjeremylt           ierr = CeedVectorGetArray(impl->evecsin[i], CEED_MEM_HOST,
23189c6efa4Sjeremylt                                     &impl->edata[i]); CeedChk(ierr);
23289c6efa4Sjeremylt           ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST,
23389c6efa4Sjeremylt                                     CEED_USE_POINTER,
23489c6efa4Sjeremylt                                     impl->edata[i]); CeedChk(ierr);
23589c6efa4Sjeremylt           ierr = CeedVectorRestoreArray(impl->evecsin[i],
23689c6efa4Sjeremylt                                         &impl->edata[i]); CeedChk(ierr);
23789c6efa4Sjeremylt         }
23889c6efa4Sjeremylt       }
23989c6efa4Sjeremylt       // Get evec
24089c6efa4Sjeremylt       ierr = CeedVectorGetArrayRead(impl->evecs[i], CEED_MEM_HOST,
24189c6efa4Sjeremylt                                     (const CeedScalar **) &impl->edata[i]);
24289c6efa4Sjeremylt       CeedChk(ierr);
24389c6efa4Sjeremylt     }
24489c6efa4Sjeremylt   }
2451d102b48SJeremy L Thompson   return 0;
2461d102b48SJeremy L Thompson }
24789c6efa4Sjeremylt 
248f10650afSjeremylt //------------------------------------------------------------------------------
249f10650afSjeremylt // Input Basis Action
250f10650afSjeremylt //------------------------------------------------------------------------------
2511d102b48SJeremy L Thompson static inline int CeedOperatorInputBasis_Opt(CeedInt e, CeedInt Q,
2521d102b48SJeremy L Thompson     CeedQFunctionField *qfinputfields, CeedOperatorField *opinputfields,
2531d102b48SJeremy L Thompson     CeedInt numinputfields, CeedInt blksize, CeedVector invec, bool skipactive,
2541d102b48SJeremy L Thompson     CeedOperator_Opt *impl, CeedRequest *request) {
2551d102b48SJeremy L Thompson   CeedInt ierr;
2561d102b48SJeremy L Thompson   CeedInt dim, elemsize, size;
2571d102b48SJeremy L Thompson   CeedElemRestriction Erestrict;
2581d102b48SJeremy L Thompson   CeedEvalMode emode;
2591d102b48SJeremy L Thompson   CeedBasis basis;
2601d102b48SJeremy L Thompson   CeedVector vec;
2611d102b48SJeremy L Thompson   CeedTransposeMode lmode;
26289c6efa4Sjeremylt 
26389c6efa4Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
2641d102b48SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
2651d102b48SJeremy L Thompson     // Skip active input
2661d102b48SJeremy L Thompson     if (skipactive) {
2671d102b48SJeremy L Thompson       if (vec == CEED_VECTOR_ACTIVE)
2681d102b48SJeremy L Thompson         continue;
2691d102b48SJeremy L Thompson     }
2701d102b48SJeremy L Thompson 
27189c6efa4Sjeremylt     CeedInt activein = 0;
2724d537eeaSYohann     // Get elemsize, emode, size
27389c6efa4Sjeremylt     ierr = CeedOperatorFieldGetElemRestriction(opinputfields[i], &Erestrict);
27489c6efa4Sjeremylt     CeedChk(ierr);
27589c6efa4Sjeremylt     ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
27689c6efa4Sjeremylt     CeedChk(ierr);
27789c6efa4Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
27889c6efa4Sjeremylt     CeedChk(ierr);
2794d537eeaSYohann     ierr = CeedQFunctionFieldGetSize(qfinputfields[i], &size); CeedChk(ierr);
28089c6efa4Sjeremylt     // Restrict block active input
28189c6efa4Sjeremylt     if (vec == CEED_VECTOR_ACTIVE) {
28289c6efa4Sjeremylt       ierr = CeedOperatorFieldGetLMode(opinputfields[i], &lmode);
28389c6efa4Sjeremylt       CeedChk(ierr);
28489c6efa4Sjeremylt       ierr = CeedElemRestrictionApplyBlock(impl->blkrestr[i], e/blksize,
28589c6efa4Sjeremylt                                            CEED_NOTRANSPOSE, lmode, invec,
28689c6efa4Sjeremylt                                            impl->evecsin[i], request);
28789c6efa4Sjeremylt       CeedChk(ierr);
28889c6efa4Sjeremylt       activein = 1;
28989c6efa4Sjeremylt     }
29089c6efa4Sjeremylt     // Basis action
29189c6efa4Sjeremylt     switch(emode) {
29289c6efa4Sjeremylt     case CEED_EVAL_NONE:
29389c6efa4Sjeremylt       if (!activein) {
29489c6efa4Sjeremylt         ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST,
29589c6efa4Sjeremylt                                   CEED_USE_POINTER,
2964d537eeaSYohann                                   &impl->edata[i][e*Q*size]); CeedChk(ierr);
29789c6efa4Sjeremylt       }
29889c6efa4Sjeremylt       break;
29989c6efa4Sjeremylt     case CEED_EVAL_INTERP:
30089c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis);
30189c6efa4Sjeremylt       CeedChk(ierr);
30289c6efa4Sjeremylt       if (!activein) {
30389c6efa4Sjeremylt         ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
30489c6efa4Sjeremylt                                   CEED_USE_POINTER,
3054d537eeaSYohann                                   &impl->edata[i][e*elemsize*size]);
30689c6efa4Sjeremylt         CeedChk(ierr);
30789c6efa4Sjeremylt       }
30889c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
30989c6efa4Sjeremylt                             CEED_EVAL_INTERP, impl->evecsin[i],
31089c6efa4Sjeremylt                             impl->qvecsin[i]); CeedChk(ierr);
31189c6efa4Sjeremylt       break;
31289c6efa4Sjeremylt     case CEED_EVAL_GRAD:
31389c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis);
31489c6efa4Sjeremylt       CeedChk(ierr);
31589c6efa4Sjeremylt       if (!activein) {
3164d537eeaSYohann         ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
31789c6efa4Sjeremylt         ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
31889c6efa4Sjeremylt                                   CEED_USE_POINTER,
3194d537eeaSYohann                                   &impl->edata[i][e*elemsize*size/dim]);
32089c6efa4Sjeremylt         CeedChk(ierr);
32189c6efa4Sjeremylt       }
32289c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
32389c6efa4Sjeremylt                             CEED_EVAL_GRAD, impl->evecsin[i],
32489c6efa4Sjeremylt                             impl->qvecsin[i]); CeedChk(ierr);
32589c6efa4Sjeremylt       break;
32689c6efa4Sjeremylt     case CEED_EVAL_WEIGHT:
32789c6efa4Sjeremylt       break;  // No action
32889c6efa4Sjeremylt     case CEED_EVAL_DIV:
3291d102b48SJeremy L Thompson     case CEED_EVAL_CURL: {
3301d102b48SJeremy L Thompson       // LCOV_EXCL_START
3311d102b48SJeremy L Thompson       ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis);
33289c6efa4Sjeremylt       CeedChk(ierr);
3331d102b48SJeremy L Thompson       Ceed ceed;
3341d102b48SJeremy L Thompson       ierr = CeedBasisGetCeed(basis, &ceed); CeedChk(ierr);
3351d102b48SJeremy L Thompson       return CeedError(ceed, 1, "Ceed evaluation mode not implemented");
3361d102b48SJeremy L Thompson       // LCOV_EXCL_STOP
3371d102b48SJeremy L Thompson       break; // Not implemented
3381d102b48SJeremy L Thompson     }
3391d102b48SJeremy L Thompson     }
3401d102b48SJeremy L Thompson   }
3411d102b48SJeremy L Thompson   return 0;
3421d102b48SJeremy L Thompson }
34389c6efa4Sjeremylt 
344f10650afSjeremylt //------------------------------------------------------------------------------
345f10650afSjeremylt // Output Basis Action
346f10650afSjeremylt //------------------------------------------------------------------------------
3471d102b48SJeremy L Thompson static inline int CeedOperatorOutputBasis_Opt(CeedInt e, CeedInt Q,
3481d102b48SJeremy L Thompson     CeedQFunctionField *qfoutputfields, CeedOperatorField *opoutputfields,
3491d102b48SJeremy L Thompson     CeedInt blksize, CeedInt numinputfields, CeedInt numoutputfields,
3501d102b48SJeremy L Thompson     CeedOperator op, CeedVector outvec, CeedOperator_Opt *impl,
3511d102b48SJeremy L Thompson     CeedRequest *request) {
3521d102b48SJeremy L Thompson   CeedInt ierr;
3531d102b48SJeremy L Thompson   CeedElemRestriction Erestrict;
3541d102b48SJeremy L Thompson   CeedEvalMode emode;
3551d102b48SJeremy L Thompson   CeedBasis basis;
3561d102b48SJeremy L Thompson   CeedVector vec;
3571d102b48SJeremy L Thompson   CeedTransposeMode lmode;
3581d102b48SJeremy L Thompson 
35989c6efa4Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
3604d537eeaSYohann     // Get elemsize, emode, size
36189c6efa4Sjeremylt     ierr = CeedOperatorFieldGetElemRestriction(opoutputfields[i], &Erestrict);
36289c6efa4Sjeremylt     CeedChk(ierr);
36389c6efa4Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
36489c6efa4Sjeremylt     CeedChk(ierr);
36589c6efa4Sjeremylt     // Basis action
36689c6efa4Sjeremylt     switch(emode) {
36789c6efa4Sjeremylt     case CEED_EVAL_NONE:
36889c6efa4Sjeremylt       break; // No action
36989c6efa4Sjeremylt     case CEED_EVAL_INTERP:
37089c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
37189c6efa4Sjeremylt       CeedChk(ierr);
37289c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
37389c6efa4Sjeremylt                             CEED_EVAL_INTERP, impl->qvecsout[i],
37489c6efa4Sjeremylt                             impl->evecsout[i]); CeedChk(ierr);
37589c6efa4Sjeremylt       break;
37689c6efa4Sjeremylt     case CEED_EVAL_GRAD:
37789c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
37889c6efa4Sjeremylt       CeedChk(ierr);
37989c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
38089c6efa4Sjeremylt                             CEED_EVAL_GRAD, impl->qvecsout[i],
38189c6efa4Sjeremylt                             impl->evecsout[i]); CeedChk(ierr);
38289c6efa4Sjeremylt       break;
38389c6efa4Sjeremylt     case CEED_EVAL_WEIGHT: {
384c042f62fSJeremy L Thompson       // LCOV_EXCL_START
38589c6efa4Sjeremylt       Ceed ceed;
38689c6efa4Sjeremylt       ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
3871d102b48SJeremy L Thompson       return CeedError(ceed, 1, "CEED_EVAL_WEIGHT cannot be an output "
3881d102b48SJeremy L Thompson                        "evaluation mode");
389c042f62fSJeremy L Thompson       // LCOV_EXCL_STOP
39089c6efa4Sjeremylt       break; // Should not occur
39189c6efa4Sjeremylt     }
39289c6efa4Sjeremylt     case CEED_EVAL_DIV:
3931d102b48SJeremy L Thompson     case CEED_EVAL_CURL: {
3941d102b48SJeremy L Thompson       // LCOV_EXCL_START
3951d102b48SJeremy L Thompson       Ceed ceed;
3961d102b48SJeremy L Thompson       ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
3971d102b48SJeremy L Thompson       return CeedError(ceed, 1, "Ceed evaluation mode not implemented");
3981d102b48SJeremy L Thompson       // LCOV_EXCL_STOP
3991d102b48SJeremy L Thompson       break; // Not implemented
4001d102b48SJeremy L Thompson     }
40189c6efa4Sjeremylt     }
40289c6efa4Sjeremylt     // Restrict output block
40389c6efa4Sjeremylt     // Get output vector
40489c6efa4Sjeremylt     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
40589c6efa4Sjeremylt     if (vec == CEED_VECTOR_ACTIVE)
40689c6efa4Sjeremylt       vec = outvec;
40789c6efa4Sjeremylt     // Restrict
40889c6efa4Sjeremylt     ierr = CeedOperatorFieldGetLMode(opoutputfields[i], &lmode);
40989c6efa4Sjeremylt     CeedChk(ierr);
41089c6efa4Sjeremylt     ierr = CeedElemRestrictionApplyBlock(impl->blkrestr[i+impl->numein],
41189c6efa4Sjeremylt                                          e/blksize, CEED_TRANSPOSE,
41289c6efa4Sjeremylt                                          lmode, impl->evecsout[i],
41389c6efa4Sjeremylt                                          vec, request); CeedChk(ierr);
41489c6efa4Sjeremylt   }
4151d102b48SJeremy L Thompson   return 0;
41689c6efa4Sjeremylt }
41789c6efa4Sjeremylt 
418f10650afSjeremylt //------------------------------------------------------------------------------
419f10650afSjeremylt // Restore Input Vectors
420f10650afSjeremylt //------------------------------------------------------------------------------
4211d102b48SJeremy L Thompson static inline int CeedOperatorRestoreInputs_Opt(CeedInt numinputfields,
4221d102b48SJeremy L Thompson     CeedQFunctionField *qfinputfields, CeedOperatorField *opinputfields,
4231d102b48SJeremy L Thompson     bool skipactive, CeedOperator_Opt *impl) {
4241d102b48SJeremy L Thompson   CeedInt ierr;
4251d102b48SJeremy L Thompson   CeedEvalMode emode;
4261d102b48SJeremy L Thompson 
42789c6efa4Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
4281d102b48SJeremy L Thompson     // Skip active inputs
4291d102b48SJeremy L Thompson     if (skipactive) {
4301d102b48SJeremy L Thompson       CeedVector vec;
4311d102b48SJeremy L Thompson       ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
4321d102b48SJeremy L Thompson       if (vec == CEED_VECTOR_ACTIVE)
4331d102b48SJeremy L Thompson         continue;
4341d102b48SJeremy L Thompson     }
43589c6efa4Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
43689c6efa4Sjeremylt     CeedChk(ierr);
43789c6efa4Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
43889c6efa4Sjeremylt     } else {
43989c6efa4Sjeremylt       ierr = CeedVectorRestoreArrayRead(impl->evecs[i],
44089c6efa4Sjeremylt                                         (const CeedScalar **) &impl->edata[i]);
44189c6efa4Sjeremylt       CeedChk(ierr);
44289c6efa4Sjeremylt     }
44389c6efa4Sjeremylt   }
4441d102b48SJeremy L Thompson   return 0;
4451d102b48SJeremy L Thompson }
4461d102b48SJeremy L Thompson 
447f10650afSjeremylt //------------------------------------------------------------------------------
448f10650afSjeremylt // Operator Apply
449f10650afSjeremylt //------------------------------------------------------------------------------
4501d102b48SJeremy L Thompson static int CeedOperatorApply_Opt(CeedOperator op, CeedVector invec,
4511d102b48SJeremy L Thompson                                  CeedVector outvec, CeedRequest *request) {
4521d102b48SJeremy L Thompson   int ierr;
4531d102b48SJeremy L Thompson   Ceed ceed;
4541d102b48SJeremy L Thompson   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
4551d102b48SJeremy L Thompson   Ceed_Opt *ceedimpl;
4561d102b48SJeremy L Thompson   ierr = CeedGetData(ceed, (void *)&ceedimpl); CeedChk(ierr);
4571d102b48SJeremy L Thompson   CeedInt blksize = ceedimpl->blksize;
4581d102b48SJeremy L Thompson   CeedOperator_Opt *impl;
4591d102b48SJeremy L Thompson   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
4601d102b48SJeremy L Thompson   CeedInt Q, numinputfields, numoutputfields, numelements;
4611d102b48SJeremy L Thompson   ierr = CeedOperatorGetNumElements(op, &numelements); CeedChk(ierr);
4621d102b48SJeremy L Thompson   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
4631d102b48SJeremy L Thompson   CeedInt nblks = (numelements/blksize) + !!(numelements%blksize);
4641d102b48SJeremy L Thompson   CeedQFunction qf;
4651d102b48SJeremy L Thompson   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
4661d102b48SJeremy L Thompson   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
4671d102b48SJeremy L Thompson   CeedChk(ierr);
4681d102b48SJeremy L Thompson   CeedOperatorField *opinputfields, *opoutputfields;
4691d102b48SJeremy L Thompson   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
4701d102b48SJeremy L Thompson   CeedChk(ierr);
4711d102b48SJeremy L Thompson   CeedQFunctionField *qfinputfields, *qfoutputfields;
4721d102b48SJeremy L Thompson   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
4731d102b48SJeremy L Thompson   CeedChk(ierr);
4741d102b48SJeremy L Thompson   CeedEvalMode emode;
4751d102b48SJeremy L Thompson 
4761d102b48SJeremy L Thompson   // Setup
4771d102b48SJeremy L Thompson   ierr = CeedOperatorSetup_Opt(op); CeedChk(ierr);
4781d102b48SJeremy L Thompson 
4791d102b48SJeremy L Thompson   // Input Evecs and Restriction
4801d102b48SJeremy L Thompson   ierr = CeedOperatorSetupInputs_Opt(numinputfields, qfinputfields,
48116911fdaSjeremylt                                      opinputfields, invec, impl, request);
48216911fdaSjeremylt   CeedChk(ierr);
4831d102b48SJeremy L Thompson 
4841d102b48SJeremy L Thompson   // Output Lvecs, Evecs, and Qvecs
4851d102b48SJeremy L Thompson   for (CeedInt i=0; i<numoutputfields; i++) {
4861d102b48SJeremy L Thompson     // Set Qvec if needed
4871d102b48SJeremy L Thompson     ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
4881d102b48SJeremy L Thompson     CeedChk(ierr);
4891d102b48SJeremy L Thompson     if (emode == CEED_EVAL_NONE) {
4901d102b48SJeremy L Thompson       // Set qvec to single block evec
4911d102b48SJeremy L Thompson       ierr = CeedVectorGetArray(impl->evecsout[i], CEED_MEM_HOST,
4921d102b48SJeremy L Thompson                                 &impl->edata[i + numinputfields]);
4931d102b48SJeremy L Thompson       CeedChk(ierr);
4941d102b48SJeremy L Thompson       ierr = CeedVectorSetArray(impl->qvecsout[i], CEED_MEM_HOST,
4951d102b48SJeremy L Thompson                                 CEED_USE_POINTER,
4961d102b48SJeremy L Thompson                                 impl->edata[i + numinputfields]); CeedChk(ierr);
4971d102b48SJeremy L Thompson       ierr = CeedVectorRestoreArray(impl->evecsout[i],
4981d102b48SJeremy L Thompson                                     &impl->edata[i + numinputfields]);
4991d102b48SJeremy L Thompson       CeedChk(ierr);
5001d102b48SJeremy L Thompson     }
5011d102b48SJeremy L Thompson   }
5021d102b48SJeremy L Thompson 
5031d102b48SJeremy L Thompson   // Loop through elements
5041d102b48SJeremy L Thompson   for (CeedInt e=0; e<nblks*blksize; e+=blksize) {
5051d102b48SJeremy L Thompson     // Input basis apply
5061d102b48SJeremy L Thompson     ierr = CeedOperatorInputBasis_Opt(e, Q, qfinputfields, opinputfields,
5071d102b48SJeremy L Thompson                                       numinputfields, blksize, invec, false,
5081d102b48SJeremy L Thompson                                       impl, request); CeedChk(ierr);
5091d102b48SJeremy L Thompson 
5101d102b48SJeremy L Thompson     // Q function
51116911fdaSjeremylt     if (!impl->identityqf) {
5121d102b48SJeremy L Thompson       ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout);
5131d102b48SJeremy L Thompson       CeedChk(ierr);
51416911fdaSjeremylt     }
5151d102b48SJeremy L Thompson 
5161d102b48SJeremy L Thompson     // Output basis apply and restrict
5171d102b48SJeremy L Thompson     ierr = CeedOperatorOutputBasis_Opt(e, Q, qfoutputfields, opoutputfields,
5187f823360Sjeremylt                                        blksize, numinputfields, numoutputfields,
5197f823360Sjeremylt                                        op, outvec, impl, request);
5207f823360Sjeremylt     CeedChk(ierr);
5211d102b48SJeremy L Thompson   }
5221d102b48SJeremy L Thompson 
5231d102b48SJeremy L Thompson   // Restore input arrays
5241d102b48SJeremy L Thompson   ierr = CeedOperatorRestoreInputs_Opt(numinputfields, qfinputfields,
5251d102b48SJeremy L Thompson                                        opinputfields, false, impl);
5261d102b48SJeremy L Thompson   CeedChk(ierr);
52789c6efa4Sjeremylt 
52889c6efa4Sjeremylt   return 0;
52989c6efa4Sjeremylt }
53089c6efa4Sjeremylt 
531f10650afSjeremylt //------------------------------------------------------------------------------
5321d102b48SJeremy L Thompson // Assemble Linear QFunction
533f10650afSjeremylt //------------------------------------------------------------------------------
5341d102b48SJeremy L Thompson static int CeedOperatorAssembleLinearQFunction_Opt(CeedOperator op,
5351d102b48SJeremy L Thompson     CeedVector *assembled, CeedElemRestriction *rstr, CeedRequest *request) {
5361d102b48SJeremy L Thompson   int ierr;
5371d102b48SJeremy L Thompson   Ceed ceed;
5381d102b48SJeremy L Thompson   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
5391d102b48SJeremy L Thompson   Ceed_Opt *ceedimpl;
5401d102b48SJeremy L Thompson   ierr = CeedGetData(ceed, (void *)&ceedimpl); CeedChk(ierr);
5411d102b48SJeremy L Thompson   const CeedInt blksize = ceedimpl->blksize;
5421d102b48SJeremy L Thompson   CeedOperator_Opt *impl;
5431d102b48SJeremy L Thompson   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
5441d102b48SJeremy L Thompson   CeedInt Q, numinputfields, numoutputfields, numelements, size;
5451d102b48SJeremy L Thompson   ierr = CeedOperatorGetNumElements(op, &numelements); CeedChk(ierr);
5461d102b48SJeremy L Thompson   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
5471d102b48SJeremy L Thompson   CeedInt nblks = (numelements/blksize) + !!(numelements%blksize);
5481d102b48SJeremy L Thompson   CeedQFunction qf;
5491d102b48SJeremy L Thompson   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
5501d102b48SJeremy L Thompson   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
5511d102b48SJeremy L Thompson   CeedChk(ierr);
5521d102b48SJeremy L Thompson   CeedOperatorField *opinputfields, *opoutputfields;
5531d102b48SJeremy L Thompson   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
5541d102b48SJeremy L Thompson   CeedChk(ierr);
5551d102b48SJeremy L Thompson   CeedQFunctionField *qfinputfields, *qfoutputfields;
5561d102b48SJeremy L Thompson   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
5571d102b48SJeremy L Thompson   CeedChk(ierr);
5581d102b48SJeremy L Thompson   CeedVector vec, lvec;
5591d102b48SJeremy L Thompson   CeedInt numactivein = 0, numactiveout = 0;
56042ea3801Sjeremylt   CeedVector *activein = NULL;
5611d102b48SJeremy L Thompson   CeedScalar *a, *tmp;
5621d102b48SJeremy L Thompson 
5631d102b48SJeremy L Thompson   // Setup
5641d102b48SJeremy L Thompson   ierr = CeedOperatorSetup_Opt(op); CeedChk(ierr);
5651d102b48SJeremy L Thompson 
56616911fdaSjeremylt   // Check for identity
56716911fdaSjeremylt   if (impl->identityqf)
56816911fdaSjeremylt     // LCOV_EXCL_START
56967db23e4Sjeremylt     return CeedError(ceed, 1, "Assembling identity qfunctions not supported");
57016911fdaSjeremylt   // LCOV_EXCL_STOP
57116911fdaSjeremylt 
5721d102b48SJeremy L Thompson   // Input Evecs and Restriction
5731d102b48SJeremy L Thompson   ierr = CeedOperatorSetupInputs_Opt(numinputfields, qfinputfields,
5741d102b48SJeremy L Thompson                                      opinputfields, NULL, impl, request);
5751d102b48SJeremy L Thompson   CeedChk(ierr);
5761d102b48SJeremy L Thompson 
5771d102b48SJeremy L Thompson   // Count number of active input fields
5781d102b48SJeremy L Thompson   for (CeedInt i=0; i<numinputfields; i++) {
5791d102b48SJeremy L Thompson     // Get input vector
5801d102b48SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
5811d102b48SJeremy L Thompson     // Check if active input
5821d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
5831d102b48SJeremy L Thompson       ierr = CeedQFunctionFieldGetSize(qfinputfields[i], &size); CeedChk(ierr);
5841d102b48SJeremy L Thompson       ierr = CeedVectorSetValue(impl->qvecsin[i], 0.0); CeedChk(ierr);
5851d102b48SJeremy L Thompson       ierr = CeedVectorGetArray(impl->qvecsin[i], CEED_MEM_HOST, &tmp);
5861d102b48SJeremy L Thompson       CeedChk(ierr);
5871d102b48SJeremy L Thompson       ierr = CeedRealloc(numactivein + size, &activein); CeedChk(ierr);
5881d102b48SJeremy L Thompson       for (CeedInt field=0; field<size; field++) {
58942ea3801Sjeremylt         ierr = CeedVectorCreate(ceed, Q*blksize, &activein[numactivein+field]);
59042ea3801Sjeremylt         CeedChk(ierr);
59142ea3801Sjeremylt         ierr = CeedVectorSetArray(activein[numactivein+field], CEED_MEM_HOST,
59242ea3801Sjeremylt                                   CEED_USE_POINTER, &tmp[field*Q*blksize]);
593112e3f70Sjeremylt         CeedChk(ierr);
5941d102b48SJeremy L Thompson       }
5951d102b48SJeremy L Thompson       numactivein += size;
5961d102b48SJeremy L Thompson       ierr = CeedVectorRestoreArray(impl->qvecsin[i], &tmp); CeedChk(ierr);
5971d102b48SJeremy L Thompson     }
59889c6efa4Sjeremylt   }
59989c6efa4Sjeremylt 
6001d102b48SJeremy L Thompson   // Count number of active output fields
6011d102b48SJeremy L Thompson   for (CeedInt i=0; i<numoutputfields; i++) {
6021d102b48SJeremy L Thompson     // Get output vector
6031d102b48SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
6041d102b48SJeremy L Thompson     // Check if active output
6051d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
6061d102b48SJeremy L Thompson       ierr = CeedQFunctionFieldGetSize(qfoutputfields[i], &size); CeedChk(ierr);
6071d102b48SJeremy L Thompson       numactiveout += size;
6081d102b48SJeremy L Thompson     }
6091d102b48SJeremy L Thompson   }
6101d102b48SJeremy L Thompson 
6111d102b48SJeremy L Thompson   // Check sizes
6121d102b48SJeremy L Thompson   if (!numactivein || !numactiveout)
6131d102b48SJeremy L Thompson     // LCOV_EXCL_START
6141d102b48SJeremy L Thompson     return CeedError(ceed, 1, "Cannot assemble QFunction without active inputs "
6151d102b48SJeremy L Thompson                      "and outputs");
6161d102b48SJeremy L Thompson   // LCOV_EXCL_STOP
6171d102b48SJeremy L Thompson 
6181d102b48SJeremy L Thompson   // Setup lvec
6191d102b48SJeremy L Thompson   ierr = CeedVectorCreate(ceed, nblks*blksize*Q*numactivein*numactiveout,
6201d102b48SJeremy L Thompson                           &lvec); CeedChk(ierr);
6211d102b48SJeremy L Thompson   ierr = CeedVectorGetArray(lvec, CEED_MEM_HOST, &a); CeedChk(ierr);
6221d102b48SJeremy L Thompson 
6231d102b48SJeremy L Thompson   // Create output restriction
6241d102b48SJeremy L Thompson   ierr = CeedElemRestrictionCreateIdentity(ceed, numelements, Q,
6257f823360Sjeremylt          numelements*Q, numactivein*numactiveout, rstr); CeedChk(ierr);
6261d102b48SJeremy L Thompson   // Create assembled vector
6271d102b48SJeremy L Thompson   ierr = CeedVectorCreate(ceed, numelements*Q*numactivein*numactiveout,
6281d102b48SJeremy L Thompson                           assembled); CeedChk(ierr);
6291d102b48SJeremy L Thompson 
6301d102b48SJeremy L Thompson   // Loop through elements
6311d102b48SJeremy L Thompson   for (CeedInt e=0; e<nblks*blksize; e+=blksize) {
6321d102b48SJeremy L Thompson     // Input basis apply
6331d102b48SJeremy L Thompson     ierr = CeedOperatorInputBasis_Opt(e, Q, qfinputfields, opinputfields,
6341d102b48SJeremy L Thompson                                       numinputfields, blksize, NULL, true,
6351d102b48SJeremy L Thompson                                       impl, request); CeedChk(ierr);
6361d102b48SJeremy L Thompson 
6371d102b48SJeremy L Thompson     // Assemble QFunction
6381d102b48SJeremy L Thompson     for (CeedInt in=0; in<numactivein; in++) {
6391d102b48SJeremy L Thompson       // Set Inputs
64042ea3801Sjeremylt       ierr = CeedVectorSetValue(activein[in], 1.0); CeedChk(ierr);
64142ea3801Sjeremylt       if (numactivein > 1) {
64242ea3801Sjeremylt         ierr = CeedVectorSetValue(activein[(in+numactivein-1)%numactivein],
64342ea3801Sjeremylt                                   0.0); CeedChk(ierr);
64442ea3801Sjeremylt       }
6451d102b48SJeremy L Thompson       // Set Outputs
6461d102b48SJeremy L Thompson       for (CeedInt out=0; out<numoutputfields; out++) {
6471d102b48SJeremy L Thompson         // Get output vector
6481d102b48SJeremy L Thompson         ierr = CeedOperatorFieldGetVector(opoutputfields[out], &vec);
6491d102b48SJeremy L Thompson         CeedChk(ierr);
6501d102b48SJeremy L Thompson         // Check if active output
6511d102b48SJeremy L Thompson         if (vec == CEED_VECTOR_ACTIVE) {
6521d102b48SJeremy L Thompson           CeedVectorSetArray(impl->qvecsout[out], CEED_MEM_HOST,
6531d102b48SJeremy L Thompson                              CEED_USE_POINTER, a); CeedChk(ierr);
6541d102b48SJeremy L Thompson           ierr = CeedQFunctionFieldGetSize(qfoutputfields[out], &size);
6551d102b48SJeremy L Thompson           CeedChk(ierr);
6561d102b48SJeremy L Thompson           a += size*Q*blksize; // Advance the pointer by the size of the output
6571d102b48SJeremy L Thompson         }
6581d102b48SJeremy L Thompson       }
6591d102b48SJeremy L Thompson       // Apply QFunction
6601d102b48SJeremy L Thompson       ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout);
6611d102b48SJeremy L Thompson       CeedChk(ierr);
6621d102b48SJeremy L Thompson     }
6631d102b48SJeremy L Thompson   }
6641d102b48SJeremy L Thompson 
6651d102b48SJeremy L Thompson   // Un-set output Qvecs to prevent accidental overwrite of Assembled
6661d102b48SJeremy L Thompson   for (CeedInt out=0; out<numoutputfields; out++) {
6671d102b48SJeremy L Thompson     // Get output vector
6681d102b48SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opoutputfields[out], &vec);
6691d102b48SJeremy L Thompson     CeedChk(ierr);
6701d102b48SJeremy L Thompson     // Check if active output
6711d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
6721d102b48SJeremy L Thompson       CeedVectorSetArray(impl->qvecsout[out], CEED_MEM_HOST, CEED_COPY_VALUES,
6731d102b48SJeremy L Thompson                          NULL); CeedChk(ierr);
6741d102b48SJeremy L Thompson     }
6751d102b48SJeremy L Thompson   }
6761d102b48SJeremy L Thompson 
6771d102b48SJeremy L Thompson   // Restore input arrays
6781d102b48SJeremy L Thompson   ierr = CeedOperatorRestoreInputs_Opt(numinputfields, qfinputfields,
6791d102b48SJeremy L Thompson                                        opinputfields, true, impl);
6801d102b48SJeremy L Thompson   CeedChk(ierr);
6811d102b48SJeremy L Thompson 
6821d102b48SJeremy L Thompson   // Output blocked restriction
6831d102b48SJeremy L Thompson   ierr = CeedVectorRestoreArray(lvec, &a); CeedChk(ierr);
6841d102b48SJeremy L Thompson   ierr = CeedVectorSetValue(*assembled, 0.0); CeedChk(ierr);
6851d102b48SJeremy L Thompson   CeedElemRestriction blkrstr;
6861d102b48SJeremy L Thompson   ierr = CeedElemRestrictionCreateBlocked(ceed, numelements, Q, blksize,
6871d102b48SJeremy L Thompson                                           numelements*Q,
6881d102b48SJeremy L Thompson                                           numactivein*numactiveout,
6891d102b48SJeremy L Thompson                                           CEED_MEM_HOST, CEED_COPY_VALUES,
6901d102b48SJeremy L Thompson                                           NULL, &blkrstr); CeedChk(ierr);
6911d102b48SJeremy L Thompson   ierr = CeedElemRestrictionApply(blkrstr, CEED_TRANSPOSE, CEED_NOTRANSPOSE,
6921d102b48SJeremy L Thompson                                   lvec, *assembled, request); CeedChk(ierr);
6931d102b48SJeremy L Thompson 
6941d102b48SJeremy L Thompson   // Cleanup
69542ea3801Sjeremylt   for (CeedInt i=0; i<numactivein; i++) {
69642ea3801Sjeremylt     ierr = CeedVectorDestroy(&activein[i]); CeedChk(ierr);
69742ea3801Sjeremylt   }
6981d102b48SJeremy L Thompson   ierr = CeedFree(&activein); CeedChk(ierr);
6991d102b48SJeremy L Thompson   ierr = CeedVectorDestroy(&lvec); CeedChk(ierr);
7001d102b48SJeremy L Thompson   ierr = CeedElemRestrictionDestroy(&blkrstr); CeedChk(ierr);
7011d102b48SJeremy L Thompson 
7021d102b48SJeremy L Thompson   return 0;
70389c6efa4Sjeremylt }
70489c6efa4Sjeremylt 
705f10650afSjeremylt //------------------------------------------------------------------------------
706f10650afSjeremylt // Operator Destroy
707f10650afSjeremylt //------------------------------------------------------------------------------
708f10650afSjeremylt static int CeedOperatorDestroy_Opt(CeedOperator op) {
709f10650afSjeremylt   int ierr;
710f10650afSjeremylt   CeedOperator_Opt *impl;
711f10650afSjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
712f10650afSjeremylt 
713f10650afSjeremylt   for (CeedInt i=0; i<impl->numein+impl->numeout; i++) {
714f10650afSjeremylt     ierr = CeedElemRestrictionDestroy(&impl->blkrestr[i]); CeedChk(ierr);
715f10650afSjeremylt     ierr = CeedVectorDestroy(&impl->evecs[i]); CeedChk(ierr);
716f10650afSjeremylt   }
717f10650afSjeremylt   ierr = CeedFree(&impl->blkrestr); CeedChk(ierr);
718f10650afSjeremylt   ierr = CeedFree(&impl->evecs); CeedChk(ierr);
719f10650afSjeremylt   ierr = CeedFree(&impl->edata); CeedChk(ierr);
720f10650afSjeremylt   ierr = CeedFree(&impl->inputstate); CeedChk(ierr);
721f10650afSjeremylt 
722f10650afSjeremylt   for (CeedInt i=0; i<impl->numein; i++) {
723f10650afSjeremylt     ierr = CeedVectorDestroy(&impl->evecsin[i]); CeedChk(ierr);
724f10650afSjeremylt     ierr = CeedVectorDestroy(&impl->qvecsin[i]); CeedChk(ierr);
725f10650afSjeremylt   }
726f10650afSjeremylt   ierr = CeedFree(&impl->evecsin); CeedChk(ierr);
727f10650afSjeremylt   ierr = CeedFree(&impl->qvecsin); CeedChk(ierr);
728f10650afSjeremylt 
729f10650afSjeremylt   for (CeedInt i=0; i<impl->numeout; i++) {
730f10650afSjeremylt     ierr = CeedVectorDestroy(&impl->evecsout[i]); CeedChk(ierr);
731f10650afSjeremylt     ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChk(ierr);
732f10650afSjeremylt   }
733f10650afSjeremylt   ierr = CeedFree(&impl->evecsout); CeedChk(ierr);
734f10650afSjeremylt   ierr = CeedFree(&impl->qvecsout); CeedChk(ierr);
735f10650afSjeremylt 
736f10650afSjeremylt   ierr = CeedFree(&impl); CeedChk(ierr);
737f10650afSjeremylt   return 0;
738f10650afSjeremylt }
739f10650afSjeremylt 
740f10650afSjeremylt //------------------------------------------------------------------------------
741f10650afSjeremylt // Operator Create
742f10650afSjeremylt //------------------------------------------------------------------------------
74389c6efa4Sjeremylt int CeedOperatorCreate_Opt(CeedOperator op) {
74489c6efa4Sjeremylt   int ierr;
74589c6efa4Sjeremylt   Ceed ceed;
74689c6efa4Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
74789c6efa4Sjeremylt   Ceed_Opt *ceedimpl;
74889c6efa4Sjeremylt   ierr = CeedGetData(ceed, (void *)&ceedimpl); CeedChk(ierr);
74989c6efa4Sjeremylt   CeedInt blksize = ceedimpl->blksize;
75089c6efa4Sjeremylt   CeedOperator_Opt *impl;
75189c6efa4Sjeremylt 
75289c6efa4Sjeremylt   ierr = CeedCalloc(1, &impl); CeedChk(ierr);
75389c6efa4Sjeremylt   ierr = CeedOperatorSetData(op, (void *)&impl); CeedChk(ierr);
75489c6efa4Sjeremylt 
755*82946b17Sjeremylt   if (blksize != 1 && blksize != 8)
756*82946b17Sjeremylt     // LCOV_EXCL_START
757*82946b17Sjeremylt     return CeedError(ceed, 1, "Opt backend cannot use blocksize: %d", blksize);
758*82946b17Sjeremylt   // LCOV_EXCL_STOP
759*82946b17Sjeremylt 
7601d102b48SJeremy L Thompson   ierr = CeedSetBackendFunction(ceed, "Operator", op, "AssembleLinearQFunction",
7611d102b48SJeremy L Thompson                                 CeedOperatorAssembleLinearQFunction_Opt);
7621d102b48SJeremy L Thompson   CeedChk(ierr);
763cae8b89aSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "ApplyAdd",
7641d102b48SJeremy L Thompson                                 CeedOperatorApply_Opt); CeedChk(ierr);
76589c6efa4Sjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Destroy",
76689c6efa4Sjeremylt                                 CeedOperatorDestroy_Opt); CeedChk(ierr);
76789c6efa4Sjeremylt   return 0;
76889c6efa4Sjeremylt }
769f10650afSjeremylt //------------------------------------------------------------------------------
770